<?xml version="1.0" encoding="UTF-8"?>
<!DOCTYPE article
  PUBLIC "-//NLM//DTD Journal Publishing DTD v3.0 20080202//EN" "http://dtd.nlm.nih.gov/publishing/3.0/journalpublishing3.dtd">
<article xmlns:mml="http://www.w3.org/1998/Math/MathML" xmlns:xlink="http://www.w3.org/1999/xlink" article-type="research-article" dtd-version="3.0" xml:lang="en">
  <front>
    <journal-meta>
      <journal-id journal-id-type="nlm-ta">PLoS ONE</journal-id>
      <journal-id journal-id-type="publisher-id">plos</journal-id>
      <journal-id journal-id-type="pmc">plosone</journal-id>
      <journal-title-group>
        <journal-title>PLoS ONE</journal-title>
      </journal-title-group>
      <issn pub-type="epub">1932-6203</issn>
      <publisher>
        <publisher-name>Public Library of Science</publisher-name>
        <publisher-loc>San Francisco, USA</publisher-loc>
      </publisher>
    </journal-meta>
    <article-meta>
      <article-id pub-id-type="publisher-id">PONE-D-12-23793</article-id>
      <article-id pub-id-type="doi">10.1371/journal.pone.0055864</article-id>
      <article-categories>
        <subj-group subj-group-type="heading">
          <subject>Research Article</subject>
        </subj-group>
        <subj-group subj-group-type="Discipline-v2">
          <subject>Biology</subject>
          <subj-group>
            <subject>Biotechnology</subject>
          </subj-group>
          <subj-group>
            <subject>Computational biology</subject>
            <subj-group>
              <subject>Genomics</subject>
              <subj-group>
                <subject>Genome analysis tools</subject>
                <subject>Genome complexity</subject>
                <subject>Genome sequencing</subject>
              </subj-group>
            </subj-group>
          </subj-group>
          <subj-group>
            <subject>Genomics</subject>
            <subj-group>
              <subject>Genome analysis tools</subject>
              <subject>Genome complexity</subject>
              <subject>Genome sequencing</subject>
            </subj-group>
          </subj-group>
          <subj-group>
            <subject>Plant science</subject>
            <subj-group>
              <subject>Plant biotechnology</subject>
              <subj-group>
                <subject>Plant genomics</subject>
              </subj-group>
            </subj-group>
            <subj-group>
              <subject>Plant genomics</subject>
            </subj-group>
          </subj-group>
        </subj-group>
        <subj-group subj-group-type="Discipline-v2">
          <subject>Engineering</subject>
          <subj-group>
            <subject>Systems engineering</subject>
            <subj-group>
              <subject>Technology assessment</subject>
            </subj-group>
          </subj-group>
          <subj-group>
            <subject>Bioengineering</subject>
            <subj-group>
              <subject>Medical devices</subject>
            </subj-group>
          </subj-group>
        </subj-group>
        <subj-group subj-group-type="Discipline">
          <subject>Genetics and Genomics</subject>
          <subject>Plant Biology</subject>
          <subject>Biotechnology</subject>
          <subject>Computational Biology</subject>
        </subj-group>
      </article-categories>
      <title-group>
        <article-title>Rapid Genome Mapping in Nanochannel Arrays for Highly Complete and Accurate <italic>De Novo</italic> Sequence Assembly of the Complex <italic>Aegilops tauschii</italic> Genome</article-title>
        <alt-title alt-title-type="running-head">Rapid Genome Mapping in Nanochannel Arrays</alt-title>
      </title-group>
      <contrib-group>
        <contrib contrib-type="author" xlink:type="simple">
          <name name-style="western">
            <surname>Hastie</surname>
            <given-names>Alex R.</given-names>
          </name>
          <xref ref-type="aff" rid="aff1">
            <sup>1</sup>
          </xref>
        </contrib>
        <contrib contrib-type="author" xlink:type="simple">
          <name name-style="western">
            <surname>Dong</surname>
            <given-names>Lingli</given-names>
          </name>
          <xref ref-type="aff" rid="aff2">
            <sup>2</sup>
          </xref>
          <xref ref-type="aff" rid="aff3">
            <sup>3</sup>
          </xref>
        </contrib>
        <contrib contrib-type="author" xlink:type="simple">
          <name name-style="western">
            <surname>Smith</surname>
            <given-names>Alexis</given-names>
          </name>
          <xref ref-type="aff" rid="aff1">
            <sup>1</sup>
          </xref>
        </contrib>
        <contrib contrib-type="author" xlink:type="simple">
          <name name-style="western">
            <surname>Finklestein</surname>
            <given-names>Jeff</given-names>
          </name>
          <xref ref-type="aff" rid="aff1">
            <sup>1</sup>
          </xref>
        </contrib>
        <contrib contrib-type="author" xlink:type="simple">
          <name name-style="western">
            <surname>Lam</surname>
            <given-names>Ernest T.</given-names>
          </name>
          <xref ref-type="aff" rid="aff1">
            <sup>1</sup>
          </xref>
        </contrib>
        <contrib contrib-type="author" xlink:type="simple">
          <name name-style="western">
            <surname>Huo</surname>
            <given-names>Naxin</given-names>
          </name>
          <xref ref-type="aff" rid="aff2">
            <sup>2</sup>
          </xref>
          <xref ref-type="aff" rid="aff3">
            <sup>3</sup>
          </xref>
        </contrib>
        <contrib contrib-type="author" xlink:type="simple">
          <name name-style="western">
            <surname>Cao</surname>
            <given-names>Han</given-names>
          </name>
          <xref ref-type="aff" rid="aff1">
            <sup>1</sup>
          </xref>
        </contrib>
        <contrib contrib-type="author" xlink:type="simple">
          <name name-style="western">
            <surname>Kwok</surname>
            <given-names>Pui-Yan</given-names>
          </name>
          <xref ref-type="aff" rid="aff4">
            <sup>4</sup>
          </xref>
        </contrib>
        <contrib contrib-type="author" xlink:type="simple">
          <name name-style="western">
            <surname>Deal</surname>
            <given-names>Karin R.</given-names>
          </name>
          <xref ref-type="aff" rid="aff3">
            <sup>3</sup>
          </xref>
        </contrib>
        <contrib contrib-type="author" xlink:type="simple">
          <name name-style="western">
            <surname>Dvorak</surname>
            <given-names>Jan</given-names>
          </name>
          <xref ref-type="aff" rid="aff3">
            <sup>3</sup>
          </xref>
        </contrib>
        <contrib contrib-type="author" xlink:type="simple">
          <name name-style="western">
            <surname>Luo</surname>
            <given-names>Ming-Cheng</given-names>
          </name>
          <xref ref-type="aff" rid="aff3">
            <sup>3</sup>
          </xref>
        </contrib>
        <contrib contrib-type="author" xlink:type="simple">
          <name name-style="western">
            <surname>Gu</surname>
            <given-names>Yong</given-names>
          </name>
          <xref ref-type="aff" rid="aff2">
            <sup>2</sup>
          </xref>
          <xref ref-type="aff" rid="aff3">
            <sup>3</sup>
          </xref>
          <xref ref-type="corresp" rid="cor1">
            <sup>*</sup>
          </xref>
        </contrib>
        <contrib contrib-type="author" xlink:type="simple">
          <name name-style="western">
            <surname>Xiao</surname>
            <given-names>Ming</given-names>
          </name>
          <xref ref-type="aff" rid="aff1">
            <sup>1</sup>
          </xref>
          <xref ref-type="corresp" rid="cor1">
            <sup>*</sup>
          </xref>
          <xref ref-type="fn" rid="fn1">
            <sup>¤</sup>
          </xref>
        </contrib>
      </contrib-group>
      <aff id="aff1">
        <label>1</label>
        <addr-line>BioNano Genomics, San Diego, California, United States of America</addr-line>
      </aff>
      <aff id="aff2">
        <label>2</label>
        <addr-line>Genomics and Gene Discovery Research Unit, United States Department of Agriculture - Agricultural Research Service, Albany, California, United States of America</addr-line>
      </aff>
      <aff id="aff3">
        <label>3</label>
        <addr-line>Department of Plant Sciences, University of California Davis, Davis, California, United States of America</addr-line>
      </aff>
      <aff id="aff4">
        <label>4</label>
        <addr-line>Institute for Human Genetics, University of California San Francisco, San Francisco, California, United States of America</addr-line>
      </aff>
      <contrib-group>
        <contrib contrib-type="editor" xlink:type="simple">
          <name name-style="western">
            <surname>Nelson</surname>
            <given-names>James C.</given-names>
          </name>
          <role>Editor</role>
          <xref ref-type="aff" rid="edit1"/>
        </contrib>
      </contrib-group>
      <aff id="edit1">
        <addr-line>Kansas State University, United States of America</addr-line>
      </aff>
      <author-notes>
        <corresp id="cor1">* E-mail: <email xlink:type="simple">ming.xiao@drexel.edu</email> (MX); <email xlink:type="simple">yong.gu@ars.usda.gov</email> (YG)</corresp>
        <fn fn-type="conflict">
          <p>The authors have the following interests. Alex R Hastie, Alexis Smith, Jeff Finklestein, Ernest Lam, Han Cao and Ming Xiao were employees of BioNano Genomics at the time of the study, and they own company stock options. Alex R Hastie, Han Cao, and Ming Xiao are inventors of several patents owned by BioNano Genomics. The authors have a patent application with inventors Ming Xiao and Alex R Hastie (application number 13/606,819) relating to this work, “PHYSICAL MAP CONSTRUCTION OF WHOLE GENOME AND POOLED CLONE MAPPING IN NANOCHANNEL ARRAY. The technology platform (BioNano Genomics IrysTM) described in this paper was developed by BioNano Genomics. There are no further patents, products in development or marketed products to declare. This does not alter the authors' adherence to all the PLOS ONE policies on sharing data and materials, as detailed online in the guide for authors.</p>
        </fn>
        <fn fn-type="con">
          <p>Conceived and designed the experiments: AH HC MCL YG MX. Performed the experiments: AH LD AS JF NH KRD. Analyzed the data: AH LD JD MCL YG MX. Contributed reagents/materials/analysis tools: ETL PYK. Wrote the paper: AH LD ETL KRD JD HC MCL YG MX.</p>
        </fn>
        <fn id="fn1" fn-type="current-aff">
          <label>¤</label>
          <p>Current address: Drexel University, Philadelphia, Pennsylvania, United States of America</p>
        </fn>
      </author-notes>
      <pub-date pub-type="collection">
        <year>2013</year>
      </pub-date>
      <pub-date pub-type="epub">
        <day>6</day>
        <month>2</month>
        <year>2013</year>
      </pub-date>
      <volume>8</volume>
      <issue>2</issue>
      <elocation-id>e55864</elocation-id>
      <history>
        <date date-type="received">
          <day>9</day>
          <month>8</month>
          <year>2012</year>
        </date>
        <date date-type="accepted">
          <day>3</day>
          <month>1</month>
          <year>2013</year>
        </date>
      </history>
      <permissions>
        <copyright-year>2013</copyright-year>
        <copyright-holder>Hastie et al</copyright-holder>
        <license xlink:type="simple">
          <license-p>This is an open-access article distributed under the terms of the Creative Commons Attribution License, which permits unrestricted use, distribution, and reproduction in any medium, provided the original author and source are credited.</license-p>
        </license>
      </permissions>
      <abstract>
        <p>Next-generation sequencing (NGS) technologies have enabled high-throughput and low-cost generation of sequence data; however, <italic>de novo</italic> genome assembly remains a great challenge, particularly for large genomes. NGS short reads are often insufficient to create large contigs that span repeat sequences and to facilitate unambiguous assembly. Plant genomes are notorious for containing high quantities of repetitive elements, which combined with huge genome sizes, makes accurate assembly of these large and complex genomes intractable thus far. Using two-color genome mapping of tiling bacterial artificial chromosomes (BAC) clones on nanochannel arrays, we completed high-confidence assembly of a 2.1-Mb, highly repetitive region in the large and complex genome of <italic>Aegilops tauschii</italic>, the D-genome donor of hexaploid wheat (<italic>Triticum aestivum</italic>). Genome mapping is based on direct visualization of sequence motifs on single DNA molecules hundreds of kilobases in length. With the genome map as a scaffold, we anchored unplaced sequence contigs, validated the initial draft assembly, and resolved instances of misassembly, some involving contigs &lt;2 kb long, to dramatically improve the assembly from 75% to 95% complete.</p>
      </abstract>
      <funding-group>
        <funding-statement>This research is supported in part by US National Institutes of Health (NIH) award to P.-Y.K. and M.X. (R01 HG005946). The funders had no role in study design, data collection and analysis, decision to publish, or preparation of the manuscript. No additional external funding was received for this study.</funding-statement>
      </funding-group>
      <counts>
        <page-count count="10"/>
      </counts>
    </article-meta>
  </front>
  <body>
    <sec id="s1">
      <title>Introduction</title>
      <p>Accurate <italic>de novo</italic> assembly of sequence reads represents the weak link in genome projects despite advances in high-throughput sequencing <xref ref-type="bibr" rid="pone.0055864-Blakesley1">[1]</xref>, <xref ref-type="bibr" rid="pone.0055864-Chain1">[2]</xref> . There are two general steps in genome sequence assembly: generation of sequence contigs and scaffolds, and their anchoring on genome-wide, lower resolution maps. NGS platforms generate sequence reads ranging from 25 to more than 500 bases <xref ref-type="bibr" rid="pone.0055864-Lee1">[3]</xref>, while reads of up to 1000 bases can be obtained by Sanger sequencing with high accuracy. NGS reads are often too short for unambiguous assembly. Paired-end reads can bridge contigs into scaffolds, but there are often gaps within the scaffolds. To order contigs and scaffolds, high-resolution genomic maps from an independent technology platform are needed. They may be of chromosomal scale, i.e., genetic maps, or regional scale, i.e., contigs of bacterial artificial chromosomes (BACs) or fosmids <xref ref-type="bibr" rid="pone.0055864-Green1">[4]</xref>. Contigs and scaffolds may be difficult to map if they are too short compared to the map resolution. For example, maps may have a resolution of 50–150 kb while many contigs and scaffolds may only span a few kilobases. Additionally, there are errors in the contigs and scaffolds themselves, often due to misassembly of repeat sequences. Typical medium to large genomes contain 40–85% repetitive sequences <xref ref-type="bibr" rid="pone.0055864-McPherson1">[5]</xref>–<xref ref-type="bibr" rid="pone.0055864-Zuccolo1">[8]</xref>, dramatically hindering effective <italic>de novo</italic> sequence assembly.</p>
      <p>Genome finishing has relied on guidance of a physical map for large and complex genomes, including human, arabidopsis <xref ref-type="bibr" rid="pone.0055864-Initiative1">[9]</xref>, rice <xref ref-type="bibr" rid="pone.0055864-Project1">[10]</xref> and maize <xref ref-type="bibr" rid="pone.0055864-Zhou1">[11]</xref>, <xref ref-type="bibr" rid="pone.0055864-Schnable1">[12]</xref>. BAC-based restriction fragment physical mapping of complex genomes is fairly robust because even in the presence of interspersed repeat sequences along the BAC inserts (typically 100–220 kb long) a unique pattern of restriction fragments is generated. The state of the art technologies for physical map construction include SNaPshot <xref ref-type="bibr" rid="pone.0055864-Luo1">13</xref>,<xref ref-type="bibr" rid="pone.0055864-Paux1">14</xref>, whole-genome profiling <xref ref-type="bibr" rid="pone.0055864-Philippe1">[15]</xref>, <xref ref-type="bibr" rid="pone.0055864-vanOeveren1">[16]</xref>, optical mapping <xref ref-type="bibr" rid="pone.0055864-Schwartz1">[17]</xref>, <xref ref-type="bibr" rid="pone.0055864-Teague1">[18]</xref> and genome mapping <xref ref-type="bibr" rid="pone.0055864-Lam1">[19]</xref>. SNaPshot is a restriction fingerprinting method which uses one or more restriction enzymes and fluorescent labels followed by separation of fragments by capillary electrophoresis. SNaPshot has been used for physical mapping of wheat and other genomes <xref ref-type="bibr" rid="pone.0055864-Paux1">[14]</xref>, <xref ref-type="bibr" rid="pone.0055864-Mun1">[20]</xref>. Optical mapping provides an additional layer of information by retaining the physical order of restriction sites along DNA molecules immobilized on a surface <xref ref-type="bibr" rid="pone.0055864-Teague1">[18]</xref>. It has been applied to the maize and the rice genome <xref ref-type="bibr" rid="pone.0055864-Zhou1">[11]</xref>, <xref ref-type="bibr" rid="pone.0055864-Zhou2">[21]</xref>. One can validate a sequence assembly by comparing <italic>in silico</italic> sequence motif maps to consensus optical maps <xref ref-type="bibr" rid="pone.0055864-Nagarajan1">[22]</xref>–<xref ref-type="bibr" rid="pone.0055864-Lin1">[25]</xref>. However, information density for optical maps is only about one site per 20 kb, and the technology is limited in utility by high error-rates, non-uniform DNA linearization, and low throughput. Therefore, a high-resolution (&lt;5 kb), DNA sequencing-independent mapping method that can overcome these constraints of optical mapping is much needed.</p>
      <p>Genome mapping on nanochannel arrays at the single-molecule level overcomes many of the limitations of preexisting technologies and has recently been described in depth <xref ref-type="bibr" rid="pone.0055864-Lam1">[19]</xref>. This technology uses nicking enzymes to create sequence-specific nicks that are subsequently labeled by a fluorescent nucleotide analog <xref ref-type="bibr" rid="pone.0055864-Xiao1">[26]</xref>. The nick-labeled DNA is stained with the intercalating dye YOYO-1, loaded onto the nanofluidic chip by an electric field, and imaged with a CCD camera. The DNA is linearized by confinement in a nanochannel array <xref ref-type="bibr" rid="pone.0055864-Das1">[27]</xref>, resulting in uniform linearization and allowing precise and accurate measurement of the distance between nick-labels on DNA molecules comprising a signature pattern. Also, the DNA loading and imaging cycle can be repeated many times in a completely automated fashion; data can be obtained at a high throughput of ∼5 Gb/hour. Genome mapping was previously used to map the 4.7-Mb, highly variable human MHC region. It was able to distinguish haplotype differences, identify a segmental duplication, and identify errors in the reference assembly <xref ref-type="bibr" rid="pone.0055864-Lam1">[19]</xref>.</p>
      <p>The tribe Triticeae includes wheat and the closely related genus Aegilops, the source of two of the three wheat genomes. Diploid Triticeae genomes range from &lt;4 to &gt;8 Gb, and approximately 90% of their genomes is comprised of repetitive sequences <xref ref-type="bibr" rid="pone.0055864-Dvorak1">[28]</xref>. For example, <italic>Ae. tauschii</italic> contains 91% repetitive elements, 2.5% known genes and 6% low-copy sequence <xref ref-type="bibr" rid="pone.0055864-Li1">[29]</xref>. Shotgun genome sequencing of several Triticeae species has been attempted, but the resulting assemblies remain incomplete and coverage is uneven, limiting analysis to mostly gene-rich regions <xref ref-type="bibr" rid="pone.0055864-Cassidy1">[30]</xref>–<xref ref-type="bibr" rid="pone.0055864-Brenchley1">[33]</xref>. The current consensus is that these genomes can only be tackled with an ordered clone sequencing approach. This approach is being adopted by the International Wheat Genome Sequencing Consortium (IWGSC, <ext-link ext-link-type="uri" xlink:href="http://www.wheatgenome.org" xlink:type="simple">www.wheatgenome.org</ext-link>) and other genome sequencing projects involving species in this tribe. For ordering and selecting BAC clones, a physical map is constructed with SNaPshot fingerprinting <xref ref-type="bibr" rid="pone.0055864-Luo1">[13]</xref>, <xref ref-type="bibr" rid="pone.0055864-Paux1">[14]</xref> or whole-genome profiling <xref ref-type="bibr" rid="pone.0055864-Philippe1">[15]</xref>, <xref ref-type="bibr" rid="pone.0055864-vanOeveren1">[16]</xref>. A set of minimal tiling path (MTP) BAC clones is selected to maximize coverage while controlling for redundancy. The MTP BAC clones are sequenced as pools, and the sequences are scaffolded with the ultimate goal of creating a complete and high-fidelity <italic>de novo</italic> genome assembly.</p>
      <p>We have further extended genome mapping for use with a second nicking enzyme and a second label color. This new strategy greatly improves information density of genome mapping. We have applied it to a 2.1-Mb prolamin gene family region from the genome of <italic>Ae. tauschii</italic>, the D genome donor of hexaploid bread wheat. This region is rich in syntelogs, and its assembly is exceptionally challenging. With the improved genome mapping technology, we have constructed a high-resolution physical map and used it to correct the physical map of the region and used it to validate and correct the physical map generated by SNaPshot fingerprinting technology. We then used the genome map to facilitate <italic>de novo</italic> sequence assembly of the region by anchoring sequence scaffolds, validating correctly assembled regions and correcting inaccuracies in the scaffolds, and producing a highly accurate and complete sequence assembly.</p>
    </sec>
    <sec id="s2">
      <title>Results</title>
      <sec id="s2a">
        <title>Generation of a two-color genome map with two nicking enzymes</title>
        <p>We constructed a genome map using two nicking enzymes, Nt.BbvCI and Nt.BspQI, whose nick motifs were labeled with red and green dyes, respectively, across 27 BACs making up an MTP of a 2.1-Mb region containing the prolamin multigene family in the <italic>Ae. tauschii</italic> genome. <xref ref-type="fig" rid="pone-0055864-g001">Figure 1A</xref> shows the layout of the IrysChip. The YOYO-stained DNA was loaded into the port, unwound within the pillar structures, and linearized inside the 45 nm nanochannels (<xref ref-type="fig" rid="pone-0055864-g001">Figure 1B</xref>). After image processing, individual BAC molecules with red and green labels distributed at sequence-specific locations were compared and clustered into a pools with similar map patterns (<xref ref-type="fig" rid="pone-0055864-g001">Figure 1C</xref>, top). Density plots for the BAC clones were generated to determine the consensus peak locations (<xref ref-type="fig" rid="pone-0055864-g001">Figure 1C</xref>, bottom). These consensus maps of individual BAC clones were aligned based on overlaps of consensus maps of adjacent BACs (<xref ref-type="fig" rid="pone-0055864-g001">Figure 1D</xref>) to create a genome map of the entire region. The two-color labeling strategy resulted in an average information density of one label per 4.8 kb (437 labels in 2.1 Mb). Since each motif was marked by its own color, peaks of different motifs could be distinguished from each other even if their peaks were almost overlapping (arrow in <xref ref-type="fig" rid="pone-0055864-g001">Figure 1D</xref>). Peaks of the same motif (the same color) could be resolved when they were at least ∼1.5 kb apart, as previously established <xref ref-type="bibr" rid="pone.0055864-Lam1">[19]</xref>. Taking advantage of the combination of long molecule lengths (∼140 kb average), high-resolution, accurate length measurement, and multiple sequence motifs, we generated a high-quality genome map of the 2.1-Mb region for scaffold assembly.</p>
        <fig id="pone-0055864-g001" position="float">
          <object-id pub-id-type="doi">10.1371/journal.pone.0055864.g001</object-id>
          <label>Figure 1</label>
          <caption>
            <title>Two-color genome mapping with two enzymes.</title>
            <p>A. The DNA backbone is stained with YOYO-1 and loaded into the port of a nanochannel array chip. The DNA molecules are introduced into the region with pillars and micron-scale relaxation channels by an electric field where they unwind and linearize. Finally, they are moved into the 45 nm nanochannels, where they stretch uniformly to 85% of the length of perfectly linear B-DNA. B. Linearized BAC DNA molecules in nanochannels. The DNA molecule is stained with YOYO-1, and Nt.BspQI and Nt.BbvCI nicks are labeled with green and red dyes, respectively. C. Molecule length and nick locations are extracted from the images by custom image-analysis software. By clustering individual molecules with high similarity of green label patterns, distinct patterns are extracted (top panel). The locations of the red labels are then overlaid on the green label patterns (middle pattern). A histogram plot of the above clusters is shown in the bottom panel. The peaks represent the location of each sequence motif (GCTCTTC and CCTCAGC) along the linearized DNA molecules. D. Consensus maps for individual BAC clones are shown. Consensus maps are combined by using overlapping patterns, and the final genome map is shown at the bottom. E. The clone map from genome mapping is shown at the top and the full genome map as a grey bar with Nt.BspQI and Nt.BbvCI motif locations in green and red. Below the genome map, in blue, is the physical map from SNaPshot fingerprinting. The total length of the genome map is 2.1 Mb.</p>
          </caption>
          <graphic mimetype="image" xlink:href="info:doi/10.1371/journal.pone.0055864.g001" position="float" xlink:type="simple"/>
        </fig>
      </sec>
      <sec id="s2b">
        <title>Validation of SNaPshot-based physical map by genome mapping</title>
        <p>The entire genome map for the 2.1-Mb region is shown in <xref ref-type="fig" rid="pone-0055864-g001">Figure 1E</xref>. The predicted positions of BAC clones based on the genome map are above the map. There were no overlapping nick motifs between clones 10 and 13 and between clones 13 and 9. However, sequence overlap between the three clones confirmed tiling and allowed for assembly of the region. BAC clone positions in the 2.1-Mb region were also determined by SNaPshot fingerprinting and shown for comparison (bottom of <xref ref-type="fig" rid="pone-0055864-g001">Figure 1E</xref>). The genome map-based and the SNaPshot-base clone positions and overlaps were concordant for most clones. SNaPshot does not provide precise clone sizes or overlap lengths because it generates an incomplete collection of short and unordered fragments for a BAC. Therefore, the clone positions and boundaries are much more accurate on the genome map than on a physical map generated by SNaPshot fingerprinting.</p>
        <p>Of the 27 BAC clones analyzed by genome mapping, 25 had the same placement as in the SNaPshot map. Genome mapping suggested that clones 5 and 20 did not belong in this 2.1-Mb region. Clone 5 was reanalyzed by SNaPshot and found to differ from the clone originally used for MTP construction, implicating contamination during clone picking for sequencing. SNaPshot fingerprinting of clone 20 was consistent with the original SNaPshot result. However, upon reevaluation, the clone was found to have a weak score for placement on the SNaPshot physical map and, in agreement with the genome map result, was most likely misplaced during BAC contig assembly.</p>
      </sec>
      <sec id="s2c">
        <title>Genome map-based assessment of <italic>de novo</italic> sequence assembly</title>
        <p>The minimal tiling path for this region contained 23 BAC clones, and they were initially sequenced using the 454 platform. For constructing a contiguous physical map using genome mapping, four additional clones were also analyzed but were not sequenced. The sequence reads were assembled and scaffolded with 3-kb paired-end reads. Several clones (5, 7, 9, 14 and 20) were not included in the paired-end sequencing because they were selected later for additional coverage. BAC end-sequences and genetic markers were also used to improve scaffolding, and this resulted in sequence scaffolds that were ordered along the genetic map, covering 2.1 Mb. A total of 254 contigs were joined primarily by paired-end reads into a scaffold. Several small scaffolds could not be placed within the main scaffold. <xref ref-type="fig" rid="pone-0055864-g002">Figure 2</xref> shows the <italic>in silico</italic> map for Nt.BspQI and Nt.BbvCI motifs in the main scaffold. Regions numbered 1 to 13 (<xref ref-type="fig" rid="pone-0055864-g002">Figure 2</xref>) showed discrepancies. About 75% of the genome map could be aligned with sequence scaffolds (region 12 was not included in this calculation; see Results section “Genome map guided <italic>de novo</italic> sequence assembly” for explanation).</p>
        <fig id="pone-0055864-g002" position="float">
          <object-id pub-id-type="doi">10.1371/journal.pone.0055864.g002</object-id>
          <label>Figure 2</label>
          <caption>
            <title>Comparison of the sequence assembly scaffold to the genome map.</title>
            <p>The sequence assembly is shown with dark grey boxes representing sequence contigs. Contigs were bridged by paired-end sequence reads. The genome map is represented by light grey boxes. Shaded boxes around regions of both maps denote regions where the sequence assembly matches the genome map well. Regions where there are significant discrepancies are numbered and discussed in the results section. The two gaps in the genome map are denoted with asterisks.</p>
          </caption>
          <graphic mimetype="image" xlink:href="info:doi/10.1371/journal.pone.0055864.g002" position="float" xlink:type="simple"/>
        </fig>
      </sec>
      <sec id="s2d">
        <title>Genome map guided sequence assembly: removing incorrectly scaffolded sequence contigs in <italic>de novo</italic> assemblies</title>
        <p>The most common and easily identified type of discrepancy between the genome map and the scaffold came from regions where sequence contigs were incorrectly inserted into the scaffold. In these cases, motif sites were present in the scaffold but absent in the genome map while the flanking regions matched well (as in <xref ref-type="fig" rid="pone-0055864-g002">Figure 2</xref>, #2, 4, 5, 8 and 9). This type of discrepancy can also be identified based on length measurements alone, as in <xref ref-type="fig" rid="pone-0055864-g002">Figure 2</xref>, #3. <xref ref-type="fig" rid="pone-0055864-g003">Figure 3</xref> shows the scaffolding results generated with gsAssembler using the paired-end reads. Two of the three contigs with red bars beneath them contained Nt.BspQI sites that were absent in the genome map and the length of the three together is equal to the total discrepancy between the genome map and the assembly. Additionally, the contigs had weak paired-end coverage; therefore, they were likely incorrectly scaffolded and were removed from the assembly. The resulting scaffold mapped correctly to the consensus genome map.</p>
        <fig id="pone-0055864-g003" position="float">
          <object-id pub-id-type="doi">10.1371/journal.pone.0055864.g003</object-id>
          <label>Figure 3</label>
          <caption>
            <title>Deletion of incorrect contigs in genome map-guided <italic>de novo</italic> sequence assembly.</title>
            <p>The original assembly contained two Nt.BspQI sites and ∼8 kb of sequence that were absent from the genome map. The top image is output from gsAssembler and shows the scaffolding of contigs using paired-end reads. The green line represents the sequence coverage for each region. Paired-end reads are represented by pink (high coverage) and aqua (low coverage) carrots (&amp;squ;). The three contigs with red bars beneath them contain the extra sequence motifs and total sequence consistent with the predicted incorrect scaffold. They also contain weak paired-end data indicating that the contigs are misplaced. The bottom line shows the sequence assembly after deletion of the three contigs with red bars.</p>
          </caption>
          <graphic mimetype="image" xlink:href="info:doi/10.1371/journal.pone.0055864.g003" position="float" xlink:type="simple"/>
        </fig>
      </sec>
      <sec id="s2e">
        <title>Genome map guided sequence assembly: filling gaps in <italic>de novo</italic> assemblies</title>
        <p>A number of the inconsistencies between the genome map and the scaffold were due to gaps in the sequence scaffold (<xref ref-type="fig" rid="pone-0055864-g002">Figure 2</xref>, #1, 6, 7, 10). The largest gap was 85 kb (<xref ref-type="fig" rid="pone-0055864-g002">Figure 2</xref>, #10; <xref ref-type="fig" rid="pone-0055864-g004">Figure 4</xref>), which belonged to a region with incomplete paired-end information (BACs 7 and 9 did not have paired-end reads). In order to fill this gap, we used the genome map directly as a scaffold. We searched for the missing sequence in unplaced sequence contigs and scaffolds from the original sequence assembly. An 85-kb scaffold that did not contain a BAC end sequence matched the genome map in the gap; we placed this scaffold in the gap denoted with a triangle in the original scaffold (<xref ref-type="fig" rid="pone-0055864-g004">Figure 4</xref>). The final assembly is shown in the bottom line with the genome map guided insertion marked with a box in the middle. The right junction of the insertion was confirmed by PCR (data not shown).</p>
        <fig id="pone-0055864-g004" position="float">
          <object-id pub-id-type="doi">10.1371/journal.pone.0055864.g004</object-id>
          <label>Figure 4</label>
          <caption>
            <title>Insertion of unassembled contigs into gaps in genome map-guided <italic>de novo</italic> sequence assembly.</title>
            <p>The top line shows the <italic>in silico</italic> map for the original sequence assembly; the genome map is shown in the middle. An 85-kb segment of DNA was missing from the sequence assembly. The corresponding DNA sequence was added to the assembly by anchoring a previously unplaced scaffold on the genome map. The resulting corrected assembly is shown in the bottom row.</p>
          </caption>
          <graphic mimetype="image" xlink:href="info:doi/10.1371/journal.pone.0055864.g004" position="float" xlink:type="simple"/>
        </fig>
      </sec>
      <sec id="s2f">
        <title>Genome map guided sequence assembly: identifying misassembled contigs in <italic>de novo</italic> assemblies</title>
        <p>In addition to scaffolding errors, we observed errors in the sequence contig assembly (<xref ref-type="fig" rid="pone-0055864-g002">Figure 2</xref>, #2 and 11). <xref ref-type="fig" rid="pone-0055864-g005">Figure 5</xref> represents a zoomed-in view of example 11 from <xref ref-type="fig" rid="pone-0055864-g001">Figure 1</xref>. The first ∼40 kb and the last ∼30 kb of this assembled region matched the genome map (boxed regions in <xref ref-type="fig" rid="pone-0055864-g005">Figure 5</xref>). This region was covered by a single, 79-kb sequence contig. Based on the contig sequence, the distance between the marked Nt.BbvCI site and the adjacent Nt.BspQI site was 20,763 bp. However, it was measured to be 17.7 kb in the genome map. The histogram for the consensus genome map is shown. It has robust peaks giving high confidence in the distance measurement. This region is primarily made up of LTR elements and is likely prone to assembly errors.</p>
        <fig id="pone-0055864-g005" position="float">
          <object-id pub-id-type="doi">10.1371/journal.pone.0055864.g005</object-id>
          <label>Figure 5</label>
          <caption>
            <title>Contig assembly error identification through genome map comparison.</title>
            <p>The top line represents the <italic>in silico</italic> map for the original sequence assembly, the majority of which is covered by a single sequence contig. The genome map matches on the left and right sides of the contig (shown with shaded boxes). ∼3 kb of sequence was incorrectly inserted into the contig during assembly.</p>
          </caption>
          <graphic mimetype="image" xlink:href="info:doi/10.1371/journal.pone.0055864.g005" position="float" xlink:type="simple"/>
        </fig>
      </sec>
      <sec id="s2g">
        <title>Genome map guided sequence assembly: identification and reassembly of satellite repeated sequences</title>
        <p>At position 1 in <xref ref-type="fig" rid="pone-0055864-g002">Figure 2</xref>, there was a 28-kb gap in the scaffold. The consensus genome map showed two blocks of DNA, each containing a very high density of Nt.BspQI sites beyond our optical resolution. <xref ref-type="fig" rid="pone-0055864-g006">Figure 6A</xref> shows the strip diagram for one of the clones that covers the region; each line represents a different molecule and the green spots are locations of the Nt.BspQI motif. The region marked “label repeats” (<xref ref-type="fig" rid="pone-0055864-g006">Figure 6</xref>) did not cluster into discrete peaks because not all close labels were detected as separate. This high-density nick motif commonly represents a tandem sequence repeat (satellite) and in this case, we found two blocks of potential tandem repeats. In order to appropriately assemble the sequence, we found the repeat sequence in an unplaced contig. The repeat was 670 bp long with 95–99% identity and contained an Nt.BspQI site. The original assembly had only a single occurrence of the 670-bp sequence cassette. By extracting all of the sequence reads that contained the cassette, incorporating additional reads generated by Sanger sequencing and separately reassembling the region using Consed, we were able to assemble two regions containing direct repeats of lengths 9.8 kb and 8 kb, in good agreement with the measured lengths of the high density Nt.BspQI regions. <xref ref-type="fig" rid="pone-0055864-g006">Figure 6B</xref> shows a pairwise alignment of the repeat structure with two blocks of direct tandem repeats which are inversely oriented with respect to each other. The genome map-independent sequence assembly, where it is consistent with the consensus genome map, is shown by shaded boxes in <xref ref-type="fig" rid="pone-0055864-g006">Figure 6C</xref>. The label repeats are noted in the genome map and absent in the scaffold. After insertion of 28 kb of reassembled contigs, the assembly was in good agreement with the genome map.</p>
        <fig id="pone-0055864-g006" position="float">
          <object-id pub-id-type="doi">10.1371/journal.pone.0055864.g006</object-id>
          <label>Figure 6</label>
          <caption>
            <title>Identification and assembly of repeat sequences in genome map-guided <italic>de novo</italic> sequence assembly.</title>
            <p>Panel A shows the strip diagram for one of the clones that covers the high density region, each line represents a different molecule and the spots are the location of green (Nt.BspQI) labels. Two high-density label regions are marked “label repeats,” and they do not cluster into discrete peaks. Panel B shows a pairwise alignment of the high-density region after reassembly based on the genome map. The alignment shows two blocks of direct repeats which are inverted with respect to one another. Panel C shows the original assembly on the top, the genome map in the middle and the final assembly on the bottom, containing the repeat sequence as predicted by the genome map.</p>
          </caption>
          <graphic mimetype="image" xlink:href="info:doi/10.1371/journal.pone.0055864.g006" position="float" xlink:type="simple"/>
        </fig>
      </sec>
      <sec id="s2h">
        <title>Genome map guided <italic>de novo</italic> sequence assembly</title>
        <p>After correcting the sequence scaffold (<xref ref-type="fig" rid="pone-0055864-g007">Figure 7</xref>, #1–6, 8–13), as discussed in the examples above, and excluding the gap at position 12 (which resulted from missing sequenced clone coverage), we have generated a high-confidence assembly that covered 95% of the consensus genome map. We were unable to fill the gap at position 7 because the region contained no sequence motifs; however, we measured this gap to be 12.79 kb. Position 12 is in the region that clone 5 should have filled, based on SNaPshot (<xref ref-type="fig" rid="pone-0055864-g001">Figure 1</xref>). The genome mapped clone 5 was deemed contaminated, but it was unclear if the correct or contaminated BAC was sequenced. Much of the sequence at this region (position 12) matched the clone 5 genome map, suggesting that the contaminated clone 5 was also sequenced and then misassembled into the scaffold. Since clone 27 actually overlaps with clones 4 and 7 but was not sequenced, properly filling this gap is not possible without additional sequencing. Discrepancy 13 was partially corrected but remains incomplete. Filling the remaining gaps will require sequencing of additional BACs that were not originally part of the MTP in this region.</p>
        <fig id="pone-0055864-g007" position="float">
          <object-id pub-id-type="doi">10.1371/journal.pone.0055864.g007</object-id>
          <label>Figure 7</label>
          <caption>
            <title>Comparison of the final genome map guided sequence assembly to the genome map.</title>
            <p>The final sequence assembly almost completely spans the genome map. Gaps in the genome map are denoted with asterisks.</p>
          </caption>
          <graphic mimetype="image" xlink:href="info:doi/10.1371/journal.pone.0055864.g007" position="float" xlink:type="simple"/>
        </fig>
      </sec>
    </sec>
    <sec id="s3">
      <title>Discussion</title>
      <p>Having a high-quality reference genome assembly for an organism is critical to the understanding of its biology and evolutionary relationship with other organisms. With a good reference sequence assembly, resequencing with NGS can be inexpensive and useful for studies of genetic variation. <italic>De novo</italic> assembly of additional genomes may allow more comprehensive surveys of variation, especially structural variation, as it relieves some of the biases associated with reference-dependent alignment and variant calling <xref ref-type="bibr" rid="pone.0055864-Li2">[34]</xref>. However, whole genome assembly is often difficult and cost-prohibitive. To date, very few high-quality assemblies are available for large and complex genomes <xref ref-type="bibr" rid="pone.0055864-Chain1">[2]</xref>, and additional <italic>de novo</italic> assemblies beyond a single reference assembly are uncommon. Generated using BioNano Genomics’ nanochannel array mapping technology, genome maps can guide sequence assembly. It fills a void in <italic>de novo</italic> assembly strategies by economically providing a high-resolution physical map that can be used for anchoring contigs and scaffolds (<xref ref-type="fig" rid="pone-0055864-g004">Figure 4</xref>) and capable of identifying misassembled contigs (<xref ref-type="fig" rid="pone-0055864-g003">Figures 3</xref>, <xref ref-type="fig" rid="pone-0055864-g005">5</xref> and <xref ref-type="fig" rid="pone-0055864-g006">6</xref>), thus dramatically improving the fidelity of the final assembly (<xref ref-type="fig" rid="pone-0055864-g007">Figure 7</xref>).</p>
      <p>Genome mapping has several advantages over SNaPshot fingerprinting for physical map construction. With SNaPshot fingerprinting, the sized fragments used for the physical map construction are only a portion of all fragments generated from a simultaneous, 5-enzyme digestion. Small (&lt;70 bp), large (&gt;1000 bp), those with non-labeling ends (HaeIII – HaeIII cleavage) and those from high-copy repeats are not sized and therefore not included in the fingerprint. An additional impediment is that the order of the fragment along a molecule is not known. In genome mapping, clone sizes are directly measured using the YOYO-stained DNA backbone. Since the DNA molecules are not cut but labeled at nick sites, it provides relative locations and linear order of sequence motifs on the DNA molecule. As a result, genome mapping produces accurate size estimates and facilitates tiling as shown in our ability to improve (and in some instances correct) the SNaPshot physical map. Additionally, multiple BACs can be analyzed in pools to increase throughput. Whole-genome profiling is a sequence tag based approach that is high throughput but since it is based on sets of sequence tags adjacent to restriction sites, it is unable to provide any length measurements. Optical mapping of static substrate-affixed DNA molecules has been used to construct physical maps for large and complex organisms <xref ref-type="bibr" rid="pone.0055864-Teague1">[18]</xref>, <xref ref-type="bibr" rid="pone.0055864-Zhou2">[21]</xref>. It suffers from low throughput, high error-rates and inconsistent DNA stretching. It is also a highly specialized technique that is difficult to master and therefore used by few labs. Genome mapping uses a nanochannel array to reproducibly and uniformly linearize DNA. In addition to the improved noise characteristics, by virtue of keeping DNA in solution rather than affixed, the system can perform cycles of channel-loading and imaging to generate throughputs of at least 5 Gb of DNA per hour. As shown here, genome mapping has the fundamental advantage that multiple motifs can be labeled with different colors, significantly increasing the information density.</p>
      <p>An important advantage of BAC-by-BAC sequencing strategies is that a physical map, with BAC positional information, is available prior to the start of genome sequencing. Physical maps provide very long and rigid scaffolds with BAC resolution and allow sequence contigs and paired-end scaffolds, from a limited number of pooled BACs, to be placed in a defined region. Distance information from physical map restriction fragments and BAC overlap have been used to guide sequence assembly by incorporating length constraints on the contig assembly <xref ref-type="bibr" rid="pone.0055864-Soderlund1">[35]</xref>, <xref ref-type="bibr" rid="pone.0055864-Warren1">[36]</xref>. This improves assembly, but these restriction maps are unordered and thus provide limited information. None of these impediments are present in genome mapping, which can therefore be used to guide and validate <italic>de novo</italic> assembly of sequencing data.</p>
      <p>DNA sequence contigs from shotgun sequencing for <italic>de novo</italic> assembly are usually small due to the presence of repetitive elements (such as satellites, tandem repeats and retroelements) and other sequence duplications. Assembly of contigs into scaffolds can be aided by the use of paired-end reads. However, we have shown even with a physical map generated by SNaPshot, relatively long sequence reads and short paired-end reads, the sequence assembly was only about 75% complete. By incorporating genome mapping, we improved the assembly to about 95% complete. The improvement stemmed from a number of factors: genome mapping is DNA sequencing-independent, it is high resolution (1.5 kb), has accurate length measurements (within 1 kb <xref ref-type="bibr" rid="pone.0055864-Lam1">[19]</xref>) and high information density (one site/5 kb). Genome mapping is ideally suited for validating, improving and facilitating anchoring of sequence contigs and paired-end scaffolds for full genome assembly.</p>
      <p>One of the most difficult challenges in sequence assembly is caused by highly repetitive sequences. Repetitive sequences often collapse in assembly due to the high sequence identity. Therefore, it is often impossible to assemble them correctly with sequence information alone. Genome maps can span long repeat regions not containing nick motifs. Even if sequence assemblies cannot be anchored directly to the repeat region, the repeat unit identity can be identified by use of adjacent mapped sequence contigs and the length of the repeat containing region can be measured on the genome map. By doing so, unambiguous contigs on each side of the repeat stretch can be linked with the aid of a genome map. Additionally, some repeat sequences contain nick motif sites and can be predicted by a characteristic genome map pattern such as a high density of labels, as in <xref ref-type="fig" rid="pone-0055864-g006">Figure 6</xref> or a repetitive pattern such as may be seen in a segmental duplication. This information can guide more accurate assembly of difficult and complex regions such as tandem gene duplication regions that are present in eukaryotic genomes.</p>
      <p>In summary, we envision adoption of genome mapping in whole genome <italic>de novo</italic> sequence assembly projects as it fills two critical needs, providing and/or correcting the physical map-based minimal tiling path and guiding <italic>de novo</italic> sequence assembly. Genome mapping can be used for <italic>ab initio</italic> physical map generation from BACs or from genomic DNA. BAC pools or genomic DNA can be sequenced with NGS, and assemblies can be guided with the genome map. This process would greatly improve the fidelity of the assembly process by detecting and correcting incorrect assemblies at an early stage. The genome map can also be used to recognize certain structural elements such as tandem repeats, which cause problems during assembly, and guide their assembly. Based on our current throughput of ∼5 Gb per hour, we expect to be able to collect data for 20x coverage of the <italic>Ae. tauschii</italic> genome in less than one day. BACs can be accomplished in a greatly reduced time frame than with the current methods. We expect this workflow to provide an unprecedented level of completion and accuracy in <italic>de novo</italic> genome sequence assembly.</p>
    </sec>
    <sec id="s4" sec-type="materials|methods">
      <title>Materials and Methods</title>
      <sec id="s4a">
        <title>Sample preparation and data collection</title>
        <p>A total of 663,000 BAC clones were fingerprinted with the SNaPshot high-information-content-fingerprinting (HICF) technology <xref ref-type="bibr" rid="pone.0055864-Luo1">[13]</xref>. Minimum tiling path BACs (27 total) were selected to cover a ∼2-Mb region; clones 1–27 are: RI628H11, RI339A11, HD254F14, HD330o06, MI263M23, HI297K16, MI236G21, RI313E10, TCM018D08, HI219K02, HD251I24, RI339P08, RI374K13, MI305M04, RI591E12, HD057J22, RI346M11, RI575G17, HD525B13, N_BB038E03, HI050E23, HD470O12, RI543P09, HI242N01, RI524D20, RI549B21, HD451L17, HD321F15.</p>
        <p>All DNA samples used in the study were prepared using the Qiagen Large-Construct Kit. To prepare BAC mixtures, we grew 250 mL cultures of each BAC in LB containing 20 µg/mL chloramphenicol or 12.5 µg/mL tetracycline overnight and combined the separate cultures before proceeding with DNA extraction of the BACs as a pool. The DNA samples were quantified using Quant-iTdsDNA Assay Kit (Invitrogen/Molecular Probes) and their quality assessed using pulsed-field gel electrophoresis. One microgram of BAC DNA was nicked with the nicking endonuclease, Nt.BspQI (New England BioLabs, NEB). Nicked DNA was labeled with Alexa546-dUTP (Invitrogen) and Taq polymerase (NEB). After labeling, the nick was ligated by adding dNTPs and T4 ligase (NEB). DNA was purified. The second nick labeling step was performed the same as the first except Nt.BbvCI was used and the nick was labeled with Alexa647-dUTP. DNA was linearized with Cre recombinase and LoxP containing dsDNA oligonucleotides. Cre was removed with Qiagen protease (Qiagen). The backbone of fluorescently tagged DNA was stained with YOYO-1 (Invitrogen).</p>
        <p>DNA was loaded in BioNano Genomics nanochannel array chips by electrophoresis of DNA, automated by the Irys system. Twelve volts were applied to concentrate the DNA, 30 V were applied to move DNA into the nanochannels, and 10 V was applied to distribute the DNA in the nanochannels. Linearized DNA molecules were imaged using blue, green and red lasers for YOYO-1, Alexa546, and Alexa647 on the BioNano Genomics Irys system.</p>
      </sec>
      <sec id="s4b">
        <title>Genome map analysis</title>
        <p>The entire DNA molecule (YOYO-1) and locations of fluorescent labels (Alexa546 and Alexa647) along each molecule were detected using the in house software package, IrysView. A set of label locations of each DNA molecule comprises the individual DNA molecular map.</p>
        <p>Single-molecule Nt.BspQI maps were clustered, as previously described <xref ref-type="bibr" rid="pone.0055864-Lam1">[19]</xref>. Briefly, all single-molecule maps were scored for similarity to one another in a pairwise comparison, and a Euclidian distance matrix was built. Maps were then clustered with the R package, <italic>fastcluster</italic>. From the clusters, the label locations were plotted as histograms and the peaks identified by fitting a Gaussian curve and used to make consensus maps. The Nt.BbvCI label positions were overlaid on the Nt.BspQI clusters, histograms were plotted, and peaks were fitted to produce a two-color map.</p>
      </sec>
      <sec id="s4c">
        <title>Next-generation sequencing of BAC clones</title>
        <p>BAC clones 1–23 were sequenced using the Roche/454 Titanium platform (using GS FLX or GS FLX+ chemistry) in barcoded pools with an average of five BACs per pool. The average read length was 570 bases with 20x coverage depth. gsAssembler was used for de novo assembly; after testing several stringency parameters to optimize assembly, we used reads with 40-bp overlap and a 95% identity threshold. For the GS FLX method, the average contig size is 5116 bp, N50 is 12937 bp and the largest contig size is 47797 bp. For the GS FLX+ method, the average contig size is 15758 bp, N50 is 40861 bp, and the largest contig size is 93822 bp. Eighteen of the BACs (1–4, 6, 8, 10–13, 15–19, 21–23) were pooled together with an additional ∼300 non-contiguous BACs to make one paired-end library with an average insert size of 3 kb and then sequenced by Roche/454 for 10x average coverage per BAC. The region of the MTP not covered with paired-end reads is minimal. We used these paired-end reads to scaffold the contigs with a 98% identity threshold, using gsAssembler. A total of 254 contigs and 13 scaffolds were generated. BAC end sequences were used to orient the contigs and scaffolds on the physical map.</p>
      </sec>
      <sec id="s4d">
        <title>DATA ACCESS</title>
        <p>The NCBI accession number for the final sequence (after genome map assisted assembly) is: JX295577.</p>
      </sec>
    </sec>
  </body>
  <back>
    <ack>
      <p>The authors thank H. VanSteenhouse for providing critical comments.</p>
    </ack>
    <ref-list>
      <title>References</title>
      <ref id="pone.0055864-Blakesley1">
        <label>1</label>
        <mixed-citation publication-type="journal" xlink:type="simple"><name name-style="western"><surname>Blakesley</surname><given-names>R</given-names></name>, <name name-style="western"><surname>Hansen</surname><given-names>N</given-names></name>, <name name-style="western"><surname>Gupta</surname><given-names>J</given-names></name>, <name name-style="western"><surname>McDowell</surname><given-names>J</given-names></name>, <name name-style="western"><surname>Maskeri</surname><given-names>B</given-names></name>, <etal>et al</etal>. (<year>2010</year>) <article-title>Effort required to finish shotgun-generated genome sequences differs significantly among vertebrates</article-title>. <source>BMC Genomics</source> <volume>11</volume>: <fpage>21</fpage>.</mixed-citation>
      </ref>
      <ref id="pone.0055864-Chain1">
        <label>2</label>
        <mixed-citation publication-type="journal" xlink:type="simple"><name name-style="western"><surname>Chain</surname><given-names>PSG</given-names></name>, <name name-style="western"><surname>Grafham</surname><given-names>DV</given-names></name>, <name name-style="western"><surname>Fulton</surname><given-names>RS</given-names></name>, <name name-style="western"><surname>FitzGerald</surname><given-names>MG</given-names></name>, <name name-style="western"><surname>Hostetler</surname><given-names>J</given-names></name>, <etal>et al</etal>. (<year>2009</year>) <article-title>Genome Project Standards in a New Era of Sequencing</article-title>. <source>Science</source> <volume>326</volume>: <fpage>236</fpage>–<lpage>237</lpage>.</mixed-citation>
      </ref>
      <ref id="pone.0055864-Lee1">
        <label>3</label>
        <mixed-citation publication-type="journal" xlink:type="simple"><name name-style="western"><surname>Lee</surname><given-names>H</given-names></name>, <name name-style="western"><surname>Tang</surname><given-names>H</given-names></name> (<year>2012</year>) <article-title>Next-generation sequencing technologies and fragment assembly algorithms</article-title>. <source>Methods Mol Biol</source> <volume>855</volume>: <fpage>155</fpage>–<lpage>174</lpage>.</mixed-citation>
      </ref>
      <ref id="pone.0055864-Green1">
        <label>4</label>
        <mixed-citation publication-type="journal" xlink:type="simple"><name name-style="western"><surname>Green</surname><given-names>ED</given-names></name> (<year>2001</year>) <article-title>Strategies for the systematic sequencing of complex genomes</article-title>. <source>Nat Rev Genet</source> <volume>2</volume>: <fpage>573</fpage>–<lpage>583</lpage>.</mixed-citation>
      </ref>
      <ref id="pone.0055864-McPherson1">
        <label>5</label>
        <mixed-citation publication-type="journal" xlink:type="simple"><name name-style="western"><surname>McPherson</surname><given-names>TIHGMCJD</given-names></name> (<year>2001</year>) <article-title>A physical map of the human genome</article-title>. <source>Nature</source> <volume>409</volume>: <fpage>934</fpage>–<lpage>941</lpage>.</mixed-citation>
      </ref>
      <ref id="pone.0055864-Smith1">
        <label>6</label>
        <mixed-citation publication-type="journal" xlink:type="simple"><name name-style="western"><surname>Smith</surname><given-names>DB</given-names></name>, <name name-style="western"><surname>Flavell</surname><given-names>RB</given-names></name> (<year>1975</year>) <article-title>Characterisation of the wheat genome by renaturation kinetics</article-title>. <source>Chromosoma</source> <volume>50</volume>: <fpage>223</fpage>–<lpage>242</lpage>.</mixed-citation>
      </ref>
      <ref id="pone.0055864-Venter1">
        <label>7</label>
        <mixed-citation publication-type="journal" xlink:type="simple"><name name-style="western"><surname>Venter</surname><given-names>JC</given-names></name>, <name name-style="western"><surname>Adams</surname><given-names>MD</given-names></name>, <name name-style="western"><surname>Myers</surname><given-names>EW</given-names></name>, <name name-style="western"><surname>Li</surname><given-names>PW</given-names></name>, <name name-style="western"><surname>Mural</surname><given-names>RJ</given-names></name>, <etal>et al</etal>. (<year>2001</year>) <article-title>The Sequence of the Human Genome</article-title>. <source>Science</source> <volume>291</volume>: <fpage>1304</fpage>–<lpage>1351</lpage>.</mixed-citation>
      </ref>
      <ref id="pone.0055864-Zuccolo1">
        <label>8</label>
        <mixed-citation publication-type="journal" xlink:type="simple"><name name-style="western"><surname>Zuccolo</surname><given-names>A</given-names></name>, <name name-style="western"><surname>Sebastian</surname><given-names>A</given-names></name>, <name name-style="western"><surname>Talag</surname><given-names>J</given-names></name>, <name name-style="western"><surname>Yu</surname><given-names>Y</given-names></name>, <name name-style="western"><surname>Kim</surname><given-names>H</given-names></name>, <etal>et al</etal>. (<year>2007</year>) <article-title>Transposable element distribution, abundance and role in genome size variation in the genus Oryza</article-title>. <source>BMC Evolutionary Biology</source> <volume>7</volume>: <fpage>152</fpage>.</mixed-citation>
      </ref>
      <ref id="pone.0055864-Initiative1">
        <label>9</label>
        <mixed-citation publication-type="journal" xlink:type="simple"><name name-style="western"><surname>Initiative</surname><given-names>TAG</given-names></name> (<year>2000</year>) <article-title>Analysis of the genome sequence of the flowering plant Arabidopsis thaliana</article-title>. <source>Nature</source> <volume>408</volume>: <fpage>796</fpage>–<lpage>815</lpage>.</mixed-citation>
      </ref>
      <ref id="pone.0055864-Project1">
        <label>10</label>
        <mixed-citation publication-type="journal" xlink:type="simple"><collab xlink:type="simple">Project IRGS</collab> (<year>2005</year>) <article-title>The map-based sequence of the rice genome</article-title>. <source>Nature</source> <volume>436</volume>: <fpage>793</fpage>–<lpage>800</lpage>.</mixed-citation>
      </ref>
      <ref id="pone.0055864-Zhou1">
        <label>11</label>
        <mixed-citation publication-type="journal" xlink:type="simple"><name name-style="western"><surname>Zhou</surname><given-names>S</given-names></name>, <name name-style="western"><surname>Wei</surname><given-names>F</given-names></name>, <name name-style="western"><surname>Nguyen</surname><given-names>J</given-names></name>, <name name-style="western"><surname>Bechner</surname><given-names>M</given-names></name>, <name name-style="western"><surname>Potamousis</surname><given-names>K</given-names></name>, <etal>et al</etal>. (<year>2009</year>) <article-title>A single molecule scaffold for the maize genome</article-title>. <source>PLoS Genet</source> <volume>5</volume>: <fpage>e1000711</fpage>.</mixed-citation>
      </ref>
      <ref id="pone.0055864-Schnable1">
        <label>12</label>
        <mixed-citation publication-type="journal" xlink:type="simple"><name name-style="western"><surname>Schnable</surname><given-names>PS</given-names></name>, <name name-style="western"><surname>Ware</surname><given-names>D</given-names></name>, <name name-style="western"><surname>Fulton</surname><given-names>RS</given-names></name>, <name name-style="western"><surname>Stein</surname><given-names>JC</given-names></name>, <name name-style="western"><surname>Wei</surname><given-names>F</given-names></name>, <etal>et al</etal>. (<year>2009</year>) <article-title>The B73 maize genome: complexity, diversity, and dynamics</article-title>. <source>Science</source> <volume>326</volume>: <fpage>1112</fpage>–<lpage>1115</lpage>.</mixed-citation>
      </ref>
      <ref id="pone.0055864-Luo1">
        <label>13</label>
        <mixed-citation publication-type="journal" xlink:type="simple"><name name-style="western"><surname>Luo</surname><given-names>MC</given-names></name>, <name name-style="western"><surname>Thomas</surname><given-names>C</given-names></name>, <name name-style="western"><surname>You</surname><given-names>FM</given-names></name>, <name name-style="western"><surname>Hsiao</surname><given-names>J</given-names></name>, <name name-style="western"><surname>Ouyang</surname><given-names>S</given-names></name>, <etal>et al</etal>. (<year>2003</year>) <article-title>High-throughput fingerprinting of bacterial artificial chromosomes using the snapshot labeling kit and sizing of restriction fragments by capillary electrophoresis</article-title>. <source>Genomics</source> <volume>82</volume>: <fpage>378</fpage>–<lpage>389</lpage>.</mixed-citation>
      </ref>
      <ref id="pone.0055864-Paux1">
        <label>14</label>
        <mixed-citation publication-type="journal" xlink:type="simple"><name name-style="western"><surname>Paux</surname><given-names>E</given-names></name>, <name name-style="western"><surname>Sourdille</surname><given-names>P</given-names></name>, <name name-style="western"><surname>Salse</surname><given-names>Jrm</given-names></name>, <name name-style="western"><surname>Saintenac</surname><given-names>C</given-names></name>, <name name-style="western"><surname>Choulet</surname><given-names>Fdr</given-names></name>, <etal>et al</etal>. (<year>2008</year>) <article-title>A Physical Map of the 1-Gigabase Bread Wheat Chromosome 3B</article-title>. <source>Science</source> <volume>322</volume>: <fpage>101</fpage>–<lpage>104</lpage>.</mixed-citation>
      </ref>
      <ref id="pone.0055864-Philippe1">
        <label>15</label>
        <mixed-citation publication-type="journal" xlink:type="simple"><name name-style="western"><surname>Philippe</surname><given-names>R</given-names></name>, <name name-style="western"><surname>Choulet</surname><given-names>F</given-names></name>, <name name-style="western"><surname>Paux</surname><given-names>E</given-names></name>, <name name-style="western"><surname>van Oeveren</surname><given-names>J</given-names></name>, <name name-style="western"><surname>Tang</surname><given-names>J</given-names></name>, <etal>et al</etal>. (<year>2012</year>) <article-title>Whole Genome Profiling provides a robust framework for physical mapping and sequencing in the highly complex and repetitive wheat genome</article-title>. <source>BMC Genomics</source> <volume>13</volume>: <fpage>47</fpage>.</mixed-citation>
      </ref>
      <ref id="pone.0055864-vanOeveren1">
        <label>16</label>
        <mixed-citation publication-type="journal" xlink:type="simple"><name name-style="western"><surname>van Oeveren</surname><given-names>J</given-names></name>, <name name-style="western"><surname>de Ruiter</surname><given-names>M</given-names></name>, <name name-style="western"><surname>Jesse</surname><given-names>T</given-names></name>, <name name-style="western"><surname>van der Poel</surname><given-names>H</given-names></name>, <name name-style="western"><surname>Tang</surname><given-names>J</given-names></name>, <etal>et al</etal>. (<year>2011</year>) <article-title>Sequence-based physical mapping of complex genomes by whole genome profiling</article-title>. <source>Genome Research</source> <volume>21(4)</volume>: <fpage>618</fpage>–<lpage>625</lpage>.</mixed-citation>
      </ref>
      <ref id="pone.0055864-Schwartz1">
        <label>17</label>
        <mixed-citation publication-type="journal" xlink:type="simple"><name name-style="western"><surname>Schwartz</surname><given-names>DC</given-names></name>, <name name-style="western"><surname>Li</surname><given-names>X</given-names></name>, <name name-style="western"><surname>Hernandez</surname><given-names>LI</given-names></name>, <name name-style="western"><surname>Ramnarain</surname><given-names>SP</given-names></name>, <name name-style="western"><surname>Huff</surname><given-names>EJ</given-names></name>, <etal>et al</etal>. (<year>1993</year>) <article-title>Ordered restriction maps of Saccharomyces cerevisiae chromosomes constructed by optical mapping</article-title>. <source>Science</source> <volume>262</volume>: <fpage>110</fpage>–<lpage>114</lpage>.</mixed-citation>
      </ref>
      <ref id="pone.0055864-Teague1">
        <label>18</label>
        <mixed-citation publication-type="journal" xlink:type="simple"><name name-style="western"><surname>Teague</surname><given-names>B</given-names></name>, <name name-style="western"><surname>Waterman</surname><given-names>MS</given-names></name>, <name name-style="western"><surname>Goldstein</surname><given-names>S</given-names></name>, <name name-style="western"><surname>Potamousis</surname><given-names>K</given-names></name>, <name name-style="western"><surname>Zhou</surname><given-names>S</given-names></name>, <etal>et al</etal>. (<year>2010</year>) <article-title>High-resolution human genome structure by single-molecule analysis</article-title>. <source>Proc Natl Acad Sci U S A</source> <volume>107</volume>: <fpage>10848</fpage>–<lpage>10853</lpage>.</mixed-citation>
      </ref>
      <ref id="pone.0055864-Lam1">
        <label>19</label>
        <mixed-citation publication-type="journal" xlink:type="simple"><name name-style="western"><surname>Lam</surname><given-names>ET</given-names></name>, <name name-style="western"><surname>Hastie</surname><given-names>A</given-names></name>, <name name-style="western"><surname>Lin</surname><given-names>C</given-names></name>, <name name-style="western"><surname>Ehrlich</surname><given-names>D</given-names></name>, <name name-style="western"><surname>Das</surname><given-names>SK</given-names></name>, <etal>et al</etal>. (<year>2012</year>) <article-title>Genome mapping on nanochannel arrays for structural variation analysis and sequence assembly</article-title>. <source>Nat Biotechnol</source> <volume>30</volume>: <fpage>771</fpage>–<lpage>776</lpage>.</mixed-citation>
      </ref>
      <ref id="pone.0055864-Mun1">
        <label>20</label>
        <mixed-citation publication-type="journal" xlink:type="simple"><name name-style="western"><surname>Mun</surname><given-names>JH</given-names></name>, <name name-style="western"><surname>Kwon</surname><given-names>SJ</given-names></name>, <name name-style="western"><surname>Yang</surname><given-names>TJ</given-names></name>, <name name-style="western"><surname>Kim</surname><given-names>HS</given-names></name>, <name name-style="western"><surname>Choi</surname><given-names>BS</given-names></name>, <etal>et al</etal>. (<year>2008</year>) <article-title>The first generation of a BAC-based physical map of Brassica rapa</article-title>. <source>BMC Genomics</source> <volume>9</volume>: <fpage>280</fpage>.</mixed-citation>
      </ref>
      <ref id="pone.0055864-Zhou2">
        <label>21</label>
        <mixed-citation publication-type="journal" xlink:type="simple"><name name-style="western"><surname>Zhou</surname><given-names>S</given-names></name>, <name name-style="western"><surname>Bechner</surname><given-names>MC</given-names></name>, <name name-style="western"><surname>Place</surname><given-names>M</given-names></name>, <name name-style="western"><surname>Churas</surname><given-names>CP</given-names></name>, <name name-style="western"><surname>Pape</surname><given-names>L</given-names></name>, <etal>et al</etal>. (<year>2007</year>) <article-title>Validation of rice genome sequence by optical mapping</article-title>. <source>BMC Genomics</source> <volume>8</volume>: <fpage>278</fpage>.</mixed-citation>
      </ref>
      <ref id="pone.0055864-Nagarajan1">
        <label>22</label>
        <mixed-citation publication-type="journal" xlink:type="simple"><name name-style="western"><surname>Nagarajan</surname><given-names>N</given-names></name>, <name name-style="western"><surname>Read</surname><given-names>TD</given-names></name>, <name name-style="western"><surname>Pop</surname><given-names>M</given-names></name> (<year>2008</year>) <article-title>Scaffolding and validation of bacterial genome assemblies using optical restriction maps</article-title>. <source>Bioinformatics</source> <volume>24</volume>: <fpage>1229</fpage>–<lpage>1235</lpage>.</mixed-citation>
      </ref>
      <ref id="pone.0055864-Howden1">
        <label>23</label>
        <mixed-citation publication-type="journal" xlink:type="simple"><name name-style="western"><surname>Howden</surname><given-names>BP</given-names></name>, <name name-style="western"><surname>Seemann</surname><given-names>T</given-names></name>, <name name-style="western"><surname>Harrison</surname><given-names>PF</given-names></name>, <name name-style="western"><surname>McEvoy</surname><given-names>CR</given-names></name>, <name name-style="western"><surname>Stanton</surname><given-names>JA</given-names></name>, <etal>et al</etal>. (<year>2010</year>) <article-title>Complete genome sequence of Staphylococcus aureus strain JKD6008, an ST239 clone of methicillin-resistant Staphylococcus aureus with intermediate-level vancomycin resistance</article-title>. <source>J Bacteriol</source> <volume>192</volume>: <fpage>5848</fpage>–<lpage>5849</lpage>.</mixed-citation>
      </ref>
      <ref id="pone.0055864-Riley1">
        <label>24</label>
        <mixed-citation publication-type="journal" xlink:type="simple"><name name-style="western"><surname>Riley</surname><given-names>MC</given-names></name>, <name name-style="western"><surname>Lee</surname><given-names>JE</given-names></name>, <name name-style="western"><surname>Lesho</surname><given-names>E</given-names></name>, <name name-style="western"><surname>Kirkup</surname><given-names>BC</given-names><suffix>Jr</suffix></name> (<year>2011</year>) <article-title>Optically mapping multiple bacterial genomes simultaneously in a single run</article-title>. <source>PLoS One</source> <volume>6</volume>: <fpage>e27085</fpage>.</mixed-citation>
      </ref>
      <ref id="pone.0055864-Lin1">
        <label>25</label>
        <mixed-citation publication-type="journal" xlink:type="simple"><name name-style="western"><surname>Lin</surname><given-names>HC</given-names></name>, <name name-style="western"><surname>Goldstein</surname><given-names>S</given-names></name>, <name name-style="western"><surname>Mendelowitz</surname><given-names>L</given-names></name>, <name name-style="western"><surname>Zhou</surname><given-names>S</given-names></name>, <name name-style="western"><surname>Wetzel</surname><given-names>J</given-names></name>, <etal>et al</etal>. (<year>2012</year>) <article-title>AGORA: Assembly Guided by Optical Restriction Alignment</article-title>. <source>BMC Bioinformatics</source> <volume>13</volume>: <fpage>189</fpage>.</mixed-citation>
      </ref>
      <ref id="pone.0055864-Xiao1">
        <label>26</label>
        <mixed-citation publication-type="journal" xlink:type="simple"><name name-style="western"><surname>Xiao</surname><given-names>M</given-names></name>, <name name-style="western"><surname>Phong</surname><given-names>A</given-names></name>, <name name-style="western"><surname>Ha</surname><given-names>C</given-names></name>, <name name-style="western"><surname>Chan</surname><given-names>T-F</given-names></name>, <name name-style="western"><surname>Cai</surname><given-names>D</given-names></name>, <etal>et al</etal>. (<year>2007</year>) <article-title>Rapid DNA mapping by fluorescent single molecule detection</article-title>. <source>Nucleic Acids Research</source> <volume>35</volume>: <fpage>e16</fpage>.</mixed-citation>
      </ref>
      <ref id="pone.0055864-Das1">
        <label>27</label>
        <mixed-citation publication-type="journal" xlink:type="simple"><name name-style="western"><surname>Das</surname><given-names>SK</given-names></name>, <name name-style="western"><surname>Austin</surname><given-names>MD</given-names></name>, <name name-style="western"><surname>Akana</surname><given-names>MC</given-names></name>, <name name-style="western"><surname>Deshpande</surname><given-names>P</given-names></name>, <name name-style="western"><surname>Cao</surname><given-names>H</given-names></name>, <etal>et al</etal>. (<year>2010</year>) <article-title>Single molecule linear analysis of DNA in nano-channel labeled with sequence specific fluorescent probes</article-title>. <source>Nucleic Acids Research</source> <volume>38</volume>: <fpage>e177</fpage>.</mixed-citation>
      </ref>
      <ref id="pone.0055864-Dvorak1">
        <label>28</label>
        <mixed-citation publication-type="book" xlink:type="simple">Dvorak J (2009) Triticeae Genome Structure and Evolution. Genetics and Genomics of the Triticeae Springer Science.</mixed-citation>
      </ref>
      <ref id="pone.0055864-Li1">
        <label>29</label>
        <mixed-citation publication-type="journal" xlink:type="simple"><name name-style="western"><surname>Li</surname><given-names>W</given-names></name>, <name name-style="western"><surname>Zhang</surname><given-names>P</given-names></name>, <name name-style="western"><surname>Fellers</surname><given-names>JP</given-names></name>, <name name-style="western"><surname>Friebe</surname><given-names>B</given-names></name>, <name name-style="western"><surname>Gill</surname><given-names>BS</given-names></name> (<year>2004</year>) <article-title>Sequence composition, organization, and evolution of the core Triticeae genome</article-title>. <source>Plant J</source> <volume>40</volume>: <fpage>500</fpage>–<lpage>511</lpage>.</mixed-citation>
      </ref>
      <ref id="pone.0055864-Cassidy1">
        <label>30</label>
        <mixed-citation publication-type="journal" xlink:type="simple"><name name-style="western"><surname>Cassidy</surname><given-names>BG</given-names></name>, <name name-style="western"><surname>Dvorak</surname><given-names>J</given-names></name> (<year>1991</year>) <article-title>Molecular Characterization of a Low-Molecular-Weight Glutenin Cdna Clone from Triticum-Durum</article-title>. <source>Theoretical and Applied Genetics</source> <volume>81</volume>: <fpage>653</fpage>–<lpage>660</lpage>.</mixed-citation>
      </ref>
      <ref id="pone.0055864-Hernandez1">
        <label>31</label>
        <mixed-citation publication-type="journal" xlink:type="simple"><name name-style="western"><surname>Hernandez</surname><given-names>P</given-names></name>, <name name-style="western"><surname>Martis</surname><given-names>M</given-names></name>, <name name-style="western"><surname>Dorado</surname><given-names>G</given-names></name>, <name name-style="western"><surname>Pfeifer</surname><given-names>M</given-names></name>, <name name-style="western"><surname>Galvez</surname><given-names>S</given-names></name>, <etal>et al</etal>. (<year>2012</year>) <article-title>Next-generation sequencing and syntenic integration of flow-sorted arms of wheat chromosome 4A exposes the chromosome structure and gene content</article-title>. <source>Plant J</source> <volume>69</volume>: <fpage>377</fpage>–<lpage>386</lpage>.</mixed-citation>
      </ref>
      <ref id="pone.0055864-Leroy1">
        <label>32</label>
        <mixed-citation publication-type="journal" xlink:type="simple"><name name-style="western"><surname>Leroy</surname><given-names>P</given-names></name>, <name name-style="western"><surname>Guilhot</surname><given-names>N</given-names></name>, <name name-style="western"><surname>Sakai</surname><given-names>H</given-names></name>, <name name-style="western"><surname>Bernard</surname><given-names>A</given-names></name>, <name name-style="western"><surname>Choulet</surname><given-names>F</given-names></name>, <etal>et al</etal>. (<year>2012</year>) <article-title>TriAnnot: A Versatile and High Performance Pipeline for the Automated Annotation of Plant Genomes</article-title>. <source>Front Plant Sci</source> <volume>3</volume>: <fpage>5</fpage>.</mixed-citation>
      </ref>
      <ref id="pone.0055864-Brenchley1">
        <label>33</label>
        <mixed-citation publication-type="journal" xlink:type="simple"><name name-style="western"><surname>Brenchley</surname><given-names>R</given-names></name>, <name name-style="western"><surname>Spannagl</surname><given-names>M</given-names></name>, <name name-style="western"><surname>Pfeifer</surname><given-names>M</given-names></name>, <name name-style="western"><surname>Barker</surname><given-names>GL</given-names></name>, <name name-style="western"><surname>D'Amore</surname><given-names>R</given-names></name>, <etal>et al</etal>. (<year>2012</year>) <article-title>Analysis of the bread wheat genome using whole-genome shotgun sequencing</article-title>. <source>Nature</source> <volume>491</volume>: <fpage>705</fpage>–<lpage>710</lpage>.</mixed-citation>
      </ref>
      <ref id="pone.0055864-Li2">
        <label>34</label>
        <mixed-citation publication-type="journal" xlink:type="simple"><name name-style="western"><surname>Li</surname><given-names>Y</given-names></name>, <name name-style="western"><surname>Zheng</surname><given-names>H</given-names></name>, <name name-style="western"><surname>Luo</surname><given-names>R</given-names></name>, <name name-style="western"><surname>Wu</surname><given-names>H</given-names></name>, <name name-style="western"><surname>Zhu</surname><given-names>H</given-names></name>, <etal>et al</etal>. (<year>2011</year>) <article-title>Structural variation in two human genomes mapped at single-nucleotide resolution by whole genome de novo assembly</article-title>. <source>Nat Biotechnol</source> <volume>29</volume>: <fpage>723</fpage>–<lpage>730</lpage>.</mixed-citation>
      </ref>
      <ref id="pone.0055864-Soderlund1">
        <label>35</label>
        <mixed-citation publication-type="journal" xlink:type="simple"><name name-style="western"><surname>Soderlund</surname><given-names>C</given-names></name>, <name name-style="western"><surname>Longden</surname><given-names>I</given-names></name>, <name name-style="western"><surname>Mott</surname><given-names>R</given-names></name> (<year>1997</year>) <article-title>FPC: a system for building contigs from restriction fingerprinted clones</article-title>. <source>Comput Appl Biosci</source> <volume>13</volume>: <fpage>523</fpage>–<lpage>535</lpage>.</mixed-citation>
      </ref>
      <ref id="pone.0055864-Warren1">
        <label>36</label>
        <mixed-citation publication-type="journal" xlink:type="simple"><name name-style="western"><surname>Warren</surname><given-names>RL</given-names></name>, <name name-style="western"><surname>Varabei</surname><given-names>D</given-names></name>, <name name-style="western"><surname>Platt</surname><given-names>D</given-names></name>, <name name-style="western"><surname>Huang</surname><given-names>X</given-names></name>, <name name-style="western"><surname>Messina</surname><given-names>D</given-names></name>, <etal>et al</etal>. (<year>2006</year>) <article-title>Physical map-assisted whole-genome shotgun sequence assemblies</article-title>. <source>Genome Res</source> <volume>16</volume>: <fpage>768</fpage>–<lpage>775</lpage>.</mixed-citation>
      </ref>
    </ref-list>
  </back>
</article>