<?xml version="1.0" encoding="UTF-8"?><!DOCTYPE article PUBLIC "-//NLM//DTD JATS (Z39.96) Journal Publishing DTD v1.2 20190208//EN" "http://jats.nlm.nih.gov/publishing/1.2/JATS-journalpublishing1.dtd"><article xmlns:mml="http://www.w3.org/1998/Math/MathML" xmlns:xlink="http://www.w3.org/1999/xlink" article-type="review-article" dtd-version="1.2" xml:lang="en">
    <front>
        <journal-meta>
            <journal-id journal-id-type="pmc">F1000Research</journal-id>
            <journal-title-group>
                <journal-title>F1000Research</journal-title>
            </journal-title-group>
            <issn pub-type="epub">2046-1402</issn>
            <publisher>
                <publisher-name>F1000 Research Limited</publisher-name>
                <publisher-loc>London, UK</publisher-loc>
            </publisher>
        </journal-meta>
        <article-meta>
            <article-id pub-id-type="doi">10.12688/f1000research.13980.1</article-id>
            <article-categories>
                <subj-group subj-group-type="heading">
                    <subject>Review</subject>
                </subj-group>
                <subj-group>
                    <subject>Articles</subject>
                </subj-group>
            </article-categories>
            <title-group>
                <article-title>Recent advances in the detection of repeat expansions with short-read next-generation sequencing</article-title>
                <fn-group content-type="pub-status">
                    <fn>
                        <p>[version 1; peer review: 3 approved]</p>
                    </fn>
                </fn-group>
            </title-group>
            <contrib-group>
                <contrib contrib-type="author" corresp="yes">
                    <name>
                        <surname>Bahlo</surname>
                        <given-names>Melanie</given-names>
                    </name>
                    <role content-type="http://credit.niso.org/">Conceptualization</role>
                    <role content-type="http://credit.niso.org/">Funding Acquisition</role>
                    <role content-type="http://credit.niso.org/">Investigation</role>
                    <role content-type="http://credit.niso.org/">Methodology</role>
                    <role content-type="http://credit.niso.org/">Project Administration</role>
                    <role content-type="http://credit.niso.org/">Resources</role>
                    <role content-type="http://credit.niso.org/">Software</role>
                    <role content-type="http://credit.niso.org/">Supervision</role>
                    <role content-type="http://credit.niso.org/">Writing &#x2013; Original Draft Preparation</role>
                    <role content-type="http://credit.niso.org/">Writing &#x2013; Review &amp; Editing</role>
                    <uri content-type="orcid">https://orcid.org/0000-0001-5132-0774</uri>
                    <xref ref-type="corresp" rid="c1">a</xref>
                    <xref ref-type="aff" rid="a1">1</xref>
                    <xref ref-type="aff" rid="a2">2</xref>
                </contrib>
                <contrib contrib-type="author" corresp="no">
                    <name>
                        <surname>Bennett</surname>
                        <given-names>Mark F</given-names>
                    </name>
                    <role content-type="http://credit.niso.org/">Formal Analysis</role>
                    <role content-type="http://credit.niso.org/">Investigation</role>
                    <role content-type="http://credit.niso.org/">Methodology</role>
                    <role content-type="http://credit.niso.org/">Resources</role>
                    <role content-type="http://credit.niso.org/">Software</role>
                    <role content-type="http://credit.niso.org/">Visualization</role>
                    <role content-type="http://credit.niso.org/">Writing &#x2013; Review &amp; Editing</role>
                    <xref ref-type="aff" rid="a1">1</xref>
                    <xref ref-type="aff" rid="a2">2</xref>
                    <xref ref-type="aff" rid="a3">3</xref>
                </contrib>
                <contrib contrib-type="author" corresp="no">
                    <name>
                        <surname>Degorski</surname>
                        <given-names>Peter</given-names>
                    </name>
                    <role content-type="http://credit.niso.org/">Formal Analysis</role>
                    <role content-type="http://credit.niso.org/">Investigation</role>
                    <role content-type="http://credit.niso.org/">Methodology</role>
                    <role content-type="http://credit.niso.org/">Resources</role>
                    <role content-type="http://credit.niso.org/">Software</role>
                    <role content-type="http://credit.niso.org/">Visualization</role>
                    <role content-type="http://credit.niso.org/">Writing &#x2013; Review &amp; Editing</role>
                    <xref ref-type="aff" rid="a1">1</xref>
                </contrib>
                <contrib contrib-type="author" corresp="no">
                    <name>
                        <surname>Tankard</surname>
                        <given-names>Rick M</given-names>
                    </name>
                    <role content-type="http://credit.niso.org/">Data Curation</role>
                    <role content-type="http://credit.niso.org/">Formal Analysis</role>
                    <role content-type="http://credit.niso.org/">Investigation</role>
                    <role content-type="http://credit.niso.org/">Methodology</role>
                    <role content-type="http://credit.niso.org/">Resources</role>
                    <role content-type="http://credit.niso.org/">Software</role>
                    <role content-type="http://credit.niso.org/">Validation</role>
                    <role content-type="http://credit.niso.org/">Visualization</role>
                    <role content-type="http://credit.niso.org/">Writing &#x2013; Review &amp; Editing</role>
                    <xref ref-type="aff" rid="a4">4</xref>
                </contrib>
                <contrib contrib-type="author" corresp="no">
                    <name>
                        <surname>Delatycki</surname>
                        <given-names>Martin B</given-names>
                    </name>
                    <role content-type="http://credit.niso.org/">Conceptualization</role>
                    <role content-type="http://credit.niso.org/">Investigation</role>
                    <role content-type="http://credit.niso.org/">Project Administration</role>
                    <role content-type="http://credit.niso.org/">Writing &#x2013; Review &amp; Editing</role>
                    <xref ref-type="aff" rid="a5">5</xref>
                    <xref ref-type="aff" rid="a6">6</xref>
                    <xref ref-type="aff" rid="a7">7</xref>
                </contrib>
                <contrib contrib-type="author" corresp="no">
                    <name>
                        <surname>Lockhart</surname>
                        <given-names>Paul J</given-names>
                    </name>
                    <role content-type="http://credit.niso.org/">Conceptualization</role>
                    <role content-type="http://credit.niso.org/">Investigation</role>
                    <role content-type="http://credit.niso.org/">Methodology</role>
                    <role content-type="http://credit.niso.org/">Project Administration</role>
                    <role content-type="http://credit.niso.org/">Validation</role>
                    <role content-type="http://credit.niso.org/">Writing &#x2013; Review &amp; Editing</role>
                    <xref ref-type="aff" rid="a5">5</xref>
                    <xref ref-type="aff" rid="a7">7</xref>
                </contrib>
                <aff id="a1">
                    <label>1</label>Population Health and Immunity Division, The Walter and Eliza Hall Institute of Medical Research, Parkville, Victoria, Australia</aff>
                <aff id="a2">
                    <label>2</label>Department of Medical Biology, The University of Melbourne, Parkville, Victoria, Australia</aff>
                <aff id="a3">
                    <label>3</label>Epilepsy Research Centre, Department of Medicine, The University of Melbourne, Heidelberg, Victoria, Australia</aff>
                <aff id="a4">
                    <label>4</label>Mathematics and Statistics, Murdoch University, Murdoch, Australia</aff>
                <aff id="a5">
                    <label>5</label>Bruce Lefroy Centre for Genetic Health Research, Murdoch Children&#x2019;s Research Institute, Royal Children&#x2019;s Hospital, Parkville, Victoria, Australia</aff>
                <aff id="a6">
                    <label>6</label>Victorian Clinical Genetics Services, Parkville, Victoria, Australia</aff>
                <aff id="a7">
                    <label>7</label>Department of Paediatrics, The University of Melbourne, Parkville, Victoria, Australia</aff>
            </contrib-group>
            <author-notes>
                <corresp id="c1">
                    <label>a</label>
                    <email xlink:href="mailto:bahlo@wehi.edu.au">bahlo@wehi.edu.au</email>
                </corresp>
                <fn fn-type="con">
                    <p>MB wrote the manuscript. MFB, PD, RMT, PJL, and MBD edited the manuscript.</p>
                </fn>
                <fn fn-type="conflict">
                    <p>No competing interests were disclosed.</p>
                </fn>
            </author-notes>
            <pub-date pub-type="epub">
                <day>13</day>
                <month>6</month>
                <year>2018</year>
            </pub-date>
            <pub-date pub-type="collection">
                <year>2018</year>
            </pub-date>
            <volume>7</volume>
            <elocation-id>F1000 Faculty Rev-736</elocation-id>
            <history>
                <date date-type="accepted">
                    <day>7</day>
                    <month>6</month>
                    <year>2018</year>
                </date>
            </history>
            <permissions>
                <copyright-statement>Copyright: &#x00a9; 2018 Bahlo M et al.</copyright-statement>
                <copyright-year>2018</copyright-year>
                <license xlink:href="https://creativecommons.org/licenses/by/4.0/">
                    <license-p>This is an open access article distributed under the terms of the Creative Commons Attribution Licence, which permits unrestricted use, distribution, and reproduction in any medium, provided the original work is properly cited.</license-p>
                </license>
            </permissions>
            <self-uri content-type="pdf" xlink:href="https://f1000research.com/articles/7-736/pdf"/>
            <abstract>
                <p>Short tandem repeats (STRs), also known as microsatellites, are commonly defined as consisting of tandemly repeated nucleotide motifs of 2&#x2013;6 base pairs in length. STRs appear throughout the human genome, and about 239,000 are documented in the Simple Repeats Track available from the University of California, Santa Cruz genome browser. STRs vary in size, producing highly polymorphic markers commonly used as genetic markers. A small fraction of STRs (about 30 loci) have been associated with human disease whereby one or both alleles exceed an STR-specific threshold in size, leading to disease. Detection of repeat expansions is currently performed with polymerase chain reaction&#x2013;based assays or with Southern blots for large expansions. The tests are expensive and time-consuming and are not always conclusive, leading to lengthy diagnostic journeys for patients, potentially including missed diagnoses. The advent of whole exome and whole genome sequencing has identified the genetic cause of many genetic disorders; however, analysis pipelines are focused primarily on the detection of short nucleotide variations and short insertions and deletions. Until recently, repeat expansions, with the exception of the smallest expansion (SCA6), were not detectable in next-generation short-read sequencing datasets and would have been ignored in most analyses. In the last two years, four analysis methods with accompanying software (ExpansionHunter, exSTRa, STRetch, and TREDPARSE) have been released. Although a comprehensive comparative analysis of the performance of these methods across all known repeat expansions is still lacking, it is clear that these methods are a valuable addition to any existing analysis pipeline. Here, we detail how to assess short-read data for evidence of expansions, reviewing all four methods and outlining their strengths and weaknesses. Implementation of these methods should lead to increased diagnostic yield of repeat expansion disorders for known STR loci and has the potential to detect novel repeat expansions.</p>
            </abstract>
            <kwd-group kwd-group-type="author">
                <kwd>short-read sequencing</kwd>
                <kwd>short tandem repeats</kwd>
                <kwd>repeat expansion disorders</kwd>
                <kwd>bioinformatics</kwd>
            </kwd-group>
            <funding-group>
                <award-group id="fund-1">
                    <funding-source>NHMRC Program Grant</funding-source>
                    <award-id>1054618</award-id>
                </award-group>
                <award-group id="fund-2">
                    <funding-source>NHMRC Senior Research Fellowship</funding-source>
                    <award-id>110297</award-id>
                </award-group>
                <award-group id="fund-3">
                    <funding-source>Australian Government NHMRC IRIIS</funding-source>
                </award-group>
                <award-group id="fund-4">
                    <funding-source>Victorian Government&#x2019;s Operational Infrastructure Support Program</funding-source>
                </award-group>
                <funding-statement>This work was supported by the Victorian Government&#x2019;s Operational Infrastructure Support Program and Australian Government National Health and Medical Research Council (NHMRC) Independent Medical Research Institutes Infrastructure Support Scheme. MB is funded by NHMRC Senior Research Fellowship 110297 and NHMRC Program Grant 1054618. </funding-statement>
                <funding-statement>
                    <italic>The funders had no role in study design, data collection and analysis, decision to publish, or preparation of the manuscript.</italic>
                </funding-statement>
            </funding-group>
        </article-meta>
        <notes>
            <sec sec-type="editor-note">
                <title>Editorial Note on the Review Process</title>
                <p>
                    <ext-link ext-link-type="uri" xlink:href="http://f1000research.com/browse/faculty-reviews">F1000 Faculty Reviews</ext-link> are commissioned from members of the prestigious
                    <ext-link ext-link-type="uri" xlink:href="http://f1000.com/prime/thefaculty">F1000 Faculty</ext-link> and are edited as a service to readers. In order to make these reviews as comprehensive and accessible as possible, the referees provide input before publication and only the final, revised version is published. The referees who approved the final version are listed with their names and affiliations but without their reports on earlier versions (any comments will already have been addressed in the published version).</p>
                <p>The referees who approved this article are: </p>
                <list list-content="reviewer-list" list-type="simple">
                    <list-item>
                        <p>
                            <named-content content-type="reviewer-name">Sergei Mirkin</named-content>, Department of Biology, Tufts University, Medford, USA
                            <fn fn-type="conflict">
                                <p>No competing interests were disclosed.</p>
                            </fn>
                        </p>
                    </list-item>
                    <list-item>
                        <p>
                            <named-content content-type="reviewer-name">Samuel S Chong</named-content>, Department of Paediatrics, Yong Loo Lin School of Medicine, National University of Singapore, Singapore, Singapore
                            <fn fn-type="conflict">
                                <p>No competing interests were disclosed.</p>
                            </fn>
                        </p>
                    </list-item>
                    <list-item>
                        <p>
                            <named-content content-type="reviewer-name">Tetsuo Ashizawa</named-content>, Stanley H. Appel Department of Neurology, Houston Methodist Neurological and Research Institutes, Weill Cornell Medical College, Texas, USA
                            <fn fn-type="conflict">
                                <p>No competing interests were disclosed.</p>
                            </fn>
                        </p>
                    </list-item>
                </list>
            </sec>
        </notes>
    </front>
    <body>
        <sec sec-type="intro">
            <title>Introduction</title>
            <p>Expansions of known short tandem repeats (STRs) have been identified as the sole cause of disease for several orphan diseases but also can contribute substantially to the pathogenic variant burden in polygenic disease. Fragile X syndrome (OMIM #300624), the most common inherited cause of intellectual disability and autism, is caused by expansions of a CGG repeat in the 5&#x2032; untranslated region of the gene encoding fragile X mental retardation 1 (
                <italic toggle="yes">FMR1</italic>) on the X chromosome. Unaffected individuals usually have STR alleles with a repeat motif number between 6 and 54. Affected male individuals have more than 200 copies of the motif. Huntington&#x2019;s disease (OMIM #143100), one of the most common dominant disorders in Caucasians (the prevalence is 5 out of 100,000)
                <sup>
                    <xref ref-type="bibr" rid="ref-1">1</xref>
                </sup>, is caused by an expansion of a CAG repeat in the coding sequence of the huntingtin gene (
                <italic toggle="yes">HTT</italic>). Unaffected individuals have between 6 and 35 motif copy numbers in their genomic sequence, and affected individuals have more than 35 motifs. Further examples include an expansion of the hexamer GGGGCC in the intron of 
                <italic toggle="yes">C9orf72</italic>, which can cause both amyotrophic lateral sclerosis and fronto-temporal dementia (FTDALS1, OMIM #105550) and contributes the highest genetic risk burden of any single locus to both of these disorders. Currently, there are about 30 known repeat expansions that cause human diseases and that vary in terms of supporting literature. Twenty-one of these, which cause neurological disorders, have well-documented normal and pathogenic allele size ranges and are summarized in 
                <xref ref-type="table" rid="T1">Table 1</xref>. The table includes several important non-neurological repeat expansion disorders, including the CTG expansion in TCF4, which causes the complex eye disorder Fuchs&#x2019; endothelial corneal dystrophy
                <sup>
                    <xref ref-type="bibr" rid="ref-2">2</xref>
                </sup>. In a recent discovery, the cause of FAME1 was found to be a complex pentamer repeat, situated in the gene 
                <italic toggle="yes">SAMD12</italic>
                <sup>
                    <xref ref-type="bibr" rid="ref-3">3</xref>
                </sup>. This repeat is not present in normal individuals (
                <xref ref-type="table" rid="T1">Table 1</xref>).</p>
            <table-wrap id="T1" orientation="portrait" position="anchor">
                <label>Table 1. </label>
                <caption>
                    <title>Detailed short tandem repeat loci information for neurological disorders.</title>
                </caption>
                <table content-type="article-table" frame="hsides">
                    <thead>
                        <tr>
                            <th align="left" colspan="1" rowspan="1" valign="top">Disease</th>
                            <th align="left" colspan="1" rowspan="1" valign="top">Symbol</th>
                            <th align="left" colspan="1" rowspan="1" valign="top">OMIM</th>
                            <th align="left" colspan="1" rowspan="1" valign="top">Inheritance</th>
                            <th align="left" colspan="1" rowspan="1" valign="top">Gene</th>
                            <th align="left" colspan="1" rowspan="1" valign="top">Cytogenetic
                                <break/>location</th>
                            <th align="left" colspan="1" rowspan="1" valign="top">Type</th>
                            <th align="left" colspan="1" rowspan="1" valign="top">Repeat motif</th>
                            <th align="left" colspan="1" rowspan="1" valign="top">Normal
                                <break/>range</th>
                            <th align="left" colspan="1" rowspan="1" valign="top">Expansion
                                <break/>range</th>
                            <th align="left" colspan="1" rowspan="1" valign="top">Strand</th>
                            <th align="left" colspan="1" rowspan="1" valign="top">Start hg19</th>
                            <th align="left" colspan="1" rowspan="1" valign="top">Reference
                                <break/>repeat
                                <break/>number</th>
                            <th align="left" colspan="1" rowspan="1" valign="top">TRF
                                <break/>match,
                                <break/>%</th>
                            <th align="left" colspan="1" rowspan="1" valign="top">TRF
                                <break/>indel,
                                <break/>%</th>
                            <th align="left" colspan="1" rowspan="1" valign="top">Reference
                                <break/>STR size,
                                <break/>base pairs</th>
                        </tr>
                    </thead>
                    <tbody>
                        <tr>
                            <td colspan="1" rowspan="1" valign="top">Huntington
                                <break/>disease</td>
                            <td colspan="1" rowspan="1" valign="top">HD</td>
                            <td colspan="1" rowspan="1" valign="top">143100</td>
                            <td colspan="1" rowspan="1" valign="top">AD</td>
                            <td colspan="1" rowspan="1" valign="top">
								
                                <italic toggle="yes">HTT</italic>
							</td>
                            <td colspan="1" rowspan="1" valign="top">4p16.3</td>
                            <td colspan="1" rowspan="1" valign="top">Coding</td>
                            <td colspan="1" rowspan="1" valign="top">CAG</td>
                            <td colspan="1" rowspan="1" valign="top">6&#x2013;34</td>
                            <td colspan="1" rowspan="1" valign="top">36&#x2013;100+</td>
                            <td colspan="1" rowspan="1" valign="top">+</td>
                            <td colspan="1" rowspan="1" valign="top">3,076,604</td>
                            <td colspan="1" rowspan="1" valign="top">21.3</td>
                            <td colspan="1" rowspan="1" valign="top">96</td>
                            <td colspan="1" rowspan="1" valign="top">0</td>
                            <td colspan="1" rowspan="1" valign="top">64</td>
                        </tr>
                        <tr>
                            <td colspan="1" rowspan="1" valign="top">Kennedy disease</td>
                            <td colspan="1" rowspan="1" valign="top">SBMA</td>
                            <td colspan="1" rowspan="1" valign="top">313200</td>
                            <td colspan="1" rowspan="1" valign="top">X</td>
                            <td colspan="1" rowspan="1" valign="top">
								
                                <italic toggle="yes">AR</italic>
							</td>
                            <td colspan="1" rowspan="1" valign="top">Xq12</td>
                            <td colspan="1" rowspan="1" valign="top">Coding</td>
                            <td colspan="1" rowspan="1" valign="top">CAG</td>
                            <td colspan="1" rowspan="1" valign="top">9&#x2013;35</td>
                            <td colspan="1" rowspan="1" valign="top">38&#x2013;62</td>
                            <td colspan="1" rowspan="1" valign="top">+</td>
                            <td colspan="1" rowspan="1" valign="top">66,765,159</td>
                            <td colspan="1" rowspan="1" valign="top">33.3</td>
                            <td colspan="1" rowspan="1" valign="top">86</td>
                            <td colspan="1" rowspan="1" valign="top">9</td>
                            <td colspan="1" rowspan="1" valign="top">103</td>
                        </tr>
                        <tr>
                            <td colspan="1" rowspan="1" valign="top">Spinocerebellar
                                <break/>ataxia 1</td>
                            <td colspan="1" rowspan="1" valign="top">SCA1</td>
                            <td colspan="1" rowspan="1" valign="top">164400</td>
                            <td colspan="1" rowspan="1" valign="top">AD</td>
                            <td colspan="1" rowspan="1" valign="top">
								
                                <italic toggle="yes">ATXN1</italic>
							</td>
                            <td colspan="1" rowspan="1" valign="top">6p23</td>
                            <td colspan="1" rowspan="1" valign="top">Coding</td>
                            <td colspan="1" rowspan="1" valign="top">CAG</td>
                            <td colspan="1" rowspan="1" valign="top">6&#x2013;38</td>
                            <td colspan="1" rowspan="1" valign="top">39&#x2013;82</td>
                            <td colspan="1" rowspan="1" valign="top">&#x2212;</td>
                            <td colspan="1" rowspan="1" valign="top">16,327,865</td>
                            <td colspan="1" rowspan="1" valign="top">30.3</td>
                            <td colspan="1" rowspan="1" valign="top">95</td>
                            <td colspan="1" rowspan="1" valign="top">0</td>
                            <td colspan="1" rowspan="1" valign="top">91</td>
                        </tr>
                        <tr>
                            <td colspan="1" rowspan="1" valign="top">Spinocerebellar
                                <break/>ataxia 2</td>
                            <td colspan="1" rowspan="1" valign="top">SCA2</td>
                            <td colspan="1" rowspan="1" valign="top">183090</td>
                            <td colspan="1" rowspan="1" valign="top">AD</td>
                            <td colspan="1" rowspan="1" valign="top">
								
                                <italic toggle="yes">ATXN2</italic>
							</td>
                            <td colspan="1" rowspan="1" valign="top">12q24</td>
                            <td colspan="1" rowspan="1" valign="top">Coding</td>
                            <td colspan="1" rowspan="1" valign="top">CAG</td>
                            <td colspan="1" rowspan="1" valign="top">15&#x2013;24</td>
                            <td colspan="1" rowspan="1" valign="top">32&#x2013;200</td>
                            <td colspan="1" rowspan="1" valign="top">&#x2212;</td>
                            <td colspan="1" rowspan="1" valign="top">112,036,754</td>
                            <td colspan="1" rowspan="1" valign="top">23.3</td>
                            <td colspan="1" rowspan="1" valign="top">97</td>
                            <td colspan="1" rowspan="1" valign="top">0</td>
                            <td colspan="1" rowspan="1" valign="top">70</td>
                        </tr>
                        <tr>
                            <td colspan="1" rowspan="1" valign="top">Machado-Joseph
                                <break/>disease</td>
                            <td colspan="1" rowspan="1" valign="top">SCA3</td>
                            <td colspan="1" rowspan="1" valign="top">109150</td>
                            <td colspan="1" rowspan="1" valign="top">AD</td>
                            <td colspan="1" rowspan="1" valign="top">
								
                                <italic toggle="yes">ATXN3</italic>
							</td>
                            <td colspan="1" rowspan="1" valign="top">14q32.1</td>
                            <td colspan="1" rowspan="1" valign="top">Coding</td>
                            <td colspan="1" rowspan="1" valign="top">CAG</td>
                            <td colspan="1" rowspan="1" valign="top">13&#x2013;36</td>
                            <td colspan="1" rowspan="1" valign="top">61&#x2013;84</td>
                            <td colspan="1" rowspan="1" valign="top">&#x2212;</td>
                            <td colspan="1" rowspan="1" valign="top">92,537,355</td>
                            <td colspan="1" rowspan="1" valign="top">14</td>
                            <td colspan="1" rowspan="1" valign="top">84</td>
                            <td colspan="1" rowspan="1" valign="top">0</td>
                            <td colspan="1" rowspan="1" valign="top">42</td>
                        </tr>
                        <tr>
                            <td colspan="1" rowspan="1" valign="top">Spinocerebellar
                                <break/>ataxia 6</td>
                            <td colspan="1" rowspan="1" valign="top">SCA6</td>
                            <td colspan="1" rowspan="1" valign="top">183086</td>
                            <td colspan="1" rowspan="1" valign="top">AD</td>
                            <td colspan="1" rowspan="1" valign="top">
								
                                <italic toggle="yes">CACNA1A</italic>
							</td>
                            <td colspan="1" rowspan="1" valign="top">19p13</td>
                            <td colspan="1" rowspan="1" valign="top">Coding</td>
                            <td colspan="1" rowspan="1" valign="top">CAG</td>
                            <td colspan="1" rowspan="1" valign="top">4&#x2013;7</td>
                            <td colspan="1" rowspan="1" valign="top">21&#x2013;33</td>
                            <td colspan="1" rowspan="1" valign="top">&#x2212;</td>
                            <td colspan="1" rowspan="1" valign="top">13,318,673</td>
                            <td colspan="1" rowspan="1" valign="top">13.3</td>
                            <td colspan="1" rowspan="1" valign="top">100</td>
                            <td colspan="1" rowspan="1" valign="top">0</td>
                            <td colspan="1" rowspan="1" valign="top">40</td>
                        </tr>
                        <tr>
                            <td colspan="1" rowspan="1" valign="top">Spinocerebellar
                                <break/>ataxia 7</td>
                            <td colspan="1" rowspan="1" valign="top">SCA7</td>
                            <td colspan="1" rowspan="1" valign="top">164500</td>
                            <td colspan="1" rowspan="1" valign="top">AD</td>
                            <td colspan="1" rowspan="1" valign="top">
								
                                <italic toggle="yes">ATXN7</italic>
							</td>
                            <td colspan="1" rowspan="1" valign="top">3p14.1</td>
                            <td colspan="1" rowspan="1" valign="top">Coding</td>
                            <td colspan="1" rowspan="1" valign="top">CAG</td>
                            <td colspan="1" rowspan="1" valign="top">4&#x2013;35</td>
                            <td colspan="1" rowspan="1" valign="top">37&#x2013;306</td>
                            <td colspan="1" rowspan="1" valign="top">+</td>
                            <td colspan="1" rowspan="1" valign="top">63,898,361</td>
                            <td colspan="1" rowspan="1" valign="top">10.7</td>
                            <td colspan="1" rowspan="1" valign="top">100</td>
                            <td colspan="1" rowspan="1" valign="top">0</td>
                            <td colspan="1" rowspan="1" valign="top">32</td>
                        </tr>
                        <tr>
                            <td colspan="1" rowspan="1" valign="top">Spinocerebellar
                                <break/>ataxia 17</td>
                            <td colspan="1" rowspan="1" valign="top">SCA17</td>
                            <td colspan="1" rowspan="1" valign="top">607136</td>
                            <td colspan="1" rowspan="1" valign="top">AD</td>
                            <td colspan="1" rowspan="1" valign="top">
								
                                <italic toggle="yes">TBP</italic>
							</td>
                            <td colspan="1" rowspan="1" valign="top">6q27</td>
                            <td colspan="1" rowspan="1" valign="top">coding</td>
                            <td colspan="1" rowspan="1" valign="top">CAG</td>
                            <td colspan="1" rowspan="1" valign="top">25&#x2013;42</td>
                            <td colspan="1" rowspan="1" valign="top">47&#x2013;63</td>
                            <td colspan="1" rowspan="1" valign="top">+</td>
                            <td colspan="1" rowspan="1" valign="top">170,870,995</td>
                            <td colspan="1" rowspan="1" valign="top">37</td>
                            <td colspan="1" rowspan="1" valign="top">94</td>
                            <td colspan="1" rowspan="1" valign="top">0</td>
                            <td colspan="1" rowspan="1" valign="top">111</td>
                        </tr>
                        <tr>
                            <td colspan="1" rowspan="1" valign="top">Dentatorubral-
                                <break/>pallidoluysian
                                <break/>atrophy</td>
                            <td colspan="1" rowspan="1" valign="top">DRPLA</td>
                            <td colspan="1" rowspan="1" valign="top">125370</td>
                            <td colspan="1" rowspan="1" valign="top">AD</td>
                            <td colspan="1" rowspan="1" valign="top">
								
                                <italic toggle="yes">DRPLA/</italic>
                                <break/>
                                <italic toggle="yes">ATN1</italic>
							</td>
                            <td colspan="1" rowspan="1" valign="top">12p13.31</td>
                            <td colspan="1" rowspan="1" valign="top">Coding</td>
                            <td colspan="1" rowspan="1" valign="top">CAG</td>
                            <td colspan="1" rowspan="1" valign="top">7&#x2013;34</td>
                            <td colspan="1" rowspan="1" valign="top">49&#x2013;88</td>
                            <td colspan="1" rowspan="1" valign="top">+</td>
                            <td colspan="1" rowspan="1" valign="top">7,045,880</td>
                            <td colspan="1" rowspan="1" valign="top">19.7</td>
                            <td colspan="1" rowspan="1" valign="top">92</td>
                            <td colspan="1" rowspan="1" valign="top">0</td>
                            <td colspan="1" rowspan="1" valign="top">59</td>
                        </tr>
                        <tr>
                            <td colspan="1" rowspan="1" valign="top">Huntington
                                <break/>disease-like 2</td>
                            <td colspan="1" rowspan="1" valign="top">HDL2</td>
                            <td colspan="1" rowspan="1" valign="top">606438</td>
                            <td colspan="1" rowspan="1" valign="top">AD</td>
                            <td colspan="1" rowspan="1" valign="top">
								
                                <italic toggle="yes">JPH3</italic>
							</td>
                            <td colspan="1" rowspan="1" valign="top">16q24.3</td>
                            <td colspan="1" rowspan="1" valign="top">Exon</td>
                            <td colspan="1" rowspan="1" valign="top">CTG</td>
                            <td colspan="1" rowspan="1" valign="top">7&#x2013;28</td>
                            <td colspan="1" rowspan="1" valign="top">66&#x2013;78</td>
                            <td colspan="1" rowspan="1" valign="top">+</td>
                            <td colspan="1" rowspan="1" valign="top">87,637,889</td>
                            <td colspan="1" rowspan="1" valign="top">15.3</td>
                            <td colspan="1" rowspan="1" valign="top">95</td>
                            <td colspan="1" rowspan="1" valign="top">4</td>
                            <td colspan="1" rowspan="1" valign="top">47</td>
                        </tr>
                        <tr>
                            <td colspan="1" rowspan="1" valign="top">Fragile-X site A</td>
                            <td colspan="1" rowspan="1" valign="top">FRAXA</td>
                            <td colspan="1" rowspan="1" valign="top">300624</td>
                            <td colspan="1" rowspan="1" valign="top">X</td>
                            <td colspan="1" rowspan="1" valign="top">
								
                                <italic toggle="yes">FMR1</italic>
							</td>
                            <td colspan="1" rowspan="1" valign="top">Xq27.3</td>
                            <td colspan="1" rowspan="1" valign="top">5&#x2032; UTR</td>
                            <td colspan="1" rowspan="1" valign="top">CGG</td>
                            <td colspan="1" rowspan="1" valign="top">6&#x2013;54</td>
                            <td colspan="1" rowspan="1" valign="top">200&#x2013;1,000+</td>
                            <td colspan="1" rowspan="1" valign="top">+</td>
                            <td colspan="1" rowspan="1" valign="top">146,993,555</td>
                            <td colspan="1" rowspan="1" valign="top">25</td>
                            <td colspan="1" rowspan="1" valign="top">90</td>
                            <td colspan="1" rowspan="1" valign="top">5</td>
                            <td colspan="1" rowspan="1" valign="top">75</td>
                        </tr>
                        <tr>
                            <td colspan="1" rowspan="1" valign="top">Fragile-X site E</td>
                            <td colspan="1" rowspan="1" valign="top">FRAXE</td>
                            <td colspan="1" rowspan="1" valign="top">309548</td>
                            <td colspan="1" rowspan="1" valign="top">X</td>
                            <td colspan="1" rowspan="1" valign="top">
								
                                <italic toggle="yes">FMR2</italic>
							</td>
                            <td colspan="1" rowspan="1" valign="top">Xq28</td>
                            <td colspan="1" rowspan="1" valign="top">5&#x2032; UTR</td>
                            <td colspan="1" rowspan="1" valign="top">CCG</td>
                            <td colspan="1" rowspan="1" valign="top">4&#x2013;39</td>
                            <td colspan="1" rowspan="1" valign="top">200&#x2013;900</td>
                            <td colspan="1" rowspan="1" valign="top">+</td>
                            <td colspan="1" rowspan="1" valign="top">147,582,159</td>
                            <td colspan="1" rowspan="1" valign="top">15.3</td>
                            <td colspan="1" rowspan="1" valign="top">100</td>
                            <td colspan="1" rowspan="1" valign="top">0</td>
                            <td colspan="1" rowspan="1" valign="top">46</td>
                        </tr>
                        <tr>
                            <td colspan="1" rowspan="1" valign="top">Myotonic
                                <break/>dystrophy 1</td>
                            <td colspan="1" rowspan="1" valign="top">DM1</td>
                            <td colspan="1" rowspan="1" valign="top">160900</td>
                            <td colspan="1" rowspan="1" valign="top">AD</td>
                            <td colspan="1" rowspan="1" valign="top">
								
                                <italic toggle="yes">DMPK</italic>
							</td>
                            <td colspan="1" rowspan="1" valign="top">19q13</td>
                            <td colspan="1" rowspan="1" valign="top">3&#x2032; UTR</td>
                            <td colspan="1" rowspan="1" valign="top">CTG</td>
                            <td colspan="1" rowspan="1" valign="top">5&#x2013;37</td>
                            <td colspan="1" rowspan="1" valign="top">50&#x2013;10,000</td>
                            <td colspan="1" rowspan="1" valign="top">&#x2212;</td>
                            <td colspan="1" rowspan="1" valign="top">46,273,463</td>
                            <td colspan="1" rowspan="1" valign="top">20.7</td>
                            <td colspan="1" rowspan="1" valign="top">100</td>
                            <td colspan="1" rowspan="1" valign="top">0</td>
                            <td colspan="1" rowspan="1" valign="top">62</td>
                        </tr>
                        <tr>
                            <td colspan="1" rowspan="1" valign="top">Friedreich ataxia</td>
                            <td colspan="1" rowspan="1" valign="top">FRDA</td>
                            <td colspan="1" rowspan="1" valign="top">229300</td>
                            <td colspan="1" rowspan="1" valign="top">AR</td>
                            <td colspan="1" rowspan="1" valign="top">
								
                                <italic toggle="yes">FXN</italic>
							</td>
                            <td colspan="1" rowspan="1" valign="top">9q13</td>
                            <td colspan="1" rowspan="1" valign="top">Intron</td>
                            <td colspan="1" rowspan="1" valign="top">GAA</td>
                            <td colspan="1" rowspan="1" valign="top">6&#x2013;32</td>
                            <td colspan="1" rowspan="1" valign="top">200&#x2013;1,700</td>
                            <td colspan="1" rowspan="1" valign="top">+</td>
                            <td colspan="1" rowspan="1" valign="top">71,652,201</td>
                            <td colspan="1" rowspan="1" valign="top">6.7</td>
                            <td colspan="1" rowspan="1" valign="top">100</td>
                            <td colspan="1" rowspan="1" valign="top">0</td>
                            <td colspan="1" rowspan="1" valign="top">20</td>
                        </tr>
                        <tr>
                            <td colspan="1" rowspan="1" valign="top">Myotonic
                                <break/>dystrophy 2</td>
                            <td colspan="1" rowspan="1" valign="top">DM2</td>
                            <td colspan="1" rowspan="1" valign="top">602668</td>
                            <td colspan="1" rowspan="1" valign="top">AD</td>
                            <td colspan="1" rowspan="1" valign="top">
								
                                <italic toggle="yes">ZNF9/CNBP</italic>
							</td>
                            <td colspan="1" rowspan="1" valign="top">3q21.3</td>
                            <td colspan="1" rowspan="1" valign="top">Intron</td>
                            <td colspan="1" rowspan="1" valign="top">CCTG</td>
                            <td colspan="1" rowspan="1" valign="top">10&#x2013;26</td>
                            <td colspan="1" rowspan="1" valign="top">75&#x2013;11,000</td>
                            <td colspan="1" rowspan="1" valign="top">&#x2212;</td>
                            <td colspan="1" rowspan="1" valign="top">128,891,420</td>
                            <td colspan="1" rowspan="1" valign="top">20.8</td>
                            <td colspan="1" rowspan="1" valign="top">92</td>
                            <td colspan="1" rowspan="1" valign="top">0</td>
                            <td colspan="1" rowspan="1" valign="top">83</td>
                        </tr>
                        <tr>
                            <td colspan="1" rowspan="1" valign="top">Frontotemporal
                                <break/>dementia and/or
                                <break/>amyotrophic
                                <break/>lateral sclerosis 1</td>
                            <td colspan="1" rowspan="1" valign="top">FTDALS1</td>
                            <td colspan="1" rowspan="1" valign="top">105550</td>
                            <td colspan="1" rowspan="1" valign="top">AD</td>
                            <td colspan="1" rowspan="1" valign="top">
								
                                <italic toggle="yes">C9orf72</italic>
							</td>
                            <td colspan="1" rowspan="1" valign="top">9p21</td>
                            <td colspan="1" rowspan="1" valign="top">Intron</td>
                            <td colspan="1" rowspan="1" valign="top">GGGGCC</td>
                            <td colspan="1" rowspan="1" valign="top">2&#x2013;19</td>
                            <td colspan="1" rowspan="1" valign="top">250&#x2013;1,600</td>
                            <td colspan="1" rowspan="1" valign="top">&#x2212;</td>
                            <td colspan="1" rowspan="1" valign="top">27,573,483</td>
                            <td colspan="1" rowspan="1" valign="top">10.8</td>
                            <td colspan="1" rowspan="1" valign="top">74</td>
                            <td colspan="1" rowspan="1" valign="top">8</td>
                            <td colspan="1" rowspan="1" valign="top">62</td>
                        </tr>
                        <tr>
                            <td colspan="1" rowspan="1" valign="top">Spinocerebellar
                                <break/>ataxia 36</td>
                            <td colspan="1" rowspan="1" valign="top">SCA36</td>
                            <td colspan="1" rowspan="1" valign="top">614153</td>
                            <td colspan="1" rowspan="1" valign="top">AD</td>
                            <td colspan="1" rowspan="1" valign="top">
								
                                <italic toggle="yes">NOP56</italic>
							</td>
                            <td colspan="1" rowspan="1" valign="top">20p13</td>
                            <td colspan="1" rowspan="1" valign="top">Intron</td>
                            <td colspan="1" rowspan="1" valign="top">GGCCTG</td>
                            <td colspan="1" rowspan="1" valign="top">3&#x2013;8</td>
                            <td colspan="1" rowspan="1" valign="top">1500&#x2013;2,500</td>
                            <td colspan="1" rowspan="1" valign="top">+</td>
                            <td colspan="1" rowspan="1" valign="top">2,633,379</td>
                            <td colspan="1" rowspan="1" valign="top">7.2</td>
                            <td colspan="1" rowspan="1" valign="top">97</td>
                            <td colspan="1" rowspan="1" valign="top">0</td>
                            <td colspan="1" rowspan="1" valign="top">43</td>
                        </tr>
                        <tr>
                            <td colspan="1" rowspan="1" valign="top">Spinocerebellar
                                <break/>ataxia 10</td>
                            <td colspan="1" rowspan="1" valign="top">SCA10</td>
                            <td colspan="1" rowspan="1" valign="top">603516</td>
                            <td colspan="1" rowspan="1" valign="top">AD</td>
                            <td colspan="1" rowspan="1" valign="top">
								
                                <italic toggle="yes">ATXN10</italic>
							</td>
                            <td colspan="1" rowspan="1" valign="top">22q13.31</td>
                            <td colspan="1" rowspan="1" valign="top">Intron</td>
                            <td colspan="1" rowspan="1" valign="top">ATTCT</td>
                            <td colspan="1" rowspan="1" valign="top">10&#x2013;20</td>
                            <td colspan="1" rowspan="1" valign="top">500&#x2013;4,500</td>
                            <td colspan="1" rowspan="1" valign="top">+</td>
                            <td colspan="1" rowspan="1" valign="top">46,191,235</td>
                            <td colspan="1" rowspan="1" valign="top">14</td>
                            <td colspan="1" rowspan="1" valign="top">100</td>
                            <td colspan="1" rowspan="1" valign="top">0</td>
                            <td colspan="1" rowspan="1" valign="top">70</td>
                        </tr>
                        <tr>
                            <td colspan="1" rowspan="1" valign="top">Myoclonic
                                <break/>epilepsy of
                                <break/>Unverricht and
                                <break/>Lundborg</td>
                            <td colspan="1" rowspan="1" valign="top">EPM1</td>
                            <td colspan="1" rowspan="1" valign="top">254800</td>
                            <td colspan="1" rowspan="1" valign="top">AR</td>
                            <td colspan="1" rowspan="1" valign="top">
								
                                <italic toggle="yes">CSTB</italic>
							</td>
                            <td colspan="1" rowspan="1" valign="top">21q22.3</td>
                            <td colspan="1" rowspan="1" valign="top">Promoter</td>
                            <td colspan="1" rowspan="1" valign="top">CCCCGCCCCGCG</td>
                            <td colspan="1" rowspan="1" valign="top">2&#x2013;3</td>
                            <td colspan="1" rowspan="1" valign="top">40&#x2013;80</td>
                            <td colspan="1" rowspan="1" valign="top">&#x2212;</td>
                            <td colspan="1" rowspan="1" valign="top">45,196,324</td>
                            <td colspan="1" rowspan="1" valign="top">3.1</td>
                            <td colspan="1" rowspan="1" valign="top">100</td>
                            <td colspan="1" rowspan="1" valign="top">0</td>
                            <td colspan="1" rowspan="1" valign="top">37</td>
                        </tr>
                        <tr>
                            <td colspan="1" rowspan="1" valign="top">Spinocerebellar
                                <break/>ataxia 12</td>
                            <td colspan="1" rowspan="1" valign="top">SCA12</td>
                            <td colspan="1" rowspan="1" valign="top">604326</td>
                            <td colspan="1" rowspan="1" valign="top">AD</td>
                            <td colspan="1" rowspan="1" valign="top">
								
                                <italic toggle="yes">PPP2R2B</italic>
							</td>
                            <td colspan="1" rowspan="1" valign="top">5q32</td>
                            <td colspan="1" rowspan="1" valign="top">Promoter</td>
                            <td colspan="1" rowspan="1" valign="top">CAG</td>
                            <td colspan="1" rowspan="1" valign="top">7&#x2013;45</td>
                            <td colspan="1" rowspan="1" valign="top">55&#x2013;78</td>
                            <td colspan="1" rowspan="1" valign="top">&#x2212;</td>
                            <td colspan="1" rowspan="1" valign="top">146,258,291</td>
                            <td colspan="1" rowspan="1" valign="top">10.7</td>
                            <td colspan="1" rowspan="1" valign="top">100</td>
                            <td colspan="1" rowspan="1" valign="top">0</td>
                            <td colspan="1" rowspan="1" valign="top">32</td>
                        </tr>
                        <tr>
                            <td colspan="1" rowspan="1" valign="top">Spinocerebellar
                                <break/>ataxia 8</td>
                            <td colspan="1" rowspan="1" valign="top">SCA8</td>
                            <td colspan="1" rowspan="1" valign="top">608768</td>
                            <td colspan="1" rowspan="1" valign="top">AD</td>
                            <td colspan="1" rowspan="1" valign="top">
								
                                <italic toggle="yes">ATXN8OS/</italic>
                                <break/>
                                <italic toggle="yes">ATXN8</italic>
							</td>
                            <td colspan="1" rowspan="1" valign="top">13q21</td>
                            <td colspan="1" rowspan="1" valign="top">utRNA</td>
                            <td colspan="1" rowspan="1" valign="top">CTG</td>
                            <td colspan="1" rowspan="1" valign="top">16&#x2013;34</td>
                            <td colspan="1" rowspan="1" valign="top">74+</td>
                            <td colspan="1" rowspan="1" valign="top">+</td>
                            <td colspan="1" rowspan="1" valign="top">70,713,516</td>
                            <td colspan="1" rowspan="1" valign="top">15.3</td>
                            <td colspan="1" rowspan="1" valign="top">100</td>
                            <td colspan="1" rowspan="1" valign="top">0</td>
                            <td colspan="1" rowspan="1" valign="top">46</td>
                        </tr>
                        <tr>
                            <td colspan="1" rowspan="1" valign="top">Spinocerebellar
                                <break/>ataxia 31</td>
                            <td colspan="1" rowspan="1" valign="top">SCA31</td>
                            <td colspan="1" rowspan="1" valign="top">117210</td>
                            <td colspan="1" rowspan="1" valign="top">AD</td>
                            <td colspan="1" rowspan="1" valign="top">
								
                                <italic toggle="yes">BEAN1/TK2</italic>
							</td>
                            <td colspan="1" rowspan="1" valign="top">16q21</td>
                            <td colspan="1" rowspan="1" valign="top">Intron</td>
                            <td colspan="1" rowspan="1" valign="top">TGGAA
                                <sup>
                                    <xref ref-type="other" rid="fn1">a</xref>
                                </sup>
                            </td>
                            <td colspan="1" rowspan="1" valign="top">N/A</td>
                            <td colspan="1" rowspan="1" valign="top">2.5&#x2013;3.8 kb
                                <sup>
                                    <xref ref-type="other" rid="fn2">b</xref>
                                </sup>
                            </td>
                            <td colspan="1" rowspan="1" valign="top">+</td>
                            <td colspan="1" rowspan="1" valign="top">66,524,302</td>
                            <td colspan="1" rowspan="1" valign="top">0</td>
                            <td colspan="1" rowspan="1" valign="top">N/A</td>
                            <td colspan="1" rowspan="1" valign="top">N/A</td>
                            <td colspan="1" rowspan="1" valign="top">N/A</td>
                        </tr>
                        <tr>
                            <td colspan="1" rowspan="1" valign="top">Spinocerebellar
                                <break/>ataxia 37</td>
                            <td colspan="1" rowspan="1" valign="top">SCA37</td>
                            <td colspan="1" rowspan="1" valign="top">615945</td>
                            <td colspan="1" rowspan="1" valign="top">AD</td>
                            <td colspan="1" rowspan="1" valign="top">
								
                                <italic toggle="yes">DAB1</italic>
							</td>
                            <td colspan="1" rowspan="1" valign="top">1p32.2</td>
                            <td colspan="1" rowspan="1" valign="top">Intron</td>
                            <td colspan="1" rowspan="1" valign="top">ATTTC
                                <sup>
                                    <xref ref-type="other" rid="fn1">a</xref>
                                </sup>
                            </td>
                            <td colspan="1" rowspan="1" valign="top">0</td>
                            <td colspan="1" rowspan="1" valign="top">31&#x2013;75</td>
                            <td colspan="1" rowspan="1" valign="top">&#x2212;</td>
                            <td colspan="1" rowspan="1" valign="top">57,832,716
                                <sup>
                                    <xref ref-type="other" rid="fn3">c</xref>
                                </sup>
                            </td>
                            <td colspan="1" rowspan="1" valign="top">0</td>
                            <td colspan="1" rowspan="1" valign="top">N/A</td>
                            <td colspan="1" rowspan="1" valign="top">N/A</td>
                            <td colspan="1" rowspan="1" valign="top">N/A</td>
                        </tr>
                        <tr>
                            <td colspan="1" rowspan="1" valign="top">Familial adult
                                <break/>myoclonic
                                <break/>epilepsy 1
                                <sup>
                                    <xref ref-type="other" rid="fn4">d</xref>
                                </sup>
                            </td>
                            <td colspan="1" rowspan="1" valign="top">FAME1</td>
                            <td colspan="1" rowspan="1" valign="top">601068</td>
                            <td colspan="1" rowspan="1" valign="top">AD</td>
                            <td colspan="1" rowspan="1" valign="top">
								
                                <italic toggle="yes">SAMD12</italic>
							</td>
                            <td colspan="1" rowspan="1" valign="top">8q24</td>
                            <td colspan="1" rowspan="1" valign="top">Intron</td>
                            <td colspan="1" rowspan="1" valign="top">TTTCA
                                <sup>
                                    <xref ref-type="other" rid="fn1">a</xref>
                                </sup>
                            </td>
                            <td colspan="1" rowspan="1" valign="top">0</td>
                            <td colspan="1" rowspan="1" valign="top">440&#x2013;3,680
                                <sup>
                                    <xref ref-type="other" rid="fn5">e</xref>
                                </sup>
                            </td>
                            <td colspan="1" rowspan="1" valign="top">&#x2212;</td>
                            <td colspan="1" rowspan="1" valign="top">119,379,055
                                <sup>
                                    <xref ref-type="other" rid="fn3">c</xref>
                                </sup>
                            </td>
                            <td colspan="1" rowspan="1" valign="top">0</td>
                            <td colspan="1" rowspan="1" valign="top">N/A</td>
                            <td colspan="1" rowspan="1" valign="top">N/A</td>
                            <td colspan="1" rowspan="1" valign="top">N/A</td>
                        </tr>
                        <tr>
                            <td colspan="1" rowspan="1" valign="top">Fuchs endothelial
                                <break/>corneal
                                <break/>dystrophy 3</td>
                            <td colspan="1" rowspan="1" valign="top">FECD3</td>
                            <td colspan="1" rowspan="1" valign="top">613267</td>
                            <td colspan="1" rowspan="1" valign="top">AD</td>
                            <td colspan="1" rowspan="1" valign="top">
								
                                <italic toggle="yes">TCF4</italic>
							</td>
                            <td colspan="1" rowspan="1" valign="top">18q21.2</td>
                            <td colspan="1" rowspan="1" valign="top">Intron</td>
                            <td colspan="1" rowspan="1" valign="top">CTG</td>
                            <td colspan="1" rowspan="1" valign="top">10&#x2013;40</td>
                            <td colspan="1" rowspan="1" valign="top">50&#x2013;150+</td>
                            <td colspan="1" rowspan="1" valign="top">&#x2212;</td>
                            <td colspan="1" rowspan="1" valign="top">53,253,385</td>
                            <td colspan="1" rowspan="1" valign="top">25.3</td>
                            <td colspan="1" rowspan="1" valign="top">100</td>
                            <td colspan="1" rowspan="1" valign="top">0</td>
                            <td colspan="1" rowspan="1" valign="top">76</td>
                        </tr>
                        <tr>
                            <td colspan="1" rowspan="1" valign="top">Oculopharyngeal
                                <break/>muscular
                                <break/>dystrophy</td>
                            <td colspan="1" rowspan="1" valign="top">OPMD</td>
                            <td colspan="1" rowspan="1" valign="top">164300</td>
                            <td colspan="1" rowspan="1" valign="top">AD</td>
                            <td colspan="1" rowspan="1" valign="top">
								
                                <italic toggle="yes">PABPN1</italic>
							</td>
                            <td colspan="1" rowspan="1" valign="top">14q11.2</td>
                            <td colspan="1" rowspan="1" valign="top">Coding</td>
                            <td colspan="1" rowspan="1" valign="top">GCG</td>
                            <td colspan="1" rowspan="1" valign="top">6&#x2013;7</td>
                            <td colspan="1" rowspan="1" valign="top">8&#x2013;13</td>
                            <td colspan="1" rowspan="1" valign="top">+</td>
                            <td colspan="1" rowspan="1" valign="top">23,790,682</td>
                            <td colspan="1" rowspan="1" valign="top">6.7</td>
                            <td colspan="1" rowspan="1" valign="top">100</td>
                            <td colspan="1" rowspan="1" valign="top">0</td>
                            <td colspan="1" rowspan="1" valign="top">20</td>
                        </tr>
                        <tr>
                            <td colspan="1" rowspan="1" valign="top">Early infantile
                                <break/>epileptic
                                <break/>encephalopathy 1
                                <sup>
                                    <xref ref-type="other" rid="fn6">f</xref>
                                </sup>
                            </td>
                            <td colspan="1" rowspan="1" valign="top">EIEE1</td>
                            <td colspan="1" rowspan="1" valign="top">308350</td>
                            <td colspan="1" rowspan="1" valign="top">X</td>
                            <td colspan="1" rowspan="1" valign="top">
								
                                <italic toggle="yes">ARX</italic>
							</td>
                            <td colspan="1" rowspan="1" valign="top">Xp21.3</td>
                            <td colspan="1" rowspan="1" valign="top">Coding</td>
                            <td colspan="1" rowspan="1" valign="top">GCG</td>
                            <td colspan="1" rowspan="1" valign="top">7&#x2013;12
                                <sup>
                                    <xref ref-type="other" rid="fn6">f</xref>
                                </sup>
                            </td>
                            <td colspan="1" rowspan="1" valign="top">17&#x2013;20
                                <sup>
                                    <xref ref-type="other" rid="fn6">f</xref>
                                </sup>
                            </td>
                            <td colspan="1" rowspan="1" valign="top">&#x2212;</td>
                            <td colspan="1" rowspan="1" valign="top">25,031,771</td>
                            <td colspan="1" rowspan="1" valign="top">14.7</td>
                            <td colspan="1" rowspan="1" valign="top">90</td>
                            <td colspan="1" rowspan="1" valign="top">0</td>
                            <td colspan="1" rowspan="1" valign="top">44</td>
                        </tr>
                    </tbody>
                </table>
                <table-wrap-foot>
                    <fn>
                        <p>Detailed short tandem repeat (STR) loci information for disorders associated with repeat expansions. Tandem Repeats Finder (TRF) (Benson
                            <sup>
                                <xref ref-type="bibr" rid="ref-16">16</xref>
                            </sup> 1999) match and TRF indel describe the purity of the repeat. AD, autosomal dominant; AR, autosomal recessive; N/A, not applicable; UTR, untranslated region; X, X-linked.</p>
                        <p id="fn1">
							
                            <sup>a</sup>As these repeats are insertions, the motifs do not appear in the reference at the respective locus.</p>
                        <p id="fn2">
							
                            <sup>b</sup>SCA31 is caused by the insertion of a complex repeat containing (TGGAA)
                            <sub>n</sub>; thus, the base-pair length of expanded repeats is given instead of repeat number.</p>
                        <p id="fn3">
							
                            <sup>c</sup>The SCA37 position is given at the reference (ATTTT)
                            <sub>n</sub> repeat, of which affected individuals have (ATTTC)
                            <sub>n</sub> inserted. The FAME1 position is given at the reference (TTTTA)
                            <sub>n</sub> repeat, of which affected individuals have (TTTCA)
                            <sub>n</sub> inserted.</p>
                        <p id="fn4">
							
                            <sup>d</sup>Ishiura 
                            <italic toggle="yes">et al</italic>.
                            <sup>
                                <xref ref-type="bibr" rid="ref-3">3</xref>
                            </sup> identified similar expansions associated with FAME6 and FAME7 but only in single families. The same TTTCA repeat insertion was observed in the intronic region of 
                            <italic toggle="yes">TNRC6A</italic> and 
                            <italic toggle="yes">RAPGEF2</italic>, respectively.</p>
                        <p id="fn5">
							
                            <sup>e</sup>The size of the FAME1 repeat is the estimated combined size of the expanded (TTTCA)
                            <sub>n</sub> insertion and (TTTTA)
                            <sub>n</sub> reference repeat.</p>
                        <p id="fn6">
							
                            <sup>f</sup>Different polyalanine expansions in the gene can be expanded.</p>
                    </fn>
                </table-wrap-foot>
            </table-wrap>
            <p>Repeat expansion tests are instigated by clinicians in response to a suspected clinical diagnosis. Detection of repeat expansions is performed by using methods such as polymerase chain reaction (PCR) for the shorter repeat expansions or Southern blot for longer repeats. Repeat expansion locus-specific PCR methods, such as repeat-primed PCR
                <sup>
                    <xref ref-type="bibr" rid="ref-4">4</xref>
                </sup>, have also been developed by individual laboratories and represent an active area of research in diagnostic methods
                <sup>
                    <xref ref-type="bibr" rid="ref-5">5</xref>
                </sup>. These methods are also able to accurately size repeat expansions.</p>
            <p>Genetic laboratories conduct a large number of tests for repeat expansion disorders, but the detection rate is low. Turnaround times are of the order of weeks or months. No comprehensive panel or testing method exists that simultaneously tests for all known repeat expansions using the current gold-standard detection methods of PCR and Southern blot.</p>
            <p>Next-generation sequencing (NGS), with either whole exome (WES) or whole genome (WGS) sequencing, is now a standard test for many individuals with a suspected genetic disorder. DNA sequencing analysis is a highly streamlined process that can be outsourced to one of the many clinically accredited sequencing laboratories worldwide. The analysis is performed by using sophisticated pipelines
                <sup>
                    <xref ref-type="bibr" rid="ref-6">6</xref>
                </sup>. Even with outsourced data, results are often delivered within a few weeks. If analysis is performed in-house, turnaround could be as fast as a week for WGS and a few days for WES, depending on the computer capacity available.</p>
            <p>Analysis of WES and WGS data is very efficient in the identification of single-nucleotide variants and indels but also can examine structural variation, such as copy number variation. Standard variant pipelines report mismatches of up to 50 base pairs (bp)
                <sup>
                    <xref ref-type="bibr" rid="ref-7">7</xref>
                </sup> and thus can identify only short STR alleles
                <sup>
                    <xref ref-type="bibr" rid="ref-8">8</xref>
                </sup>. Furthermore, these are often poorly described in the variant call format output files. Some improvements in the identification of STR variants came from larger indel detection methods such as DINDEL
                <sup>
                    <xref ref-type="bibr" rid="ref-9">9</xref>
                </sup> and PINDEL
                <sup>
                    <xref ref-type="bibr" rid="ref-10">10</xref>
                </sup>. Since 2000, several methods have also sought to specifically identify the lengths of STR alleles from short-read NGS data. One of the most recent methods is HipSTR
                <sup>
                    <xref ref-type="bibr" rid="ref-11">11</xref>
                </sup>, which uses an Expectation Maximization (EM) algorithm to determine the set of STR alleles present at a locus. The EM algorithm is combined with a local realignment step, and was found to outperform existing methods. However, all of these methods are constrained to STR alleles with repeat lengths smaller than the read length employed in the sequencing. Standard WGS short-read sequencing for the highest throughput sequencing platform, the Illumina HiSeq X Ten, uses a paired-end protocol with reads of 150 bp in length. WES is now also performed by using paired-end sequencing with reads of 150 bp; however, some data sets&#x2014;in particular, older data sets&#x2014;have shorter read lengths. Hence, many of the repeat expansion alleles that cause disease remain undetectable by these standard pipeline variant-calling methods. The ability to detect known and possibly novel repeat expansions with short-read sequencing data would be a valuable addition to any clinical genomics or diagnostic sequencing pipeline.</p>
            <p>Four new methods to detect repeat expansions have recently been described: ExpansionHunter
                <sup>
                    <xref ref-type="bibr" rid="ref-12">12</xref>
                </sup>, exSTRa
                <sup>
                    <xref ref-type="bibr" rid="ref-13">13</xref>
                </sup>, STRetch
                <sup>
                    <xref ref-type="bibr" rid="ref-14">14</xref>
                </sup>, and TREDPARSE
                <sup>
                    <xref ref-type="bibr" rid="ref-15">15</xref>
                </sup>. All four have demonstrated the ability to detect repeat expansions where the expanded allele size is greater than the length of standard short-read sequencing reads and even the read pair fragment length. In this review, we briefly outline the principles behind these methods, comparing their approaches. By introducing these methods, we hope to encourage researchers and clinical genomics facilities to incorporate them into their pipelines, as we believe it will improve molecular genetic diagnosis with the greatest impact to be expected for neurological disorders. We also discuss applications of these approaches beyond clinical genomics and finish with some comments regarding the potential of the developing long-read sequencing technologies for the detection of expanded alleles.</p>
        </sec>
        <sec>
            <title>How to detect repeat expansions with short-read data</title>
            <p>The repeat expansion detection methods discussed here all require paired-end sequencing data. Standard paired-end sequencing provides a pair of reads that flank a DNA fragment of about 350 bp in length. Library preparations can vary this DNA fragment size, and larger fragments are known to be advantageous for applications such as genome assembly, which could also be potentially useful for expansion detection. The two reads that comprise a read pair are sequenced in opposite directions, toward each other. Between the read pairs, there is typically a short sequence of DNA (of about 50 bp in length) that is not sequenced. The key to repeat expansion detection is to assess reads that are found to lie partially, or entirely, in an STR for their repeat content. This can be done heuristically (Expansion, exSTRa, and STRetch) or can be integrated into a likelihood model (TREDPARSE). Expanded alleles at an STR will contribute reads with more repeat content and more reads in total than reads stemming from the normal, unexpanded, allele (
                <xref ref-type="fig" rid="f1">Figure 1</xref>).</p>
            <fig fig-type="figure" id="f1" orientation="portrait" position="float">
                <label>Figure 1. </label>
                <caption>
                    <title>Detecting repeat expansions with short-read sequencing data.</title>
                    <p>Depicted are three scenarios: (
                        <bold>I</bold>) a short repeat expansion where the repeat expansion is less than 150 base pairs (bp), or smaller than a read; (
                        <bold>II</bold>) a medium-size repeat expansion where the repeat expansion is between 150 and 350 bp; and (
                        <bold>III</bold>) a large repeat expansion, where the repeat expansion is greater than 350 bp. For each of the three panels, 
                        <bold>I</bold>&#x2013;
                        <bold>III</bold>, the top line of DNA sequence depicts the reference sequence, and the bottom line depicts the (not known) repeat expansion size sequence. Red segments in reads signify repeat sequence. Evidence from reads varies according to the repeat size. For all three scenarios, there is information in reads that map into the repeat (A) but for scenario 
                        <bold>I</bold> occasional reads span the expanded allele (B), giving information about the size of the expanded allele. In scenario 
                        <bold>II</bold>, some read fragments can span expanded allele and can also be used for inference. For large expansions, some read fragments stem entirely from the expanded alleles. These may not be unambiguously mapped and are exploited only by ExpansionHunter (large motifs only) and TREDPARSE (based on fragment size information).</p>
                </caption>
                <graphic orientation="portrait" position="float" xlink:href="https://f1000research-files.f1000.com/manuscripts/15195/3a42f0b4-f995-4fa7-8174-b339e0544bbe_figure1.gif"/>
            </fig>
            <p>Key factors that will influence the ability to detect repeat expansions are (i) the library preparation protocol, (ii) the read length and likely also the DNA fragment length, and (iii) the depth of sequencing employed for the sample. These factors influence the number of reads that cover each STR locus. Tankard 
                <italic toggle="yes">et al</italic>.
                <sup>
                    <xref ref-type="bibr" rid="ref-13">13</xref>
                </sup> compared several library preparation protocols over the 21 known neurological STR loci, showing locus-specific effects for these factors (Supplementary Figure 1). In general, PCR-free WGS library preparation protocols yield the best data to allow repeat expansion detection, but even WES data could be successfully interrogated for repeat expansions for most of the known repeat expansion loci captured during library preparation
                <sup>
                    <xref ref-type="bibr" rid="ref-13">13</xref>
                </sup>.</p>
            <p>exSTRa and ExpansionHunter determine the repeat content of all reads mapped to a particular STR locus. This then forms the source data for their respective analyses. TREDPARSE includes the repeat content into its likelihood model and estimates the repeat motif number, similarly to HipSTR and lobSTR
                <sup>
                    <xref ref-type="bibr" rid="ref-17">17</xref>
                </sup>. For large expanded repeats, it is possible that entire DNA fragments lie within the STR. The paired-end reads that capture only repeat content either map to other regions in the genome where longer copies of this repeat are present in the genomic reference or remain unmapped for both reads of the read pair (
                <xref ref-type="fig" rid="f1">Figure 1</xref>). ExpansionHunter labels these reads as in-repeat reads, or IRRs, whereas exSTRa, STRetch, and TREDPARSE discard them from analysis. If the motif is long and sufficiently under-represented in the genome, as is the case with the hexamer GGGGCC 
                <italic toggle="yes">C9orf72</italic> expansion (
                <xref ref-type="table" rid="T1">Table 1</xref>), then there will be only a small number of alternate locations where reads containing the expanded allele could preferentially map to instead of the original locus. ExpansionHunter assesses the additional 29 sites with larger copy numbers of this hexamer repeat and incorporates this information into its likelihood to estimate the allele sizes of the individual.</p>
            <p>STRetch employs a different approach, whereby a new reference genome is proposed with additional decoy chromosomes containing artificially long versions of all repeat motif combinations. The decoy chromosomes provide an alternative mapping location for the expanded reads. This method requires the initially computationally expensive step of realignment to a new reference genome but results in a very natural statistical testing framework where relative read alignment between a candidate STR and its decoy are compared in a likelihood ratio test. However, it requires that the decoy chromosomes encode repeat motif representations that ensure that the reads of the expanded STR allele preferentially map there. Short repeat expansions, such as SCA6, where expanded alleles have as few as 21 repeat motifs, may preferentially map as insertions at their original location rather than to the (longer) expansion in the decoy chromosome, remaining undetected. In contrast, older, within-read only detection methods, such as lobSTR, can detect such expansions.</p>
            <p>ExpansionHunter and exSTRa do not require additional alignment steps. Instead, they interrogate existing alignments. STRetch requires re-alignment to the augmented reference genome, although alignment can potentially be performed in an 
                <italic toggle="yes">ad hoc</italic> manner by taking reads that have failed to align with the standard genome reference and aligning these solely to the set of decoy chromosomes. The effects of this approach have not yet been evaluated. TREDPARSE also has a potentially time-consuming local realignment step similar to HipSTR
                <sup>
                    <xref ref-type="bibr" rid="ref-11">11</xref>
                </sup>, which has yet to be evaluated in a genome-wide analysis. We refer the reader to each of the four articles for depictions of the types of read evidence that are used in each of the algorithms.</p>
            <p>The possibility of an expansion is assessed differently for each of the methods. exSTRa and STRetch rely on the availability of controls to allow them to determine whether an individual is an outlier with respect to their statistical measures. ExpansionHunter and TREDPARSE can be used on a single sample for known loci, making use of known thresholds and empirical distribution properties in the STR allele size to identify individuals with expansions. For novel repeat expansion loci, appropriate thresholds are not known and will require 
                <italic toggle="yes">post hoc</italic> testing of the allele size distribution in a control cohort to assess likely outlier individuals. This latter test is not currently implemented in ExpansionHunter or TREDPARSE.</p>
            <p>TREDPARSE makes use of a highly parameterized likelihood framework with a stuttering model and a local realignment step to infer allele sizes and determine the likelihood of pathogenicity. ExpansionHunter employs a much simpler likelihood model, which infers allele sizes and then uses allele thresholds to determine significance. STRetch applies a likelihood ratio test comparing the relative mapping of the reads for a known repeat to the expected genomic location or to its decoy chromosome containing that repeat. exSTRa uses a simple summary statistic combined with an outlier detection method and, like STRetch, uses a set of controls to apply permutation testing to assess the significance of the findings. Despite the variety of evaluation frameworks and statistical approaches, ExpansionHunter, exSTRa, and STRetch were able to detect almost all of the known repeat expansions that they were tested on whereas TREDPARSE was found to produce results that were validated with alternative methods such as long-read sequencing. We refer readers to the respective articles for details of the variety of performance evaluations that have been employed by these methods. We summarize the properties of the four algorithms in 
                <xref ref-type="table" rid="T2">Table 2</xref>.</p>
            <table-wrap id="T2" orientation="portrait" position="anchor">
                <label>Table 2. </label>
                <caption>
                    <title>Summary of computational methods, evaluation framework, and limitations for ExpansionHunter, exSTRa, STRetch, and TREDPARSE.</title>
                </caption>
                <table content-type="article-table" frame="hsides">
                    <thead>
                        <tr>
                            <th align="left" colspan="1" rowspan="1" valign="top">Software</th>
                            <th align="left" colspan="1" rowspan="1" valign="top">Publication</th>
                            <th align="left" colspan="1" rowspan="1" valign="top">Computational
                                <break/>burden
                                <sup>
                                    <xref ref-type="other" rid="fn7">a</xref>
                                </sup>:
                                <break/>known loci/
                                <break/>genome-wide</th>
                            <th align="left" colspan="1" rowspan="1" valign="top">Statistical test</th>
                            <th align="left" colspan="1" rowspan="1" valign="top">Reported
                                <break/>WGS/WES
                                <break/>analysis
                                <break/>capability</th>
                            <th align="left" colspan="1" rowspan="1" valign="top">Software
                                <break/>ease of
                                <break/>use</th>
                            <th align="left" colspan="1" rowspan="1" valign="top">Ability to
                                <break/>search
                                <break/>genome-
                                <break/>wide</th>
                            <th align="left" colspan="1" rowspan="1" valign="top">Graphical
                                <break/>output</th>
                            <th align="left" colspan="1" rowspan="1" valign="top">Length of STR
                                <break/>expansion
                                <break/>detection bias</th>
                        </tr>
                    </thead>
                    <tbody>
                        <tr>
                            <td colspan="1" rowspan="1" valign="top">ExpansionHunter</td>
                            <td colspan="1" rowspan="1" valign="top">Dolzhenko
                                <break/>
                                <italic toggle="yes">et al</italic>.
                                <sup>
                                    <xref ref-type="bibr" rid="ref-12">12</xref>
                                </sup>,
                                <break/>
                                <italic toggle="yes">Genome</italic>
                                <break/>
                                <italic toggle="yes">Research</italic>
                                <break/>2017</td>
                            <td colspan="1" rowspan="1" valign="top">Low/Low</td>
                            <td colspan="1" rowspan="1" valign="top">None &#x2013; estimates
                                <break/>allele sizes.
                                <break/>Significance
                                <break/>determined on
                                <break/>the basis of
                                <break/>thresholds
                                <sup>
                                    <xref ref-type="other" rid="fn7">b</xref>
                                </sup>.</td>
                            <td colspan="1" rowspan="1" valign="top">WGS</td>
                            <td colspan="1" rowspan="1" valign="top">High</td>
                            <td colspan="1" rowspan="1" valign="top">Possible</td>
                            <td colspan="1" rowspan="1" valign="top">No</td>
                            <td colspan="1" rowspan="1" valign="top">Repeats with
                                <break/>long motifs (e.g.,
                                <break/>c9orf72
                                <sup>
                                    <xref ref-type="other" rid="fn7">c</xref>
                                </sup>) gain
                                <break/>extra evidence
                                <break/>for expansion
                                <break/>with usage of
                                <break/>in-repeat reads
                                <break/>(IRRs)</td>
                        </tr>
                        <tr>
                            <td colspan="1" rowspan="1" valign="top">exSTRa</td>
                            <td colspan="1" rowspan="1" valign="top">Tankard
                                <break/>
                                <italic toggle="yes">et al</italic>.
                                <sup>
                                    <xref ref-type="bibr" rid="ref-13">13</xref>
                                </sup>,
                                <break/>bioRxiv, 2017</td>
                            <td colspan="1" rowspan="1" valign="top">Low/Medium</td>
                            <td colspan="1" rowspan="1" valign="top">Permutation
                                <break/>based outlier
                                <break/>detection test</td>
                            <td colspan="1" rowspan="1" valign="top">WGS and
                                <break/>WES</td>
                            <td colspan="1" rowspan="1" valign="top">Medium</td>
                            <td colspan="1" rowspan="1" valign="top">Possible</td>
                            <td colspan="1" rowspan="1" valign="top">Yes</td>
                            <td colspan="1" rowspan="1" valign="top">No known bias</td>
                        </tr>
                        <tr>
                            <td colspan="1" rowspan="1" valign="top">STRetch</td>
                            <td colspan="1" rowspan="1" valign="top">Dashnow
                                <break/>
                                <italic toggle="yes">et al</italic>.
                                <sup>
                                    <xref ref-type="bibr" rid="ref-14">14</xref>
                                </sup>,
                                <break/>bioRxiv, 2017</td>
                            <td colspan="1" rowspan="1" valign="top">High/Medium</td>
                            <td colspan="1" rowspan="1" valign="top">Likelihood ratio
                                <break/>test with reads
                                <break/>mapping to
                                <break/>decoy. Estimates
                                <break/>allele sizes.</td>
                            <td colspan="1" rowspan="1" valign="top">WGS</td>
                            <td colspan="1" rowspan="1" valign="top">Low</td>
                            <td colspan="1" rowspan="1" valign="top">Easy</td>
                            <td colspan="1" rowspan="1" valign="top">No</td>
                            <td colspan="1" rowspan="1" valign="top">Short expansions
                                <break/>may not map
                                <break/>to the decoy
                                <break/>chromosomes
                                <break/>and remain
                                <break/>undetected (e.g.,
                                <break/>SCA6
                                <sup>
                                    <xref ref-type="other" rid="fn7">d</xref>
                                </sup>)</td>
                        </tr>
                        <tr>
                            <td colspan="1" rowspan="1" valign="top">TREDPARSE</td>
                            <td colspan="1" rowspan="1" valign="top">Tang 
                                <italic toggle="yes">et al</italic>.
                                <sup>
                                    <xref ref-type="bibr" rid="ref-15">15</xref>
                                </sup>,
                                <break/>AJHG, 2017</td>
                            <td colspan="1" rowspan="1" valign="top">Low/Unknown</td>
                            <td colspan="1" rowspan="1" valign="top">Likelihood of
                                <break/>pathogenicity,
                                <break/>genetic model,
                                <break/>estimates allele
                                <break/>sizes
                                <sup>
                                    <xref ref-type="other" rid="fn7">b</xref>
                                </sup>
                            </td>
                            <td colspan="1" rowspan="1" valign="top">WGS</td>
                            <td colspan="1" rowspan="1" valign="top">High</td>
                            <td colspan="1" rowspan="1" valign="top">Possible</td>
                            <td colspan="1" rowspan="1" valign="top">Yes</td>
                            <td colspan="1" rowspan="1" valign="top">Does not detect
                                <break/>expansions
                                <break/>that exceed
                                <break/>its detection
                                <break/>threshold (300
                                <break/>repeats)</td>
                        </tr>
                    </tbody>
                </table>
                <table-wrap-foot>
                    <fn>
                        <p id="fn7">
							
                            <sup>a</sup>Computational burden has been split into two components: known loci&#x2014;a small subset of all short tandem repeat (STR) loci&#x2014;and genome-wide, representing thousands of STR loci. 
                            <sup>b</sup>Requires prior information for STR in terms of allele size to aid statistical test. 
                            <sup>c</sup>The C9orf72 repeat expansion is a hexamer repeat. 
                            <sup>d</sup>SCA6 is the smallest repeat expansion currently known. WES, whole exome sequencing; WGS, whole genome sequencing.</p>
                    </fn>
                </table-wrap-foot>
            </table-wrap>
        </sec>
        <sec>
            <title>The role of known repeat expansions in related disorders such as epilepsy</title>
            <p>Several genes that contain disease-causing repeat expansions that cause ataxias have been implicated in other disorders such as epilepsy and migraine. For example, the clinical spectrum of 
                <italic toggle="yes">C9orf72</italic> has broadened to encompass Huntington&#x2019;s disease-like disorder
                <sup>
                    <xref ref-type="bibr" rid="ref-18">18</xref>
                </sup>. The new repeat expansion detection methods described here permit an investigation of the role of all known repeat expansions in cohorts of individuals with related phenotypes, such as epilepsy and migraine. These can be examined using existing short-read data with the new repeat expansion detection methods, potentially providing new insights.</p>
        </sec>
        <sec>
            <title>Evaluating the variation of repeat expansion short tandem repeats in cohorts</title>
            <p>STRs vary in frequency and length distributions between different ethnic groups because of founder effects. Willems 
                <italic toggle="yes">et al</italic>.
                <sup>
                    <xref ref-type="bibr" rid="ref-11">11</xref>
                </sup> used lobSTR to generate an online STR catalogue (
                <ext-link ext-link-type="uri" xlink:href="http://strcat.teamerlich.org">http://strcat.teamerlich.org</ext-link>) of about 700,000 STRs, which displays STR repeat number distributions and, where possible, stratification by the 14 ethnicities represented in the 1000 Genomes project
                <sup>
                    <xref ref-type="bibr" rid="ref-19">19</xref>
                </sup>. Many STR loci were found to display ethnicity-specific distributions. Several repeat expansion STRs also show ethnicity effects. For example, the CAG repeat in the ataxin 7 gene (
                <italic toggle="yes">ATXN7</italic>) displays multiple founder events in Scandinavia, Mexico, and South Africa/Zimbabwe
                <sup>
                    <xref ref-type="bibr" rid="ref-20">20</xref>
                </sup>, and multiple founder events have also been documented for Huntington&#x2019;s disease
                <sup>
                    <xref ref-type="bibr" rid="ref-1">1</xref>
                </sup>.</p>
            <p>The remarkable reduction in the price of short-read sequencing has led to the sequencing of greater numbers of individuals and new study cohorts. By making use of the methods reviewed here, analyses of the genetic composition of pathogenic STR loci in hitherto unexamined cohorts will be possible. An understanding of the natural variation of both normal and repeat expansion alleles in different populations will also be helpful to refine the statistical tests of the methods, including providing more accurate information for the determination of significance thresholds and prior information required for testing.</p>
            <p>Prospective and retrospective analysis of sequencing data sets with the repeat expansion detection methods described here should provide clinically actionable outcomes. Also exciting are the research opportunities in our understanding of STRs. For example, spinocerebellar ataxia-8 (SCA8, OMIM #608768) is one of several poorly understood disorders caused by a repeat expansion
                <sup>
                    <xref ref-type="bibr" rid="ref-21">21</xref>
                </sup>. The repeat is bidirectionally transcribed
                <sup>
                    <xref ref-type="bibr" rid="ref-22">22</xref>
                </sup>. Additionally, its clinical implications are still uncertain, and the understanding of its clinical spectrum and penetrance is incomplete. Using large population-based and disease-ascertained cohorts containing thousands of individuals, we will be able to gather hundreds of detected repeat expansions for these repeats, allowing a more precise determination of penetrance and potential co-morbidities. To determine proof of principle, Tang 
                <italic toggle="yes">et al</italic>.
                <sup>
                    <xref ref-type="bibr" rid="ref-15">15</xref>
                </sup> profiled 12,632 individuals, identifying 132 individuals with larger-than-normal-range STR alleles at 15 different known repeat expansion STR loci.</p>
        </sec>
        <sec>
            <title>Detecting novel repeat expansions</title>
            <p>It is likely that novel repeat expansion loci are awaiting discovery. In OMIM, there are several reported SCA loci, such as SCA32 (OMIM %613909, 7q32-q33), that as yet have no determined genetic cause. Families with linkage to SCA25
                <sup>
                    <xref ref-type="bibr" rid="ref-23">23</xref>,
                    <xref ref-type="bibr" rid="ref-24">24</xref>
                </sup> (OMIM %608703, 2p21-p13) furthermore report the phenomenon of anticipation. Anticipation is a hallmark of repeat expansions since these can become more unstable (and usually larger) with subsequent meioses, after the initial expansion step, thus leading to earlier ages of onset or more severe symptoms (or both) in affected individuals from more recent generations in the pedigree.</p>
            <p>ExpansionHunter, TREDPARSE, exSTRa, and STRetch are all able to detect novel repeat expansions but require that the putative expansion STRs be explicitly specified. Hence, all methods rely on 
                <italic toggle="yes">a priori</italic> knowledge of STR loci to be examined. STR sets of interest can be assembled by using annotation of STRs from Tandem Repeats Finder results
                <sup>
                    <xref ref-type="bibr" rid="ref-16">16</xref>
                </sup> and appropriate search parameters. Relevant parameters such as motif length, existing repeats, and purity of the repeat will determine the number of STRs detected in the reference genome being examined. The expansion detection performance of methods will be influenced by the genomic composition of the STRs, with complex STRs, with features such as impure repeats or multiple repeat motifs comprising a single STR likely to be more difficult to detect.</p>
        </sec>
        <sec>
            <title>Implementation limitations</title>
            <p>Although all four of the methods discussed (TREDPARSE, exSTRa, ExpansionHunter, and STRetch) will benefit from further development, we advocate the immediate implementation of these methods to any existing analysis pipelines for WES, WGS, or even suitable gene panels. The initial benefit will be through the examination of retrospective and prospectively sequenced individuals for all known repeat expansion loci to prevent a missed diagnosis due to a known expansion
                <sup>
                    <xref ref-type="bibr" rid="ref-18">18</xref>,
                    <xref ref-type="bibr" rid="ref-25">25</xref>
                </sup>. These missed molecular diagnoses are an important contributor to increasing diagnostic yield in clinical genomic sequencing
                <sup>
                    <xref ref-type="bibr" rid="ref-26">26</xref>
                </sup>. Individuals detected to have a repeat expansion with one, or more, of ExpansionHunter, TREDPARSE, exSTRa, or STRetch should undergo the gold-standard assays at a certified laboratory, when possible, or at a research laboratory specializing in the repeat expansion detection of that STR locus.</p>
            <p>Although all four publications describe the application of the methods to a variety of known repeat expansions, none of them encompasses a complete list of known loci. Indeed, at the moment, there are several repeat expansion loci that have never been tested with any of the four computational approaches. These include EPM1 (
                <italic toggle="yes">CSTB</italic>, OMIM #254800), HDL2 (
                <italic toggle="yes">JPH3</italic>, OMIM #606438), SCA10 (
                <italic toggle="yes">ATXN10</italic>, OMIM #603516), and SCA12 (
                <italic toggle="yes">PPP2R2B</italic>, OMIM #604326). It is likely that the algorithms will be able to efficiently interrogate most or all of these loci, similar to the majority of other repeat expansion loci that have been tested. All four of these loci achieve good coverage with PCR-free WGS protocols
                <sup>
                    <xref ref-type="bibr" rid="ref-13">13</xref>
                </sup>.</p>
            <p>Some STR loci such as FRAXA (
                <italic toggle="yes">FMR1</italic>, OMIM #300624) are highly adversely affected by PCR amplification bias introduced during library preparation and could be assessed only with a PCR-free library preparation protocol
                <sup>
                    <xref ref-type="bibr" rid="ref-13">13</xref>
                </sup>. FRAXE (
                <italic toggle="yes">FMR2</italic>, OMIM #309548) remains refractory to capture with short-read sequencing, regardless of the protocol used, and is not currently assessable with any of the repeat expansion detection methods.</p>
        </sec>
        <sec>
            <title>The continuing role of gold-standard repeat detection methods such as Southern blots and repeat-primed polymerase chain reaction</title>
            <p>Repeat expansion detection methods such as Southern blots, PCRs, and repeat primed PCR will not be supplanted soon, even with these developments in the detection of repeat expansions with NGS. First, the latter should be seen as a screening method, requiring validation with the gold-standard methods. Second, the NGS-based methods cannot, as yet, accurately and reliably size repeat expansions. Furthermore, it is unlikely that the short-read methods will be able to do so, even in the future, since they rely on imperfect relationships between read numbers and repeat allele length, which is more difficult for larger repeats.</p>
            <p>Comprehensive prospective studies will also be needed to compare the cost and efficacy of the NGS-based methods for screening for repeat expansions before NGS-based screening approaches are adopted.</p>
            <p>Long repeat expansion alleles (&gt;500 bp), such as those found in DM1 (
                <italic toggle="yes">DMPK</italic>, OMIM #160900), FRAXA, FRAXE, and SCA10 patients, are difficult to detect with standard diagnostic tests. The use of NGS-based detection of repeat expansions could improve overall diagnostic yield for these ultra-long expansions since these methods have been demonstrated to perform well for loci such as DM1 and FRDA (
                <italic toggle="yes">FXN</italic>, OMIM #229300). NGS-based repeat expansion detection may also be more accessible for some patients than current gold-standard methods because NGS is a commonly used, robust platform, which has seen a continuing drop in costs and even wider availability.</p>
        </sec>
        <sec>
            <title>The impact of long-read sequencing</title>
            <p>Although advances in repeat expansion detection with short-read sequencing are exciting, the next wave of discovery, owing to the increased quality and rapidly decreasing costs of long-read sequencing, is already upon us. Long-read sequencing technologies such as PacBio and Nanopore sequencing are rapidly gaining popularity and attracting significant bioinformatics interest to improve analysis pipelines. The reported read lengths are in the tens of thousands of base pairs rather than the hundreds. As such, these long-read sequencing platforms will sequence through STR loci for both normal and expanded alleles. This will be particularly useful for complex expanded alleles, where the repeat may be interrupted multiple times. Neither NGS-based methods nor current diagnostics methods do well in these cases. Sequencing error rates are currently still much higher for long-read sequencing than for short-read sequencing and will require further work to be able to reliably determine repeat lengths
                <sup>
                    <xref ref-type="bibr" rid="ref-27">27</xref>
                </sup>. Nanopore sequencing has the additional advantage of having no GC coverage bias because there is no DNA polymerization step
                <sup>
                    <xref ref-type="bibr" rid="ref-28">28</xref>
                </sup>. GC bias in repeats or flanking regions (or both) can lead to a bias in allele amplification with bias observed both for and against the expanded allele
                <sup>
                    <xref ref-type="bibr" rid="ref-13">13</xref>
                </sup>.</p>
            <p>Further novel repeat expansions are doubtlessly awaiting discovery. Their discovery will be aided by novel analytical approaches such as those reviewed here. They will likely require support from other sequencing methods, including long-read sequencing and RNA-seq
                <sup>
                    <xref ref-type="bibr" rid="ref-26">26</xref>
                </sup>. The biological mechanisms underpinning these diseases are a separately fascinating and rapidly broadening field of research. Additionally, new technologies are leading to renewed hope for potential treatments. A recent publication described the elimination of the toxic effect of the CTG expansion in DM1 (OMIM #160900) with RNA targeting Cas9 excision
                <sup>
                    <xref ref-type="bibr" rid="ref-29">29</xref>
                </sup>. This is an exciting time for research in repeat expansion disorders.</p>
        </sec>
    </body>
    <back>
        <ref-list>
            <ref id="ref-1">
                <label>1</label>
                <mixed-citation publication-type="journal">
                    <person-group person-group-type="author">
				
                        <name name-style="western">
                            <surname>Warby</surname>
                            <given-names>SC</given-names>
                        </name>
				
                        <name name-style="western">
                            <surname>Visscher</surname>
                            <given-names>H</given-names>
                        </name>
				
                        <name name-style="western">
                            <surname>Collins</surname>
                            <given-names>JA</given-names>
                        </name>
				
                        <etal/>
			</person-group>:
                    <article-title>HTT haplotypes contribute to differences in Huntington disease prevalence between Europe and East Asia.</article-title>
                    <source>
				
                        <italic toggle="yes">Eur J Hum Genet.</italic>
			</source>
                    <year>2011</year>;<volume>19</volume>(<issue>5</issue>):<fpage>561</fpage>&#x2013;<lpage>6</lpage>.
                    <pub-id pub-id-type="pmid">21248742</pub-id>
                    <pub-id pub-id-type="doi">10.1038/ejhg.2010.229</pub-id>
                    <pub-id pub-id-type="pmcid">3083615</pub-id>
                </mixed-citation>
            </ref>
            <ref id="ref-2">
                <label>2</label>
                <mixed-citation publication-type="journal">
                    <person-group person-group-type="author">
				
                        <name name-style="western">
                            <surname>Mootha</surname>
                            <given-names>VV</given-names>
                        </name>
				
                        <name name-style="western">
                            <surname>Gong</surname>
                            <given-names>X</given-names>
                        </name>
				
                        <name name-style="western">
                            <surname>Ku</surname>
                            <given-names>HC</given-names>
                        </name>
				
                        <etal/>
			</person-group>:
                    <article-title>Association and familial segregation of CTG18.1 trinucleotide repeat expansion of 
                        <italic toggle="yes">TCF4</italic> gene in Fuchs' endothelial corneal dystrophy.</article-title>
                    <source>
				
                        <italic toggle="yes">Invest Ophthalmol Vis Sci.</italic>
			</source>
                    <year>2014</year>;<volume>55</volume>(<issue>1</issue>):<fpage>33</fpage>&#x2013;<lpage>42</lpage>.
                    <pub-id pub-id-type="pmid">24255041</pub-id>
                    <pub-id pub-id-type="doi">10.1167/iovs.13-12611</pub-id>
                    <pub-id pub-id-type="pmcid">3880006</pub-id>
                </mixed-citation>
            </ref>
            <ref id="ref-3">
                <label>3</label>
                <mixed-citation publication-type="journal">
                    <person-group person-group-type="author">
				
                        <name name-style="western">
                            <surname>Ishiura</surname>
                            <given-names>H</given-names>
                        </name>
				
                        <name name-style="western">
                            <surname>Doi</surname>
                            <given-names>K</given-names>
                        </name>
				
                        <name name-style="western">
                            <surname>Mitsui</surname>
                            <given-names>J</given-names>
                        </name>
				
                        <etal/>
			</person-group>:
                    <article-title>Expansions of intronic TTTCA and TTTTA repeats in benign adult familial myoclonic epilepsy.</article-title>
                    <source>
				
                        <italic toggle="yes">Nat Genet.</italic>
			</source>
                    <year>2018</year>;<volume>50</volume>(<issue>4</issue>):<fpage>581</fpage>&#x2013;<lpage>90</lpage>.
                    <pub-id pub-id-type="pmid">29507423</pub-id>
                    <pub-id pub-id-type="doi">10.1038/s41588-018-0067-2</pub-id>
                </mixed-citation>
                <note>
                    <p>
                        <ext-link ext-link-type="uri" xlink:href="https://f1000.com/prime/732796809">F1000 Recommendation</ext-link>
                    </p>
                </note>
            </ref>
            <ref id="ref-4">
                <label>4</label>
                <mixed-citation publication-type="journal">
                    <person-group person-group-type="author">
				
                        <name name-style="western">
                            <surname>Warner</surname>
                            <given-names>JP</given-names>
                        </name>
				
                        <name name-style="western">
                            <surname>Barron</surname>
                            <given-names>LH</given-names>
                        </name>
				
                        <name name-style="western">
                            <surname>Goudie</surname>
                            <given-names>D</given-names>
                        </name>
				
                        <etal/>
			</person-group>:
                    <article-title>A general method for the detection of large CAG repeat expansions by fluorescent PCR.</article-title>
                    <source>
				
                        <italic toggle="yes">J Med Genet.</italic>
			</source>
                    <year>1996</year>;<volume>33</volume>(<issue>12</issue>):<fpage>1022</fpage>&#x2013;<lpage>6</lpage>.
                    <pub-id pub-id-type="pmid">9004136</pub-id>
                    <pub-id pub-id-type="doi">10.1136/jmg.33.12.1022</pub-id>
                    <pub-id pub-id-type="pmcid">1050815</pub-id>
                </mixed-citation>
            </ref>
            <ref id="ref-5">
                <label>5</label>
                <mixed-citation publication-type="journal">
                    <person-group person-group-type="author">
				
                        <name name-style="western">
                            <surname>Zhao</surname>
                            <given-names>M</given-names>
                        </name>
				
                        <name name-style="western">
                            <surname>Cheah</surname>
                            <given-names>FSH</given-names>
                        </name>
				
                        <name name-style="western">
                            <surname>Chen</surname>
                            <given-names>M</given-names>
                        </name>
				
                        <etal/>
			</person-group>:
                    <article-title>Improved high sensitivity screen for Huntington disease using a one-step triplet-primed PCR and melting curve assay.</article-title>
                    <source>
				
                        <italic toggle="yes">PLoS One.</italic>
			</source>
                    <year>2017</year>;<volume>12</volume>(<issue>7</issue>):<fpage>e0180984</fpage>.
                    <pub-id pub-id-type="pmid">28700716</pub-id>
                    <pub-id pub-id-type="doi">10.1371/journal.pone.0180984</pub-id>
                    <pub-id pub-id-type="pmcid">5507316</pub-id>
                </mixed-citation>
            </ref>
            <ref id="ref-6">
                <label>6</label>
                <mixed-citation publication-type="journal">
                    <person-group person-group-type="author">
				
                        <name name-style="western">
                            <surname>Sadedin</surname>
                            <given-names>SP</given-names>
                        </name>
				
                        <name name-style="western">
                            <surname>Dashnow</surname>
                            <given-names>H</given-names>
                        </name>
				
                        <name name-style="western">
                            <surname>James</surname>
                            <given-names>PA</given-names>
                        </name>
				
                        <etal/>
			</person-group>:
                    <article-title>Cpipe: a shared variant detection pipeline designed for diagnostic settings.</article-title>
                    <source>
				
                        <italic toggle="yes">Genome Med.</italic>
			</source>
                    <year>2015</year>;<volume>7</volume>(<issue>1</issue>):<fpage>68</fpage>.
                    <pub-id pub-id-type="pmid">26217397</pub-id>
                    <pub-id pub-id-type="doi">10.1186/s13073-015-0191-x</pub-id>
                    <pub-id pub-id-type="pmcid">4515933</pub-id>
                </mixed-citation>
            </ref>
            <ref id="ref-7">
                <label>7</label>
                <mixed-citation publication-type="journal">
                    <person-group person-group-type="author">
				
                        <name name-style="western">
                            <surname>Hasan</surname>
                            <given-names>MS</given-names>
                        </name>
				
                        <name name-style="western">
                            <surname>Wu</surname>
                            <given-names>X</given-names>
                        </name>
				
                        <name name-style="western">
                            <surname>Zhang</surname>
                            <given-names>L</given-names>
                        </name>
			</person-group>:
                    <article-title>Performance evaluation of indel calling tools using real short-read data.</article-title>
                    <source>
				
                        <italic toggle="yes">Hum Genomics.</italic>
			</source>
                    <year>2015</year>;<volume>9</volume>:<fpage>20</fpage>.
                    <pub-id pub-id-type="pmid">26286629</pub-id>
                    <pub-id pub-id-type="doi">10.1186/s40246-015-0042-2</pub-id>
                    <pub-id pub-id-type="pmcid">4545535</pub-id>
                </mixed-citation>
            </ref>
            <ref id="ref-8">
                <label>8</label>
                <mixed-citation publication-type="journal">
                    <person-group person-group-type="author">
				
                        <name name-style="western">
                            <surname>McKenna</surname>
                            <given-names>A</given-names>
                        </name>
				
                        <name name-style="western">
                            <surname>Hanna</surname>
                            <given-names>M</given-names>
                        </name>
				
                        <name name-style="western">
                            <surname>Banks</surname>
                            <given-names>E</given-names>
                        </name>
				
                        <etal/>
			</person-group>:
                    <article-title>The Genome Analysis Toolkit: a MapReduce framework for analyzing next-generation DNA sequencing data.</article-title>
                    <source>
				
                        <italic toggle="yes">Genome Res.</italic>
			</source>
                    <year>2010</year>;<volume>20</volume>(<issue>9</issue>):<fpage>1297</fpage>&#x2013;<lpage>303</lpage>.
                    <pub-id pub-id-type="pmid">20644199</pub-id>
                    <pub-id pub-id-type="doi">10.1101/gr.107524.110</pub-id>
                    <pub-id pub-id-type="pmcid">2928508</pub-id>
                </mixed-citation>
            </ref>
            <ref id="ref-9">
                <label>9</label>
                <mixed-citation publication-type="journal">
                    <person-group person-group-type="author">
				
                        <name name-style="western">
                            <surname>Albers</surname>
                            <given-names>CA</given-names>
                        </name>
				
                        <name name-style="western">
                            <surname>Lunter</surname>
                            <given-names>G</given-names>
                        </name>
				
                        <name name-style="western">
                            <surname>MacArthur</surname>
                            <given-names>DG</given-names>
                        </name>
				
                        <etal/>
			</person-group>:
                    <article-title>Dindel: accurate indel calls from short-read data.</article-title>
                    <source>
				
                        <italic toggle="yes">Genome Res.</italic>
			</source>
                    <year>2011</year>;<volume>21</volume>(<issue>6</issue>):<fpage>961</fpage>&#x2013;<lpage>73</lpage>.
                    <pub-id pub-id-type="pmid">20980555</pub-id>
                    <pub-id pub-id-type="doi">10.1101/gr.112326.110</pub-id>
                    <pub-id pub-id-type="pmcid">3106329</pub-id>
                </mixed-citation>
            </ref>
            <ref id="ref-10">
                <label>10</label>
                <mixed-citation publication-type="journal">
                    <person-group person-group-type="author">
				
                        <name name-style="western">
                            <surname>Ye</surname>
                            <given-names>K</given-names>
                        </name>
				
                        <name name-style="western">
                            <surname>Schulz</surname>
                            <given-names>MH</given-names>
                        </name>
				
                        <name name-style="western">
                            <surname>Long</surname>
                            <given-names>Q</given-names>
                        </name>
				
                        <etal/>
			</person-group>:
                    <article-title>Pindel: a pattern growth approach to detect break points of large deletions and medium sized insertions from paired-end short reads.</article-title>
                    <source>
				
                        <italic toggle="yes">Bioinformatics.</italic>
			</source>
                    <year>2009</year>;<volume>25</volume>(<issue>21</issue>):<fpage>2865</fpage>&#x2013;<lpage>71</lpage>.
                    <pub-id pub-id-type="pmid">19561018</pub-id>
                    <pub-id pub-id-type="doi">10.1093/bioinformatics/btp394</pub-id>
                    <pub-id pub-id-type="pmcid">2781750</pub-id>
                </mixed-citation>
            </ref>
            <ref id="ref-11">
                <label>11</label>
                <mixed-citation publication-type="journal">
                    <person-group person-group-type="author">
				
                        <name name-style="western">
                            <surname>Willems</surname>
                            <given-names>T</given-names>
                        </name>
				
                        <name name-style="western">
                            <surname>Zielinski</surname>
                            <given-names>D</given-names>
                        </name>
				
                        <name name-style="western">
                            <surname>Yuan</surname>
                            <given-names>J</given-names>
                        </name>
				
                        <etal/>
			</person-group>:
                    <article-title>Genome-wide profiling of heritable and 
                        <italic toggle="yes">de novo</italic> STR variations.</article-title>
                    <source>
				
                        <italic toggle="yes">Nat Methods.</italic>
			</source>
                    <year>2017</year>;<volume>14</volume>(<issue>6</issue>):<fpage>590</fpage>&#x2013;<lpage>2</lpage>.
                    <pub-id pub-id-type="pmid">28436466</pub-id>
                    <pub-id pub-id-type="doi">10.1038/nmeth.4267</pub-id>
                    <pub-id pub-id-type="pmcid">5482724</pub-id>
                </mixed-citation>
                <note>
                    <p>
                        <ext-link ext-link-type="uri" xlink:href="https://f1000.com/prime/727556430">F1000 Recommendation</ext-link>
                    </p>
                </note>
            </ref>
            <ref id="ref-12">
                <label>12</label>
                <mixed-citation publication-type="journal">
                    <person-group person-group-type="author">
				
                        <name name-style="western">
                            <surname>Dolzhenko</surname>
                            <given-names>E</given-names>
                        </name>
				
                        <name name-style="western">
                            <surname>van Vugt</surname>
                            <given-names>JJFA</given-names>
                        </name>
					
                        <name name-style="western">
                            <surname>Shaw</surname>
                            <given-names>RJ</given-names>
                        </name>
				
                        <etal/>
			</person-group>:
                    <article-title>Detection of long repeat expansions from PCR-free whole-genome sequence data.</article-title>
                    <source>
				
                        <italic toggle="yes">Genome Res.</italic>
			</source>
                    <year>2017</year>;<volume>27</volume>(<issue>11</issue>):<fpage>1895</fpage>&#x2013;<lpage>903</lpage>.
                    <pub-id pub-id-type="pmid">28887402</pub-id>
                    <pub-id pub-id-type="doi">10.1101/gr.225672.117</pub-id>
                    <pub-id pub-id-type="pmcid">5668946</pub-id>
                </mixed-citation>
                <note>
                    <p>
                        <ext-link ext-link-type="uri" xlink:href="https://f1000.com/prime/730933359">F1000 Recommendation</ext-link>
                    </p>
                </note>
            </ref>
            <ref id="ref-13">
                <label>13</label>
                <mixed-citation publication-type="journal">
                    <person-group person-group-type="author">
				
                        <name name-style="western">
                            <surname>Tankard</surname>
                            <given-names>RM</given-names>
                        </name>
				
                        <name name-style="western">
                            <surname>Delatycki</surname>
                            <given-names>MB</given-names>
                        </name>
				
                        <name name-style="western">
                            <surname>Lockhart</surname>
                            <given-names>PJ</given-names>
                        </name>
				
                        <etal/>
			</person-group>:
                    <article-title>Detecting known repeat expansions with standard protocol next generation sequencing, towards developing a single screening test for neurological repeat expansion disorders.</article-title>
                    <source>
				
                        <italic toggle="yes">bioRxiv.</italic>
			</source>
                    <year>2017</year>.
                    <pub-id pub-id-type="doi">10.1101/157792</pub-id>
                </mixed-citation>
            </ref>
            <ref id="ref-14">
                <label>14</label>
                <mixed-citation publication-type="journal">
                    <person-group person-group-type="author">
				
                        <name name-style="western">
                            <surname>Dashnow</surname>
                            <given-names>H</given-names>
                        </name>
				
                        <name name-style="western">
                            <surname>Lek</surname>
                            <given-names>M</given-names>
                        </name>
				
                        <name name-style="western">
                            <surname>Phipson</surname>
                            <given-names>B</given-names>
                        </name>
				
                        <etal/>
			</person-group>:
                    <article-title>STRetch: detecting and discovering pathogenic short tandem repeats expansions.</article-title>
                    <source>
				
                        <italic toggle="yes">bioRxiv.</italic>
			</source>
                    <year>2017</year>.
                    <pub-id pub-id-type="doi">10.1101/159228</pub-id>
                </mixed-citation>
            </ref>
            <ref id="ref-15">
                <label>15</label>
                <mixed-citation publication-type="journal">
                    <person-group person-group-type="author">
				
                        <name name-style="western">
                            <surname>Tang</surname>
                            <given-names>H</given-names>
                        </name>
				
                        <name name-style="western">
                            <surname>Kirkness</surname>
                            <given-names>EF</given-names>
                        </name>
				
                        <name name-style="western">
                            <surname>Lippert</surname>
                            <given-names>C</given-names>
                        </name>
				
                        <etal/>
			</person-group>:
                    <article-title>Profiling of Short-Tandem-Repeat Disease Alleles in 12,632 Human Whole Genomes.</article-title>
                    <source>
				
                        <italic toggle="yes">Am J Hum Genet.</italic>
			</source>
                    <year>2017</year>;<volume>101</volume>(<issue>5</issue>):<fpage>700</fpage>&#x2013;<lpage>15</lpage>.
                    <pub-id pub-id-type="pmid">29100084</pub-id>
                    <pub-id pub-id-type="doi">10.1016/j.ajhg.2017.09.013</pub-id>
                    <pub-id pub-id-type="pmcid">5673627</pub-id>
                </mixed-citation>
                <note>
                    <p>
                        <ext-link ext-link-type="uri" xlink:href="https://f1000.com/prime/732071174">F1000 Recommendation</ext-link>
                    </p>
                </note>
            </ref>
            <ref id="ref-16">
                <label>16</label>
                <mixed-citation publication-type="journal">
                    <person-group person-group-type="author">
				
                        <name name-style="western">
                            <surname>Benson</surname>
                            <given-names>G</given-names>
                        </name>
			</person-group>:
                    <article-title>Tandem repeats finder: a program to analyze DNA sequences.</article-title>
                    <source>
				
                        <italic toggle="yes">Nucleic Acids Res.</italic>
			</source>
                    <year>1999</year>;<volume>27</volume>(<issue>2</issue>):<fpage>573</fpage>&#x2013;<lpage>80</lpage>.
                    <pub-id pub-id-type="pmid">9862982</pub-id>
                    <pub-id pub-id-type="doi">10.1093/nar/27.2.573</pub-id>
                    <pub-id pub-id-type="pmcid">148217</pub-id>
                </mixed-citation>
            </ref>
            <ref id="ref-17">
                <label>17</label>
                <mixed-citation publication-type="journal">
                    <person-group person-group-type="author">
				
                        <name name-style="western">
                            <surname>Gymrek</surname>
                            <given-names>M</given-names>
                        </name>
				
                        <name name-style="western">
                            <surname>Golan</surname>
                            <given-names>D</given-names>
                        </name>
				
                        <name name-style="western">
                            <surname>Rosset</surname>
                            <given-names>S</given-names>
                        </name>
				
                        <etal/>
			</person-group>:
                    <article-title>lobSTR: A short tandem repeat profiler for personal genomes.</article-title>
                    <source>
				
                        <italic toggle="yes">Genome Res.</italic>
			</source>
                    <year>2012</year>;<volume>22</volume>(<issue>6</issue>):<fpage>1154</fpage>&#x2013;<lpage>62</lpage>.
                    <pub-id pub-id-type="pmid">22522390</pub-id>
                    <pub-id pub-id-type="doi">10.1101/gr.135780.111</pub-id>
                    <pub-id pub-id-type="pmcid">3371701</pub-id>
                </mixed-citation>
                <note>
                    <p>
                        <ext-link ext-link-type="uri" xlink:href="https://f1000.com/prime/717957214">F1000 Recommendation</ext-link>
                    </p>
                </note>
            </ref>
            <ref id="ref-18">
                <label>18</label>
                <mixed-citation publication-type="journal">
                    <person-group person-group-type="author">
				
                        <name name-style="western">
                            <surname>Hensman Moss</surname>
                            <given-names>DJ</given-names>
                        </name>
				
                        <name name-style="western">
                            <surname>Poulter</surname>
                            <given-names>M</given-names>
                        </name>
				
                        <name name-style="western">
                            <surname>Beck</surname>
                            <given-names>J</given-names>
                        </name>
				
                        <etal/>
			</person-group>:
                    <article-title>
                        <italic toggle="yes">C9orf72</italic> expansions are the most common genetic cause of Huntington disease phenocopies.</article-title>
                    <source>
				
                        <italic toggle="yes">Neurology.</italic>
			</source>
                    <year>2014</year>;<volume>82</volume>(<issue>4</issue>):<fpage>292</fpage>&#x2013;<lpage>9</lpage>.
                    <pub-id pub-id-type="pmid">24363131</pub-id>
                    <pub-id pub-id-type="doi">10.1212/WNL.0000000000000061</pub-id>
                    <pub-id pub-id-type="pmcid">3929197</pub-id>
                </mixed-citation>
                <note>
                    <p>
                        <ext-link ext-link-type="uri" xlink:href="https://f1000.com/prime/718215353">F1000 Recommendation</ext-link>
                    </p>
                </note>
            </ref>
            <ref id="ref-19">
                <label>19</label>
                <mixed-citation publication-type="journal">
                    <person-group person-group-type="author">
				
                        <collab>1000 Genomes Project Consortium, </collab>
						
                        <name name-style="western">
                            <surname>Abecasis</surname>
                            <given-names>GR</given-names>
                        </name>
				
                        <name name-style="western">
                            <surname>Auton</surname>
                            <given-names>A</given-names>
                        </name>
				
                        <etal/>
			</person-group>:
                    <article-title>An integrated map of genetic variation from 1,092 human genomes.</article-title>
                    <source>
				
                        <italic toggle="yes">Nature.</italic>
			</source>
                    <year>2012</year>;<volume>491</volume>(<issue>7422</issue>):<fpage>56</fpage>&#x2013;<lpage>65</lpage>.
                    <pub-id pub-id-type="pmid">23128226</pub-id>
                    <pub-id pub-id-type="doi">10.1038/nature11632</pub-id>
                    <pub-id pub-id-type="pmcid">3498066</pub-id>
                </mixed-citation>
                <note>
                    <p>
                        <ext-link ext-link-type="uri" xlink:href="https://f1000.com/prime/717971074">F1000 Recommendation</ext-link>
                    </p>
                </note>
            </ref>
            <ref id="ref-20">
                <label>20</label>
                <mixed-citation publication-type="journal">
                    <person-group person-group-type="author">
				
                        <name name-style="western">
                            <surname>Smith</surname>
                            <given-names>DC</given-names>
                        </name>
				
                        <name name-style="western">
                            <surname>Atadzhanov</surname>
                            <given-names>M</given-names>
                        </name>
				
                        <name name-style="western">
                            <surname>Mwaba</surname>
                            <given-names>M</given-names>
                        </name>
				
                        <etal/>
			</person-group>:
                    <article-title>Evidence for a common founder effect amongst South African and Zambian individuals with Spinocerebellar ataxia type 7.</article-title>
                    <source>
				
                        <italic toggle="yes">J Neurol Sci.</italic>
			</source>
                    <year>2015</year>;<volume>354</volume>(<issue>1&#x2013;2</issue>):<fpage>75</fpage>&#x2013;<lpage>8</lpage>.
                    <pub-id pub-id-type="pmid">26003224</pub-id>
                    <pub-id pub-id-type="doi">10.1016/j.jns.2015.04.053</pub-id>
                </mixed-citation>
                <note>
                    <p>
                        <ext-link ext-link-type="uri" xlink:href="https://f1000.com/prime/725512547">F1000 Recommendation</ext-link>
                    </p>
                </note>
            </ref>
            <ref id="ref-21">
                <label>21</label>
                <mixed-citation publication-type="journal">
                    <person-group person-group-type="author">
				
                        <name name-style="western">
                            <surname>Brusco</surname>
                            <given-names>A</given-names>
                        </name>
				
                        <name name-style="western">
                            <surname>Gellera</surname>
                            <given-names>C</given-names>
                        </name>
				
                        <name name-style="western">
                            <surname>Cagnoli</surname>
                            <given-names>C</given-names>
                        </name>
				
                        <etal/>
			</person-group>:
                    <article-title>Molecular genetics of hereditary spinocerebellar ataxia: mutation analysis of spinocerebellar ataxia genes and CAG/CTG repeat expansion detection in 225 Italian families.</article-title>
                    <source>
				
                        <italic toggle="yes">Arch Neurol.</italic>
			</source>
                    <year>2004</year>;<volume>61</volume>(<issue>5</issue>):<fpage>727</fpage>&#x2013;<lpage>33</lpage>.
                    <pub-id pub-id-type="pmid">15148151</pub-id>
                    <pub-id pub-id-type="doi">10.1001/archneur.61.5.727</pub-id>
                </mixed-citation>
            </ref>
            <ref id="ref-22">
                <label>22</label>
                <mixed-citation publication-type="journal">
                    <person-group person-group-type="author">
				
                        <name name-style="western">
                            <surname>Moseley</surname>
                            <given-names>ML</given-names>
                        </name>
				
                        <name name-style="western">
                            <surname>Zu</surname>
                            <given-names>T</given-names>
                        </name>
				
                        <name name-style="western">
                            <surname>Ikeda</surname>
                            <given-names>Y</given-names>
                        </name>
				
                        <etal/>
			</person-group>:
                    <article-title>Bidirectional expression of CUG and CAG expansion transcripts and intranuclear polyglutamine inclusions in spinocerebellar ataxia type 8.</article-title>
                    <source>
				
                        <italic toggle="yes">Nat Genet.</italic>
			</source>
                    <year>2006</year>;<volume>38</volume>(<issue>7</issue>):<fpage>758</fpage>&#x2013;<lpage>69</lpage>.
                    <pub-id pub-id-type="pmid">16804541</pub-id>
                    <pub-id pub-id-type="doi">10.1038/ng1827</pub-id>
                </mixed-citation>
            </ref>
            <ref id="ref-23">
                <label>23</label>
                <mixed-citation publication-type="journal">
                    <person-group person-group-type="author">
				
                        <name name-style="western">
                            <surname>Stevanin</surname>
                            <given-names>G</given-names>
                        </name>
				
                        <name name-style="western">
                            <surname>Bouslam</surname>
                            <given-names>N</given-names>
                        </name>
				
                        <name name-style="western">
                            <surname>Thobois</surname>
                            <given-names>S</given-names>
                        </name>
				
                        <etal/>
			</person-group>:
                    <article-title>Spinocerebellar ataxia with sensory neuropathy (SCA25) maps to chromosome 2p.</article-title>
                    <source>
				
                        <italic toggle="yes">Ann Neurol.</italic>
			</source>
                    <year>2004</year>;<volume>55</volume>(<issue>1</issue>):<fpage>97</fpage>&#x2013;<lpage>104</lpage>.
                    <pub-id pub-id-type="pmid">14705117</pub-id>
                    <pub-id pub-id-type="doi">10.1002/ana.10798</pub-id>
                </mixed-citation>
            </ref>
            <ref id="ref-24">
                <label>24</label>
                <mixed-citation publication-type="journal">
                    <person-group person-group-type="author">
				
                        <name name-style="western">
                            <surname>Tankard</surname>
                            <given-names>RM</given-names>
                        </name>
			</person-group>:
                    <article-title>Identifying disease-causing short tandem repeat expansions in massively parallel sequencing data, with a focus on ataxias</article-title>. PhD thesis. The University of Melbourne.<year>2017</year>.
                    <ext-link ext-link-type="uri" xlink:href="https://minerva-access.unimelb.edu.au/bitstream/handle/11343/197796/TankardRickM_thesis_2018_01_12.pdf?sequence=1&amp;isAllowed=y">Reference Source</ext-link>
                </mixed-citation>
            </ref>
            <ref id="ref-25">
                <label>25</label>
                <mixed-citation publication-type="journal">
                    <person-group person-group-type="author">
				
                        <name name-style="western">
                            <surname>Cooper-Knock</surname>
                            <given-names>J</given-names>
                        </name>
				
                        <name name-style="western">
                            <surname>Shaw</surname>
                            <given-names>PJ</given-names>
                        </name>
				
                        <name name-style="western">
                            <surname>Kirby</surname>
                            <given-names>J</given-names>
                        </name>
			</person-group>:
                    <article-title>The widening spectrum of 
                        <italic toggle="yes">C9ORF72</italic>-related disease; genotype/phenotype correlations and potential modifiers of clinical phenotype.</article-title>
                    <source>
				
                        <italic toggle="yes">Acta Neuropathol.</italic>
			</source>
                    <year>2014</year>;<volume>127</volume>(<issue>3</issue>):<fpage>333</fpage>&#x2013;<lpage>45</lpage>.
                    <pub-id pub-id-type="pmid">24493408</pub-id>
                    <pub-id pub-id-type="doi">10.1007/s00401-014-1251-9</pub-id>
                    <pub-id pub-id-type="pmcid">3925297</pub-id>
                </mixed-citation>
                <note>
                    <p>
                        <ext-link ext-link-type="uri" xlink:href="https://f1000.com/prime/718264712">F1000 Recommendation</ext-link>
                    </p>
                </note>
            </ref>
            <ref id="ref-26">
                <label>26</label>
                <mixed-citation publication-type="journal">
                    <person-group person-group-type="author">
				
                        <name name-style="western">
                            <surname>Cummings</surname>
                            <given-names>BB</given-names>
                        </name>
				
                        <name name-style="western">
                            <surname>Marshall</surname>
                            <given-names>JL</given-names>
                        </name>
				
                        <name name-style="western">
                            <surname>Tukiainen</surname>
                            <given-names>T</given-names>
                        </name>
				
                        <etal/>
			</person-group>:
                    <article-title>Improving genetic diagnosis in Mendelian disease with transcriptome sequencing.</article-title>
                    <source>
				
                        <italic toggle="yes">Sci Transl Med.</italic>
			</source>
                    <year>2017</year>;<volume>9</volume>(<issue>386</issue>):  pii: eaal5209.
                    <pub-id pub-id-type="pmid">28424332</pub-id>
                    <pub-id pub-id-type="doi">10.1126/scitranslmed.aal5209</pub-id>
                    <pub-id pub-id-type="pmcid">5548421</pub-id>
                </mixed-citation>
                <note>
                    <p>
                        <ext-link ext-link-type="uri" xlink:href="https://f1000.com/prime/727520084">F1000 Recommendation</ext-link>
                    </p>
                </note>
            </ref>
            <ref id="ref-27">
                <label>27</label>
                <mixed-citation publication-type="journal">
                    <person-group person-group-type="author">
				
                        <name name-style="western">
                            <surname>Ganesamoorthy</surname>
                            <given-names>D</given-names>
                        </name>
					
                        <name name-style="western">
                            <surname>Cao</surname>
                            <given-names>MD</given-names>
                        </name>
				
                        <name name-style="western">
                            <surname>Duarte</surname>
                            <given-names>T</given-names>
                        </name>
				
                        <etal/>
			</person-group>:
                    <article-title>GtTR: Bayesian estimation of absolute tandem repeat copy number using sequence capture and high throughput sequencing.</article-title>
                    <source>
				
                        <italic toggle="yes">bioRxiv.</italic>
			</source>
                    <year>2018</year>.
                    <pub-id pub-id-type="doi">10.1101/246108</pub-id>
                </mixed-citation>
            </ref>
            <ref id="ref-28">
                <label>28</label>
                <mixed-citation publication-type="journal">
                    <person-group person-group-type="author">
				
                        <name name-style="western">
                            <surname>McGinty</surname>
                            <given-names>RJ</given-names>
                        </name>
				
                        <name name-style="western">
                            <surname>Rubinstein</surname>
                            <given-names>RG</given-names>
                        </name>
				
                        <name name-style="western">
                            <surname>Neil</surname>
                            <given-names>AJ</given-names>
                        </name>
				
                        <etal/>
			</person-group>:
                    <article-title>Nanopore sequencing of complex genomic rearrangements in yeast reveals mechanisms of repeat-mediated double-strand break repair.</article-title>
                    <source>
				
                        <italic toggle="yes">Genome Res.</italic>
			</source>
                    <year>2017</year>;<volume>27</volume>(<issue>12</issue>):<fpage>2072</fpage>&#x2013;<lpage>82</lpage>.
                    <pub-id pub-id-type="pmid">29113982</pub-id>
                    <pub-id pub-id-type="doi">10.1101/gr.228148.117</pub-id>
                    <pub-id pub-id-type="pmcid">5741057</pub-id>
                </mixed-citation>
                <note>
                    <p>
                        <ext-link ext-link-type="uri" xlink:href="https://f1000.com/prime/732092181">F1000 Recommendation</ext-link>
                    </p>
                </note>
            </ref>
            <ref id="ref-29">
                <label>29</label>
                <mixed-citation publication-type="journal">
                    <person-group person-group-type="author">
				
                        <name name-style="western">
                            <surname>Batra</surname>
                            <given-names>R</given-names>
                        </name>
				
                        <name name-style="western">
                            <surname>Nelles</surname>
                            <given-names>DA</given-names>
                        </name>
				
                        <name name-style="western">
                            <surname>Pirie</surname>
                            <given-names>E</given-names>
                        </name>
				
                        <etal/>
			</person-group>:
                    <article-title>Elimination of Toxic Microsatellite Repeat Expansion RNA by RNA-Targeting Cas9.</article-title>
                    <source>
				
                        <italic toggle="yes">Cell.</italic>
			</source>
                    <year>2017</year>;<volume>170</volume>(<issue>5</issue>):<fpage>899</fpage>&#x2013;<lpage>912.e10</lpage>.
                    <pub-id pub-id-type="pmid">28803727</pub-id>
                    <pub-id pub-id-type="doi">10.1016/j.cell.2017.07.010</pub-id>
                    <pub-id pub-id-type="pmcid">5873302</pub-id>
                </mixed-citation>
                <note>
                    <p>
                        <ext-link ext-link-type="uri" xlink:href="https://f1000.com/prime/728639308">F1000 Recommendation</ext-link>
                    </p>
                </note>
            </ref>
        </ref-list>
    </back>
</article>
