<?xml version="1.0" encoding="UTF-8"?>
<?xml-model https://raw.githubusercontent.com/i-d-e/ride/master/schema/ride.rng application/xml http://relaxng.org/ns/structure/1.0?>
<?xml-model https://raw.githubusercontent.com/i-d-e/ride/master/schema/ride.rng application/xml http://purl.oclc.org/dsdl/schematron?>
<TEI xmlns="http://www.tei-c.org/ns/1.0" xml:id="ride.9.3">
   <teiHeader>
      <fileDesc>
         <titleStmt>
            <title>Review of “ShakespearePlaysPlus Text Corpus”</title>
            <author>
               <name>
                  <forename>Katharina</forename>
                  <surname>Mahler</surname>
               </name>
               <affiliation>
                  <orgName>CCeH, University of Cologne</orgName>
                  <placeName ref="https://www.geonames.org/2886242">Cologne, Germany</placeName>
               </affiliation>
               <email>english.textworks@gmail.com</email>
            </author>
         </titleStmt>
         <publicationStmt>
            <publisher>Institut für Dokumentologie und Editorik e. V.</publisher>
            <date when="2018-11">2018</date>
            <idno type="URI">http://ride.i-d-e.de/issues/issue-9/shakespeare-plays</idno>
            <idno type="DOI">10.18716/ride.a.9.3</idno>
            <idno type="archive">https://github.com/i-d-e/ride/raw/master/issues/issue09/shakespeare-plays/shakespeare-plays.pdf</idno>
            <availability>
               <licence target="http://creativecommons.org/licenses/by/4.0/"/>
            </availability>
         </publicationStmt>
         <seriesStmt>
            <title level="j">RIDE - A review journal for digital editions and resources</title>
            <editor ref="https://orcid.org/0000-0003-2852-065X">Ulrike Henny-Krahmer</editor>
            <editor ref="https://orcid.org/0000-0001-8279-9298">Frederike Neuber</editor>
            <editor ref="https://orcid.org/0000-0002-6457-0913" role="managing">Philipp
               Steinkrüger</editor>
            <editor ref="http://viaf.org/viaf/80243768" role="technical">Bernhard Assmann</editor>
            <editor ref="https://orcid.org/0000-0003-2852-065X" role="technical">Ulrike
               Henny-Krahmer</editor>
            <editor ref="https://orcid.org/0000-0001-8279-9298" role="technical">Frederike
               Neuber</editor>
            <biblScope unit="issue" n="9">Digital Text Collections</biblScope>
            <idno type="URI">http://ride.i-d-e.de/issues/issue-9</idno>
            <idno type="DOI">10.18716/ride.a.9</idno>
         </seriesStmt>
         <notesStmt>
            <relatedItem type="reviewed_resource">
               <bibl>
                  <title>Shakespeare Corpus: ShakespearePlaysPlus</title>
                  <editor>Mike Scott</editor>
                  <respStmt>
                     <resp>Editor</resp>
                     <persName>
                        <name>Mike Scott</name>
                     </persName>
                  </respStmt>
                  <respStmt>
                     <resp>Encoder</resp>
                     <persName>
                        <name>Mike Scott</name>
                     </persName>
                  </respStmt>
                  <respStmt>
                     <resp>Administrator</resp>
                     <persName>
                        <name>Mike Scott</name>
                     </persName>
                  </respStmt>
                  <date type="publication">2006</date>
                  <idno type="URI">http://lexically.net/wordsmith/support/shakespeare.html</idno>
                  <date type="accessed">2018-07-03</date>
               </bibl>
            </relatedItem>
            <relatedItem type="reviewing_criteria">
               <bibl>
                  <ref target="http://www.i-d-e.de/publikationen/weitereschriften/criteria-text-collections-version-1-0/">Criteria for Reviewing Digital Text Collections</ref>
               </bibl>
            </relatedItem>
         </notesStmt>
         <sourceDesc>
            <p>born digital</p>
         </sourceDesc>
      </fileDesc>
      <encodingDesc>
         <classDecl>
            <taxonomy xml:base="https://www.i-d-e.de/publikationen/weitereschriften/criteria-text-collections-version-1-0/">
               <category xml:id="general_information">
                  <category xml:id="bibl_desc">
                     <catDesc>
                        <ref target="#K1.1">cf. Catalogue 1.1</ref>
                     </catDesc>
                     <catDesc>Can the text collection be identified in terms similar to traditional
                        bibliographic descriptions (title, responsible editors, institution, date(s)
                        of publication, identifier/address)?</catDesc>
                     <catDesc>
                        <num type="boolean" value="1"/>
                     </catDesc>
                  </category>
                  <category xml:id="contributors">
                     <catDesc>
                        <ref target="#K1.3">cf. Catalogue 1.3</ref>
                     </catDesc>
                     <catDesc>Are the contributors (editors, institutions, associates) of the
                        project documented?</catDesc>
                     <catDesc>
                        <num type="boolean" value="1"/>
                     </catDesc>
                  </category>
                  <category xml:id="contacts">
                     <catDesc>
                        <ref target="#K1.4">cf. Catalogue 1.4</ref>
                     </catDesc>
                     <catDesc>Is contact information given?</catDesc>
                     <catDesc>
                        <num type="boolean" value="1"/>
                     </catDesc>
                  </category>
               </category>
               <category xml:id="aims">
                  <category xml:id="doc_contents">
                     <catDesc>
                        <ref target="#K2.1">cf. Catalogue 2.1</ref>
                     </catDesc>
                     <catDesc>Is there a description of the aims and contents of the text
                        collection?</catDesc>
                     <catDesc>
                        <num type="boolean" value="1"/>
                     </catDesc>
                  </category>
                  <category xml:id="purpose">
                     <catDesc>
                        <ref target="#K2.2">cf. Catalogue 2.2</ref>
                     </catDesc>
                     <catDesc>What is the purpose of the text collection?</catDesc>
                     <category xml:id="research">
                        <catDesc>
                           <num type="boolean" value="1"/>
                        </catDesc>
                     </category>
                     <category xml:id="teaching">
                        <catDesc>
                           <num type="boolean" value="0"/>
                        </catDesc>
                     </category>
                     <category xml:id="general_purpose">
                        <catDesc>
                           <num type="boolean" value="1"/>
                        </catDesc>
                     </category>
                     <category corresp="#other">
                        <catDesc>
                           <num type="boolean" value="0"/>
                        </catDesc>
                        <category corresp="#free">
                           <desc/>
                        </category>
                     </category>
                  </category>
                  <category xml:id="research_type">
                     <catDesc>
                        <ref target="#K3.1.8">cf. Catalogue 3.1.8</ref>
                     </catDesc>
                     <catDesc>What kind of research does the collection allow to conduct
                        primarily?</catDesc>
                     <category xml:id="qualitative">
                        <catDesc>
                           <num type="boolean" value="0"/>
                        </catDesc>
                     </category>
                     <category xml:id="quantitative">
                        <catDesc>
                           <num type="boolean" value="1"/>
                        </catDesc>
                     </category>
                     <category corresp="#none">
                        <catDesc>
                           <num type="boolean" value="0"/>
                        </catDesc>
                     </category>
                  </category>
                  <category xml:id="classification">
                     <catDesc>
                        <ref target="#K2.3">cf. Catalogue 2.3</ref>
                     </catDesc>
                     <catDesc>How does the text collection classify itself (e.g. in its title or
                        documentation)?</catDesc>
                     <category xml:id="collection">
                        <catDesc>
                           <num type="boolean" value="0"/>
                        </catDesc>
                     </category>
                     <category xml:id="corpus">
                        <catDesc>
                           <num type="boolean" value="1"/>
                        </catDesc>
                     </category>
                     <category xml:id="digital_archive">
                        <catDesc>
                           <num type="boolean" value="0"/>
                        </catDesc>
                     </category>
                     <category xml:id="digital_library">
                        <catDesc>
                           <num type="boolean" value="0"/>
                        </catDesc>
                     </category>
                     <category xml:id="digital_edition">
                        <catDesc>
                           <num type="boolean" value="0"/>
                        </catDesc>
                     </category>
                     <category xml:id="portal">
                        <catDesc>
                           <num type="boolean" value="0"/>
                        </catDesc>
                     </category>
                     <category xml:id="database">
                        <catDesc>
                           <num type="boolean" value="0"/>
                        </catDesc>
                     </category>
                     <category xml:id="no_classification">
                        <catDesc>
                           <num type="boolean" value="0"/>
                        </catDesc>
                     </category>
                     <category corresp="#other">
                        <catDesc>
                           <num type="boolean" value="0"/>
                        </catDesc>
                        <category corresp="#free">
                           <desc/>
                        </category>
                     </category>
                  </category>
                  <category xml:id="research_fields">
                     <catDesc>
                        <ref target="#K2.2">cf. Catalogue 2.2</ref>
                     </catDesc>
                     <catDesc>To which field(s) of research does the text collection
                        contribute?</catDesc>
                     <category xml:id="field_history">
                        <catDesc>
                           <num type="boolean" value="0"/>
                        </catDesc>
                     </category>
                     <category xml:id="field_literary_studies">
                        <catDesc>
                           <num type="boolean" value="1"/>
                        </catDesc>
                     </category>
                     <category xml:id="field_linguistics">
                        <catDesc>
                           <num type="boolean" value="1"/>
                        </catDesc>
                     </category>
                     <category xml:id="field_musicology">
                        <catDesc>
                           <num type="boolean" value="0"/>
                        </catDesc>
                     </category>
                     <category xml:id="field_art_history">
                        <catDesc>
                           <num type="boolean" value="0"/>
                        </catDesc>
                     </category>
                     <category xml:id="field_archaeology">
                        <catDesc>
                           <num type="boolean" value="0"/>
                        </catDesc>
                     </category>
                     <category xml:id="field_philosophy">
                        <catDesc>
                           <num type="boolean" value="0"/>
                        </catDesc>
                     </category>
                     <category xml:id="field_religious_studies">
                        <catDesc>
                           <num type="boolean" value="0"/>
                        </catDesc>
                     </category>
                     <category xml:id="field_sociology">
                        <catDesc>
                           <num type="boolean" value="0"/>
                        </catDesc>
                     </category>
                     <category corresp="#other">
                        <catDesc>
                           <num type="boolean" value="0"/>
                        </catDesc>
                        <category corresp="#free">
                           <desc/>
                        </category>
                     </category>
                     <category corresp="#none">
                        <catDesc>
                           <num type="boolean" value="0"/>
                        </catDesc>
                     </category>
                  </category>
               </category>
               <category xml:id="content">
                  <category xml:id="era">
                     <catDesc>
                        <ref target="#K2.5">cf. Catalogue 2.5</ref>
                     </catDesc>
                     <catDesc>What era(s) do the texts belong to?</catDesc>
                     <category xml:id="era_classics">
                        <catDesc>
                           <num type="boolean" value="1"/>
                        </catDesc>
                     </category>
                     <category xml:id="era_medieval">
                        <catDesc>
                           <num type="boolean" value="0"/>
                        </catDesc>
                     </category>
                     <category xml:id="era_early_modern">
                        <catDesc>
                           <num type="boolean" value="1"/>
                        </catDesc>
                     </category>
                     <category xml:id="era_modern">
                        <catDesc>
                           <num type="boolean" value="0"/>
                        </catDesc>
                     </category>
                     <category xml:id="era_contemporary">
                        <catDesc>
                           <num type="boolean" value="0"/>
                        </catDesc>
                     </category>
                  </category>
                  <category xml:id="language">
                     <catDesc>
                        <ref target="#K2.5">cf. Catalogue 2.5</ref>
                     </catDesc>
                     <catDesc>What languages are the texts in?</catDesc>
                     <category xml:id="arabic">
                        <catDesc>
                           <num type="boolean" value="0"/>
                        </catDesc>
                     </category>
                     <category xml:id="chinese">
                        <catDesc>
                           <num type="boolean" value="0"/>
                        </catDesc>
                     </category>
                     <category xml:id="danish">
                        <catDesc>
                           <num type="boolean" value="0"/>
                        </catDesc>
                     </category>
                     <category xml:id="english">
                        <catDesc>
                           <num type="boolean" value="1"/>
                        </catDesc>
                     </category>
                     <category xml:id="finnish">
                        <catDesc>
                           <num type="boolean" value="0"/>
                        </catDesc>
                     </category>
                     <category xml:id="french">
                        <catDesc>
                           <num type="boolean" value="0"/>
                        </catDesc>
                     </category>
                     <category xml:id="german">
                        <catDesc>
                           <num type="boolean" value="0"/>
                        </catDesc>
                     </category>
                     <category xml:id="greek">
                        <catDesc>
                           <num type="boolean" value="0"/>
                        </catDesc>
                     </category>
                     <category xml:id="hebrew">
                        <catDesc>
                           <num type="boolean" value="0"/>
                        </catDesc>
                     </category>
                     <category xml:id="hindi">
                        <catDesc>
                           <num type="boolean" value="0"/>
                        </catDesc>
                     </category>
                     <category xml:id="italian">
                        <catDesc>
                           <num type="boolean" value="0"/>
                        </catDesc>
                     </category>
                     <category xml:id="japanese">
                        <catDesc>
                           <num type="boolean" value="0"/>
                        </catDesc>
                     </category>
                     <category xml:id="latin">
                        <catDesc>
                           <num type="boolean" value="0"/>
                        </catDesc>
                     </category>
                     <category xml:id="norwegian">
                        <catDesc>
                           <num type="boolean" value="0"/>
                        </catDesc>
                     </category>
                     <category xml:id="polish">
                        <catDesc>
                           <num type="boolean" value="0"/>
                        </catDesc>
                     </category>
                     <category xml:id="portuguese">
                        <catDesc>
                           <num type="boolean" value="0"/>
                        </catDesc>
                     </category>
                     <category xml:id="russian">
                        <catDesc>
                           <num type="boolean" value="0"/>
                        </catDesc>
                     </category>
                     <category xml:id="spanish">
                        <catDesc>
                           <num type="boolean" value="0"/>
                        </catDesc>
                     </category>
                     <category xml:id="swedish">
                        <catDesc>
                           <num type="boolean" value="0"/>
                        </catDesc>
                     </category>
                     <category xml:id="turkish">
                        <catDesc>
                           <num type="boolean" value="0"/>
                        </catDesc>
                     </category>
                     <category corresp="#other">
                        <catDesc>
                           <num type="boolean" value="0"/>
                        </catDesc>
                        <category corresp="#free">
                           <desc/>
                        </category>
                     </category>
                     <category corresp="#free" xml:id="language_note">
                        <desc/>
                     </category>
                  </category>
                  <category xml:id="text_type">
                     <catDesc>
                        <ref target="#K2.5">cf. Catalogue 2.5</ref>
                     </catDesc>
                     <catDesc>What kind of texts are in the collection?</catDesc>
                     <category xml:id="literary_works">
                        <catDesc>
                           <num type="boolean" value="1"/>
                        </catDesc>
                     </category>
                     <category xml:id="private_documents">
                        <catDesc>
                           <num type="boolean" value="0"/>
                        </catDesc>
                     </category>
                     <category xml:id="essays">
                        <catDesc>
                           <num type="boolean" value="0"/>
                        </catDesc>
                     </category>
                     <category xml:id="newspaper_articles">
                        <catDesc>
                           <num type="boolean" value="0"/>
                        </catDesc>
                     </category>
                     <category xml:id="charters">
                        <catDesc>
                           <num type="boolean" value="0"/>
                        </catDesc>
                     </category>
                     <category xml:id="inscriptions">
                        <catDesc>
                           <num type="boolean" value="0"/>
                        </catDesc>
                     </category>
                     <category xml:id="files_records">
                        <catDesc>
                           <num type="boolean" value="0"/>
                        </catDesc>
                     </category>
                     <category xml:id="protocols">
                        <catDesc>
                           <num type="boolean" value="0"/>
                        </catDesc>
                     </category>
                     <category xml:id="scientific_papers">
                        <catDesc>
                           <num type="boolean" value="0"/>
                        </catDesc>
                     </category>
                     <category xml:id="speech_transcripts">
                        <catDesc>
                           <num type="boolean" value="0"/>
                        </catDesc>
                     </category>
                     <category corresp="#other">
                        <catDesc>
                           <num type="boolean" value="0"/>
                        </catDesc>
                        <category corresp="#free">
                           <desc/>
                        </category>
                     </category>
                  </category>
                  <category xml:id="add_information">
                     <catDesc>
                        <ref target="#K2.5">cf. Catalogue 2.5</ref>
                     </catDesc>
                     <catDesc>What kind of information is published in addition to the
                        texts?</catDesc>
                     <category xml:id="introduction">
                        <catDesc>
                           <num type="boolean" value="0"/>
                        </catDesc>
                     </category>
                     <category xml:id="commentary">
                        <catDesc>
                           <num type="boolean" value="0"/>
                        </catDesc>
                     </category>
                     <category xml:id="context_material">
                        <catDesc>
                           <num type="boolean" value="0"/>
                        </catDesc>
                     </category>
                     <category xml:id="bibliography">
                        <catDesc>
                           <num type="boolean" value="0"/>
                        </catDesc>
                     </category>
                     <category xml:id="facsimile">
                        <catDesc>
                           <num type="boolean" value="0"/>
                        </catDesc>
                     </category>
                     <category corresp="#other">
                        <catDesc>
                           <num type="boolean" value="0"/>
                        </catDesc>
                        <category corresp="#free">
                           <desc/>
                        </category>
                     </category>
                     <category corresp="#none">
                        <catDesc>
                           <num type="boolean" value="1"/>
                        </catDesc>
                     </category>
                  </category>
               </category>
               <category xml:id="composition">
                  <category xml:id="documentation_methods">
                     <catDesc>
                        <ref target="#K3.1.1">cf. Catalogue 3.1.1-3.1.3</ref>
                     </catDesc>
                     <catDesc>Are the principles and decisions regarding the design of the text
                        collection, its composition and the selection of texts documented?</catDesc>
                     <catDesc>
                        <num type="boolean" value="1"/>
                     </catDesc>
                  </category>
                  <category xml:id="selection">
                     <catDesc>
                        <ref target="#K3.1">cf. Catalogue 3.1</ref>
                     </catDesc>
                     <catDesc>What selection criteria have been chosen for the text
                        collection?</catDesc>
                     <category xml:id="selection_language">
                        <catDesc>
                           <num type="boolean" value="0"/>
                        </catDesc>
                     </category>
                     <category xml:id="selection_author">
                        <catDesc>
                           <num type="boolean" value="1"/>
                        </catDesc>
                     </category>
                     <category xml:id="selection_country">
                        <catDesc>
                           <num type="boolean" value="0"/>
                        </catDesc>
                     </category>
                     <category xml:id="selection_epoch">
                        <catDesc>
                           <num type="boolean" value="0"/>
                        </catDesc>
                     </category>
                     <category xml:id="selection_genre">
                        <catDesc>
                           <num type="boolean" value="1"/>
                        </catDesc>
                     </category>
                     <category xml:id="selection_topic">
                        <catDesc>
                           <num type="boolean" value="0"/>
                        </catDesc>
                     </category>
                     <category xml:id="selection_style">
                        <catDesc>
                           <num type="boolean" value="0"/>
                        </catDesc>
                     </category>
                     <category xml:id="selection_linguistic_characteristics">
                        <catDesc>
                           <num type="boolean" value="0"/>
                        </catDesc>
                     </category>
                     <category corresp="#other">
                        <catDesc>
                           <num type="boolean" value="0"/>
                        </catDesc>
                        <category corresp="#free">
                           <desc/>
                        </category>
                     </category>
                  </category>
                  <category xml:id="size">
                     <category xml:id="size_texts">
                        <catDesc>
                           <ref target="#K3.1.4">cf. Catalogue 3.1.4</ref>
                        </catDesc>
                        <catDesc>How large is the text collection in number of
                           texts/records?</catDesc>
                        <category xml:id="texts_le10">
                           <catDesc>
                              <num type="boolean" value="0"/>
                           </catDesc>
                        </category>
                        <category xml:id="texts_11-50">
                           <catDesc>
                              <num type="boolean" value="0"/>
                           </catDesc>
                        </category>
                        <category xml:id="texts_51-100">
                           <catDesc>
                              <num type="boolean" value="0"/>
                           </catDesc>
                        </category>
                        <category xml:id="texts_gt100_">
                           <catDesc>
                              <num type="boolean" value="0"/>
                           </catDesc>
                        </category>
                        <category xml:id="texts_gt1000">
                           <catDesc>
                              <num type="boolean" value="1"/>
                           </catDesc>
                        </category>
                        <category corresp="#unknown">
                           <catDesc>
                              <num type="boolean" value="0"/>
                           </catDesc>
                        </category>
                     </category>
                     <category xml:id="size_tokens">
                        <catDesc>
                           <ref target="#K3.1.4">cf. Catalogue 3.1.4</ref>
                        </catDesc>
                        <catDesc>How large is the text collection in number of tokens?</catDesc>
                        <category xml:id="tokens_lt100.000">
                           <catDesc>
                              <num type="boolean" value="0"/>
                           </catDesc>
                        </category>
                        <category xml:id="tokens_100.000-1mio">
                           <catDesc>
                              <num type="boolean" value="1"/>
                           </catDesc>
                        </category>
                        <category xml:id="tokens_gt1mio">
                           <catDesc>
                              <num type="boolean" value="0"/>
                           </catDesc>
                        </category>
                        <category xml:id="tokens_gt10mio">
                           <catDesc>
                              <num type="boolean" value="0"/>
                           </catDesc>
                        </category>
                        <category corresp="#unknown">
                           <catDesc>
                              <num type="boolean" value="0"/>
                           </catDesc>
                        </category>
                     </category>
                  </category>
                  <category xml:id="structure">
                     <catDesc>
                        <ref target="#K3.1.5">cf. Catalogue 3.1.5</ref>
                     </catDesc>
                     <catDesc>Does the text collection have identifiable sub-collections or
                        components?</catDesc>
                     <catDesc>
                        <num type="boolean" value="1"/>
                     </catDesc>
                  </category>
                  <category xml:id="acquisition">
                     <category xml:id="text_recording">
                        <catDesc>
                           <ref target="#K3.1.6">cf. Catalogue 3.1.6</ref>
                        </catDesc>
                        <catDesc>Does the text collection record or transcribe the textual data for
                           the first time?</catDesc>
                        <catDesc>
                           <num type="boolean" value="0"/>
                        </catDesc>
                     </category>
                     <category xml:id="text_integration">
                        <catDesc>
                           <ref target="#K3.1.6">cf. Catalogue 3.1.6</ref>
                        </catDesc>
                        <catDesc>What kind of material has been taken over from other
                           sources?</catDesc>
                        <category xml:id="full_text">
                           <catDesc>
                              <num type="boolean" value="1"/>
                           </catDesc>
                        </category>
                        <category xml:id="reuse_metadata">
                           <catDesc>
                              <num type="boolean" value="0"/>
                           </catDesc>
                        </category>
                        <category xml:id="reuse_annotation">
                           <catDesc>
                              <num type="boolean" value="1"/>
                           </catDesc>
                        </category>
                        <category corresp="#other">
                           <catDesc>
                              <num type="boolean" value="0"/>
                           </catDesc>
                           <category corresp="#free">
                              <desc/>
                           </category>
                        </category>
                        <category corresp="#none">
                           <catDesc>
                              <num type="boolean" value="0"/>
                           </catDesc>
                        </category>
                     </category>
                  </category>
                  <category xml:id="quality">
                     <catDesc>
                        <ref target="#K3.1.7">cf. Catalogue 3.1.7</ref>
                     </catDesc>
                     <catDesc>Has the quality of the data (transcriptions, metadata, annotations,
                        etc.) been checked?</catDesc>
                     <catDesc>
                        <num type="boolean" value="1"/>
                     </catDesc>
                  </category>
                  <category xml:id="typology">
                     <catDesc>
                        <ref target="#K3.1.8">cf. Catalogue 3.1.8</ref>
                     </catDesc>
                     <catDesc>Considering aims and methods of the text collection, how would you
                        classify it further? For definitions please consider the
                        help-texts.</catDesc>
                     <category xml:id="typology_general_purpose">
                        <catDesc>
                           <num type="boolean" value="0"/>
                        </catDesc>
                     </category>
                     <category xml:id="typology_corpus">
                        <catDesc>
                           <num type="boolean" value="1"/>
                        </catDesc>
                     </category>
                     <category xml:id="typology_collection_records">
                        <catDesc>
                           <num type="boolean" value="0"/>
                        </catDesc>
                     </category>
                     <category xml:id="typology_canon">
                        <catDesc>
                           <num type="boolean" value="1"/>
                        </catDesc>
                     </category>
                     <category xml:id="typology_oeuvre">
                        <catDesc>
                           <num type="boolean" value="1"/>
                        </catDesc>
                     </category>
                     <category xml:id="typology_reference_corpus">
                        <catDesc>
                           <num type="boolean" value="0"/>
                        </catDesc>
                     </category>
                     <category xml:id="typology_contrastive_corpus">
                        <catDesc>
                           <num type="boolean" value="0"/>
                        </catDesc>
                     </category>
                     <category xml:id="typology_parallel_corpus">
                        <catDesc>
                           <num type="boolean" value="0"/>
                        </catDesc>
                     </category>
                     <category xml:id="typology_diachronic_corpus">
                        <catDesc>
                           <num type="boolean" value="0"/>
                        </catDesc>
                     </category>
                     <category corresp="#other">
                        <catDesc>
                           <num type="boolean" value="0"/>
                        </catDesc>
                        <category corresp="#free">
                           <desc/>
                        </category>
                     </category>
                  </category>
               </category>
               <category xml:id="data_modelling">
                  <category xml:id="text_treatment">
                     <catDesc>
                        <ref target="#K3.2.1">cf. Catalogue 3.2.1</ref>
                     </catDesc>
                     <catDesc>How are the textual sources represented in the digital
                        collection?</catDesc>
                     <category xml:id="normalized_transcription">
                        <catDesc>
                           <num type="boolean" value="1"/>
                        </catDesc>
                     </category>
                     <category xml:id="orthographic_transcription">
                        <catDesc>
                           <num type="boolean" value="0"/>
                        </catDesc>
                     </category>
                     <category xml:id="phonetic_transcription">
                        <catDesc>
                           <num type="boolean" value="0"/>
                        </catDesc>
                     </category>
                     <category xml:id="diplomatic_transcription">
                        <catDesc>
                           <num type="boolean" value="0"/>
                        </catDesc>
                     </category>
                     <category xml:id="transliteration">
                        <catDesc>
                           <num type="boolean" value="0"/>
                        </catDesc>
                     </category>
                     <category xml:id="edited_text">
                        <catDesc>
                           <num type="boolean" value="0"/>
                        </catDesc>
                     </category>
                     <category xml:id="translated_text">
                        <catDesc>
                           <num type="boolean" value="0"/>
                        </catDesc>
                     </category>
                     <category xml:id="summarized_text">
                        <catDesc>
                           <num type="boolean" value="0"/>
                        </catDesc>
                     </category>
                     <category xml:id="sampled_text">
                        <catDesc>
                           <num type="boolean" value="0"/>
                        </catDesc>
                     </category>
                     <category corresp="#other">
                        <catDesc>
                           <num type="boolean" value="0"/>
                        </catDesc>
                        <category corresp="#free">
                           <desc/>
                        </category>
                     </category>
                  </category>
                  <category xml:id="basic_format">
                     <catDesc>
                        <ref target="#K3.2.4">cf. Catalogue 3.2.4</ref>
                     </catDesc>
                     <catDesc>In which basic format are the texts encoded?</catDesc>
                     <category xml:id="plain_text">
                        <catDesc>
                           <num type="boolean" value="1"/>
                        </catDesc>
                     </category>
                     <category xml:id="xml">
                        <catDesc>
                           <num type="boolean" value="0"/>
                        </catDesc>
                     </category>
                     <category xml:id="html">
                        <catDesc>
                           <num type="boolean" value="0"/>
                        </catDesc>
                     </category>
                     <category corresp="#other">
                        <catDesc>
                           <num type="boolean" value="0"/>
                        </catDesc>
                        <category corresp="#free">
                           <desc/>
                        </category>
                     </category>
                  </category>
                  <category xml:id="annotation">
                     <category xml:id="annotation_type">
                        <catDesc>
                           <ref target="#K3.2.2">cf. Catalogue 3.2.2</ref>
                        </catDesc>
                        <catDesc>With what information are the texts further enriched?</catDesc>
                        <category xml:id="semantic_annotations">
                           <catDesc>
                              <num type="boolean" value="0"/>
                           </catDesc>
                        </category>
                        <category xml:id="linguistic_annotations">
                           <catDesc>
                              <num type="boolean" value="0"/>
                           </catDesc>
                        </category>
                        <category xml:id="editorial_annotations">
                           <catDesc>
                              <num type="boolean" value="0"/>
                           </catDesc>
                        </category>
                        <category xml:id="structural_information">
                           <catDesc>
                              <num type="boolean" value="1"/>
                           </catDesc>
                        </category>
                        <category corresp="#other">
                           <catDesc>
                              <num type="boolean" value="0"/>
                           </catDesc>
                           <category corresp="#free">
                              <desc/>
                           </category>
                        </category>
                        <category corresp="#none">
                           <catDesc>
                              <num type="boolean" value="0"/>
                           </catDesc>
                        </category>
                     </category>
                     <category xml:id="annotation_integration">
                        <catDesc>
                           <ref target="#K3.2.2">cf. Catalogue 3.2.2</ref>
                        </catDesc>
                        <catDesc>How are the annotations linked to the texts themselves?</catDesc>
                        <category xml:id="embedded">
                           <catDesc>
                              <num type="boolean" value="1"/>
                           </catDesc>
                        </category>
                        <category xml:id="stand-off">
                           <catDesc>
                              <num type="boolean" value="0"/>
                           </catDesc>
                        </category>
                        <category corresp="#other">
                           <catDesc>
                              <num type="boolean" value="0"/>
                           </catDesc>
                           <category corresp="#free">
                              <desc/>
                           </category>
                        </category>
                        <category corresp="#not_applicable">
                           <catDesc>
                              <num type="boolean" value="0"/>
                           </catDesc>
                        </category>
                     </category>
                  </category>
                  <category xml:id="metadata">
                     <category xml:id="metadata_type">
                        <catDesc>
                           <ref target="#K3.2.3">cf. Catalogue 3.2.3</ref>
                        </catDesc>
                        <catDesc>What kind of metadata are included in the text
                           collection?</catDesc>
                        <category xml:id="descriptive">
                           <catDesc>
                              <num type="boolean" value="1"/>
                           </catDesc>
                        </category>
                        <category xml:id="structural">
                           <catDesc>
                              <num type="boolean" value="0"/>
                           </catDesc>
                        </category>
                        <category xml:id="administrative">
                           <catDesc>
                              <num type="boolean" value="0"/>
                           </catDesc>
                        </category>
                        <category corresp="#other">
                           <catDesc>
                              <num type="boolean" value="0"/>
                           </catDesc>
                           <category corresp="#free">
                              <desc/>
                           </category>
                        </category>
                        <category corresp="#none">
                           <catDesc>
                              <num type="boolean" value="0"/>
                           </catDesc>
                        </category>
                     </category>
                     <category xml:id="metadata_level">
                        <catDesc>
                           <ref target="#K3.2.2">cf. Catalogue 3.2.2</ref>
                        </catDesc>
                        <catDesc>On which level are the metadata included?</catDesc>
                        <category xml:id="whole_collection">
                           <catDesc>
                              <num type="boolean" value="0"/>
                           </catDesc>
                        </category>
                        <category xml:id="collection_parts">
                           <catDesc>
                              <num type="boolean" value="0"/>
                           </catDesc>
                        </category>
                        <category xml:id="individual_texts">
                           <catDesc>
                              <num type="boolean" value="1"/>
                           </catDesc>
                        </category>
                        <category corresp="#other">
                           <catDesc>
                              <num type="boolean" value="0"/>
                           </catDesc>
                           <category corresp="#free">
                              <desc/>
                           </category>
                        </category>
                        <category corresp="#not_applicable">
                           <catDesc>
                              <num type="boolean" value="0"/>
                           </catDesc>
                        </category>
                     </category>
                  </category>
                  <category xml:id="data_standards">
                     <category xml:id="data_schema">
                        <catDesc>
                           <ref target="#K3.2.4">cf. Catalogue 3.2.4</ref>
                        </catDesc>
                        <catDesc>What kind of data/metadata/annotation schemas are used for the text
                           collection?</catDesc>
                        <category xml:id="standardized_schema">
                           <catDesc>
                              <num type="boolean" value="0"/>
                           </catDesc>
                        </category>
                        <category xml:id="customized_standard_schema">
                           <catDesc>
                              <num type="boolean" value="0"/>
                           </catDesc>
                        </category>
                        <category xml:id="project_specific">
                           <catDesc>
                              <num type="boolean" value="1"/>
                           </catDesc>
                        </category>
                        <category corresp="#none">
                           <catDesc>
                              <num type="boolean" value="0"/>
                           </catDesc>
                        </category>
                        <category corresp="#unknown">
                           <catDesc>
                              <num type="boolean" value="0"/>
                           </catDesc>
                        </category>
                     </category>
                     <category xml:id="standard_format">
                        <catDesc>
                           <ref target="#K3.2.4">cf. Catalogue 3.2.4</ref>
                        </catDesc>
                        <catDesc>Which standards for text encoding, metadata and annotation are used
                           in the text collection?</catDesc>
                        <category xml:id="tei">
                           <catDesc>
                              <num type="boolean" value="0"/>
                           </catDesc>
                        </category>
                        <category xml:id="cei">
                           <catDesc>
                              <num type="boolean" value="0"/>
                           </catDesc>
                        </category>
                        <category xml:id="ead">
                           <catDesc>
                              <num type="boolean" value="0"/>
                           </catDesc>
                        </category>
                        <category xml:id="xces">
                           <catDesc>
                              <num type="boolean" value="0"/>
                           </catDesc>
                        </category>
                        <category xml:id="dublin_core">
                           <catDesc>
                              <num type="boolean" value="0"/>
                           </catDesc>
                        </category>
                        <category xml:id="edm">
                           <catDesc>
                              <num type="boolean" value="0"/>
                           </catDesc>
                        </category>
                        <category xml:id="mets">
                           <catDesc>
                              <num type="boolean" value="0"/>
                           </catDesc>
                        </category>
                        <category xml:id="mods">
                           <catDesc>
                              <num type="boolean" value="0"/>
                           </catDesc>
                        </category>
                        <category xml:id="skos">
                           <catDesc>
                              <num type="boolean" value="0"/>
                           </catDesc>
                        </category>
                        <category xml:id="owl">
                           <catDesc>
                              <num type="boolean" value="0"/>
                           </catDesc>
                        </category>
                        <category xml:id="imdi">
                           <catDesc>
                              <num type="boolean" value="0"/>
                           </catDesc>
                        </category>
                        <category xml:id="cmdi">
                           <catDesc>
                              <num type="boolean" value="0"/>
                           </catDesc>
                        </category>
                        <category xml:id="tcf">
                           <catDesc>
                              <num type="boolean" value="0"/>
                           </catDesc>
                        </category>
                        <category xml:id="olac">
                           <catDesc>
                              <num type="boolean" value="0"/>
                           </catDesc>
                        </category>
                        <category xml:id="eagles">
                           <catDesc>
                              <num type="boolean" value="0"/>
                           </catDesc>
                        </category>
                        <category xml:id="pos_tagsets">
                           <catDesc>
                              <num type="boolean" value="0"/>
                           </catDesc>
                        </category>
                        <category corresp="#other">
                           <catDesc>
                              <num type="boolean" value="1"/>
                           </catDesc>
                           <category corresp="#free">
                              <desc/>
                           </category>
                        </category>
                        <category corresp="#none">
                           <catDesc>
                              <num type="boolean" value="0"/>
                           </catDesc>
                        </category>
                     </category>
                  </category>
               </category>
               <category xml:id="provision">
                  <category xml:id="basic_data_accessible">
                     <catDesc>
                        <ref target="#K4.1">cf. Catalogue 4.1</ref>
                     </catDesc>
                     <catDesc>Is the textual data accessible in a source format (e.g. XML,
                        TXT)?</catDesc>
                     <catDesc>
                        <num type="boolean" value="1"/>
                     </catDesc>
                  </category>
                  <category xml:id="download">
                     <catDesc>
                        <ref target="#K4.2">cf. Catalogue 4.2</ref>
                     </catDesc>
                     <catDesc>Can the entire raw data of the project be downloaded (as a
                        whole)?</catDesc>
                     <catDesc>
                        <num type="boolean" value="1"/>
                     </catDesc>
                  </category>
                  <category xml:id="technical_interfaces">
                     <catDesc>
                        <ref target="#K4.2">cf. Catalogue 4.2</ref>
                     </catDesc>
                     <catDesc>Are there technical interfaces which allow the reuse of the data of
                        the text collection in other contexts?</catDesc>
                     <category xml:id="OAI-PMH">
                        <catDesc>
                           <num type="boolean" value="0"/>
                        </catDesc>
                     </category>
                     <category xml:id="REST">
                        <catDesc>
                           <num type="boolean" value="0"/>
                        </catDesc>
                     </category>
                     <category xml:id="SPARQL_endpoint">
                        <catDesc>
                           <num type="boolean" value="0"/>
                        </catDesc>
                     </category>
                     <category xml:id="general_API">
                        <catDesc>
                           <num type="boolean" value="0"/>
                        </catDesc>
                     </category>
                     <category corresp="#other">
                        <catDesc>
                           <num type="boolean" value="0"/>
                        </catDesc>
                        <category corresp="#free">
                           <desc/>
                        </category>
                     </category>
                     <category corresp="#none">
                        <catDesc>
                           <num type="boolean" value="1"/>
                        </catDesc>
                     </category>
                  </category>
                  <category xml:id="analytical_data">
                     <catDesc>
                        <ref target="#K4.3">cf. Catalogue 4.3</ref>
                     </catDesc>
                     <catDesc>Besides the textual data, does the project provide analytical data
                        (e.g. statistics) to download or harvest?</catDesc>
                     <catDesc>
                        <num type="boolean" value="0"/>
                     </catDesc>
                  </category>
                  <category xml:id="reuse">
                     <catDesc>
                        <ref target="#K4.4">cf. Catalogue 4.4</ref>
                     </catDesc>
                     <catDesc>Can you use the data with other tools useful for this kind of
                        content?</catDesc>
                     <catDesc>
                        <num type="boolean" value="1"/>
                     </catDesc>
                  </category>
               </category>
               <category xml:id="user_interface">
                  <category xml:id="interface_provision">
                     <catDesc>
                        <ref target="#K5.1">cf. Catalogue 5.1</ref>
                     </catDesc>
                     <catDesc>Does the text collection have a dedicated user interface designed for
                        the collection at hand in which the texts of the collection are represented
                        and/or in which the data is analyzable?</catDesc>
                     <catDesc>
                        <num type="boolean" value="0"/>
                     </catDesc>
                  </category>
               </category>
               <category xml:id="preservation">
                  <category xml:id="documentation_project">
                     <catDesc>
                        <ref target="#K6.1">cf. Catalogue 6.1</ref>
                     </catDesc>
                     <catDesc>Does the text collection provide sufficient documentation about the
                        project in general as well as about the aims, contents and methods of the
                        text collection?</catDesc>
                     <catDesc>
                        <num type="boolean" value="1"/>
                     </catDesc>
                  </category>
                  <category xml:id="open_access">
                     <catDesc>
                        <ref target="#K6.2">cf. Catalogue 6.2</ref>
                     </catDesc>
                     <catDesc>Is the text collection Open Access?</catDesc>
                     <catDesc>
                        <num type="boolean" value="1"/>
                     </catDesc>
                  </category>
                  <category xml:id="rights">
                     <category xml:id="rights_declared">
                        <catDesc>
                           <ref target="#K6.2">cf. Catalogue 6.2</ref>
                        </catDesc>
                        <catDesc>Are the rights to (re)use the content declared?</catDesc>
                        <catDesc>
                           <num type="boolean" value="1"/>
                        </catDesc>
                     </category>
                     <category xml:id="rights_license">
                        <catDesc>
                           <ref target="#K6.2">cf. Catalogue 6.2</ref>
                        </catDesc>
                        <catDesc>Under what license are the contents released?</catDesc>
                        <category xml:id="CC0">
                           <catDesc>
                              <num type="boolean" value="0"/>
                           </catDesc>
                        </category>
                        <category xml:id="CC-BY_only">
                           <catDesc>
                              <num type="boolean" value="0"/>
                           </catDesc>
                        </category>
                        <category xml:id="CC-BY-ND">
                           <catDesc>
                              <num type="boolean" value="0"/>
                           </catDesc>
                        </category>
                        <category xml:id="CC-BY-NC">
                           <catDesc>
                              <num type="boolean" value="0"/>
                           </catDesc>
                        </category>
                        <category xml:id="CC-BY-SA">
                           <catDesc>
                              <num type="boolean" value="0"/>
                           </catDesc>
                        </category>
                        <category xml:id="CC-BY-NC-ND">
                           <catDesc>
                              <num type="boolean" value="0"/>
                           </catDesc>
                        </category>
                        <category xml:id="CC-BY-NC-SA">
                           <catDesc>
                              <num type="boolean" value="0"/>
                           </catDesc>
                        </category>
                        <category xml:id="PDM">
                           <catDesc>
                              <num type="boolean" value="0"/>
                           </catDesc>
                        </category>
                        <category corresp="#other">
                           <catDesc>
                              <num type="boolean" value="0"/>
                           </catDesc>
                           <category corresp="#free">
                              <desc/>
                           </category>
                        </category>
                        <category xml:id="no_license">
                           <catDesc>
                              <num type="boolean" value="1"/>
                           </catDesc>
                        </category>
                     </category>
                  </category>
                  <category xml:id="persistent_identification">
                     <catDesc>
                        <ref target="#K6.3">cf. Catalogue 6.3</ref>
                     </catDesc>
                     <catDesc>Are there persistent identifiers and an addressing system for the text
                        collection and/or parts/objects of it and which mechanism is used to that
                        end?</catDesc>
                     <category xml:id="DOI">
                        <catDesc>
                           <num type="boolean" value="0"/>
                        </catDesc>
                     </category>
                     <category xml:id="ARK">
                        <catDesc>
                           <num type="boolean" value="0"/>
                        </catDesc>
                     </category>
                     <category xml:id="URN">
                        <catDesc>
                           <num type="boolean" value="0"/>
                        </catDesc>
                     </category>
                     <category xml:id="PURL.ORG">
                        <catDesc>
                           <num type="boolean" value="0"/>
                        </catDesc>
                     </category>
                     <category corresp="#other">
                        <catDesc>
                           <num type="boolean" value="0"/>
                        </catDesc>
                        <category corresp="#free">
                           <desc/>
                        </category>
                     </category>
                     <category xml:id="persistent_URLs">
                        <catDesc>
                           <num type="boolean" value="0"/>
                        </catDesc>
                     </category>
                     <category corresp="#none">
                        <catDesc>
                           <num type="boolean" value="1"/>
                        </catDesc>
                     </category>
                  </category>
                  <category xml:id="citation">
                     <catDesc>
                        <ref target="#K6.3">cf. Catalogue 6.3</ref>
                     </catDesc>
                     <catDesc>Does the text collection supply citation guidelines?</catDesc>
                     <catDesc>
                        <num type="boolean" value="0"/>
                     </catDesc>
                  </category>
                  <category xml:id="archiving">
                     <catDesc>
                        <ref target="#K6.4">cf. Catalogue 6.4</ref>
                     </catDesc>
                     <catDesc>Does the documentation include information about the long term
                        sustainability of the basic data (archiving of the data)?</catDesc>
                     <catDesc>
                        <num type="boolean" value="0"/>
                     </catDesc>
                  </category>
                  <category xml:id="curation">
                     <catDesc>
                        <ref target="#K6.4">cf. Catalogue 6.4</ref>
                     </catDesc>
                     <catDesc>Does the project provide information about institutional support for
                        the curation and sustainability of the project?</catDesc>
                     <catDesc>
                        <num type="boolean" value="0"/>
                     </catDesc>
                  </category>
                  <category xml:id="completion">
                     <catDesc>
                        <ref target="#K6.4">cf. Catalogue 6.4</ref>
                     </catDesc>
                     <catDesc>Is the text collection completed?</catDesc>
                     <catDesc>
                        <num type="boolean" value="1"/>
                     </catDesc>
                  </category>
               </category>
            </taxonomy>
         </classDecl>
      </encodingDesc>
      <profileDesc>
         <langUsage>
            <language ident="en"/>
         </langUsage>
         <textClass>
            <keywords xml:lang="en">
               <term>corpus analysis</term>
               <term>english</term>
               <term>shakespeare</term>
               <term>text collection</term>
               <term>theatre</term>
            </keywords>
         </textClass>
      </profileDesc>
   </teiHeader>
   <text>
      <front>
         <div type="abstract">
            <p>
               <emph>ShakespearePlaysPlus</emph> is a freely available digital text corpus of
               William Shakespeare’s plays. The 37 plays were compiled from the Oxford University
               Press 1916 Edition of “The Complete Works of William Shakespeare” and annotated by
               Mike Scott for his own research in 2006. The plays are organized in three categories
               according to their type, i.e., comedies, historical plays and tragedies. The speeches
               of all characters have been extracted into separate text files. The text files are
               marked up in a pseudo-XML style and stored in Unicode. The corpus is downloadable as
               an extractable zip-file. This review presents a detailed look at the text corpus, its
               creation and composition. <emph>ShakespearePlaysPlus</emph> is compact and marked up
               with essential information, making it a durable, portable and easily reusable
               resource.</p>
         </div>
      </front>
      <body>
         <div xml:id="div1">
            <head>Introduction</head>
            <p xml:id="p1">
               <emph>ShakespearePlaysPlus</emph> is a freely available digital text corpus of
               William Shakespeare’s plays. The original source is the Oxford University Press 1916
               Edition of “The Complete Works of William Shakespeare (The Oxford Shakespeare /
                  OUP)”.<note xml:id="ftn1">Original source of the online version: William
                  Shakespeare, <emph>The Complete Works of William Shakespeare (The Oxford
                     Shakespeare)</emph>, ed. with a glossary by W. J. Craig M. A. (Oxford
                  University Press, 1916). See <ref target="https://web.archive.org/web/20180426191653/http://oll.libertyfund.org/titles/shakespeare-the-complete-works-of-william-shakespeare-the-oxford-shakespeare">https://web.archive.org/web/20180426191653/http://oll.libertyfund.org/titles/shakespeare-the-complete-works-of-william-shakespeare-the-oxford-shakespeare</ref>
                  for the online OLL version. Several formats are available for download, including
                  HTML (“Every effort has been taken to translate the unique features of the printed
                  book into the HTML medium.“), Ebook and ePub, as well as a facsimile of the
                  original book as an image-based PDF.</note> The corpus was compiled, processed and
               annotated by Dr. Mike Scott<note xml:id="ftn2">Dr. Mike Scott, Liverpool University
                  (1990-2009), Aston University (2010-), see <ref target="https://web.archive.org/web/20180126155259/http://www.aston.ac.uk/lss/staff-directory/scottmdr/">https://web.archive.org/web/20180126155259/http://www.aston.ac.uk/lss/staff-directory/scottmdr/</ref>.
                  Developer of WordSmith Tools (1996-2018). For further publications, see <ref target="https://web.archive.org/web/20180426191409/http://lexically.net/publications/publications.htm">https://web.archive.org/web/20180426191409/http://lexically.net/publications/publications.htm</ref>.</note>
               for his own research in 2006. The source texts were collected from the Online Library
               of Liberty (OLL) in digitized form, are in the public domain and “may be used freely
               for educational and academic purposes”. The corpus is available for download and
               re-use as a zip-file on the WordSmith Tools website<note xml:id="ftn3">Available on
                  the Wordsmith Homepage <ref target="https://web.archive.org/web/20180703103234/http://www.lexically.net/wordsmith/index.html">https://web.archive.org/web/20180703103234/http://www.lexically.net/wordsmith/index.html</ref>
                  under “Downloads”, then “Extras”, see <ref target="https://web.archive.org/web/20180118112105/http://lexically.net/wordsmith/support/shakespeare.html">https://web.archive.org/web/20180118112105/http://lexically.net/wordsmith/support/shakespeare.html</ref>.</note>
               under “Extra downloads for WordSmith Tools”. Information concerning the origin of the
               text corpus, its design, contents and format is presented on the
                  <emph>ShakespearePlaysPlus</emph> download page.</p>
         </div>
         <div xml:id="div2">
            <head>Aims, content and design of the corpus</head>
            <p xml:id="p2">This corpus provides a complete digital text collection of all of
               Shakespeare’s plays in standardized modern spelling, in a durable format that can
               easily be reused for further research. It can be used for analysis and research
               within the WordSmith environment or with any other lexical analysis tool. This
               primary source collection can be understood as a ‘general purpose corpus’, open to
               many possibilities for reuse and adaptation depending on the focus and interest of
               the user, ranging from literary to linguistic and quantitative as well as qualitative
               research.</p>
            <p xml:id="p3">The description of the content presented on the Wordsmith website is
               straightforward: “You get 37 plays, plus all the speeches of all the characters. That
               is, you get the whole play Hamlet, plus separately all the speeches of Prince Hamlet,
               all the speeches of Horatio, etc. There is also a list of the plays and their dates.”
               The overview of all the plays and years of publication is presented as a link to a
               new page.<note xml:id="ftn4">See <ref target="https://web.archive.org/web/20180703103115/http://lexically.net/downloads/corpus_linguistics/shakespeare_plays_dated_plain.txt">https://web.archive.org/web/20180703103115/http://lexically.net/downloads/corpus_linguistics/shakespeare_plays_dated_plain.txt</ref>.</note>
               The years of publication do not coincide with the dates presented on the OLL website,
               matching better with other Shakespeare sources.<note xml:id="ftn5">For simple
                  accessibility, only freely available online sources are cited in this review.
                  Please refer to the canon of Shakespeare research for in-depth discussion. The
                  publication dates of Shakespeare’s plays are not precisely historically
                  documented, a topic which will not be discussed here in length. See, e.g., the
                  timelines presented online by the Royal Shakespeare Company <ref target="https://web.archive.org/web/20180118115432/https://www.rsc.org.uk/shakespeares-plays/timeline">https://web.archive.org/web/20180118115432/https://www.rsc.org.uk/shakespeares-plays/timeline</ref>,
                  on Open Source Shakespeare <ref target="https://web.archive.org/web/20180118115735/https://www.opensourceshakespeare.org/views/plays/plays_date.php">https://web.archive.org/web/20180118115735/https://www.opensourceshakespeare.org/views/plays/plays_date.php</ref>,
                  or Shakespeare Online <ref target="https://web.archive.org/web/20180118120823/http://shakespeare-online.com/keydates/playchron.html">https://web.archive.org/web/20180118120823/http://shakespeare-online.com/keydates/playchron.html</ref>,
                  which also lists the presumable years of first performance vs. publication.</note>
            </p>
            <p xml:id="p4">Basic relevant documentary information regarding format and design is
               presented in a pop-up link on the same page. The text files are stored in 16-bit
               Unicode format.<note xml:id="ftn6">There are no special characters that make 16-bit
                  Unicode necessary, but, as Scott notes, it is a generally useful encoding.</note>
               The zip-file is 5.4 MB large and unzips to 24.3 MB on the computer. It contains 41
               file folders and 1387 files altogether.</p>
            <p xml:id="p5">The composition of this collection is guided by the principle of
               completeness, i.e., the complete plays of Shakespeare.<note xml:id="ftn7">The number
                  of plays written by Shakespeare is a controversial topic. The Royal Shakespeare
                  Company (RSC), e.g., writes “[i]t is believed that he wrote around 38 plays,
                  including collaborations with other writers”, see <ref target="https://web.archive.org/web/20180118115432/https://www.rsc.org.uk/shakespeares-plays/timeline">https://web.archive.org/web/20180118115432/https://www.rsc.org.uk/shakespeares-plays/timeline</ref>.
                  The RSC include the play “Two noble Kinsmen” (1613-1614) in their listing, which
                  is not included in the Oxford 1916 edition. This play is also listed on
                  Shakespeare Online <ref target="https://web.archive.org/web/20180124122014/http://www.shakespeare-online.com/keydates/playchron.html">https://web.archive.org/web/20180124122014/http://www.shakespeare-online.com/keydates/playchron.html</ref>,
                  where it is noted that “all but a few scholars believe it not to be an original
                  work by Shakespeare. The majority of the play was probably written by John
                  Fletcher, who was a prominent actor and Shakespeare's close friend” (Mabillard
                  2000).</note> The 37 plays provided in the 1916 Oxford edition are organized
               according to three main categories, namely comedies, historical and tragedies. The
               plays and the respective “character speeches” sub-folders are available at the root
               level of each category folder.</p>
            <p xml:id="p6">
               <figure xml:id="img1">
                  <graphic url="http://ride.i-d-e.de/wp-content/uploads/issue_9/shakespeare-plays/pictures/picture-1.png"/>
                  <head type="legend">Character speech folders, play files and <emph>Hamlet</emph>’s
                     character speech files.</head>
               </figure> The 1313 extracted speech files are listed alphabetically in separate,
               accordingly named, sub-folders, e.g., “Hamlet, Prince of Denmark_characters” (see
                  <ref type="crossref" target="#img1">Fig. 1</ref>). A plain text alphabetical
               listing of each play’s characters is included in the character speech folder, i.e.,
               “dramatis_personae”, providing an overview of the play’s characters. These files
               reflect the <emph>Dramatis Personae</emph> listing at the beginning of each printed
               original play, however, they differ a bit from the source version in that they do not
               include further character descriptions (e.g., Friend to Hamlet, a Soldier, Officers,
               Courtiers, etc.) and some of the characters are listed more specifically, e.g., Clown
               1 and Clown 2 instead of ‘two clowns’ as in the OUP. <figure xml:id="code1">
                  <eg>ALL<lb/> AMBASSADOR 1<lb/> BERNARDO<lb/> CAPTAIN<lb/> CLOWN 1<lb/> CLOWN
                     2<lb/> CORNELIUS<lb/> DANES<lb/> FORTINBRAS<lb/> […] </eg>
                  <head type="legend">Excerpt from <emph>Hamlet</emph>’s dramatis personae file
                     list.</head>
               </figure> The dramatis personae file can be useful for specifying parameters in text
               analysis tools, e.g., the analysis of selected character speeches and their context,
               or creating concordances of a full play text without focus on the speakers.</p>
            <p xml:id="p7">The contents of the corpus provide what they promise, as described on the
               website. The structure of the data storage is clear and easy to understand. The
               simple structure makes it easy to integrate and reuse the text files in other
               systems.</p>
            <div xml:id="div2.1">
               <head>Data modeling</head>
               <p xml:id="p8">The text files include embedded annotations in the form of “pseudo-XML
                  angle-bracketed tags”.<note xml:id="ftn8">See link ‘Format information’ on the
                     download page, <ref target="https://web.archive.org/web/20180118112105/http://lexically.net/wordsmith/support/shakespeare.html">https://web.archive.org/web/20180118112105/http://lexically.net/wordsmith/support/shakespeare.html</ref>.</note>
                  The annotations reflect the structure and contents of the play. Further analytic
                  information not included in the HTML source texts in this form has been added,
                  namely percentages indicating the relative position of the text segment<note xml:id="ftn9">The OLL HTML does include information about the numbering of the
                     text lines, but not in percentage form. For example, one sees Craig1916: 4 in
                     the front-end browser version, standing for line 4, reflecting the numbering as
                     printed in the original book, see, e.g., <ref target="https://web.archive.org/web/20180124112122/http://oll.libertyfund.org/titles/shakespeare-hamlet-prince-of-denmark--5">https://web.archive.org/web/20180124112122/http://oll.libertyfund.org/titles/shakespeare-hamlet-prince-of-denmark--5</ref>.
                     In the HTML code, each text line has a specific embedded numbering (as viewed
                     in December 2017).</note> within the entirety of the play.<note xml:id="ftn10">At the very beginning of a play it is &lt;0%&gt; and at the end of the play
                     &lt;100%&gt; (or &lt;99%&gt;, if it ends with a single long speech, e.g., in
                        <emph>The Life of King Henry V</emph>). Depending on the length of the
                     speeches, a single percent may be repeated within the markup of several
                     following character speeches (see <ref type="crossref" target="#img3">Fig.
                        3</ref>), or one long speech can cover more than a single percent, e.g., in
                        <emph>The Life of King Henry V</emph>, a King Henry speech is marked with
                     &lt;5%&gt;, and the following speech in marked with &lt;7%&gt; (Act 1, Scene
                     2).</note> The percentages are based on the number of tokens altogether.</p>
               <p xml:id="p9">Each play consists of a single text file. Following the basic markup
                  structure of HTML/XML,<note xml:id="ftn11">The format is, however, not well-formed
                     XML because (1) the tag names contain whitespace, (2) angle brackets are used
                     to surround the stage directions, percentages, and metadata, (3) metadata,
                     percentages, and markup in character speech files are not surrounded by opening
                     and closing tags, i.e., not properly nested and (4) there is no root
                     element.</note> each act, scene and character speech, as well as extra
                  angle-bracketed stage directions, are surrounded by opening and closing
                  angle-bracketed tags. Speeches are indented, making them easier to read while
                  maintaining the visual form of standard HTML/XML. Each speech tag also contains
                  the bracketed percentage information within the first line. <figure xml:id="code2">
                     <eg> &lt;ACT 1&gt;<lb/> &lt;SCENE 1&gt;<lb/> &lt;Elsinore. A Platform before
                        the Castle.&gt;<lb/> &lt;STAGE DIR&gt;<lb/> &lt;Francisco at his post. Enter
                        to him Bernardo.&gt;<lb/> &lt;/STAGE DIR&gt;<lb/> &lt;BERNARDO&gt;
                        &lt;0%&gt;<lb/> Who's there?<lb/> &lt;/BERNARDO&gt;<lb/> &lt;FRANCISCO&gt;
                        &lt;0%&gt; <lb/> Nay, answer me; stand, and unfold yourself.<lb/>
                        &lt;/FRANCISCO&gt;<lb/> &lt;BERNARDO&gt; &lt;0%&gt;<lb/> Long live the
                        king!<lb/> &lt;/BERNARDO&gt;<lb/> […]<lb/> &lt;/SCENE 1&gt; </eg>
                     <head type="legend">Excerpt from the play file at the beginning of
                           <emph>Hamlet</emph>.</head>
                  </figure>
               </p>
               <p xml:id="p10">In comparison, the character speech files are annotated without
                  opening and closing brackets, deviating from the pseudo-XML style. The annotation
                  at the beginning of each segment consists of the angle-bracketed number of the
                  speech, act and scene, as well as the relative percentage information regarding
                  the position within the play, as can be seen in the following: <figure xml:id="code3">
                     <eg>&lt;SPEECH 131&gt;&lt;ACT 3&gt;&lt;SCENE 1&gt;&lt;42%&gt;<lb/> To be, or
                        not to be: that is the question:<lb/> Whether 'tis nobler in the mind to
                        suffer<lb/> The slings and arrows of outrageous fortune,<lb/> Or to take
                        arms against a sea of troubles,<lb/> And by opposing end them? To die: to
                        sleep;<lb/> No more; and, by a sleep to say we end<lb/> The heart-ache and
                        the thousand natural shocks<lb/> That flesh is heir to, 'tis a
                        consummation<lb/> Devoutly to be wish'd. To die, to sleep;<lb/> To sleep:
                        perchance to dream: ay, there's the rub;</eg>
                     <head type="legend">Excerpt from <emph>Hamlet</emph>’s speech file.</head>
                  </figure>
               </p>
               <p xml:id="p11">Overall, the data model is simple and easy to understand; however,
                  the style of markup would make certain conversions necessary for reuse with
                  XML-technologies. It is interesting that the creator specifically chose to model
                  the texts in this manner, i.e., that the XML-version created during the
                  transformation process was not used (see below). But considering that the corpus
                  was originally created for personal research in WordSmith, by default ignoring
                  angle bracket pairs and their content,<note xml:id="ftn12">See <ref target="http://www.lexically.net/wordsmith/Handling%20BNC/tag_file.htm">http://www.lexically.net/wordsmith/Handling%20BNC/tag_file.htm</ref>, last
                     accessed May 14, 2018. Tags can be retained for the analysis with WordSmith
                     when a special tag file is created.</note> this markup style serves perfectly
                  well. This markup style privileges speeches,<note xml:id="ftn13">Compared to,
                     e.g., printed text from the source edition.</note> i.e., blending out all the
                  lines starting with angle brackets leads to a text consisting only of the
                  speeches.</p>
            </div>
            <div xml:id="div2.2">
               <head>Data acquisition and transformation</head>
               <p xml:id="p12">Besides essential information, i.e., the original book edition, the
                  source of the digital texts and the final encoding, not much is told on the
                  Wordsmith website regarding the creation and format of the original data and the
                  transformation process behind this corpus. Upon my request, Mike Scott very kindly
                  provided more detailed information, which is reused here with permission.</p>
               <p xml:id="p13">The 37 plays were downloaded individually in HTML, having “the merit
                  of having detailed style information”, as can be seen below: <figure xml:id="code4">
                     <eg>&lt;div class="sp"&gt;&lt;span
                        class="ital_speaker"&gt;Lys.&lt;/span&gt;<lb/> &lt;p style="margin-top:
                        -0.5em;"&gt;I am, my lord, as well<lb/> deriv&amp;#8217;d as
                        he,&lt;/p&gt;<lb/> &lt;p class="p-no-indent1"&gt;As well
                        possess&amp;#8217;d; my love is<lb/> more than his;&lt;span
                        class="milestone_right"<lb/>
                        title="Craig1916_line_100"&gt;100&lt;/span&gt;&lt;/p&gt;<lb/> &lt;p
                        class="p-no-indent1"&gt;My fortunes every way as fairly<lb/>
                        rank&amp;#8217;d&lt;/p&gt;<lb/> […]<lb/> &lt;div&gt; </eg>
                     <head type="legend">Sample excerpt of source HTML from OLL, <emph>A Midsummer
                           Night’s Dream</emph>.</head>
                  </figure> The original markup thus distinguishes the name of the character
                  speaking from the lines spoken (<code>class="ital_speaker"</code>) and marks the
                  line numbers in the OUP edition (<code>"milestone_right"</code>).</p>
               <p xml:id="p14">The HTML versions were converted to XML using Dreamweaver MX 2004. A
                  program was written to convert the XML into plain text (TXT) with the desired
                  markup. For this, it was necessary to <list rend="ordered">
                     <item>remove the standard headers and transform the text into Unicode with
                        easily read punctuation instead of markup such as
                           “<code>deriv&amp;#8217;d</code>” (deriv’d).</item>
                     <item>find the XML markup for Dramatis Personae and build a standardised list
                        of characters. This was problematic since in the text there are minor
                        variations of spelling and it was not always easy to identify the character.
                        In the example above, “Lys.” is easily identified as Lysander, but in some
                        cases character names were simply absent from the Dramatis Personae, stage
                        directions and the speech-wording.</item>
                     <item>identify Act and Scene numbers, mark up speech beginnings and endings so
                        that these would be easy to locate, and to export all the speeches by all
                        the characters in the plays to separate new files.</item>
                  </list> The final output is shown below: <figure xml:id="code5">
                     <eg>&lt;LYSANDER&gt; &lt;5%&gt;<lb/> I am, my lord, as well deriv'd as he,<lb/>
                        As well possess'd; my love is more than his;<lb/> My fortunes every way as
                        fairly rank'd<lb/> [...]<lb/> &lt;/LYSANDER&gt;</eg>
                     <head type="legend">TXT Extract of Fig. 5 from <emph>A Midsummer Night’s
                           dream</emph>, play file.</head>
                  </figure>
               </p>
               <p xml:id="p15">
                  <figure xml:id="img2">
                     <graphic url="http://ride.i-d-e.de/wp-content/uploads/issue_9/shakespeare-plays/pictures/picture-2.png"/>
                     <head type="legend">Lines 1 to 13 of <emph>The Merry Wives of Windsor</emph>
                        from OUP, with line numbering and italics.</head>
                  </figure>
                  <figure xml:id="img3">
                     <graphic url="http://ride.i-d-e.de/wp-content/uploads/issue_9/shakespeare-plays/pictures/picture-3.png"/>
                     <head type="legend">Lines 5 to 13 of OUP in HTML presentation, italics and line
                        numbering as in OUP.</head>
                  </figure> As can be seen above, some information was dropped from the OLL markup
                  during the transformation into TXT format, i.e., the line numbering of the OUP and
                  the markup reflecting the OUP pages (e.g., [52] for p. 52 of OUP). The cursive
                  writing in the speeches is also not marked up in the TXT versions, as can be
                  compared in the following examples (all from <emph>The Merry Wives of
                     Windsor</emph>, see <ref type="crossref" target="#img2">Fig. 2</ref>, <ref type="crossref" target="#img3">3</ref>, <ref type="crossref" target="#img4">4</ref> and <ref type="crossref" target="#code6">Code 6</ref>, <ref type="crossref" target="#code7">7</ref>).<note xml:id="ftn14">Screenshots from
                     the fascsimile of the OUP (p. 51 of OUP, p. 62 of 1390 in PDF), see <ref target="https://web.archive.org/web/20180428131356/http://lf-oll.s3.amazonaws.com/titles/1608/0612_Bk_Sm.pdf">https://web.archive.org/web/20180428131356/http://lf-oll.s3.amazonaws.com/titles/1608/0612_Bk_Sm.pdf</ref>.</note>
                  <figure xml:id="code6">
                     <eg>&lt;SLENDER&gt; &lt;1%&gt;<lb/> In the county of Gloster, justice of peace,
                        and coram.<lb/> &lt;/SLENDER&gt;<lb/> &lt;SHALLOW&gt; &lt;1%&gt;<lb/> Ay,
                        cousin Slender, and cust-alorum.<lb/> &lt;/SHALLOW&gt;<lb/> &lt;SLENDER&gt;
                        &lt;1%&gt;<lb/> Ay, and rato-lorum too; and a gentleman born, Master<lb/>
                        Parson; who writes himself armigero, in any bill,<lb/> warrant, quittance,
                        or obligation,—armigero.<lb/> &lt;/SLENDER&gt; </eg>
                     <head type="legend">Lines 5 to 13 of OUP in <emph>ShakespearePlaysPlus</emph>,
                        without line numbering and italics.</head>
                  </figure>
               </p>
               <p xml:id="p16">
                  <figure xml:id="img4">
                     <graphic url="http://ride.i-d-e.de/wp-content/uploads/issue_9/shakespeare-plays/pictures/picture-4.png"/>
                     <head type="legend">Stage directions within speeches in square brackets,
                        OUP.</head>
                  </figure> In the OUP, stage directions within speeches are marked with square
                  brackets and italics, as in <ref type="crossref" target="#img4">Fig. 4</ref>.
                  These are visually reproduced within the OLL and marked up in the HTML.</p>
               <p xml:id="p17">In the play files of <emph>ShakespearePlaysPlus</emph>, the target
                  markup replaces the square brackets with angle brackets and surrounds these with
                  &lt;STAGE DIR&gt; tags, as below. <figure xml:id="code7">
                     <eg>&lt;EVANS&gt; &lt;2%&gt;<lb/> Shall I tell you a lie? I do despise a har as
                        I do<lb/> despise one that is false; or as I despise one that is not<lb/>
                        true. The knight, Sir John, is there; and, I beseech you, be<lb/> ruled by
                        your well-willers. I will peat the door for Master<lb/> Page.<lb/> &lt;STAGE
                        DIR&gt;<lb/> &lt;Knocks.&gt;<lb/> &lt;/STAGE DIR&gt; What, hoa! Got pless
                        your house here!<lb/> &lt;/EVANS&gt;</eg>
                     <head type="legend">Stage directions within speeches in angle brackets in TXT
                        file.</head>
                  </figure>
               </p>
               <p xml:id="p18">Some deviations – which have been corrected – occurred during the
                  original transformation, however. The more relevant ones concerned 13 speech
                  segments where in-speech stage directions resulted in lines of speech text being
                  surrounded by angled brackets, enclosing 32 speech lines altogether, e.g.: <figure xml:id="code8">
                     <eg>&lt;BAPTISTA&gt; &lt;35%&gt;<lb/> A mighty man of Pisa; by report<lb/> I
                        know him well: you are very welcome, sir.<lb/> &lt;STAGE DIR&gt;<lb/> &lt;To
                        Hortensio.] Take you the lute, [To Lucentio.&gt;<lb/> &lt;/STAGE DIR&gt; and
                        you the set of books;<lb/> You shall go see your pupils presently.<lb/>
                        Holla, within! </eg>
                     <head type="legend">Speech mixed with stage directions in angle brackets in TXT
                           file.<note xml:id="ftn16">Excerpt from <emph>The Taming of the
                              Shrew</emph>, Baptista at 35%. Depending on the setting of the
                           research environment, angle-bracketed content may be ignored (as, e.g.,
                           by default in WordSmith), thus excluding these 32 lines from
                           analysis.</note>
                     </head>
                  </figure>
               </p>
               <p xml:id="p19">In the single speech files, the 32 enclosed speech lines of the play
                  files did not show up at all, thus deviating from the original OUP content and
                  resulting in a loss of information.<note xml:id="ftn17">Compared to the number of
                     lines in Shakespeare's plays altogether, the impact of the 32 enclosed speech
                     lines in the play files and the 32 corresponding missing lines in the speech
                     files may be deemed quite minimal — yet, depending on the research focus, they
                     could have distorted results to a certain degree. For example, in Rosalind’s
                     &lt;speech 192&gt; at 89 % (from <emph>As You Like It</emph>), entire nine
                     lines of speech were missing.</note> Mike Scott set about correcting these
                  deviations as soon as I informed him of my findings. The corrected version is now
                  online for future downloads.<note xml:id="ftn18">As of the 14th of June,
                     2018.</note>
               </p>
            </div>
            <div xml:id="div2.3">
               <head>Metadata</head>
               <p xml:id="p20">
                  <emph>ShakespearePlaysPlus</emph> includes embedded metadata at the beginning of
                  each play’s text file. The metadata provide basic information about the original
                  source of the texts as well as information about the collection at hand and the
                  file format: <figure xml:id="code9">
                     <eg>&lt; Shakespeare - - HAMLET, PRINCE OF DENMARK &gt;<lb/> &lt; from Online
                        Library of Liberty (http://oll.libertyfund.org) &gt;<lb/> &lt; Unicode .txt
                        version by Mike Scott (http://www.lexically.net) &gt;<lb/> &lt; from "The
                        Complete Works of William Shakespeare" &gt;<lb/> &lt; ed. with a glossary by
                        W.J. Craig M.A. &gt;<lb/> &lt; (London: Oxford University Press, 1916)
                        &gt;</eg>
                     <head type="legend">Metadata from the <emph>Hamlet</emph> play file.</head>
                  </figure> The metadata are not repeated in the single characters’ speech files. A
                  date of creation is not provided, but the entire corpus was created in 2006.<note xml:id="ftn19">Personal correspondence.</note>
               </p>
            </div>
         </div>
         <div xml:id="div3">
            <head>Further aspects: access, preservation, reuse</head>
            <p xml:id="p21">As the basic format of the text collection is 16-bit Unicode,<note xml:id="ftn20">For the Unicode Character Encoding Stability Policies, see <ref target="https://web.archive.org/web/20180118134036/http://www.unicode.org/policies/stability_policy.html">https://web.archive.org/web/20180118134036/http://www.unicode.org/policies/stability_policy.html</ref>.</note>
               it is a durable corpus<note xml:id="ftn21">Regarding encoding choices, Folger
                  Digital, e.g., notes that they offer TXT versions in ASCII-7, a subset of Unicode,
                  and that these “files are the most likely to render properly in the widest number
                  of applications and the least likely to present conversion errors when being
                  incorporated into text analysis tools.” Their unadorned TXT files are “for
                  projects and applications where simplicity and/or stability is the highest
                  priority”, see <ref target="https://web.archive.org/web/20180118123036/http://www.folgerdigitaltexts.org/download/txt.html">https://web.archive.org/web/20180118123036/http://www.folgerdigitaltexts.org/download/txt.html</ref>.</note>
               that is readable, processible and accessible for the long term. No specific user
               interface is provided or necessary. This corpus consists of plain text files, so a
               basic editor or text program can in principle suffice for human reading. For specific
               research questions, WordSmith or any other tool that works with annotated text files
               can be used as an interface. Whether the format or markup have to be adapted depends
               on the environment, e.g., for use with XML-technologies, the markup would have to be
               adapted to be well-formed (<ref target="#ftn11">see footnote 11</ref>).<note xml:id="ftn22">It may be noted that other online Shakespeare sources, e.g., Folger
                  Digital Texts, provide XML and TEI markup for download. Folger Digital provides
                  complex XML markup, including markup for line breaks, individual words and
                  punctuation marks, as well as IDs for each tag. The TEISimple version is also
                  quite complex, including linguistic annotation and single line IDs, see <ref target="https://web.archive.org/web/20180124140232/http://www.folgerdigitaltexts.org/download/">https://web.archive.org/web/20180124140232/http://www.folgerdigitaltexts.org/download/</ref>.
                  Christof Schöch (2014) notes in his review of Folger Digital that the texts were
                  marked up in an exemplary manner for maximal interoperability with other
                  resources.</note>
            </p>
            <p xml:id="p22">The text collection has brief documentation regarding source, format,
               annotation and design on the website. The original source of the collection, as well
               as creator, contact and download website address, are provided in the metadata of
               each of the plays. There are no persistent identifiers. The original source texts are
               in the public domain and no specific license is in use.</p>
            <p xml:id="p23">Due to its simple and durable design, this text corpus has good
               prospects of being available for research and reuse for a long time. Continuous
               access is provided by the WordSmith Tools Website. The website entered its 22th year
               of existence,<note xml:id="ftn23">See <ref target="https://web.archive.org/web/20180703103234/http://www.lexically.net/wordsmith/index.html">https://web.archive.org/web/20180703103234/http://www.lexically.net/wordsmith/index.html</ref>.</note>
               so one may be optimistic that it will remain online well into the future. The corpus
               is not officially mirrored or archived in an academic repository, but it is archived
               at the WayBackMachine Internet Repository.</p>
            <p xml:id="p24">Regarding reuse, at least one author has used this corpus as a basis for
               her MA thesis investigating key word clusters in the dialogue of Shakespeare’s female
                  characters.<note xml:id="ftn24">Demmen, Jane (2009). <emph>Charmed and chattering
                     tongues: Investigating the functions and effects of key word clusters in the
                     dialogue of Shakespeare’s female characters</emph>. Lancaster University, MA
                  thesis (unpublished). See <ref target="https://web.archive.org/web/20180118142818/http://lexically.net/wordsmith/corpus_linguistics_links/Jane%20Demmen_MA_word%20clusters_in_Shakespeare%27s%20plays.pdf">https://web.archive.org/web/20180118142818/http://lexically.net/wordsmith/corpus_linguistics_links/Jane%20Demmen_MA_word%20clusters_in_Shakespeare%27s%20plays.pdf</ref>.</note>
               Scott does not have additional records of people or projects using this corpus, but
               expects that “quite a lot of students will have downloaded it for their term
                  papers”.<note xml:id="ftn25">Personal correspondence.</note>
            </p>
         </div>
         <div xml:id="div4">
            <head>Conclusion</head>
            <p xml:id="p25">
               <emph>ShakespearePlaysPlus</emph> is a practical text collection that is open for
               many potential reuses. Important aspects behind the creation of the corpus were
               completeness, durability, reusability and standardized modern English spelling. The
               markup reflects the structure of the play (act, scene, stage direction, speeches of
               characters), and the relative position of the speech tags. Considering more complex
               data models and current standards of research data management, some modifications
               could possibly be undertaken, e.g., adapting the metadata (adding distinct IDs/DOIs
               and creation dates, including the metadata in the speech files) and the markup to be
               well-formed XML.</p>
            <p xml:id="p26">There is no particular theoretical stance behind the text corpus, at
               least not explicitly, but a special interest in the analysis of speeches may be
               presumed. As Mike Scott writes on the website, “I made this version for my own
               research. If you use it please let me know!” And further in personal correspondence,
               “I thought others might find it useful, no other great ambition!”.</p>
            <p xml:id="p27">
               <emph>ShakespearePlaysPlus</emph> is a durable and portable text corpus marked up
               with relevant information. Making this practical corpus created for personal use
               available to the public for further reuse and research was a collegial and friendly
               deed in the spirit of open access, adding to the pool of digital resources available
               for linguistic and literary Shakespeare research.</p>
         </div>
      </body>
      <back>
         <div type="bibliography">
            <listBibl>
               <bibl>
                  <emph>Folger Digital Texts</emph>. “Download.” <ref target="https://web.archive.org/web/20180124140232/http://www.folgerdigitaltexts.org/download/">https://web.archive.org/web/20180124140232/http://www.folgerdigitaltexts.org/download/</ref>.</bibl>
               <bibl>Mabillard, Amanda. 2000. “The Chronology of Shakespeare's Plays.”
                     <emph>Shakespeare Online</emph>. <ref target="https://web.archive.org/web/20180124122014/http://www.shakespeare-online.com/keydates/playchron.html">https://web.archive.org/web/20180124122014/http://www.shakespeare-online.com/keydates/playchron.html</ref>.</bibl>
               <bibl>
                  <emph>Online Library of Liberty (OLL)</emph>, <ref target="https://web.archive.org/web/20180118112729/http://oll.libertyfund.org">https://web.archive.org/web/20180118112729/http://oll.libertyfund.org</ref>.</bibl>
               <bibl>
                  <emph>Open Source Shakespeare</emph>. 2003-2018. “Shakespeare's plays, listed by
                  presumed date of composition.” <ref target="https://web.archive.org/web/20180118115735/https://www.opensourceshakespeare.org/views/plays/plays_date.php">https://web.archive.org/web/20180118115735/https://www.opensourceshakespeare.org/views/plays/plays_date.php</ref>.</bibl>
               <bibl>
                  <emph>Royal Shakespeare Company</emph>. “Timeline of Shakespeare's plays. A
                  chronological list of Shakespeare's plays by decade.” <ref target="https://web.archive.org/web/20180118115432/https://www.rsc.org.uk/shakespeares-plays/timeline">https://web.archive.org/web/20180118115432/https://www.rsc.org.uk/shakespeares-plays/timeline</ref>.</bibl>
               <bibl>Schöch, Christof. 2014. “Michael Poston and Rebecca Niles, eds.: Folger Digital
                  Texts. Washington: The Folger Shakespeare Library, 2012-2013.” <emph>Variants -
                     The Journal of the European Society for Textual Scholarship</emph>, pp. 16-20.
                     <ref target="https://zenodo.org/record/13745">https://zenodo.org/record/13745</ref>.</bibl>
               <bibl>Scott, Mike. 2016. “Extra Downloads for WordSmith Tools: Shakespeare Corpus.”
                     <emph>WordSmith Tools.</emph>
                  <ref target="https://web.archive.org/web/20180118112105/http://lexically.net/wordsmith/support/shakespeare.html">https://web.archive.org/web/20180118112105/http://lexically.net/wordsmith/support/shakespeare.html</ref>.</bibl>
               <bibl>
                  <emph>Shakespeare Online.</emph> 1999-2012. “The Chronology of Shakespeare's
                  Plays.” <ref target="https://web.archive.org/web/20180118120823/http://shakespeare-online.com/keydates/playchron.html">https://web.archive.org/web/20180118120823/http://shakespeare-online.com/keydates/playchron.html</ref>.</bibl>
               <bibl>Shakespeare, William. 1916. <emph>The Complete Works of William Shakespeare
                     (The Oxford Shakespeare)</emph>. Edited with a glossary by W. J. Craig M. A.
                  Oxford: Oxford University Press. HTML: <ref target="https://web.archive.org/web/20180126162402/http://oll.libertyfund.org/titles/shakespeare-the-complete-works-of-william-shakespeare-the-oxford-shakespeare">https://web.archive.org/web/20180126162402/http://oll.libertyfund.org/titles/shakespeare-the-complete-works-of-william-shakespeare-the-oxford-shakespeare</ref>.
                  Facsimile: <ref target="https://web.archive.org/web/20180428131356/http://lf-oll.s3.amazonaws.com/titles/1608/0612_Bk_Sm.pdf">https://web.archive.org/web/20180428131356/http://lf-oll.s3.amazonaws.com/titles/1608/0612_Bk_Sm.pdf</ref>.</bibl>
               <bibl>
                  <emph>WordSmith Tools.</emph> Lexical Analysis Software and Oxford University
                  Press. <ref target="https://web.archive.org/web/20180703103234/http://www.lexically.net/wordsmith/index.html">https://web.archive.org/web/20180703103234/http://www.lexically.net/wordsmith/index.html</ref>.
               </bibl>
            </listBibl>
         </div>
      </back>
   </text>
</TEI>
