[{"data":1,"prerenderedAt":-1},["ShallowReactive",2],{"doc-seo-422918-105":3,"detail-sidebar-cat-0-en-105":80,"doc-detail-422918-en":130},{"code":4,"msg":5,"data":6},0,"ok",{"site_id":7,"language":8,"slug":9,"title":10,"keywords":11,"description":12,"schema_data":13,"social_meta":73,"head_meta":75,"extra_data":77,"updated_unix":79},105,"en","text-processing-procedures-for-analysing-a-corpus-with-medievalmarian-miracle-tales-in-old-swedish","Text Processing Procedures for Analysing a Corpus with MedievalMarian Miracle Tales in Old Swedish","","A text corpus of 101 Marian miracle tales in Old Swedish, written between c.1272 and 1430, compiled digitally from three 19th-century transcriptions. Interpreting the medieval language requires specialized knowledge because vocabulary, spelling, grammar, and rich morphology differ markedly from modern Swedish. The paper preliminarily tests automated text-processing strategies including frequency-list analysis and spelling-variation methods to generate stop-word lists and reveal key terms, supporting word-form grouping, lexicon lookup, and lemmatization.",{"@graph":14,"@context":72},[15,34,55],{"@type":16,"itemListElement":17},"BreadcrumbList",[18,23,27,31],{"item":19,"name":20,"@type":21,"position":22},"https://docshare.wps.com","Home","ListItem",1,{"item":24,"name":25,"@type":21,"position":26},"https://docshare.wps.com/document/","Document",2,{"item":28,"name":29,"@type":21,"position":30},"https://docshare.wps.com/document/research-report/","Research & Report",3,{"item":32,"name":10,"@type":21,"position":33},"https://docshare.wps.com/document/text-processing-procedures-for-analysing-a-corpus-with-medievalmarian-miracle-tales-in-old-swedish/422918/",4,{"url":32,"name":10,"@type":35,"image":36,"author":41,"headline":10,"publisher":44,"fileFormat":47,"inLanguage":8,"description":12,"dateModified":48,"datePublished":49,"encodingFormat":47,"isAccessibleForFree":50,"interactionStatistic":51},"DigitalDocument",{"url":37,"@type":38,"width":39,"height":40},"https://docshare.wps.com/thumbnails/text-processing-procedures-for-analysing-a-corpus-with-medievalmarian-miracle-tales-in-old-swedish/422918.png","ImageObject",300,407,{"name":42,"@type":43},"Sarah ","Person",{"url":19,"name":45,"@type":46},"DocShare","Organization","application/pdf","2026-09-29","2026-09-28",true,{"@type":52,"interactionType":53,"userInteractionCount":26},"InteractionCounter",{"@type":54},"ViewAction",{"@type":56,"mainEntity":57},"FAQPage",[58,64,68],{"name":59,"@type":60,"acceptedAnswer":61},"What is the corpus used in the study, and where does it come from?","Question",{"text":62,"@type":63},"The study uses a corpus of 101 Marian miracle tales in Old Swedish, digitized from three transcribed sources published in the 19th century.","Answer",{"name":65,"@type":60,"acceptedAnswer":66},"Why are Old Swedish texts difficult to analyze computationally?",{"text":67,"@type":63},"The texts show substantial differences from modern Swedish, including spelling variation, inflectional word forms, and rich morphology that make automated tasks difficult.",{"name":69,"@type":60,"acceptedAnswer":70},"Which text-processing strategies are investigated to support analysis?",{"text":71,"@type":63},"The paper investigates strategies such as frequency-list analysis and methods for identifying spelling variations to build stop-word lists and expose key words, facilitating word-form grouping and lexicon lookup.","https://schema.org",{"og:url":32,"og:type":74,"og:title":10,"og:site_name":45,"og:description":12},"article",{"robots":76,"canonical":32},"index,follow",{"doc_id":78,"site_id":7},422918,1790693796,{"code":4,"msg":81,"data":82},"success",[83,87,91,95,100,105,110,114,119,122,126],{"id":22,"doc_module":4,"doc_module_name":25,"category_name":84,"show_sort_weight":85,"slug":86},"Story & Novel",90,"story-novel",{"id":26,"doc_module":4,"doc_module_name":25,"category_name":88,"show_sort_weight":89,"slug":90},"Literature",80,"literature",{"id":33,"doc_module":4,"doc_module_name":25,"category_name":92,"show_sort_weight":93,"slug":94},"Exam",70,"exam",{"id":96,"doc_module":4,"doc_module_name":25,"category_name":97,"show_sort_weight":98,"slug":99},5,"Comic",60,"comic",{"id":101,"doc_module":4,"doc_module_name":25,"category_name":102,"show_sort_weight":103,"slug":104},6,"Technology",50,"technology",{"id":106,"doc_module":4,"doc_module_name":25,"category_name":107,"show_sort_weight":108,"slug":109},7,"Healthcare",40,"healthcare",{"id":111,"doc_module":4,"doc_module_name":25,"category_name":29,"show_sort_weight":112,"slug":113},8,30,"research-report",{"id":115,"doc_module":4,"doc_module_name":25,"category_name":116,"show_sort_weight":117,"slug":118},9,"Religion & Spirituality",20,"religion-spirituality",{"id":117,"doc_module":4,"doc_module_name":25,"category_name":120,"show_sort_weight":117,"slug":121},"World Cup","world-cup",{"id":123,"doc_module":4,"doc_module_name":25,"category_name":124,"show_sort_weight":123,"slug":125},10,"Lifestyle","lifestyle",{"id":127,"doc_module":4,"doc_module_name":25,"category_name":128,"show_sort_weight":96,"slug":129},19,"General","general",{"code":4,"msg":81,"data":131},{"doc_id":78,"user_id":132,"nickname":42,"user_avatar":133,"doc_module":4,"category_id":111,"category_name":29,"doc_title":10,"doc_description":12,"doc_content":134,"file_id":135,"file_url":136,"file_type":137,"file_size":138,"view_count":26,"is_deleted":4,"is_public":22,"is_downloadable":22,"audit_status":22,"page_count":106,"language":139,"language_code":8,"site_id":7,"html_lang":8,"table_of_contents":140,"faqs":141,"seo_title":142,"seo_description":12,"update_tm":143,"read_time":144},962085320529,"https://ap-avatar.wpscdn.com/davatar_9964176cb1d06d4a9deccf72a44ae3dc","# Text Processing Procedures for Analysing a Corpus with MedievalMarian Miracle Tales in Old Swedish\n\nBengt Dahlqvist  \nDepartment of Linguistics and Philology,Uppsala University,P.O.Box 635,75126 Uppsala,Sweden  \nKeywords:Text Mining,Medieval Texts,Miracle Stories,Old Swedish,Stop Words,Word Similarity,Spelling Variations,Key Words.  \nAbstract:A text corpus of one hundred and one Marian Miracle stories in Old Swedish written between c.1272 and1430 has been digitally compiled from three transcribed sources from the 19th Century.Highly specializedknowledge is needed to interpret these texts,since the medieval variant of Swedish differs significantly fromthe modern form of the language.Both the vocabulary and spelling as well as the grammar show substantialvariances compared to modern Swedish.To advance the understanding of these texts,automated tools fortextual processing are needed.This paper preliminary investigates a number of strategies,such as frequencylist analysis and methods for identifying spelling variations for producing stop word lists and exposing thekey words of the texts.This can be a helpto understand the texts,identifying different word forms of the sameword,to ease a lexicon lookup and be a starting point for lemmatisation.  \ndeveloping more automated text processing tools foranalysing texts written in Old Swedish,primarilywith the aim to mine and deconstruct texts intoconstituents and group similar forms together.  \n## 1 INTRODUCTION\n\nTo make computer analyses of texts in Old Swedish,which was in use between 1225 and 1534,is still notan easy task.The standard lexicon for Old Swedishwas prepared in the late 19h Century by the Swedishphilologist(Soderwall,1884).An electronic versionof this exist,but gives no support for spelling variantsnor word inflections,which makes it hard to usepractically for unknown word forms.Very littlesupport exists for automated tasks.For instance,anefficient part-of-speech tagger or a parser is not to befound,even if some work in this area has been donein the last years(Adesam,2016).Neither has notmuch been done regarding the essential problems forthis type of text,which shows many features that areproblematic to handle,foremost in the area of the richmorphology and abundance of word forms.A non-specialist user,wishing to understand or eventranslate a given text,faces many problems.In thispaper,a number of text processing strategies will bediscussed and applied to a small corpus of medievalmiracle tales.  \nAside from this,the texts themselves areinteresting as witnesses of religious thinking at thetime and as evidence of the interchange andtranslation of textual material within the Europeanmedieval culture.More and better tools for theanalysis of Old Swedish as such may in this way pavethe way for more serious studies of the content ofthese types of texts,and help to facilate both literaryunderstanding and analysis.  \n## 2 THE DATA MATERIAL\n\nThe data material used in this study consists of 101medieval miracle tales in Old Swedish where theVirgin Mary figures as a saint and wonder worker.This text collection is believed to consist of allsurviving complete tales of this kind.  \nMiracle tales constitute a specific subgenre inmedieval religious writing,aside from hagiographies,visions and moral tales.The contents of miracle taleson the whole are for the most part purely apocryphal,in that they do not originate from the Bible.They alsooften take place in later times,after the death of the  \nEspecially,focus here will be given to the inherentproblems of word analysis,the elimination of wordvariants and text normalisation to be able to identifyword content with rich lexical meaning.This can beseen as a first step in an ongoing research aiming at  \n452  \nprotagonist saint.Stories of this type were commonin many European languages at the time,from Latinand Greek to French and German.Translations wereproduced into several other languages,including OldNorwegian,Danish and Swedish.  \nIn the fol","cbCaisTMSz0a68El","https://ap.wps.com/l/cbCaisTMSz0a68El","pdf",1164711,"English","# 1 INTRODUCTION\n# 2 THE DATA MATERIAL\n## 2.1 The Old Swedish Legendary\n## 2.2 Book of Miracles\n## 2.3 Solance for the","[{\"question\":\"What is the corpus used in the study, and where does it come from?\",\"answer\":\"The study uses a corpus of 101 Marian miracle tales in Old Swedish, digitized from three transcribed sources published in the 19th century.\"},{\"question\":\"Why are Old Swedish texts difficult to analyze computationally?\",\"answer\":\"The texts show substantial differences from modern Swedish, including spelling variation, inflectional word forms, and rich morphology that make automated tasks difficult.\"},{\"question\":\"Which text-processing strategies are investigated to support analysis?\",\"answer\":\"The paper investigates strategies such as frequency-list analysis and methods for identifying spelling variations to build stop-word lists and expose key words, facilitating word-form grouping and lexicon lookup.\"}]","Text Processing Procedures for Analysing a Corpus with MedievalMarian Miracle Tales in Old Swedish | PDF",1790629294,18]