[{"data":1,"prerenderedAt":-1},["ShallowReactive",2],{"detail-sidebar-cat-0-en-105":3,"doc-seo-154715-105":59,"doc-detail-154715-en":130},{"code":4,"msg":5,"data":6},0,"success",[7,13,18,23,28,33,38,43,48,51,55],{"id":8,"doc_module":4,"doc_module_name":9,"category_name":10,"show_sort_weight":11,"slug":12},1,"Document","Story & Novel",90,"story-novel",{"id":14,"doc_module":4,"doc_module_name":9,"category_name":15,"show_sort_weight":16,"slug":17},2,"Literature",80,"literature",{"id":19,"doc_module":4,"doc_module_name":9,"category_name":20,"show_sort_weight":21,"slug":22},4,"Exam",70,"exam",{"id":24,"doc_module":4,"doc_module_name":9,"category_name":25,"show_sort_weight":26,"slug":27},5,"Comic",60,"comic",{"id":29,"doc_module":4,"doc_module_name":9,"category_name":30,"show_sort_weight":31,"slug":32},6,"Technology",50,"technology",{"id":34,"doc_module":4,"doc_module_name":9,"category_name":35,"show_sort_weight":36,"slug":37},7,"Healthcare",40,"healthcare",{"id":39,"doc_module":4,"doc_module_name":9,"category_name":40,"show_sort_weight":41,"slug":42},8,"Research & Report",30,"research-report",{"id":44,"doc_module":4,"doc_module_name":9,"category_name":45,"show_sort_weight":46,"slug":47},9,"Religion & Spirituality",20,"religion-spirituality",{"id":46,"doc_module":4,"doc_module_name":9,"category_name":49,"show_sort_weight":46,"slug":50},"World Cup","world-cup",{"id":52,"doc_module":4,"doc_module_name":9,"category_name":53,"show_sort_weight":52,"slug":54},10,"Lifestyle","lifestyle",{"id":56,"doc_module":4,"doc_module_name":9,"category_name":57,"show_sort_weight":24,"slug":58},19,"General","general",{"code":4,"msg":60,"data":61},"ok",{"site_id":62,"language":63,"slug":64,"title":65,"keywords":66,"description":67,"schema_data":68,"social_meta":123,"head_meta":125,"extra_data":127,"updated_unix":129},105,"en","a-new-semantic-artifact-based-framework-for-studying-and-documenting-algospeak-and-related-phenomena","A New Semantic Artifact Based Framework for Studying and Documenting Algospeak and Related Phenomena","","A new framework is proposed to analyze, document, and publish resources on algospeak, a coded form of linguistic self-censorship used to evade social media content moderation. The framework relies on two semantic artifacts provided as SKOS semantic artifacts in RDF and a cross-lingual lexicon designed to support comparison across languages and cultural settings. The paper discusses algospeak use in Italian and Arabic based on data collected to ground the framework’s categories and taxonomy design.",{"@graph":69,"@context":122},[70,84,105],{"@type":71,"itemListElement":72},"BreadcrumbList",[73,77,79,82],{"item":74,"name":75,"@type":76,"position":8},"https://docshare.wps.com","Home","ListItem",{"item":78,"name":9,"@type":76,"position":14},"https://docshare.wps.com/document/",{"item":80,"name":40,"@type":76,"position":81},"https://docshare.wps.com/document/research-report/",3,{"item":83,"name":65,"@type":76,"position":19},"https://docshare.wps.com/document/a-new-semantic-artifact-based-framework-for-studying-and-documenting-algospeak-and-related-phenomena/154715/",{"url":83,"name":65,"@type":85,"image":86,"author":91,"headline":65,"publisher":94,"fileFormat":97,"inLanguage":63,"description":67,"dateModified":98,"datePublished":99,"encodingFormat":97,"isAccessibleForFree":100,"interactionStatistic":101},"DigitalDocument",{"url":87,"@type":88,"width":89,"height":90},"https://docshare.wps.com/thumbnails/a-new-semantic-artifact-based-framework-for-studying-and-documenting-algospeak-and-related-phenomena/154715.png","ImageObject",300,407,{"name":92,"@type":93},"Cipher","Person",{"url":74,"name":95,"@type":96},"DocShare","Organization","application/pdf","2026-10-08","2026-08-28",true,{"@type":102,"interactionType":103,"userInteractionCount":44},"InteractionCounter",{"@type":104},"ViewAction",{"@type":106,"mainEntity":107},"FAQPage",[108,114,118],{"name":109,"@type":110,"acceptedAnswer":111},"What is algospeak and what is its main function on social media?","Question",{"text":112,"@type":113},"Algospeak is coded language or linguistic self-censorship intended to bypass content moderation algorithms. It often works through abbreviations, misspellings, or substitutions of words and symbols.","Answer",{"name":115,"@type":110,"acceptedAnswer":116},"What components does the proposed framework use to study algospeak?",{"text":117,"@type":113},"The framework uses two semantic artifacts published as SKOS semantic artifacts in RDF, plus a cross-lingual lexicon following a schema to enable comparison across languages and cultural contexts.",{"name":119,"@type":110,"acceptedAnswer":120},"How does the paper apply the framework to languages beyond English?",{"text":121,"@type":113},"It discusses algospeak in Italian and Arabic, using data collection conducted by the authors to refine and expand initial categories underlying the framework.","https://schema.org",{"og:url":83,"og:type":124,"og:title":65,"og:site_name":95,"og:description":67},"article",{"robots":126,"canonical":83},"index,follow",{"doc_id":128,"site_id":62},154715,1787897455,{"code":4,"msg":5,"data":131},{"doc_id":128,"user_id":132,"nickname":92,"user_avatar":133,"doc_module":4,"category_id":39,"category_name":40,"doc_title":65,"doc_description":67,"doc_content":134,"file_id":135,"file_url":136,"file_type":137,"file_size":138,"view_count":44,"is_deleted":4,"is_public":8,"is_downloadable":8,"audit_status":8,"page_count":139,"language":140,"language_code":63,"site_id":62,"html_lang":63,"table_of_contents":141,"faqs":142,"seo_title":143,"seo_description":67,"update_tm":129,"read_time":144},687208528416,"https://ap-avatar.wpscdn.com/davatar_9964176cb1d06d4a9deccf72a44ae3dc","A New Semantic Artifact Based Framework for Studying and Documenting Algospeak and Related Phenomena  \nAnas Fahad Khan 1 , Elisa Gugliotta4 , Elisa Squadrito 1 ,2 , Maura Tarquini3 ,  \nFrancesca Frontini 1  \n1.CNR-ILC, 2 . Università degli Studi di Macerata, 3 . Università degli Studi di Sassari-DiSSUF,  \n4. Universitè Grenoble Alpes-LIDILEM  \n{fahad. khan,elisa.squadrito,[francesca.frontini}@ilc.cnr.it](francesca.frontini}@ilc.cnr.it), {mtarquini,[egugliotta}@uniss.it](egugliotta}@uniss.it)  \nAbstract  \nIn this paper we present a new framework for analysis, documenting and publishing resources about the recent linguistic phenomenon of algospeak. This proposed framework features the use of two semantic artifacts (both of which we make available as SKOS semantic artifacts in RDF), and a cross-lingual lexicon of algospeak terms which follows a schema intended to facilitate the comparison of algospeak across languages and cultural contexts. Our article also features a discussion of the use of algospeak in two non-anglophone contexts (Italian and Arabic) which resulted from a period of data collection which the authors undertook as preparation for the creation of our framework and the categories which underlie it.  \nKeywords: algospeak, semantic artifacts, arabic, italian, social media, anti-language, emoji  \n1. Introduction  \nAlgospeak is a form of coded language or linguistic self-censorship used to bypass content moderation algorithms on social media platforms1 . It is “commonly understood as abbreviating, misspelling, or substituting specific words,[...] when creating a social media post with the particular goal to circumvent a platform’s content moderation systems”(Steen et al. , 2023) . Algospeak is most commonly associated with Tiktok but has also become increasingly common on other plat-  \nAll authors contributed to the study and the linguistic resources described. The development of the two taxonomies described above was largely carried out by Khan. Material preparation and data collection were largely carried out by Khan with respect to the English dataset and by Squadrito with respect to the Italian dataset; the analysis of the Italian examples given below was carried out by Squadrito. Tarquini and Gugliotta were responsible for material prerparation and data collection of the Arabic dataset, whose empirical analysis contributed to the refinement and expansion of initial resources. Frontini provided supervision with regard to FAIR data principles, data management, and ethical considerations related to resource publication and interoperability. The structure of the manuscript was jointly developed by all of the authors. The initial draft was written collaboratively, and all of the authors contributed to reviewing and editing subsequent versions. All authors approved the final manuscript.  \n1 Content Warning: This document discusses examples of harmful content (hate, abuse, misinformation and negative stereotypes) . The authors do not support the use of harmful language, nor any of the harmful representations quoted below.  \nforms such as Reddit, Instagram and Youtube2 . Beyond this, a number of instances of algospeak have entered into wider cultural circulation and/or have begun to be used for reasons other than circumventing social media platform content moderation systems.  \nWell known examples of algospeak are p0rn and , both stand ins for the word porn; the  emoji for Palestine; and unalive for dead. Algospeak can frequently be as simple as the substitution of numbers for letters, letters for letters, or the substitution of emojis for words or morphemes, and in this it is similar to an earlier kind of internet language, namely, leetspeak. However it can often be more sophisticated than that – sophisticated at least from a linguistic point of view – by playing on phonological similarities between words or making use of speakers’ morphological and lexical knowledge, and combining these together with other (often obscur","cbCaipJ613Md4Zpz","https://ap.wps.com/l/cbCaipJ613Md4Zpz","pdf",247059,11,"English","# Introduction\n## Background on algospeak\n# Framework overview\n## Semantic artifacts and multilingual lexicon\n# Applications in Italian and Arabic\n## Formation processes and tendencies in Italian\n## Arabic studies using the framework","[{\"question\":\"What is algospeak and what is its main function on social media?\",\"answer\":\"Algospeak is coded language or linguistic self-censorship intended to bypass content moderation algorithms. It often works through abbreviations, misspellings, or substitutions of words and symbols.\"},{\"question\":\"What components does the proposed framework use to study algospeak?\",\"answer\":\"The framework uses two semantic artifacts published as SKOS semantic artifacts in RDF, plus a cross-lingual lexicon following a schema to enable comparison across languages and cultural contexts.\"},{\"question\":\"How does the paper apply the framework to languages beyond English?\",\"answer\":\"It discusses algospeak in Italian and Arabic, using data collection conducted by the authors to refine and expand initial categories underlying the framework.\"}]","A New Semantic Artifact Based Framework for Studying and Documenting Algospeak and Related Phenomena | PDF",28]