[{"data":1,"prerenderedAt":-1},["ShallowReactive",2],{"doc-seo-241620-105":3,"detail-sidebar-cat-1-en-105":84,"doc-detail-241620-en":130},{"code":4,"msg":5,"data":6},0,"ok",{"site_id":7,"language":8,"slug":9,"title":10,"keywords":11,"description":12,"schema_data":13,"social_meta":77,"head_meta":79,"extra_data":81,"updated_unix":83},105,"en","document-metadata-extraction-multimodal-241620","Document Metadata Extraction (Multimodal)","","This document outlines the critical requirements and instructions for extracting metadata from various document types, including text and images. It details the process for handling multimodal content, emphasizing language detection rules, content analysis, and structured data generation. The guide provides specific protocols for extracting titles, ensuring language consistency across all metadata fields, and formatting the output as a single JSON object. It also covers the generation of abstracts, keywords, tables of contents, and frequently asked questions, with strict adherence to length, diversity, and quality constraints. The document serves as a comprehensive manual for automated metadata extraction systems, ensuring accurate and standardized results.",{"@graph":14,"@context":76},[15,34,55],{"@type":16,"itemListElement":17},"BreadcrumbList",[18,23,27,31],{"item":19,"name":20,"@type":21,"position":22},"https://docshare.wps.com","Home","ListItem",1,{"item":24,"name":25,"@type":21,"position":26},"https://docshare.wps.com/template/","Template",2,{"item":28,"name":29,"@type":21,"position":30},"https://docshare.wps.com/template/general/","General",3,{"item":32,"name":10,"@type":21,"position":33},"https://docshare.wps.com/template/document-metadata-extraction-multimodal-241620/241620/",4,{"url":32,"name":10,"@type":35,"image":36,"author":41,"headline":10,"publisher":44,"fileFormat":47,"inLanguage":8,"description":12,"dateModified":48,"datePublished":49,"encodingFormat":47,"isAccessibleForFree":50,"interactionStatistic":51},"DigitalDocument",{"url":37,"@type":38,"width":39,"height":40},"https://docshare.wps.com/thumbnails/document-metadata-extraction-multimodal-241620/241620.png","ImageObject",442,249,{"name":42,"@type":43},"Grenda","Person",{"url":19,"name":45,"@type":46},"DocShare","Organization","application/pdf","2026-09-21","2026-09-12",true,{"@type":52,"interactionType":53,"userInteractionCount":30},"InteractionCounter",{"@type":54},"ViewAction",{"@type":56,"mainEntity":57},"FAQPage",[58,64,68,72],{"name":59,"@type":60,"acceptedAnswer":61},"What is the primary output format required for this document metadata extraction process?","Question",{"text":62,"@type":63},"The primary output format required is a single raw JSON object, starting with '{' and ending with '}', with no additional characters or markdown formatting.","Answer",{"name":65,"@type":60,"acceptedAnswer":66},"What is the language priority for detecting the document's language?",{"text":67,"@type":63},"The language priority is determined by meaningful natural language sentences in the body text and image text. If insufficient, the document title and then the file name are used. The dominant language rule prioritizes the language in which the majority of full sentences are written.",{"name":69,"@type":60,"acceptedAnswer":70},"What are the length constraints for the 'abstract' field?",{"text":71,"@type":63},"The 'abstract' field must be between a minimum of 450 characters and a maximum of 500 characters. Density is achieved through adjective precision and thematic depth, not repetition.",{"name":73,"@type":60,"acceptedAnswer":74},"How should the 'toc' (table of contents) be formatted, and when should it be omitted?",{"text":75,"@type":63},"The 'toc' should be in markdown format with at most two levels (# and ##). It should be omitted entirely if the document is too short (under 100 words) or a single flat paragraph with no discernible sub-topics.","https://schema.org",{"og:url":32,"og:type":78,"og:title":10,"og:site_name":45,"og:description":12},"article",{"robots":80,"canonical":32},"index,follow",{"doc_id":82,"site_id":7},241620,1789178227,{"code":4,"msg":85,"data":86},"success",[87,92,97,102,107,112,117,122,127],{"id":88,"doc_module":22,"doc_module_name":25,"category_name":89,"show_sort_weight":90,"slug":91},11,"Presentations",90,"presentations",{"id":93,"doc_module":22,"doc_module_name":25,"category_name":94,"show_sort_weight":95,"slug":96},12,"Resumes",80,"resumes",{"id":98,"doc_module":22,"doc_module_name":25,"category_name":99,"show_sort_weight":100,"slug":101},14,"Invoices",70,"invoices",{"id":103,"doc_module":22,"doc_module_name":25,"category_name":104,"show_sort_weight":105,"slug":106},15,"Posters",60,"posters",{"id":108,"doc_module":22,"doc_module_name":25,"category_name":109,"show_sort_weight":110,"slug":111},16,"Social Media",50,"social-media",{"id":113,"doc_module":22,"doc_module_name":25,"category_name":114,"show_sort_weight":115,"slug":116},17,"Forms",40,"forms",{"id":118,"doc_module":22,"doc_module_name":25,"category_name":119,"show_sort_weight":120,"slug":121},18,"Letters",30,"letters",{"id":123,"doc_module":22,"doc_module_name":25,"category_name":124,"show_sort_weight":125,"slug":126},21,"Paper Templates",5,"papers-templates",{"id":128,"doc_module":22,"doc_module_name":25,"category_name":29,"show_sort_weight":4,"slug":129},158,"general-158",{"code":4,"msg":85,"data":131},{"doc_id":82,"user_id":132,"nickname":42,"user_avatar":133,"doc_module":22,"category_id":128,"category_name":29,"doc_title":10,"doc_description":12,"doc_content":11,"file_id":134,"file_url":135,"file_type":136,"file_size":137,"view_count":30,"is_deleted":4,"is_public":22,"is_downloadable":22,"audit_status":22,"page_count":98,"language":138,"language_code":8,"site_id":7,"html_lang":8,"table_of_contents":139,"faqs":140,"seo_title":141,"seo_description":12,"update_tm":83,"read_time":125},7971474921005,"https://ap-avatar.wpscdn.com/davatar_29158cc5080c5b710cf443261637dec0","cbCaik081QxlXZ3c","https://ap.wps.com/l/cbCaik081QxlXZ3c","pdf",173545,"English","# Document Metadata Extraction (Multimodal)\n## Instructions\n## Language Rules","[{\"question\":\"What is the primary output format required for this document metadata extraction process?\",\"answer\":\"The primary output format required is a single raw JSON object, starting with '{' and ending with '}', with no additional characters or markdown formatting.\"},{\"question\":\"What is the language priority for detecting the document's language?\",\"answer\":\"The language priority is determined by meaningful natural language sentences in the body text and image text. If insufficient, the document title and then the file name are used. The dominant language rule prioritizes the language in which the majority of full sentences are written.\"},{\"question\":\"What are the length constraints for the 'abstract' field?\",\"answer\":\"The 'abstract' field must be between a minimum of 450 characters and a maximum of 500 characters. Density is achieved through adjective precision and thematic depth, not repetition.\"},{\"question\":\"How should the 'toc' (table of contents) be formatted, and when should it be omitted?\",\"answer\":\"The 'toc' should be in markdown format with at most two levels (# and ##). It should be omitted entirely if the document is too short (under 100 words) or a single flat paragraph with no discernible sub-topics.\"}]","Document Metadata Extraction (Multimodal) | PDF"]