[{"data":1,"prerenderedAt":-1},["ShallowReactive",2],{"detail-sidebar-cat-1-en-105":3,"doc-seo-191330-105":53,"doc-detail-191330-en":126},{"code":4,"msg":5,"data":6},0,"success",[7,14,19,24,29,34,39,44,49],{"id":8,"doc_module":9,"doc_module_name":10,"category_name":11,"show_sort_weight":12,"slug":13},11,1,"Template","Presentations",90,"presentations",{"id":15,"doc_module":9,"doc_module_name":10,"category_name":16,"show_sort_weight":17,"slug":18},12,"Resumes",80,"resumes",{"id":20,"doc_module":9,"doc_module_name":10,"category_name":21,"show_sort_weight":22,"slug":23},14,"Invoices",70,"invoices",{"id":25,"doc_module":9,"doc_module_name":10,"category_name":26,"show_sort_weight":27,"slug":28},15,"Posters",60,"posters",{"id":30,"doc_module":9,"doc_module_name":10,"category_name":31,"show_sort_weight":32,"slug":33},16,"Social Media",50,"social-media",{"id":35,"doc_module":9,"doc_module_name":10,"category_name":36,"show_sort_weight":37,"slug":38},17,"Forms",40,"forms",{"id":40,"doc_module":9,"doc_module_name":10,"category_name":41,"show_sort_weight":42,"slug":43},18,"Letters",30,"letters",{"id":45,"doc_module":9,"doc_module_name":10,"category_name":46,"show_sort_weight":47,"slug":48},21,"Paper Templates",5,"papers-templates",{"id":50,"doc_module":9,"doc_module_name":10,"category_name":51,"show_sort_weight":4,"slug":52},158,"General","general-158",{"code":4,"msg":54,"data":55},"ok",{"site_id":56,"language":57,"slug":58,"title":59,"keywords":60,"description":61,"schema_data":62,"social_meta":119,"head_meta":121,"extra_data":123,"updated_unix":125},105,"en","document-metadata-extraction-multimodal","Document Metadata Extraction (Multimodal)","","This document outlines critical instructions for extracting metadata from documents, including text and image analysis. It provides detailed rules for title generation, language detection, abstract and keyword creation, table of contents formatting, and frequently asked questions. The system emphasizes language consistency across all extracted fields and specifies strict output requirements, including a single-line JSON format, character limits for titles, and the mandatory inclusion of all eight metadata fields. Specific guidelines are provided for handling different document types and for ensuring the accuracy and SEO-friendliness of extracted information. The process mandates a rigorous self-check to confirm language adherence and compliance with all formatting and content rules before outputting the final JSON object.",{"@graph":63,"@context":118},[64,80,101],{"@type":65,"itemListElement":66},"BreadcrumbList",[67,71,74,77],{"item":68,"name":69,"@type":70,"position":9},"https://docshare.wps.com","Home","ListItem",{"item":72,"name":10,"@type":70,"position":73},"https://docshare.wps.com/template/",2,{"item":75,"name":51,"@type":70,"position":76},"https://docshare.wps.com/template/general/",3,{"item":78,"name":59,"@type":70,"position":79},"https://docshare.wps.com/template/document-metadata-extraction-multimodal/191330/",4,{"url":78,"name":59,"@type":81,"image":82,"author":87,"headline":59,"publisher":90,"fileFormat":93,"inLanguage":57,"description":61,"dateModified":94,"datePublished":95,"encodingFormat":93,"isAccessibleForFree":96,"interactionStatistic":97},"DigitalDocument",{"url":83,"@type":84,"width":85,"height":86},"https://docshare.wps.com/thumbnails/document-metadata-extraction-multimodal/191330.png","ImageObject",442,249,{"name":88,"@type":89},"Rizky","Person",{"url":68,"name":91,"@type":92},"DocShare","Organization","application/pdf","2026-09-29","2026-09-03",true,{"@type":98,"interactionType":99,"userInteractionCount":8},"InteractionCounter",{"@type":100},"ViewAction",{"@type":102,"mainEntity":103},"FAQPage",[104,110,114],{"name":105,"@type":106,"acceptedAnswer":107},"What is the primary goal of this document?","Question",{"text":108,"@type":109},"The primary goal is to provide a comprehensive set of instructions for extracting structured metadata from documents, combining text and image analysis.","Answer",{"name":111,"@type":106,"acceptedAnswer":112},"What are the key requirements for the output format?",{"text":113,"@type":109},"The output must be a single raw JSON object, starting with '{' and ending with '}', with all eight metadata fields present and strictly adhering to language and formatting rules.",{"name":115,"@type":106,"acceptedAnswer":116},"How should the dominant language of a document be determined?",{"text":117,"@type":109},"The dominant language is determined by the majority of full natural language sentences. Isolated words, code, URLs, numbers, and filler text are ignored. This also applies to other language-dependent fields like abstract, keywords, and FAQs.","https://schema.org",{"og:url":78,"og:type":120,"og:title":59,"og:site_name":91,"og:description":61},"article",{"robots":122,"canonical":78},"index,follow",{"doc_id":124,"site_id":56},191330,1788408028,{"code":4,"msg":5,"data":127},{"doc_id":124,"user_id":128,"nickname":88,"user_avatar":129,"doc_module":9,"category_id":50,"category_name":51,"doc_title":59,"doc_description":61,"doc_content":130,"file_id":131,"file_url":132,"file_type":133,"file_size":134,"view_count":8,"is_deleted":4,"is_public":9,"is_downloadable":9,"audit_status":9,"page_count":73,"language":135,"language_code":57,"site_id":56,"html_lang":57,"table_of_contents":136,"faqs":137,"seo_title":138,"seo_description":61,"update_tm":125,"read_time":9},962085564807,"https://ap-avatar.wpscdn.com/davatar_6f874abed73319feea01a86fa6f0fab8","| Salvatore Vitale | dotloop verified 12/29/23 1:10 PM CST 7JHT-TEZO-BIWK-AAE0 |\n| --- | --- |","cbCaife7hjupIF4B","https://ap.wps.com/l/cbCaife7hjupIF4B","pdf",581710,"English","# Document Metadata Extraction (Multimodal)\n## Instructions\n## Language Rules\n## Field Specifications\n### Title Rules\n### `title_en`\n### `language`\n### abstract\n### `keywords`\n### `toc`\n### `faqs`\n### `category`\n## Output Requirements (STRICT)\n## Inputs","[{\"question\":\"What is the primary goal of this document?\",\"answer\":\"The primary goal is to provide a comprehensive set of instructions for extracting structured metadata from documents, combining text and image analysis.\"},{\"question\":\"What are the key requirements for the output format?\",\"answer\":\"The output must be a single raw JSON object, starting with '{' and ending with '}', with all eight metadata fields present and strictly adhering to language and formatting rules.\"},{\"question\":\"How should the dominant language of a document be determined?\",\"answer\":\"The dominant language is determined by the majority of full natural language sentences. Isolated words, code, URLs, numbers, and filler text are ignored. This also applies to other language-dependent fields like abstract, keywords, and FAQs.\"}]","Document Metadata Extraction (Multimodal) | PDF"]