[{"data":1,"prerenderedAt":-1},["ShallowReactive",2],{"detail-sidebar-cat-1-en-105":3,"doc-seo-196932-105":53,"doc-detail-196932-en":127},{"code":4,"msg":5,"data":6},0,"success",[7,14,19,24,29,34,39,44,49],{"id":8,"doc_module":9,"doc_module_name":10,"category_name":11,"show_sort_weight":12,"slug":13},11,1,"Template","Presentations",90,"presentations",{"id":15,"doc_module":9,"doc_module_name":10,"category_name":16,"show_sort_weight":17,"slug":18},12,"Resumes",80,"resumes",{"id":20,"doc_module":9,"doc_module_name":10,"category_name":21,"show_sort_weight":22,"slug":23},14,"Invoices",70,"invoices",{"id":25,"doc_module":9,"doc_module_name":10,"category_name":26,"show_sort_weight":27,"slug":28},15,"Posters",60,"posters",{"id":30,"doc_module":9,"doc_module_name":10,"category_name":31,"show_sort_weight":32,"slug":33},16,"Social Media",50,"social-media",{"id":35,"doc_module":9,"doc_module_name":10,"category_name":36,"show_sort_weight":37,"slug":38},17,"Forms",40,"forms",{"id":40,"doc_module":9,"doc_module_name":10,"category_name":41,"show_sort_weight":42,"slug":43},18,"Letters",30,"letters",{"id":45,"doc_module":9,"doc_module_name":10,"category_name":46,"show_sort_weight":47,"slug":48},21,"Paper Templates",5,"papers-templates",{"id":50,"doc_module":9,"doc_module_name":10,"category_name":51,"show_sort_weight":4,"slug":52},158,"General","general-158",{"code":4,"msg":54,"data":55},"ok",{"site_id":56,"language":57,"slug":58,"title":59,"keywords":60,"description":61,"schema_data":62,"social_meta":120,"head_meta":122,"extra_data":124,"updated_unix":126},105,"en","document-metadata-extraction-multimodal-196932","Document Metadata Extraction \"Multimodal\"","","This document outlines the critical requirements for Multimodal Document Metadata Extraction, focusing on the robust processing of both text and image data to generate structured metadata. It details strict rules for language detection, prioritizing body and image text, and enforcing language consistency across all metadata fields. Specific instructions are provided for title generation, including handling orphaned serials, file names, and document titles, with a feature detection and formatting process that adheres to a 100-character limit. The document emphasizes the extraction of SEO-friendly abstracts, relevant keywords, a two-level table of contents, and a diverse set of FAQs, all while adhering to provided category definitions. A key aspect is the meticulous preparation of the output as a single, raw JSON object, ensuring all special characters are correctly escaped and the final output is syntactically valid and adheres to strict formatting guidelines.",{"@graph":63,"@context":119},[64,80,102],{"@type":65,"itemListElement":66},"BreadcrumbList",[67,71,74,77],{"item":68,"name":69,"@type":70,"position":9},"https://docshare.wps.com","Home","ListItem",{"item":72,"name":10,"@type":70,"position":73},"https://docshare.wps.com/template/",2,{"item":75,"name":51,"@type":70,"position":76},"https://docshare.wps.com/template/general/",3,{"item":78,"name":59,"@type":70,"position":79},"https://docshare.wps.com/template/document-metadata-extraction-multimodal-196932/196932/",4,{"url":78,"name":59,"@type":81,"image":82,"author":87,"headline":59,"publisher":90,"fileFormat":93,"inLanguage":57,"description":61,"dateModified":94,"datePublished":95,"encodingFormat":93,"isAccessibleForFree":96,"interactionStatistic":97},"DigitalDocument",{"url":83,"@type":84,"width":85,"height":86},"https://docshare.wps.com/thumbnails/document-metadata-extraction-multimodal-196932/196932.png","ImageObject",442,249,{"name":88,"@type":89},"Genevieve","Person",{"url":68,"name":91,"@type":92},"DocShare","Organization","application/pdf","2026-09-27","2026-09-03",true,{"@type":98,"interactionType":99,"userInteractionCount":101},"InteractionCounter",{"@type":100},"ViewAction",8,{"@type":103,"mainEntity":104},"FAQPage",[105,111,115],{"name":106,"@type":107,"acceptedAnswer":108},"What is the primary goal of the Document Metadata Extraction process described?","Question",{"text":109,"@type":110},"The primary goal is to extract readable text from images and combine it with document text to generate structured metadata in a specific JSON format.","Answer",{"name":112,"@type":107,"acceptedAnswer":113},"What are the key language rules for this metadata extraction process?",{"text":114,"@type":110},"The process enforces strict language detection, prioritizing body and image text, and requires all text fields in the JSON output to match the detected dominant language.",{"name":116,"@type":107,"acceptedAnswer":117},"How should the document title be formatted according to the instructions?",{"text":118,"@type":110},"Titles must be in the document's detected original language, strictly adhere to a 100-character limit, and use hyphens to separate distinct components like main titles and subtitles.","https://schema.org",{"og:url":78,"og:type":121,"og:title":59,"og:site_name":91,"og:description":61},"article",{"robots":123,"canonical":78},"index,follow",{"doc_id":125,"site_id":56},196932,1788463747,{"code":4,"msg":5,"data":128},{"doc_id":125,"user_id":129,"nickname":88,"user_avatar":130,"doc_module":9,"category_id":50,"category_name":51,"doc_title":59,"doc_description":61,"doc_content":60,"file_id":131,"file_url":132,"file_type":133,"file_size":134,"view_count":101,"is_deleted":4,"is_public":9,"is_downloadable":9,"audit_status":9,"page_count":135,"language":136,"language_code":57,"site_id":56,"html_lang":57,"table_of_contents":137,"faqs":138,"seo_title":139,"seo_description":61,"update_tm":126,"read_time":140},1374391974585,"https://ap-avatar.wpscdn.com/davatar_276721f389ce27ea32af1340a28f341c","cbCaiqZvx8zku1pT","https://ap.wps.com/l/cbCaiqZvx8zku1pT","pdf",4105478,138,"English","# Document Metadata Extraction (Multimodal)\n## Instructions\n## Language Rules\n## Field Specifications\n### Title Rules\n### `title_en`\n### `language`\n### abstract\n### `keywords`\n### `toc`\n### `faqs`\n### `category`\n## Output Requirements (STRICT)\n## Inputs","[{\"question\":\"What is the primary goal of the Document Metadata Extraction process described?\",\"answer\":\"The primary goal is to extract readable text from images and combine it with document text to generate structured metadata in a specific JSON format.\"},{\"question\":\"What are the key language rules for this metadata extraction process?\",\"answer\":\"The process enforces strict language detection, prioritizing body and image text, and requires all text fields in the JSON output to match the detected dominant language.\"},{\"question\":\"How should the document title be formatted according to the instructions?\",\"answer\":\"Titles must be in the document's detected original language, strictly adhere to a 100-character limit, and use hyphens to separate distinct components like main titles and subtitles.\"}]","Document Metadata Extraction \"Multimodal\" | PDF",48]