[{"data":1,"prerenderedAt":-1},["ShallowReactive",2],{"doc-detail-185956-en":3,"doc-seo-185956-105":31,"detail-sidebar-cat-0-en-105":92},{"code":4,"msg":5,"data":6},0,"success",{"doc_id":7,"user_id":8,"nickname":9,"user_avatar":10,"doc_module":4,"category_id":11,"category_name":12,"doc_title":13,"doc_description":14,"doc_content":15,"file_id":16,"file_url":17,"file_type":18,"file_size":19,"view_count":20,"is_deleted":4,"is_public":21,"is_downloadable":21,"audit_status":21,"page_count":22,"language":23,"language_code":24,"site_id":25,"html_lang":24,"table_of_contents":26,"faqs":27,"seo_title":28,"seo_description":14,"update_tm":29,"read_time":30},185956,5909887256941,"Mason","https://ap-avatar.wpscdn.com/davatar_9964176cb1d06d4a9deccf72a44ae3dc",19,"General","Document Metadata Extraction (Multimodal)","This document details a multimodal approach to extracting metadata, combining text analysis with image processing. It outlines a structured framework for analyzing document content, including text extraction from images, language detection, and metadata generation such as title, abstract, keywords, table of contents, and frequently asked questions. The system prioritizes accuracy and adherence to specified formatting rules, ensuring a comprehensive and standardized output for diverse documents. Special attention is given to critical language rules, title formatting, and content extraction from both textual and visual elements to provide rich metadata.","| No. Nama PDPD Kelas Nilai PT Kebutuhan PDPD |  |  |  |  |  |\n| --- | --- | --- | --- | --- | --- |\n| 1 | AFM | DKV 1 | 76 | 1. | Model pembelajaran dengan media sekolah |\n| 2 | ARA | Animasi | 75 | 2. | Penggunaan informasi tertulis |\n| 3 | AMT | DKV 1 | 64 | 1. | Model pembelajaran dengan pendekatan |\n| 4 | CF | DKV 2 | 71 |  | individual |\n| 5 | DSW | MP2 | 75 | 2. | Rentang waktu belajar lebih lama |\n\n\n| Subjek | a. Monitoring dan Evaluasi | b. Perbaikan |\n| --- | --- | --- |\n| 1. KS | 5 | 4 |\n| 2. P1 | 1 | 1 |\n| 3. P2 | 0 | 0 |\n| 4. BK1 | 1 | 0 |\n| 5. BK2 | 0 | 0 |\n| 6. GR1 | 0 | 0 |\n| 7. GR2 | 0 | 0 |\n| 8. W | 0 | 0 |\n| 9. PDPD | 0 | 0 |","cbCaifiu9dKfoKRR","https://ap.wps.com/l/cbCaifiu9dKfoKRR","pdf",708893,2,1,16,"English","en",105,"# Document Metadata Extraction (Multimodal)\n## Instructions\n## Language Rules\n## Field Specifications\n### Title Rules\n### `title_en`\n### `language`\n### abstract\n### `keywords`\n### `toc`\n### `faqs`\n### `category`\n## Output Requirements (STRICT)\n## Inputs","[{\"question\":\"What is the primary goal of this document?\",\"answer\":\"The primary goal is to describe a multimodal system for extracting metadata from documents, integrating both text and image analysis.\"},{\"question\":\"What are the key components of the metadata extraction process described?\",\"answer\":\"The process involves text extraction from images, language detection, combined analysis of text and image data, and generation of structured metadata fields like title, abstract, keywords, TOC, and FAQs.\"},{\"question\":\"What are the critical rules for language detection and application in this system?\",\"answer\":\"The system prioritizes dominant language from sentences, enforces language consistency across all text fields except title_en, and defaults to 'en' only when no other language signal is found.\"}]","Document Metadata Extraction (Multimodal) | PDF",1788370045,25,{"code":4,"msg":32,"data":33},"ok",{"site_id":25,"language":24,"slug":34,"title":13,"keywords":35,"description":14,"schema_data":36,"social_meta":87,"head_meta":89,"extra_data":91,"updated_unix":29},"document-metadata-extraction-multimodal","",{"@graph":37,"@context":86},[38,54,69],{"@type":39,"itemListElement":40},"BreadcrumbList",[41,45,48,51],{"item":42,"name":43,"@type":44,"position":21},"https://docshare.wps.com","Home","ListItem",{"item":46,"name":47,"@type":44,"position":20},"https://docshare.wps.com/document/","Document",{"item":49,"name":12,"@type":44,"position":50},"https://docshare.wps.com/document/general/",3,{"item":52,"name":13,"@type":44,"position":53},"https://docshare.wps.com/document/document-metadata-extraction-multimodal/185956/",4,{"url":52,"name":13,"@type":55,"author":56,"headline":13,"publisher":58,"fileFormat":61,"inLanguage":24,"description":14,"dateModified":62,"datePublished":63,"encodingFormat":61,"isAccessibleForFree":64,"interactionStatistic":65},"DigitalDocument",{"name":9,"@type":57},"Person",{"url":42,"name":59,"@type":60},"DocShare","Organization","application/pdf","2026-09-05","2026-09-02",true,{"@type":66,"interactionType":67,"userInteractionCount":20},"InteractionCounter",{"@type":68},"ViewAction",{"@type":70,"mainEntity":71},"FAQPage",[72,78,82],{"name":73,"@type":74,"acceptedAnswer":75},"What is the primary goal of this document?","Question",{"text":76,"@type":77},"The primary goal is to describe a multimodal system for extracting metadata from documents, integrating both text and image analysis.","Answer",{"name":79,"@type":74,"acceptedAnswer":80},"What are the key components of the metadata extraction process described?",{"text":81,"@type":77},"The process involves text extraction from images, language detection, combined analysis of text and image data, and generation of structured metadata fields like title, abstract, keywords, TOC, and FAQs.",{"name":83,"@type":74,"acceptedAnswer":84},"What are the critical rules for language detection and application in this system?",{"text":85,"@type":77},"The system prioritizes dominant language from sentences, enforces language consistency across all text fields except title_en, and defaults to 'en' only when no other language signal is found.","https://schema.org",{"og:url":52,"og:type":88,"og:title":13,"og:site_name":59,"og:description":14},"article",{"robots":90,"canonical":52},"index,follow",{"doc_id":7,"site_id":25},{"code":4,"msg":5,"data":93},[94,98,102,106,111,116,121,126,131,134,138],{"id":21,"doc_module":4,"doc_module_name":47,"category_name":95,"show_sort_weight":96,"slug":97},"Story & Novel",90,"story-novel",{"id":20,"doc_module":4,"doc_module_name":47,"category_name":99,"show_sort_weight":100,"slug":101},"Literature",80,"literature",{"id":53,"doc_module":4,"doc_module_name":47,"category_name":103,"show_sort_weight":104,"slug":105},"Exam",70,"exam",{"id":107,"doc_module":4,"doc_module_name":47,"category_name":108,"show_sort_weight":109,"slug":110},5,"Comic",60,"comic",{"id":112,"doc_module":4,"doc_module_name":47,"category_name":113,"show_sort_weight":114,"slug":115},6,"Technology",50,"technology",{"id":117,"doc_module":4,"doc_module_name":47,"category_name":118,"show_sort_weight":119,"slug":120},7,"Healthcare",40,"healthcare",{"id":122,"doc_module":4,"doc_module_name":47,"category_name":123,"show_sort_weight":124,"slug":125},8,"Research & Report",30,"research-report",{"id":127,"doc_module":4,"doc_module_name":47,"category_name":128,"show_sort_weight":129,"slug":130},9,"Religion & Spirituality",20,"religion-spirituality",{"id":129,"doc_module":4,"doc_module_name":47,"category_name":132,"show_sort_weight":129,"slug":133},"World Cup","world-cup",{"id":135,"doc_module":4,"doc_module_name":47,"category_name":136,"show_sort_weight":135,"slug":137},10,"Lifestyle","lifestyle",{"id":11,"doc_module":4,"doc_module_name":47,"category_name":12,"show_sort_weight":107,"slug":139},"general"]