{"record_type":"pith_number_record","schema_url":"https://pith.science/schemas/pith-number/v1.json","pith_number":"pith:2020:M7CEEOPAHP6SZHERGOVSOLJHSN","short_pith_number":"pith:M7CEEOPA","schema_version":"1.0","canonical_sha256":"67c44239e03bfd2c9c9133ab272d2793502de3c58426821ec458ebd164707624","source":{"kind":"arxiv","id":"2006.01038","version":3},"attestation_state":"computed","paper":{"title":"DocBank: A Benchmark Dataset for Document Layout Analysis","license":"http://creativecommons.org/licenses/by-nc-nd/4.0/","headline":"","cross_cats":[],"primary_cat":"cs.CL","authors_text":"Furu Wei, Lei Cui, Minghao Li, Ming Zhou, Shaohan Huang, Yiheng Xu, Zhoujun Li","submitted_at":"2020-06-01T16:04:30Z","abstract_excerpt":"Document layout analysis usually relies on computer vision models to understand documents while ignoring textual information that is vital to capture. Meanwhile, high quality labeled datasets with both visual and textual information are still insufficient. In this paper, we present \\textbf{DocBank}, a benchmark dataset that contains 500K document pages with fine-grained token-level annotations for document layout analysis. DocBank is constructed using a simple yet effective way with weak supervision from the \\LaTeX{} documents available on the arXiv.com. With DocBank, models from different mod"},"verification_status":{"content_addressed":true,"pith_receipt":true,"author_attested":false,"weak_author_claims":0,"strong_author_claims":0,"externally_anchored":false,"storage_verified":false,"citation_signatures":0,"replication_records":0,"graph_snapshot":true,"references_resolved":false,"formal_links_present":false},"canonical_record":{"source":{"id":"2006.01038","kind":"arxiv","version":3},"metadata":{"license":"http://creativecommons.org/licenses/by-nc-nd/4.0/","primary_cat":"cs.CL","submitted_at":"2020-06-01T16:04:30Z","cross_cats_sorted":[],"title_canon_sha256":"193538b4dd107de6a783c07b27ade42c90f4538c5299fd2a2e964aeb22a39294","abstract_canon_sha256":"694538e42d30cf32e8a2d2981fcb85cddee65e46a16cc915698556392b4bc254"},"schema_version":"1.0"},"receipt":{"kind":"pith_receipt","key_id":"pith-v1-2026-05","algorithm":"ed25519","signed_at":"2026-07-05T01:50:47.428265Z","signature_b64":"QqMgjxYFbnLA3K5bXscCnJfj14C9AMipHvm0SAawa/hJUtjJzDfurwgKym+Nou7Rcsl6x85rnlGjSCurJ4UHAQ==","signed_message":"canonical_sha256_bytes","builder_version":"pith-number-builder-2026-05-17-v1","receipt_version":"0.3","canonical_sha256":"67c44239e03bfd2c9c9133ab272d2793502de3c58426821ec458ebd164707624","last_reissued_at":"2026-07-05T01:50:47.427865Z","signature_status":"signed_v1","first_computed_at":"2026-07-05T01:50:47.427865Z","public_key_fingerprint":"8d4b5ee74e4693bcd1df2446408b0d54"},"graph_snapshot":{"paper":{"title":"DocBank: A Benchmark Dataset for Document Layout Analysis","license":"http://creativecommons.org/licenses/by-nc-nd/4.0/","headline":"","cross_cats":[],"primary_cat":"cs.CL","authors_text":"Furu Wei, Lei Cui, Minghao Li, Ming Zhou, Shaohan Huang, Yiheng Xu, Zhoujun Li","submitted_at":"2020-06-01T16:04:30Z","abstract_excerpt":"Document layout analysis usually relies on computer vision models to understand documents while ignoring textual information that is vital to capture. Meanwhile, high quality labeled datasets with both visual and textual information are still insufficient. In this paper, we present \\textbf{DocBank}, a benchmark dataset that contains 500K document pages with fine-grained token-level annotations for document layout analysis. DocBank is constructed using a simple yet effective way with weak supervision from the \\LaTeX{} documents available on the arXiv.com. With DocBank, models from different mod"},"claims":{"count":0,"items":[],"snapshot_sha256":"258153158e38e3291e3d48162225fcdb2d5a3ed65a07baac614ab91432fd4f57"},"source":{"id":"2006.01038","kind":"arxiv","version":3},"verdict":{"id":null,"model_set":{},"created_at":null,"strongest_claim":"","one_line_summary":"","pipeline_version":null,"weakest_assumption":"","pith_extraction_headline":""},"integrity":{"clean":true,"summary":{"advisory":0,"critical":0,"by_detector":{},"informational":0},"endpoint":"/pith/2006.01038/integrity.json","findings":[],"available":true,"detectors_run":[],"snapshot_sha256":"c28c3603d3b5d939e8dc4c7e95fa8dfce3d595e45f758748cecf8e644a296938"},"references":{"count":0,"sample":[],"resolved_work":0,"snapshot_sha256":"258153158e38e3291e3d48162225fcdb2d5a3ed65a07baac614ab91432fd4f57","internal_anchors":0},"formal_canon":{"evidence_count":0,"snapshot_sha256":"258153158e38e3291e3d48162225fcdb2d5a3ed65a07baac614ab91432fd4f57"},"author_claims":{"count":0,"strong_count":0,"snapshot_sha256":"258153158e38e3291e3d48162225fcdb2d5a3ed65a07baac614ab91432fd4f57"},"builder_version":"pith-number-builder-2026-05-17-v1"},"aliases":[{"alias_kind":"arxiv","alias_value":"2006.01038","created_at":"2026-07-05T01:50:47.427925+00:00"},{"alias_kind":"arxiv_version","alias_value":"2006.01038v3","created_at":"2026-07-05T01:50:47.427925+00:00"},{"alias_kind":"doi","alias_value":"10.48550/arxiv.2006.01038","created_at":"2026-07-05T01:50:47.427925+00:00"},{"alias_kind":"pith_short_12","alias_value":"M7CEEOPAHP6S","created_at":"2026-07-05T01:50:47.427925+00:00"},{"alias_kind":"pith_short_16","alias_value":"M7CEEOPAHP6SZHER","created_at":"2026-07-05T01:50:47.427925+00:00"},{"alias_kind":"pith_short_8","alias_value":"M7CEEOPA","created_at":"2026-07-05T01:50:47.427925+00:00"}],"events":[],"event_summary":{},"paper_claims":[],"inbound_citations":{"count":6,"internal_anchor_count":0,"sample":[{"citing_arxiv_id":"2606.29845","citing_title":"Bricker to BRACE: A Bracket Exposure RAW Dataset and Restoration Model for Flicker-Banding","ref_index":21,"is_internal_anchor":false},{"citing_arxiv_id":"2410.05970","citing_title":"PDF-WuKong: A Large Multimodal Model for Efficient Long PDF Reading with End-to-End Sparse Sampling","ref_index":30,"is_internal_anchor":false},{"citing_arxiv_id":"2410.21169","citing_title":"Document Parsing Unveiled: Techniques, Challenges, and Prospects for Structured Information Extraction","ref_index":119,"is_internal_anchor":false},{"citing_arxiv_id":"2605.12623","citing_title":"DocAtlas: Multilingual Document Understanding Across 80+ Languages","ref_index":57,"is_internal_anchor":false},{"citing_arxiv_id":"2605.12623","citing_title":"DocAtlas: Multilingual Document Understanding Across 80+ Languages","ref_index":63,"is_internal_anchor":false},{"citing_arxiv_id":"2604.11042","citing_title":"Improving Layout Representation Learning Across Inconsistently Annotated Datasets via Agentic Harmonization","ref_index":7,"is_internal_anchor":false}]},"formal_canon":{"evidence_count":0,"sample":[],"anchors":[]},"links":{"html":"https://pith.science/pith/M7CEEOPAHP6SZHERGOVSOLJHSN","json":"https://pith.science/pith/M7CEEOPAHP6SZHERGOVSOLJHSN.json","graph_json":"https://pith.science/api/pith-number/M7CEEOPAHP6SZHERGOVSOLJHSN/graph.json","events_json":"https://pith.science/api/pith-number/M7CEEOPAHP6SZHERGOVSOLJHSN/events.json","paper":"https://pith.science/paper/M7CEEOPA"},"agent_actions":{"view_html":"https://pith.science/pith/M7CEEOPAHP6SZHERGOVSOLJHSN","download_json":"https://pith.science/pith/M7CEEOPAHP6SZHERGOVSOLJHSN.json","view_paper":"https://pith.science/paper/M7CEEOPA","resolve_alias":"https://pith.science/api/pith-number/resolve?arxiv=2006.01038&json=true","fetch_graph":"https://pith.science/api/pith-number/M7CEEOPAHP6SZHERGOVSOLJHSN/graph.json","fetch_events":"https://pith.science/api/pith-number/M7CEEOPAHP6SZHERGOVSOLJHSN/events.json","actions":{"anchor_timestamp":"https://pith.science/pith/M7CEEOPAHP6SZHERGOVSOLJHSN/action/timestamp_anchor","attest_storage":"https://pith.science/pith/M7CEEOPAHP6SZHERGOVSOLJHSN/action/storage_attestation","attest_author":"https://pith.science/pith/M7CEEOPAHP6SZHERGOVSOLJHSN/action/author_attestation","sign_citation":"https://pith.science/pith/M7CEEOPAHP6SZHERGOVSOLJHSN/action/citation_signature","submit_replication":"https://pith.science/pith/M7CEEOPAHP6SZHERGOVSOLJHSN/action/replication_record"}},"created_at":"2026-07-05T01:50:47.427925+00:00","updated_at":"2026-07-05T01:50:47.427925+00:00"}