{"as_of":"2026-08-22T19:45:00Z","caps":{"database_statements":6,"inbound":100,"outbound":100},"context_digest":"sha256:687ac7e6874beb64338f7215ee07cd1d2120627ec6e9679b28349a19104bb7bd","coverage":[{"denominator":50,"lane":"reference_resolution","note":"Typed states for the displayed outbound observations.","records_observed":50,"source":"paper_references, paper_reference_links","source_observed_at":"2026-06-28T10:39:14.444486Z","state":"measured"},{"denominator":56,"lane":"standing_notices","note":"One-hop event checks from named stored sources.","records_observed":56,"source":"scholarly_work_events, retraction_status_cache","source_observed_at":"2026-08-22T06:32:14.747728+00:00","state":"measured"},{"denominator":6,"lane":"inbound_itemization","note":"Pith citing papers itemized under the disclosed page cap.","records_observed":6,"source":"paper_references, paper_reference_links","source_observed_at":"2026-08-15T21:04:30.139946Z","state":"measured"},{"denominator":1,"lane":"external_citation_measurements","note":"A source-named dated measurement, never combined with another source.","records_observed":0,"source":"pith","source_observed_at":"2026-08-12T00:36:24.352367Z","state":"measured"}],"external_citation_measurements":[],"inbound":[{"citation":{"cited_paper":{"arxiv_id":"2606.03264","last_updated":"2026-06-02T07:27:03Z","snapshot_observed_at":"2026-08-16T07:12:09.561259Z","submitted_at":"2026-06-02T07:27:03Z","title":"PaddleOCR-VL-1.6: Expanding the Frontier of Document Parsing with Under-Optimized Region Refinement and Progressive Post-Training","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2606.03264","snapshot_observed_at":"2026-07-11T11:58:14.396969Z","title":"PaddleOCR-VL-1.6: Expanding the frontier of document parsing with under-optimized region refinement and progressive post-training.arXiv preprint arXiv:2606.03264, 2026","venue":null,"work_id":null,"year":2026},"citing_paper":{"arxiv_id":"2607.04884","last_updated":"2026-08-06T10:55:17Z","snapshot_observed_at":"2026-08-19T06:07:12.300490Z","submitted_at":"2026-07-06T10:06:05Z","title":"HunyuanOCR-1.5: Making Lightweight OCR VLMs Faster and Better","version":1},"reference_index":76,"source":"pdf_text","source_observed_at":"2026-07-11T11:58:14.396969Z"},"links":{"cited_paper":"/paper/2606.03264","citing_paper":"/paper/2607.04884"},"observation_digest":"sha256:8f855f255a6d59009f6cc00aac01f7a650763687bc31a67a456a09cf76923a6b","observation_id":"4a43d4ac-9e39-4512-aae2-86a6b72b64b9","resolution":{"observed_at":"2026-07-11T11:58:14.396969Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2606.03264","last_updated":"2026-06-02T07:27:03Z","snapshot_observed_at":"2026-08-16T07:12:09.561259Z","submitted_at":"2026-06-02T07:27:03Z","title":"PaddleOCR-VL-1.6: Expanding the Frontier of Document Parsing with Under-Optimized Region Refinement and Progressive Post-Training","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2606.03264","snapshot_observed_at":"2026-07-14T04:45:32.682508Z","title":"Paddleocr-vl-1.6: Expanding the frontier of document parsing with under- optimized region refinement and progressive post-training.arXiv preprint arXiv:2606.03264, 2026","venue":null,"work_id":null,"year":2026},"citing_paper":{"arxiv_id":"2607.11562","last_updated":"2026-07-13T13:43:39Z","snapshot_observed_at":"2026-08-20T02:58:41.621792Z","submitted_at":"2026-07-13T13:43:39Z","title":"MonkeyOCRv2: A Visual-Text Foundation Model for Document AI","version":1},"reference_index":136,"source":"pdf_text","source_observed_at":"2026-07-14T04:45:32.682508Z"},"links":{"cited_paper":"/paper/2606.03264","citing_paper":"/paper/2607.11562"},"observation_digest":"sha256:80a61e18f87cda38762340fc8c421a94dbef0db55b44c870733fd3d5a0af2eeb","observation_id":"1f106058-a45e-487d-af66-fd1b4a519345","resolution":{"observed_at":"2026-07-14T04:45:32.682508Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2606.03264","last_updated":"2026-06-02T07:27:03Z","snapshot_observed_at":"2026-08-16T07:12:09.561259Z","submitted_at":"2026-06-02T07:27:03Z","title":"PaddleOCR-VL-1.6: Expanding the Frontier of Document Parsing with Under-Optimized Region Refinement and Progressive Post-Training","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2606.03264","snapshot_observed_at":"2026-08-02T04:40:01.558740Z","title":"PaddleOCR-VL-1.6: Expanding the frontier of document parsing with under-optimized region refinement and progressive post-training.arXiv preprint arXiv:2606.03264, 2026","venue":null,"work_id":null,"year":2026},"citing_paper":{"arxiv_id":"2607.13639","last_updated":"2026-07-15T09:32:33Z","snapshot_observed_at":"2026-08-20T14:12:11.072339Z","submitted_at":"2026-07-15T09:32:33Z","title":"OvisOCR2 Technical Report","version":1},"reference_index":3,"source":"pdf_text","source_observed_at":"2026-08-02T04:40:01.558740Z"},"links":{"cited_paper":"/paper/2606.03264","citing_paper":"/paper/2607.13639"},"observation_digest":"sha256:44c7525fe8cd4726cff2642806e3e754d24410f73eeeb07b5496223a8bac824c","observation_id":"561d56b7-659a-4b49-9a46-dfa2faf8daea","resolution":{"observed_at":"2026-08-02T04:40:01.558740Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2606.03264","last_updated":"2026-06-02T07:27:03Z","snapshot_observed_at":"2026-08-16T07:12:09.561259Z","submitted_at":"2026-06-02T07:27:03Z","title":"PaddleOCR-VL-1.6: Expanding the Frontier of Document Parsing with Under-Optimized Region Refinement and Progressive Post-Training","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2606.03264","snapshot_observed_at":"2026-08-01T14:16:08.214317Z","title":"Paddleocr-vl-1.6: Expanding the frontier of document parsing with under-optimized region refinement and progressive post-training","venue":null,"work_id":null,"year":2026},"citing_paper":{"arxiv_id":"2607.18839","last_updated":"2026-07-21T08:25:32Z","snapshot_observed_at":"2026-08-14T06:42:42.759580Z","submitted_at":"2026-07-21T08:25:32Z","title":"HPD-Parsing: Hierarchical Parallel Document Parsing","version":1},"reference_index":33,"source":"pdf_text","source_observed_at":"2026-08-01T14:16:08.214317Z"},"links":{"cited_paper":"/paper/2606.03264","citing_paper":"/paper/2607.18839"},"observation_digest":"sha256:02037bd7f21931f356871bcc63878c5c4502eb6466b5dfabdbea053b1d89a05c","observation_id":"2aae2e69-caf4-4eb6-bd6d-d933b7041a32","resolution":{"observed_at":"2026-08-01T14:16:08.214317Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2606.03264","last_updated":"2026-06-02T07:27:03Z","snapshot_observed_at":"2026-08-16T07:12:09.561259Z","submitted_at":"2026-06-02T07:27:03Z","title":"PaddleOCR-VL-1.6: Expanding the Frontier of Document Parsing with Under-Optimized Region Refinement and Progressive Post-Training","version":1},"cited_work":{"arxiv_id":"2606.03264","doi":null,"metadata_source":"pith","pith_arxiv_id":"2606.03264","snapshot_observed_at":"2026-08-12T00:36:24.352367Z","title":"PaddleOCR-VL-1.6: Expanding the Frontier of Document Parsing with Under-Optimized Region Refinement and Progressive Post-Training","venue":"cs.CV","work_id":"597a6f62-4058-4638-9f90-daa2e702e73d","year":2026},"citing_paper":{"arxiv_id":"2608.10396","last_updated":"2026-08-11T02:44:34Z","snapshot_observed_at":"2026-08-16T20:43:08.483594Z","submitted_at":"2026-08-11T02:44:34Z","title":"FormStruct-Bench:A Hierarchical and Diagnostic Benchmark for Table-Form Document Structure Recognition","version":1},"reference_index":44,"source":"pdf_text","source_observed_at":"2026-08-12T00:36:23.848604Z"},"links":{"cited_paper":"/paper/2606.03264","citing_paper":"/paper/2608.10396"},"observation_digest":"sha256:e39bcc8997849dc8e5f508819d8453359eabe3589e9a4372b7c6a31e796b7c64","observation_id":"a1013a5a-a999-43e5-8a1b-cd6febfd6262","resolution":{"observed_at":"2026-08-12T00:36:24.357936Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-22T06:32:14.747728+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-22T06:32:14.747728+00:00","source":"crossref"},{"observed_at":"2026-08-22T06:32:06.552537+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2606.03264","last_updated":"2026-06-02T07:27:03Z","snapshot_observed_at":"2026-08-16T07:12:09.561259Z","submitted_at":"2026-06-02T07:27:03Z","title":"PaddleOCR-VL-1.6: Expanding the Frontier of Document Parsing with Under-Optimized Region Refinement and Progressive Post-Training","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2606.03264","snapshot_observed_at":"2026-08-15T21:04:30.139946Z","title":"arXiv preprint arXiv:2606.03264 , year=","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2608.12898","last_updated":"2026-08-18T06:14:21Z","snapshot_observed_at":"2026-08-22T04:08:53.502592Z","submitted_at":"2026-08-13T07:34:21Z","title":"NaviDC-OCR: Navigating Document Parsing Across Digital and Camera-Captured Documents","version":1},"reference_index":70,"source":"arxiv_source","source_observed_at":"2026-08-15T21:04:30.139946Z"},"links":{"cited_paper":"/paper/2606.03264","citing_paper":"/paper/2608.12898"},"observation_digest":"sha256:2302c5dd8a7794d92ea09d8cf5c1c110704e53616c256c889a05dd77bbcfba10","observation_id":"c581d872-37f2-4249-bcb6-10756de77f52","resolution":{"observed_at":"2026-08-15T21:04:30.139946Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}}],"links":{"evidence":"/evidence","html":"/paper/2606.03264/citation-record","integrity":"/paper/2606.03264/integrity","json":"/paper/2606.03264/citation-record.json","paper":"/paper/2606.03264"},"outbound":[{"citation":{"cited_paper":{"arxiv_id":"2005.11401","last_updated":"2021-04-12T15:42:18Z","snapshot_observed_at":"2026-08-07T05:44:30.677502Z","submitted_at":"2020-05-22T21:34:34Z","title":"Retrieval-Augmented Generation for Knowledge-Intensive NLP Tasks","version":4},"cited_work":{"arxiv_id":"2005.11401","doi":"10.1145/3626772.3657717","metadata_source":"pith","pith_arxiv_id":"2005.11401","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Retrieval-Augmented Generation for Knowledge-Intensive NLP Tasks","venue":"cs.CL","work_id":"27eaec54-c105-4969-8188-da5f0fca3688","year":2020},"citing_paper":{"arxiv_id":"2606.03264","last_updated":"2026-06-02T07:27:03Z","snapshot_observed_at":"2026-08-16T07:12:09.561259Z","submitted_at":"2026-06-02T07:27:03Z","title":"PaddleOCR-VL-1.6: Expanding the Frontier of Document Parsing with Under-Optimized Region Refinement and Progressive Post-Training","version":1},"reference_index":1,"source":"pdf_text","source_observed_at":"2026-06-28T10:39:14.444486Z"},"links":{"cited_paper":"/paper/2005.11401","citing_paper":"/paper/2606.03264"},"observation_digest":"sha256:965ea34ee4b75a5fc9498a05cd4456a944dcba971e7c70cb877bd3677fa12378","observation_id":"73bf3b6b-75a7-4308-b79c-bfbe96296cda","resolution":{"observed_at":"2026-07-02T02:46:28.769701Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-22T06:32:14.747728+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-22T06:32:14.747728+00:00","source":"crossref"},{"observed_at":"2026-08-22T06:32:06.552537+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":"2506.05218","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-07-04T10:19:46.911382Z","title":"arXiv preprint arXiv:2506.05218 , year=","venue":null,"work_id":"e1783180-c7e4-4eda-92a0-bea7df4e8d95","year":2025},"citing_paper":{"arxiv_id":"2606.03264","last_updated":"2026-06-02T07:27:03Z","snapshot_observed_at":"2026-08-16T07:12:09.561259Z","submitted_at":"2026-06-02T07:27:03Z","title":"PaddleOCR-VL-1.6: Expanding the Frontier of Document Parsing with Under-Optimized Region Refinement and Progressive Post-Training","version":1},"reference_index":2,"source":"pdf_text","source_observed_at":"2026-06-28T10:39:14.444486Z"},"links":{"citing_paper":"/paper/2606.03264"},"observation_digest":"sha256:ad2dfa90278f52400fa9845fbb4942916d0dcbc2f03d22f1bd3c80c71904a030","observation_id":"19ecba42-c254-4b59-9012-71257576c0b8","resolution":{"observed_at":"2026-07-02T02:46:28.777538Z","resolver_source":"arxiv_id","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-22T06:32:14.747728+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-22T06:32:14.747728+00:00","source":"crossref"},{"observed_at":"2026-08-22T06:32:06.552537+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2509.22186","last_updated":"2025-09-29T16:41:28Z","snapshot_observed_at":"2026-08-15T07:34:12.195990Z","submitted_at":"2025-09-26T10:45:48Z","title":"MinerU2.5: A Decoupled Vision-Language Model for Efficient High-Resolution Document Parsing","version":2},"cited_work":{"arxiv_id":"2509.22186","doi":null,"metadata_source":"pith","pith_arxiv_id":"2509.22186","snapshot_observed_at":"2026-07-10T17:07:25.732459Z","title":"MinerU2.5: A Decoupled Vision-Language Model for Efficient High-Resolution Document Parsing","venue":"cs.CV","work_id":"6e06ed47-1ffa-419d-a9b6-145d5184c6cf","year":2025},"citing_paper":{"arxiv_id":"2606.03264","last_updated":"2026-06-02T07:27:03Z","snapshot_observed_at":"2026-08-16T07:12:09.561259Z","submitted_at":"2026-06-02T07:27:03Z","title":"PaddleOCR-VL-1.6: Expanding the Frontier of Document Parsing with Under-Optimized Region Refinement and Progressive Post-Training","version":1},"reference_index":3,"source":"pdf_text","source_observed_at":"2026-06-28T10:39:14.444486Z"},"links":{"cited_paper":"/paper/2509.22186","citing_paper":"/paper/2606.03264"},"observation_digest":"sha256:4af93b9165b5eca1528a9682df46628de5daed52c13a8df778d1a2d34a85cb26","observation_id":"a994715f-010f-4de8-9c72-e7859120a4ad","resolution":{"observed_at":"2026-07-02T02:46:28.765289Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-22T06:32:14.747728+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-22T06:32:14.747728+00:00","source":"crossref"},{"observed_at":"2026-08-22T06:32:06.552537+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-06-28T10:39:14.444486Z","title":"Dolphin: Document image parsing via heterogeneous anchor prompting, 2025","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2606.03264","last_updated":"2026-06-02T07:27:03Z","snapshot_observed_at":"2026-08-16T07:12:09.561259Z","submitted_at":"2026-06-02T07:27:03Z","title":"PaddleOCR-VL-1.6: Expanding the Frontier of Document Parsing with Under-Optimized Region Refinement and Progressive Post-Training","version":1},"reference_index":4,"source":"pdf_text","source_observed_at":"2026-06-28T10:39:14.444486Z"},"links":{"citing_paper":"/paper/2606.03264"},"observation_digest":"sha256:7e7fcfa2f220e4b6871d82265c9ae88d6495f7d762bb5651489cb6b79c375ae7","observation_id":"ff177cbd-f417-4dd9-8026-fc884180f813","resolution":{"observed_at":"2026-06-28T10:39:14.444486Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2509.01215","last_updated":"2025-09-01T07:54:18Z","snapshot_observed_at":"2026-08-14T02:52:32.415934Z","submitted_at":"2025-09-01T07:54:18Z","title":"POINTS-Reader: Distillation-Free Adaptation of Vision-Language Models for Document Conversion","version":1},"cited_work":{"arxiv_id":"2509.01215","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2509.01215","snapshot_observed_at":"2026-07-02T02:46:28.788373Z","title":"arXiv preprint arXiv:2509.01215 , year=","venue":null,"work_id":"0d5cc528-e64f-4b61-a453-8cffdb19a917","year":2025},"citing_paper":{"arxiv_id":"2606.03264","last_updated":"2026-06-02T07:27:03Z","snapshot_observed_at":"2026-08-16T07:12:09.561259Z","submitted_at":"2026-06-02T07:27:03Z","title":"PaddleOCR-VL-1.6: Expanding the Frontier of Document Parsing with Under-Optimized Region Refinement and Progressive Post-Training","version":1},"reference_index":5,"source":"pdf_text","source_observed_at":"2026-06-28T10:39:14.444486Z"},"links":{"cited_paper":"/paper/2509.01215","citing_paper":"/paper/2606.03264"},"observation_digest":"sha256:8e12a16153d7cef98f760aaddc85ef2b0805654d25231dfbca7cc33a46eb8162","observation_id":"466c860c-8a6d-4068-9dbb-0e3911a412f8","resolution":{"observed_at":"2026-07-02T02:46:28.790898Z","resolver_source":"arxiv_id","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-22T06:32:14.747728+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-22T06:32:14.747728+00:00","source":"crossref"},{"observed_at":"2026-08-22T06:32:06.552537+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":"2510.14528","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-07-10T17:07:25.801009Z","title":"Paddleocr-vl: Boosting multilingual document parsing via a 0.9 b ultra-compact vision-language model","venue":null,"work_id":"24ce5f62-ba1e-4587-a368-858a472f614c","year":2025},"citing_paper":{"arxiv_id":"2606.03264","last_updated":"2026-06-02T07:27:03Z","snapshot_observed_at":"2026-08-16T07:12:09.561259Z","submitted_at":"2026-06-02T07:27:03Z","title":"PaddleOCR-VL-1.6: Expanding the Frontier of Document Parsing with Under-Optimized Region Refinement and Progressive Post-Training","version":1},"reference_index":6,"source":"pdf_text","source_observed_at":"2026-06-28T10:39:14.444486Z"},"links":{"citing_paper":"/paper/2606.03264"},"observation_digest":"sha256:747cf1ac7554ece77e9ce186d3ca45763081ab4c36077559aeb93b53a9bac171","observation_id":"1e71d54d-a293-4ad5-b68c-bacebbc0a8e8","resolution":{"observed_at":"2026-07-02T02:46:28.786714Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-22T06:32:14.747728+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-22T06:32:14.747728+00:00","source":"crossref"},{"observed_at":"2026-08-22T06:32:06.552537+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-06-28T10:39:14.444486Z","title":"Deepseek-ocr: Contexts optical compression,","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2606.03264","last_updated":"2026-06-02T07:27:03Z","snapshot_observed_at":"2026-08-16T07:12:09.561259Z","submitted_at":"2026-06-02T07:27:03Z","title":"PaddleOCR-VL-1.6: Expanding the Frontier of Document Parsing with Under-Optimized Region Refinement and Progressive Post-Training","version":1},"reference_index":7,"source":"pdf_text","source_observed_at":"2026-06-28T10:39:14.444486Z"},"links":{"citing_paper":"/paper/2606.03264"},"observation_digest":"sha256:642464201df9257283af8956c8f218225e59903bd0ddfeded2e0505715e2f366","observation_id":"1a579d81-d066-49e3-ae4a-46647d58e6c1","resolution":{"observed_at":"2026-06-28T10:39:14.444486Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2510.18234","last_updated":"2025-10-21T02:41:44Z","snapshot_observed_at":"2026-08-17T02:37:33.939393Z","submitted_at":"2025-10-21T02:41:44Z","title":"DeepSeek-OCR: Contexts Optical Compression","version":1},"cited_work":{"arxiv_id":"2510.18234","doi":"10.48550/arxiv.2510.18234","metadata_source":"pith","pith_arxiv_id":"2510.18234","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"DeepSeek-OCR: Contexts Optical Compression","venue":"cs.CV","work_id":"e85551be-f243-4764-886f-76a8813ccbb8","year":2025},"citing_paper":{"arxiv_id":"2606.03264","last_updated":"2026-06-02T07:27:03Z","snapshot_observed_at":"2026-08-16T07:12:09.561259Z","submitted_at":"2026-06-02T07:27:03Z","title":"PaddleOCR-VL-1.6: Expanding the Frontier of Document Parsing with Under-Optimized Region Refinement and Progressive Post-Training","version":1},"reference_index":8,"source":"pdf_text","source_observed_at":"2026-06-28T10:39:14.444486Z"},"links":{"cited_paper":"/paper/2510.18234","citing_paper":"/paper/2606.03264"},"observation_digest":"sha256:1343d49668ce967334aeaccb1bbbd870fecbd83e28abeebb1c41ed9acb214228","observation_id":"40651e24-88d3-44fc-aef6-5ce03993e52a","resolution":{"observed_at":"2026-07-02T02:46:28.781764Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-22T06:32:14.747728+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-22T06:32:14.747728+00:00","source":"crossref"},{"observed_at":"2026-08-22T06:32:06.552537+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":"2511.19575","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-07-10T17:07:25.779583Z","title":"Hunyuanocr technical report","venue":null,"work_id":"2380b8a5-8713-4b4a-8922-7b1baf35db89","year":2025},"citing_paper":{"arxiv_id":"2606.03264","last_updated":"2026-06-02T07:27:03Z","snapshot_observed_at":"2026-08-16T07:12:09.561259Z","submitted_at":"2026-06-02T07:27:03Z","title":"PaddleOCR-VL-1.6: Expanding the Frontier of Document Parsing with Under-Optimized Region Refinement and Progressive Post-Training","version":1},"reference_index":9,"source":"pdf_text","source_observed_at":"2026-06-28T10:39:14.444486Z"},"links":{"citing_paper":"/paper/2606.03264"},"observation_digest":"sha256:378e5eedad4b224b2fcf6d882353c175d31bcb3689dda91a2e8eb1e91e88a615","observation_id":"7698eef1-14ac-454b-8127-6ffb55ca2c71","resolution":{"observed_at":"2026-07-02T02:46:28.903367Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-22T06:32:14.747728+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-22T06:32:14.747728+00:00","source":"crossref"},{"observed_at":"2026-08-22T06:32:06.552537+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-06-28T10:39:14.444486Z","title":"Omnidocbench: Benchmarking diverse pdf document parsing with comprehensive annotations","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2606.03264","last_updated":"2026-06-02T07:27:03Z","snapshot_observed_at":"2026-08-16T07:12:09.561259Z","submitted_at":"2026-06-02T07:27:03Z","title":"PaddleOCR-VL-1.6: Expanding the Frontier of Document Parsing with Under-Optimized Region Refinement and Progressive Post-Training","version":1},"reference_index":11,"source":"pdf_text","source_observed_at":"2026-06-28T10:39:14.444486Z"},"links":{"citing_paper":"/paper/2606.03264"},"observation_digest":"sha256:c7df3e7e0408f0aa4cd1db30ee038134a3e4e87a33350bd590ed92a01c3f117c","observation_id":"2da2cfc1-6a8f-4c08-b9f8-5f689c00cd32","resolution":{"observed_at":"2026-06-28T10:39:14.444486Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2402.03300","last_updated":"2024-04-27T15:25:53Z","snapshot_observed_at":"2026-08-06T14:58:42.911363Z","submitted_at":"2024-02-05T18:55:32Z","title":"DeepSeekMath: Pushing the Limits of Mathematical Reasoning in Open Language Models","version":3},"cited_work":{"arxiv_id":"2402.03300","doi":"10.1016/0004-3702(73)90011-8","metadata_source":"pith","pith_arxiv_id":"2402.03300","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"DeepSeekMath: Pushing the Limits of Mathematical Reasoning in Open Language Models","venue":"cs.CL","work_id":"c5006563-f3ec-438a-9e35-b7b484f34828","year":2024},"citing_paper":{"arxiv_id":"2606.03264","last_updated":"2026-06-02T07:27:03Z","snapshot_observed_at":"2026-08-16T07:12:09.561259Z","submitted_at":"2026-06-02T07:27:03Z","title":"PaddleOCR-VL-1.6: Expanding the Frontier of Document Parsing with Under-Optimized Region Refinement and Progressive Post-Training","version":1},"reference_index":12,"source":"pdf_text","source_observed_at":"2026-06-28T10:39:14.444486Z"},"links":{"cited_paper":"/paper/2402.03300","citing_paper":"/paper/2606.03264"},"observation_digest":"sha256:b067f1d7c5d9701cade93863ba983c30a9ebf3b2b2c0adcf460122d0bb8f7df9","observation_id":"5e75786d-f201-4bb3-9a8a-3887d1d998a1","resolution":{"observed_at":"2026-07-02T02:46:28.843075Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-22T06:32:14.747728+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-22T06:32:14.747728+00:00","source":"crossref"},{"observed_at":"2026-08-22T06:32:06.552537+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-06-28T10:39:14.444486Z","title":"Patch n’pack: Navit, a vision transformer for any aspect ratio and resolution","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2606.03264","last_updated":"2026-06-02T07:27:03Z","snapshot_observed_at":"2026-08-16T07:12:09.561259Z","submitted_at":"2026-06-02T07:27:03Z","title":"PaddleOCR-VL-1.6: Expanding the Frontier of Document Parsing with Under-Optimized Region Refinement and Progressive Post-Training","version":1},"reference_index":13,"source":"pdf_text","source_observed_at":"2026-06-28T10:39:14.444486Z"},"links":{"citing_paper":"/paper/2606.03264"},"observation_digest":"sha256:6106d55e8446c192e8af6e4c48fbc8ddcca44c57112cb105710288323988e2d8","observation_id":"b93f2d21-a85f-4f2e-aa40-f92209dda24e","resolution":{"observed_at":"2026-06-28T10:39:14.444486Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-06-28T10:39:14.444486Z","title":"Ernie 4.5 technical report, 2025","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2606.03264","last_updated":"2026-06-02T07:27:03Z","snapshot_observed_at":"2026-08-16T07:12:09.561259Z","submitted_at":"2026-06-02T07:27:03Z","title":"PaddleOCR-VL-1.6: Expanding the Frontier of Document Parsing with Under-Optimized Region Refinement and Progressive Post-Training","version":1},"reference_index":14,"source":"pdf_text","source_observed_at":"2026-06-28T10:39:14.444486Z"},"links":{"citing_paper":"/paper/2606.03264"},"observation_digest":"sha256:6b9da889808789ef23a44777beccc6ebe17303e2fe2bfd63595548cbe40c2334","observation_id":"5f163563-6bb6-4456-af48-d4e502f7bdfa","resolution":{"observed_at":"2026-06-28T10:39:14.444486Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":"2603.13398","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-07-04T10:19:46.900869Z","title":"Qianfan-ocr: A unified end-to-end model for document intelligence","venue":null,"work_id":"daabfed7-5204-4653-861d-2becb02cd711","year":2026},"citing_paper":{"arxiv_id":"2606.03264","last_updated":"2026-06-02T07:27:03Z","snapshot_observed_at":"2026-08-16T07:12:09.561259Z","submitted_at":"2026-06-02T07:27:03Z","title":"PaddleOCR-VL-1.6: Expanding the Frontier of Document Parsing with Under-Optimized Region Refinement and Progressive Post-Training","version":1},"reference_index":15,"source":"pdf_text","source_observed_at":"2026-06-28T10:39:14.444486Z"},"links":{"citing_paper":"/paper/2606.03264"},"observation_digest":"sha256:1e2bf85078aa86f7a5641b56a397d6edfc3aa71337c7d7f027c8cd09b1b796ca","observation_id":"681cb270-2960-444a-bd37-f7307f19b0f1","resolution":{"observed_at":"2026-07-02T02:46:28.825473Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-22T06:32:14.747728+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-22T06:32:14.747728+00:00","source":"crossref"},{"observed_at":"2026-08-22T06:32:06.552537+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":"2603.10910","doi":"10.48550/arxiv.2603.10910.url:http://arxiv","metadata_source":"arxiv_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-07-10T17:07:25.785152Z","title":"Glm-ocr technical report.arXiv preprint arXiv:2603.10910","venue":null,"work_id":"c18244fa-5009-4f7a-a303-76b8f016e7a0","year":2026},"citing_paper":{"arxiv_id":"2606.03264","last_updated":"2026-06-02T07:27:03Z","snapshot_observed_at":"2026-08-16T07:12:09.561259Z","submitted_at":"2026-06-02T07:27:03Z","title":"PaddleOCR-VL-1.6: Expanding the Frontier of Document Parsing with Under-Optimized Region Refinement and Progressive Post-Training","version":1},"reference_index":16,"source":"pdf_text","source_observed_at":"2026-06-28T10:39:14.444486Z"},"links":{"citing_paper":"/paper/2606.03264"},"observation_digest":"sha256:82ca36185caf867b02a6127c6930bc4ee76360a4fc52140b71ec6bd73ceb9582","observation_id":"7bad01b5-e418-40c7-9224-0d40c583d4eb","resolution":{"observed_at":"2026-07-02T02:46:28.839304Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-22T06:32:14.747728+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-22T06:32:14.747728+00:00","source":"crossref"},{"observed_at":"2026-08-22T06:32:06.552537+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-06-28T10:39:14.444486Z","title":"Mineru2.5-pro: Pushing the limits of data-centric document parsing at scale, 2026","venue":null,"work_id":null,"year":2026},"citing_paper":{"arxiv_id":"2606.03264","last_updated":"2026-06-02T07:27:03Z","snapshot_observed_at":"2026-08-16T07:12:09.561259Z","submitted_at":"2026-06-02T07:27:03Z","title":"PaddleOCR-VL-1.6: Expanding the Frontier of Document Parsing with Under-Optimized Region Refinement and Progressive Post-Training","version":1},"reference_index":17,"source":"pdf_text","source_observed_at":"2026-06-28T10:39:14.444486Z"},"links":{"citing_paper":"/paper/2606.03264"},"observation_digest":"sha256:9c1b0a811415926c683be21fd8bb9814398d521a237e06d04af54463b52e5eee","observation_id":"3736b9e3-0816-43a3-8b27-3f449ee1f183","resolution":{"observed_at":"2026-06-28T10:39:14.444486Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":"2602.04705","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-07-02T02:46:28.830041Z","title":"Ernie 5.0 technical report","venue":null,"work_id":"ffbea6f7-8129-4c8d-bcbf-1aa4cf16cf5f","year":2026},"citing_paper":{"arxiv_id":"2606.03264","last_updated":"2026-06-02T07:27:03Z","snapshot_observed_at":"2026-08-16T07:12:09.561259Z","submitted_at":"2026-06-02T07:27:03Z","title":"PaddleOCR-VL-1.6: Expanding the Frontier of Document Parsing with Under-Optimized Region Refinement and Progressive Post-Training","version":1},"reference_index":19,"source":"pdf_text","source_observed_at":"2026-06-28T10:39:14.444486Z"},"links":{"citing_paper":"/paper/2606.03264"},"observation_digest":"sha256:31b90e941fed4b1c6ccce6a23a4a67c9f63c4e2e02a437adf6831fd0380c4eca","observation_id":"3bec7abd-1e15-4645-85d3-346ea2124e05","resolution":{"observed_at":"2026-07-02T02:46:28.833162Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-22T06:32:14.747728+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-22T06:32:14.747728+00:00","source":"crossref"},{"observed_at":"2026-08-22T06:32:06.552537+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2601.21957","last_updated":"2026-04-03T06:49:11Z","snapshot_observed_at":"2026-08-13T13:13:25.351259Z","submitted_at":"2026-01-29T16:35:04Z","title":"PaddleOCR-VL-1.5: Towards a Multi-Task 0.9B VLM for Robust In-the-Wild Document Parsing","version":2},"cited_work":{"arxiv_id":"2601.21957","doi":null,"metadata_source":"pith","pith_arxiv_id":"2601.21957","snapshot_observed_at":"2026-07-04T15:09:54.670844Z","title":"PaddleOCR-VL-1.5: Towards a Multi-Task 0.9B VLM for Robust In-the-Wild Document Parsing","venue":"cs.CV","work_id":"aca6cad3-5eb4-48dc-92c1-44a5eb63772e","year":2026},"citing_paper":{"arxiv_id":"2606.03264","last_updated":"2026-06-02T07:27:03Z","snapshot_observed_at":"2026-08-16T07:12:09.561259Z","submitted_at":"2026-06-02T07:27:03Z","title":"PaddleOCR-VL-1.6: Expanding the Frontier of Document Parsing with Under-Optimized Region Refinement and Progressive Post-Training","version":1},"reference_index":21,"source":"pdf_text","source_observed_at":"2026-06-28T10:39:14.444486Z"},"links":{"cited_paper":"/paper/2601.21957","citing_paper":"/paper/2606.03264"},"observation_digest":"sha256:063441bfb2ed8753d0b7e8a98f1b2a07d5e8169ed47ad288e0ed40151f65f7b4","observation_id":"e4ad9efa-2365-4986-a30d-4028ddfdcc90","resolution":{"observed_at":"2026-07-02T02:46:28.810268Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-22T06:32:14.747728+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-22T06:32:14.747728+00:00","source":"crossref"},{"observed_at":"2026-08-22T06:32:06.552537+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2503.14476","last_updated":"2025-05-20T01:37:34Z","snapshot_observed_at":"2026-08-18T05:01:20.543826Z","submitted_at":"2025-03-18T17:49:06Z","title":"DAPO: An Open-Source LLM Reinforcement Learning System at Scale","version":2},"cited_work":{"arxiv_id":"2503.14476","doi":"10.48550/arxiv.2503.14476","metadata_source":"pith","pith_arxiv_id":"2503.14476","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"DAPO: An Open-Source LLM Reinforcement Learning System at Scale","venue":"cs.LG","work_id":"64019d00-0b11-4bbd-b173-b46c8fad0157","year":2025},"citing_paper":{"arxiv_id":"2606.03264","last_updated":"2026-06-02T07:27:03Z","snapshot_observed_at":"2026-08-16T07:12:09.561259Z","submitted_at":"2026-06-02T07:27:03Z","title":"PaddleOCR-VL-1.6: Expanding the Frontier of Document Parsing with Under-Optimized Region Refinement and Progressive Post-Training","version":1},"reference_index":22,"source":"pdf_text","source_observed_at":"2026-06-28T10:39:14.444486Z"},"links":{"cited_paper":"/paper/2503.14476","citing_paper":"/paper/2606.03264"},"observation_digest":"sha256:28bf54c979543a5253d92bd85cb799491a9de3b1e5775781bd25d31356b02896","observation_id":"eb50f87e-785b-4da5-bc60-931da5bf16f0","resolution":{"observed_at":"2026-07-02T02:46:28.773094Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-22T06:32:14.747728+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-22T06:32:14.747728+00:00","source":"crossref"},{"observed_at":"2026-08-08T16:08:20.547492+00:00","source":"crossref_status_cache"},{"observed_at":"2026-08-08T16:08:20.547492+00:00","source":"openalex_status_cache"},{"observed_at":"2026-08-22T06:32:06.552537+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2603.04205","last_updated":"2026-06-22T03:43:39Z","snapshot_observed_at":"2026-08-13T16:00:09.205622Z","submitted_at":"2026-03-04T15:49:06Z","title":"Real5-OmniDocBench: A Full-Scale Physical Reconstruction Benchmark for Robust Document Parsing in the Wild","version":2},"cited_work":{"arxiv_id":"2603.04205","doi":null,"metadata_source":"pith","pith_arxiv_id":"2603.04205","snapshot_observed_at":"2026-07-04T09:59:45.207846Z","title":"Real5-omnidocbench: A full-scale physical reconstruction benchmark for robust document parsing in the wild","venue":"cs.CV","work_id":"ee457333-fa96-4fab-8f5e-584954133d16","year":2026},"citing_paper":{"arxiv_id":"2606.03264","last_updated":"2026-06-02T07:27:03Z","snapshot_observed_at":"2026-08-16T07:12:09.561259Z","submitted_at":"2026-06-02T07:27:03Z","title":"PaddleOCR-VL-1.6: Expanding the Frontier of Document Parsing with Under-Optimized Region Refinement and Progressive Post-Training","version":1},"reference_index":23,"source":"pdf_text","source_observed_at":"2026-06-28T10:39:14.444486Z"},"links":{"cited_paper":"/paper/2603.04205","citing_paper":"/paper/2606.03264"},"observation_digest":"sha256:4e4b6fc2a0149cb3480fe00f54f0975fc170206023dfeeda80d72856c44f3b0a","observation_id":"a507985a-d69d-4d8a-9c67-944d7cadf48c","resolution":{"observed_at":"2026-07-02T02:46:28.897241Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-22T06:32:14.747728+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-22T06:32:14.747728+00:00","source":"crossref"},{"observed_at":"2026-08-22T06:32:06.552537+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-06-28T10:39:14.444486Z","title":"Image over text: Transforming formula recognition evaluation with character detection matching","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2606.03264","last_updated":"2026-06-02T07:27:03Z","snapshot_observed_at":"2026-08-16T07:12:09.561259Z","submitted_at":"2026-06-02T07:27:03Z","title":"PaddleOCR-VL-1.6: Expanding the Frontier of Document Parsing with Under-Optimized Region Refinement and Progressive Post-Training","version":1},"reference_index":24,"source":"pdf_text","source_observed_at":"2026-06-28T10:39:14.444486Z"},"links":{"citing_paper":"/paper/2606.03264"},"observation_digest":"sha256:8be132ea8485e01a973968f9475e0b2b89f04fcdb9a3868e53fca1c2fcd6c758","observation_id":"f2bcff16-567c-499b-95e0-4531412d82cf","resolution":{"observed_at":"2026-06-28T10:39:14.444486Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2508.18265","last_updated":"2025-08-27T14:39:45Z","snapshot_observed_at":"2026-08-17T12:32:16.575866Z","submitted_at":"2025-08-25T17:58:17Z","title":"InternVL3.5: Advancing Open-Source Multimodal Models in Versatility, Reasoning, and Efficiency","version":2},"cited_work":{"arxiv_id":"2508.18265","doi":"10.48550/arxiv.2508.18265","metadata_source":"pith","pith_arxiv_id":"2508.18265","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"InternVL3.5: Advancing Open-Source Multimodal Models in Versatility, Reasoning, and Efficiency","venue":"cs.CV","work_id":"b8f5e260-fff5-444e-bcf5-2c42cfefd83d","year":2025},"citing_paper":{"arxiv_id":"2606.03264","last_updated":"2026-06-02T07:27:03Z","snapshot_observed_at":"2026-08-16T07:12:09.561259Z","submitted_at":"2026-06-02T07:27:03Z","title":"PaddleOCR-VL-1.6: Expanding the Frontier of Document Parsing with Under-Optimized Region Refinement and Progressive Post-Training","version":1},"reference_index":26,"source":"pdf_text","source_observed_at":"2026-06-28T10:39:14.444486Z"},"links":{"cited_paper":"/paper/2508.18265","citing_paper":"/paper/2606.03264"},"observation_digest":"sha256:c8ada219102b9b2c700206e41d8f916c4a9d8e6c154214fc90a5a46e2b04d5c9","observation_id":"0558e9eb-d064-4bd5-ac67-d50b7626452b","resolution":{"observed_at":"2026-07-02T02:46:28.869791Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-22T06:32:14.747728+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-22T06:32:14.747728+00:00","source":"crossref"},{"observed_at":"2026-08-10T17:38:12.771444+00:00","source":"crossref_status_cache"},{"observed_at":"2026-08-10T17:38:12.771444+00:00","source":"openalex_status_cache"},{"observed_at":"2026-08-22T06:32:06.552537+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-06-28T10:39:14.444486Z","title":null,"venue":null,"work_id":null,"year":2026},"citing_paper":{"arxiv_id":"2606.03264","last_updated":"2026-06-02T07:27:03Z","snapshot_observed_at":"2026-08-16T07:12:09.561259Z","submitted_at":"2026-06-02T07:27:03Z","title":"PaddleOCR-VL-1.6: Expanding the Frontier of Document Parsing with Under-Optimized Region Refinement and Progressive Post-Training","version":1},"reference_index":27,"source":"pdf_text","source_observed_at":"2026-06-28T10:39:14.444486Z"},"links":{"citing_paper":"/paper/2606.03264"},"observation_digest":"sha256:3fd9d4275c4a7005b2215571549befb739173cc3d3e6986f89d47501adddaa3f","observation_id":"c8dfde34-8911-4ffd-83a4-8d2e77bcb7fd","resolution":{"observed_at":"2026-06-28T10:39:14.444486Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-06-28T10:39:14.444486Z","title":"Gpt-5.2 system card, 2025","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2606.03264","last_updated":"2026-06-02T07:27:03Z","snapshot_observed_at":"2026-08-16T07:12:09.561259Z","submitted_at":"2026-06-02T07:27:03Z","title":"PaddleOCR-VL-1.6: Expanding the Frontier of Document Parsing with Under-Optimized Region Refinement and Progressive Post-Training","version":1},"reference_index":28,"source":"pdf_text","source_observed_at":"2026-06-28T10:39:14.444486Z"},"links":{"citing_paper":"/paper/2606.03264"},"observation_digest":"sha256:0af631bce9ff94b671d95d9bccdd93072be74ea5027afa1efaa5eceb9a77b0e0","observation_id":"c0188602-087e-4709-935e-6501895d6a82","resolution":{"observed_at":"2026-06-28T10:39:14.444486Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2505.09388","last_updated":"2025-05-14T13:41:34Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2025-05-14T13:41:34Z","title":"Qwen3 Technical Report","version":1},"cited_work":{"arxiv_id":"2505.09388","doi":"10.1016/j.aiopen.2022.12","metadata_source":"pith","pith_arxiv_id":"2505.09388","snapshot_observed_at":"2026-07-11T11:50:26.030339Z","title":"Qwen3 Technical Report","venue":"cs.CL","work_id":"25a4e30c-1232-48e7-9925-02fa12ba7c9e","year":2025},"citing_paper":{"arxiv_id":"2606.03264","last_updated":"2026-06-02T07:27:03Z","snapshot_observed_at":"2026-08-16T07:12:09.561259Z","submitted_at":"2026-06-02T07:27:03Z","title":"PaddleOCR-VL-1.6: Expanding the Frontier of Document Parsing with Under-Optimized Region Refinement and Progressive Post-Training","version":1},"reference_index":29,"source":"pdf_text","source_observed_at":"2026-06-28T10:39:14.444486Z"},"links":{"cited_paper":"/paper/2505.09388","citing_paper":"/paper/2606.03264"},"observation_digest":"sha256:554c66707dc8f8cc255d27f7bd064f0619eeeb9f7c36fe45fae7861e7c098b34","observation_id":"5092a84c-ebb0-4c24-8d2c-2dfccd763ed3","resolution":{"observed_at":"2026-07-02T02:46:28.899870Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-22T06:32:14.747728+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-22T06:32:14.747728+00:00","source":"crossref"},{"observed_at":"2026-08-22T06:32:06.552537+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-06-28T10:39:14.444486Z","title":"Gemini 3.0","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2606.03264","last_updated":"2026-06-02T07:27:03Z","snapshot_observed_at":"2026-08-16T07:12:09.561259Z","submitted_at":"2026-06-02T07:27:03Z","title":"PaddleOCR-VL-1.6: Expanding the Frontier of Document Parsing with Under-Optimized Region Refinement and Progressive Post-Training","version":1},"reference_index":30,"source":"pdf_text","source_observed_at":"2026-06-28T10:39:14.444486Z"},"links":{"citing_paper":"/paper/2606.03264"},"observation_digest":"sha256:0ab376ae35bf0d1f39941ebc3a08b80d31e383c1b8173c04cc23519e5c864005","observation_id":"5a1c6968-1e49-4a60-be78-21748935cda9","resolution":{"observed_at":"2026-06-28T10:39:14.444486Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2405.20797","last_updated":"2024-06-17T17:51:50Z","snapshot_observed_at":"2026-08-16T13:47:26.506119Z","submitted_at":"2024-05-31T13:59:18Z","title":"Ovis: Structural Embedding Alignment for Multimodal Large Language Model","version":2},"cited_work":{"arxiv_id":"2405.20797","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2405.20797","snapshot_observed_at":"2026-07-04T19:50:10.268440Z","title":"Ovis: Structural embedding alignment for multimodal large language model","venue":null,"work_id":"6f4025f6-a13c-4549-b0dc-8f2bb5a5a954","year":2024},"citing_paper":{"arxiv_id":"2606.03264","last_updated":"2026-06-02T07:27:03Z","snapshot_observed_at":"2026-08-16T07:12:09.561259Z","submitted_at":"2026-06-02T07:27:03Z","title":"PaddleOCR-VL-1.6: Expanding the Frontier of Document Parsing with Under-Optimized Region Refinement and Progressive Post-Training","version":1},"reference_index":31,"source":"pdf_text","source_observed_at":"2026-06-28T10:39:14.444486Z"},"links":{"cited_paper":"/paper/2405.20797","citing_paper":"/paper/2606.03264"},"observation_digest":"sha256:bd6ca8d4c055c5128c6f101e943fc00c17ae34291802ced4e335f7d371d924e3","observation_id":"faa0fa58-f5bf-4400-a2d3-40f4c955cfc0","resolution":{"observed_at":"2026-07-02T02:46:28.874283Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-22T06:32:14.747728+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-22T06:32:14.747728+00:00","source":"crossref"},{"observed_at":"2026-08-22T06:32:06.552537+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2508.11737","last_updated":"2025-08-15T17:01:08Z","snapshot_observed_at":"2026-07-06T22:13:37.150049Z","submitted_at":"2025-08-15T17:01:08Z","title":"Ovis2.5 Technical Report","version":1},"cited_work":{"arxiv_id":"2508.11737","doi":null,"metadata_source":"pith","pith_arxiv_id":"2508.11737","snapshot_observed_at":"2026-07-02T22:27:25.902699Z","title":"Ovis2.5 Technical Report","venue":"cs.CV","work_id":"83cd504b-a082-4df2-8bff-919343fc7d99","year":2025},"citing_paper":{"arxiv_id":"2606.03264","last_updated":"2026-06-02T07:27:03Z","snapshot_observed_at":"2026-08-16T07:12:09.561259Z","submitted_at":"2026-06-02T07:27:03Z","title":"PaddleOCR-VL-1.6: Expanding the Frontier of Document Parsing with Under-Optimized Region Refinement and Progressive Post-Training","version":1},"reference_index":32,"source":"pdf_text","source_observed_at":"2026-06-28T10:39:14.444486Z"},"links":{"cited_paper":"/paper/2508.11737","citing_paper":"/paper/2606.03264"},"observation_digest":"sha256:e7ba58e586271ee6fc9984e293f8879fcbb5b95cd56345d7757a8ae2d5905a13","observation_id":"e2316ebd-1ddf-4df2-bf93-27dd5c0aaf45","resolution":{"observed_at":"2026-07-02T02:46:28.881273Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-22T06:32:14.747728+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-22T06:32:14.747728+00:00","source":"crossref"},{"observed_at":"2026-08-22T06:32:06.552537+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-06-28T10:39:14.444486Z","title":"Nanonets-ocr-s: A model for transforming documents into structured markdown with intelligent content recognition and semantic tagging, 2025","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2606.03264","last_updated":"2026-06-02T07:27:03Z","snapshot_observed_at":"2026-08-16T07:12:09.561259Z","submitted_at":"2026-06-02T07:27:03Z","title":"PaddleOCR-VL-1.6: Expanding the Frontier of Document Parsing with Under-Optimized Region Refinement and Progressive Post-Training","version":1},"reference_index":33,"source":"pdf_text","source_observed_at":"2026-06-28T10:39:14.444486Z"},"links":{"citing_paper":"/paper/2606.03264"},"observation_digest":"sha256:31e1f04998ff43ad94eb70fbafe36ae4632d3d5625b862f9b8f34577669c7e87","observation_id":"0d4537c8-1d4b-4d81-ad08-252f3e008ea5","resolution":{"observed_at":"2026-06-28T10:39:14.444486Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-06-28T10:39:14.444486Z","title":"Mistral-ocr","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2606.03264","last_updated":"2026-06-02T07:27:03Z","snapshot_observed_at":"2026-08-16T07:12:09.561259Z","submitted_at":"2026-06-02T07:27:03Z","title":"PaddleOCR-VL-1.6: Expanding the Frontier of Document Parsing with Under-Optimized Region Refinement and Progressive Post-Training","version":1},"reference_index":34,"source":"pdf_text","source_observed_at":"2026-06-28T10:39:14.444486Z"},"links":{"citing_paper":"/paper/2606.03264"},"observation_digest":"sha256:f752d9d5c45d858d93e10682c4fd27c29d6c216910d4b4d380d0b54ee4ea8faf","observation_id":"532b99eb-2f26-4dd2-8c6e-691520b1efda","resolution":{"observed_at":"2026-06-28T10:39:14.444486Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":"2502.18443","doi":"10.48550/arxiv.2502.18443","metadata_source":"arxiv_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"olmocr: Unlocking trillions of tokens in pdfs with vi- sion language models.arXiv preprint arXiv:2502.18443, 2025a","venue":"ArXiv.org","work_id":"bb9c92f3-5210-40cf-b3d5-7463913e3248","year":2025},"citing_paper":{"arxiv_id":"2606.03264","last_updated":"2026-06-02T07:27:03Z","snapshot_observed_at":"2026-08-16T07:12:09.561259Z","submitted_at":"2026-06-02T07:27:03Z","title":"PaddleOCR-VL-1.6: Expanding the Frontier of Document Parsing with Under-Optimized Region Refinement and Progressive Post-Training","version":1},"reference_index":35,"source":"pdf_text","source_observed_at":"2026-06-28T10:39:14.444486Z"},"links":{"citing_paper":"/paper/2606.03264"},"observation_digest":"sha256:6ab0241a9303d02122ac11bce475b1818415a83c0a6ecebc3e72586a935fc923","observation_id":"6c6a5412-3ab6-46db-b40e-98220d905705","resolution":{"observed_at":"2026-07-02T02:46:28.859589Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-22T06:32:14.747728+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-22T06:32:14.747728+00:00","source":"crossref"},{"observed_at":"2026-08-22T06:32:06.552537+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":"2601.21639","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-07-02T02:46:28.850593Z","title":"Ocrverse: Towards holistic ocr in end-to-end vision-language models","venue":null,"work_id":"882cec55-a48e-435e-aaec-5dd371cee961","year":2026},"citing_paper":{"arxiv_id":"2606.03264","last_updated":"2026-06-02T07:27:03Z","snapshot_observed_at":"2026-08-16T07:12:09.561259Z","submitted_at":"2026-06-02T07:27:03Z","title":"PaddleOCR-VL-1.6: Expanding the Frontier of Document Parsing with Under-Optimized Region Refinement and Progressive Post-Training","version":1},"reference_index":36,"source":"pdf_text","source_observed_at":"2026-06-28T10:39:14.444486Z"},"links":{"citing_paper":"/paper/2606.03264"},"observation_digest":"sha256:a513dcc230218052237293d95cf3ccbf4f9721897628f78c9d9b855f40741246","observation_id":"41377765-717b-41e3-8b2a-f3ded61d142d","resolution":{"observed_at":"2026-07-02T02:46:28.853871Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-22T06:32:14.747728+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-22T06:32:14.747728+00:00","source":"crossref"},{"observed_at":"2026-08-22T06:32:06.552537+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2512.21095","last_updated":"2026-07-11T09:04:48Z","snapshot_observed_at":"2026-08-19T23:40:02.808027Z","submitted_at":"2025-12-24T10:35:21Z","title":"UniRec-0.1B: Unified Text and Formula Recognition with 0.1B Parameters","version":2},"cited_work":{"arxiv_id":"2512.21095","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2512.21095","snapshot_observed_at":"2026-07-14T02:19:24.298662Z","title":"Unirec-0.1 b: Unified text and formula recognition with 0.1 b parameters","venue":null,"work_id":"e78d8a84-1d7c-4dc7-b8cc-2fa3a68d5d87","year":2025},"citing_paper":{"arxiv_id":"2606.03264","last_updated":"2026-06-02T07:27:03Z","snapshot_observed_at":"2026-08-16T07:12:09.561259Z","submitted_at":"2026-06-02T07:27:03Z","title":"PaddleOCR-VL-1.6: Expanding the Frontier of Document Parsing with Under-Optimized Region Refinement and Progressive Post-Training","version":1},"reference_index":37,"source":"pdf_text","source_observed_at":"2026-06-28T10:39:14.444486Z"},"links":{"cited_paper":"/paper/2512.21095","citing_paper":"/paper/2606.03264"},"observation_digest":"sha256:3941ae873f998d1df0b8dde9d93588e57107b648d0d164cb21216050bc4680d3","observation_id":"bb9c2222-ea85-4509-970d-e3e5151a9315","resolution":{"observed_at":"2026-07-14T02:19:24.298662Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-22T06:32:14.747728+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-22T06:32:14.747728+00:00","source":"crossref"},{"observed_at":"2026-08-22T06:32:06.552537+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":"2512.02498","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-07-04T09:59:45.152724Z","title":"InICASSP 2023-2023 IEEE International Conference on Acoustics, Speech and Signal Processing (ICASSP), pages 1–5","venue":null,"work_id":"cb5539a0-6843-4702-bbf3-dfe7b61e5cbb","year":2025},"citing_paper":{"arxiv_id":"2606.03264","last_updated":"2026-06-02T07:27:03Z","snapshot_observed_at":"2026-08-16T07:12:09.561259Z","submitted_at":"2026-06-02T07:27:03Z","title":"PaddleOCR-VL-1.6: Expanding the Frontier of Document Parsing with Under-Optimized Region Refinement and Progressive Post-Training","version":1},"reference_index":38,"source":"pdf_text","source_observed_at":"2026-06-28T10:39:14.444486Z"},"links":{"citing_paper":"/paper/2606.03264"},"observation_digest":"sha256:eaaaac0bb815d769a16a09f5ca26547df84c80ec124357ca14c3c9b228cf047d","observation_id":"24ea4f0a-e4bf-4cf2-9d7c-0018ba3f7c68","resolution":{"observed_at":"2026-07-02T02:46:28.885276Z","resolver_source":"arxiv_id","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-22T06:32:14.747728+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-22T06:32:14.747728+00:00","source":"crossref"},{"observed_at":"2026-08-22T06:32:06.552537+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":"2603.01840","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-07-04T10:19:46.941324Z","title":"Firered-ocr technical report","venue":null,"work_id":"36c46655-2f26-436d-a19c-295616f0ae96","year":2026},"citing_paper":{"arxiv_id":"2606.03264","last_updated":"2026-06-02T07:27:03Z","snapshot_observed_at":"2026-08-16T07:12:09.561259Z","submitted_at":"2026-06-02T07:27:03Z","title":"PaddleOCR-VL-1.6: Expanding the Frontier of Document Parsing with Under-Optimized Region Refinement and Progressive Post-Training","version":1},"reference_index":39,"source":"pdf_text","source_observed_at":"2026-06-28T10:39:14.444486Z"},"links":{"citing_paper":"/paper/2606.03264"},"observation_digest":"sha256:73bbf70f1bcfd7d8ea93f3e2decd9cd315d338a29646f5979aefd10a0cc2abb7","observation_id":"508ed194-d504-4a93-9261-0cb676365caf","resolution":{"observed_at":"2026-07-02T02:46:28.823096Z","resolver_source":"arxiv_id","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-22T06:32:14.747728+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-22T06:32:14.747728+00:00","source":"crossref"},{"observed_at":"2026-08-22T06:32:06.552537+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2603.09677","last_updated":"2026-04-08T07:57:35Z","snapshot_observed_at":"2026-08-14T19:39:31.611087Z","submitted_at":"2026-03-10T13:46:32Z","title":"Logics-Parsing-Omni Technical Report","version":3},"cited_work":{"arxiv_id":"2603.09677","doi":null,"metadata_source":"pith","pith_arxiv_id":"2603.09677","snapshot_observed_at":"2026-07-02T02:46:28.830346Z","title":"Logics-Parsing-Omni Technical Report","venue":"cs.AI","work_id":"d0941660-d2e3-4e17-8736-f411349a80b7","year":2026},"citing_paper":{"arxiv_id":"2606.03264","last_updated":"2026-06-02T07:27:03Z","snapshot_observed_at":"2026-08-16T07:12:09.561259Z","submitted_at":"2026-06-02T07:27:03Z","title":"PaddleOCR-VL-1.6: Expanding the Frontier of Document Parsing with Under-Optimized Region Refinement and Progressive Post-Training","version":1},"reference_index":40,"source":"pdf_text","source_observed_at":"2026-06-28T10:39:14.444486Z"},"links":{"cited_paper":"/paper/2603.09677","citing_paper":"/paper/2606.03264"},"observation_digest":"sha256:e8664886198234df186e878c01a4b32c1684a26e313adc4f42693ed27c04518d","observation_id":"66b8587e-86d8-46b6-b681-57367471cdee","resolution":{"observed_at":"2026-07-02T02:46:28.832260Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-22T06:32:14.747728+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-22T06:32:14.747728+00:00","source":"crossref"},{"observed_at":"2026-08-22T06:32:06.552537+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2601.20430","last_updated":"2026-07-13T14:16:50Z","snapshot_observed_at":"2026-08-21T20:22:10.786691Z","submitted_at":"2026-01-28T09:37:13Z","title":"Youtu-Parsing: Perception, Structuring and Recognition via High-Parallelism Decoding","version":2},"cited_work":{"arxiv_id":"2601.20430","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2601.20430","snapshot_observed_at":"2026-07-14T03:21:27.323268Z","title":"Youtu-parsing: Perception, structuring and recognition via high-parallelism decoding","venue":null,"work_id":"a02d784c-999f-4b21-ab04-640335260827","year":2026},"citing_paper":{"arxiv_id":"2606.03264","last_updated":"2026-06-02T07:27:03Z","snapshot_observed_at":"2026-08-16T07:12:09.561259Z","submitted_at":"2026-06-02T07:27:03Z","title":"PaddleOCR-VL-1.6: Expanding the Frontier of Document Parsing with Under-Optimized Region Refinement and Progressive Post-Training","version":1},"reference_index":41,"source":"pdf_text","source_observed_at":"2026-06-28T10:39:14.444486Z"},"links":{"cited_paper":"/paper/2601.20430","citing_paper":"/paper/2606.03264"},"observation_digest":"sha256:3eb6dc46dd7d451fe97ead3101a98dc4c68ab4b94c36b690a08d727f73e5cf43","observation_id":"4de339ab-37c2-4da1-a298-ef31081c387f","resolution":{"observed_at":"2026-07-14T03:21:27.323268Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-22T06:32:14.747728+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-22T06:32:14.747728+00:00","source":"crossref"},{"observed_at":"2026-08-22T06:32:06.552537+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-06-28T10:39:14.444486Z","title":"Omnidocbench 1.6.https://opendatalab.com/omnidocbench, 2026","venue":null,"work_id":null,"year":2026},"citing_paper":{"arxiv_id":"2606.03264","last_updated":"2026-06-02T07:27:03Z","snapshot_observed_at":"2026-08-16T07:12:09.561259Z","submitted_at":"2026-06-02T07:27:03Z","title":"PaddleOCR-VL-1.6: Expanding the Frontier of Document Parsing with Under-Optimized Region Refinement and Progressive Post-Training","version":1},"reference_index":42,"source":"pdf_text","source_observed_at":"2026-06-28T10:39:14.444486Z"},"links":{"citing_paper":"/paper/2606.03264"},"observation_digest":"sha256:71851c8cb4193e0d73a9d093c4e815a01704e30b59f21481d3290024a04144cd","observation_id":"4b739199-53e7-45ec-bfa7-328927f67877","resolution":{"observed_at":"2026-06-28T10:39:14.444486Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-06-28T10:39:14.444486Z","title":null,"venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2606.03264","last_updated":"2026-06-02T07:27:03Z","snapshot_observed_at":"2026-08-16T07:12:09.561259Z","submitted_at":"2026-06-02T07:27:03Z","title":"PaddleOCR-VL-1.6: Expanding the Frontier of Document Parsing with Under-Optimized Region Refinement and Progressive Post-Training","version":1},"reference_index":43,"source":"pdf_text","source_observed_at":"2026-06-28T10:39:14.444486Z"},"links":{"citing_paper":"/paper/2606.03264"},"observation_digest":"sha256:2ef38fe8754bcc7a2e7b0d878080c0a1c8f49d6e4c40004eef7792814e49f6f5","observation_id":"b854bf5d-6cb5-45ae-b154-110c17aec728","resolution":{"observed_at":"2026-06-28T10:39:14.444486Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2507.05595","last_updated":"2025-07-08T02:14:10Z","snapshot_observed_at":"2026-08-14T05:24:51.715708Z","submitted_at":"2025-07-08T02:14:10Z","title":"PaddleOCR 3.0 Technical Report","version":1},"cited_work":{"arxiv_id":"2507.05595","doi":"10.48550/arxiv.2507.05595.url:http://arxiv.org/","metadata_source":"pith","pith_arxiv_id":"2507.05595","snapshot_observed_at":"2026-07-11T02:17:46.362300Z","title":"PaddleOCR 3.0 Technical Report","venue":"cs.CV","work_id":"444c549c-1895-4cae-a2d3-c13d7ed5f462","year":2025},"citing_paper":{"arxiv_id":"2606.03264","last_updated":"2026-06-02T07:27:03Z","snapshot_observed_at":"2026-08-16T07:12:09.561259Z","submitted_at":"2026-06-02T07:27:03Z","title":"PaddleOCR-VL-1.6: Expanding the Frontier of Document Parsing with Under-Optimized Region Refinement and Progressive Post-Training","version":1},"reference_index":44,"source":"pdf_text","source_observed_at":"2026-06-28T10:39:14.444486Z"},"links":{"cited_paper":"/paper/2507.05595","citing_paper":"/paper/2606.03264"},"observation_digest":"sha256:844a7994d50d702bfa08d671f413a8bcc2d3577c488a792caf9ac4a068ef8ba6","observation_id":"a82c7743-a770-45f4-b63f-495fee996ed0","resolution":{"observed_at":"2026-07-02T02:46:28.878007Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-22T06:32:14.747728+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-22T06:32:14.747728+00:00","source":"crossref"},{"observed_at":"2026-08-22T06:32:06.552537+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-06-28T10:39:14.444486Z","title":"Gemini 2.5","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2606.03264","last_updated":"2026-06-02T07:27:03Z","snapshot_observed_at":"2026-08-16T07:12:09.561259Z","submitted_at":"2026-06-02T07:27:03Z","title":"PaddleOCR-VL-1.6: Expanding the Frontier of Document Parsing with Under-Optimized Region Refinement and Progressive Post-Training","version":1},"reference_index":46,"source":"pdf_text","source_observed_at":"2026-06-28T10:39:14.444486Z"},"links":{"citing_paper":"/paper/2606.03264"},"observation_digest":"sha256:ab9d2f5af7574d487fd6077485293095c682e0d3ff7dec4b6bc53afe8f8490f1","observation_id":"681fae5d-7e4f-4012-82c9-d3ba883377bc","resolution":{"observed_at":"2026-06-28T10:39:14.444486Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-06-28T10:39:14.444486Z","title":"Mineru2.0-2505-0.9b","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2606.03264","last_updated":"2026-06-02T07:27:03Z","snapshot_observed_at":"2026-08-16T07:12:09.561259Z","submitted_at":"2026-06-02T07:27:03Z","title":"PaddleOCR-VL-1.6: Expanding the Frontier of Document Parsing with Under-Optimized Region Refinement and Progressive Post-Training","version":1},"reference_index":47,"source":"pdf_text","source_observed_at":"2026-06-28T10:39:14.444486Z"},"links":{"citing_paper":"/paper/2606.03264"},"observation_digest":"sha256:0b4ac12638b3e73b99670e1999dbb7740b6a4522bb175e5ce21deed38c7144ae","observation_id":"321fecb8-65cf-4ad9-92d7-8eeac3df6281","resolution":{"observed_at":"2026-06-28T10:39:14.444486Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-06-28T10:39:14.444486Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2606.03264","last_updated":"2026-06-02T07:27:03Z","snapshot_observed_at":"2026-08-16T07:12:09.561259Z","submitted_at":"2026-06-02T07:27:03Z","title":"PaddleOCR-VL-1.6: Expanding the Frontier of Document Parsing with Under-Optimized Region Refinement and Progressive Post-Training","version":1},"reference_index":48,"source":"pdf_text","source_observed_at":"2026-06-28T10:39:14.444486Z"},"links":{"citing_paper":"/paper/2606.03264"},"observation_digest":"sha256:df92549e721ecc243561aad86b6848cae0698ccad97867a53c0effccceac56bb","observation_id":"07290425-88e9-44db-8c73-18238d28edd3","resolution":{"observed_at":"2026-06-28T10:39:14.444486Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":"2512.01248","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-07-02T02:46:28.891766Z","title":"Trivia: Self-supervised fine-tuning of vision-language models for table recognition, 2026","venue":null,"work_id":"5237c5ee-fe30-4fff-9035-da12841e8d90","year":2026},"citing_paper":{"arxiv_id":"2606.03264","last_updated":"2026-06-02T07:27:03Z","snapshot_observed_at":"2026-08-16T07:12:09.561259Z","submitted_at":"2026-06-02T07:27:03Z","title":"PaddleOCR-VL-1.6: Expanding the Frontier of Document Parsing with Under-Optimized Region Refinement and Progressive Post-Training","version":1},"reference_index":49,"source":"pdf_text","source_observed_at":"2026-06-28T10:39:14.444486Z"},"links":{"citing_paper":"/paper/2606.03264"},"observation_digest":"sha256:36e2e0cecdb04eaaaf5512608de5c2409cecac59d9748364ac98e21a84184056","observation_id":"a0ce134f-c3c4-4f5b-ac6a-1c69e3c9c459","resolution":{"observed_at":"2026-07-02T02:46:28.893949Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-22T06:32:14.747728+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-22T06:32:14.747728+00:00","source":"crossref"},{"observed_at":"2026-08-22T06:32:06.552537+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2212.10505","last_updated":"2023-05-23T18:28:39Z","snapshot_observed_at":"2026-08-16T16:07:12.861373Z","submitted_at":"2022-12-20T18:20:50Z","title":"DePlot: One-shot visual language reasoning by plot-to-table translation","version":2},"cited_work":{"arxiv_id":"2212.10505","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2212.10505","snapshot_observed_at":"2026-07-02T02:46:28.845114Z","title":"In: Findings of the 61st Annual Meeting of the Association for Computational Linguistics (2023), https://arxiv.org/abs/ 2212.10505 10","venue":null,"work_id":"5d81d551-cced-4a84-9fda-e5e7e2b9b05d","year":2023},"citing_paper":{"arxiv_id":"2606.03264","last_updated":"2026-06-02T07:27:03Z","snapshot_observed_at":"2026-08-16T07:12:09.561259Z","submitted_at":"2026-06-02T07:27:03Z","title":"PaddleOCR-VL-1.6: Expanding the Frontier of Document Parsing with Under-Optimized Region Refinement and Progressive Post-Training","version":1},"reference_index":50,"source":"pdf_text","source_observed_at":"2026-06-28T10:39:14.444486Z"},"links":{"cited_paper":"/paper/2212.10505","citing_paper":"/paper/2606.03264"},"observation_digest":"sha256:15340ed6047c424bf766d761f4be9f018f3da687633017fa41d3b3f64717ba40","observation_id":"7528ebbd-ebfd-4720-bcdf-d673ca945fe0","resolution":{"observed_at":"2026-07-02T02:46:28.848089Z","resolver_source":"arxiv_id","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-22T06:32:14.747728+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-22T06:32:14.747728+00:00","source":"crossref"},{"observed_at":"2026-08-22T06:32:06.552537+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2404.16635","last_updated":"2024-04-25T14:23:24Z","snapshot_observed_at":"2026-08-20T14:21:20.056539Z","submitted_at":"2024-04-25T14:23:24Z","title":"TinyChart: Efficient Chart Understanding with Visual Token Merging and Program-of-Thoughts Learning","version":1},"cited_work":{"arxiv_id":"2404.16635","doi":null,"metadata_source":"pith","pith_arxiv_id":"2404.16635","snapshot_observed_at":"2026-07-10T17:07:25.763363Z","title":"Tinychart: Efficient chart understanding with visual token merging and program- of-thoughts learning.arXiv preprint arXiv:2404.16635","venue":"cs.CV","work_id":"2c7f1535-0987-4081-bbaa-6eb96047b3e8","year":2024},"citing_paper":{"arxiv_id":"2606.03264","last_updated":"2026-06-02T07:27:03Z","snapshot_observed_at":"2026-08-16T07:12:09.561259Z","submitted_at":"2026-06-02T07:27:03Z","title":"PaddleOCR-VL-1.6: Expanding the Frontier of Document Parsing with Under-Optimized Region Refinement and Progressive Post-Training","version":1},"reference_index":51,"source":"pdf_text","source_observed_at":"2026-06-28T10:39:14.444486Z"},"links":{"cited_paper":"/paper/2404.16635","citing_paper":"/paper/2606.03264"},"observation_digest":"sha256:8895bbdfe651141356ab24948bc8cf3a89242dc368deb37e81bbf0486fce6321","observation_id":"08fe5135-9cdd-4322-ae39-2c808abd5889","resolution":{"observed_at":"2026-07-02T02:46:28.818593Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-22T06:32:14.747728+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-22T06:32:14.747728+00:00","source":"crossref"},{"observed_at":"2026-08-22T06:32:06.552537+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2409.01704","last_updated":"2024-09-03T08:41:31Z","snapshot_observed_at":"2026-08-18T18:59:18.095701Z","submitted_at":"2024-09-03T08:41:31Z","title":"General OCR Theory: Towards OCR-2.0 via a Unified End-to-end Model","version":1},"cited_work":{"arxiv_id":"2409.01704","doi":null,"metadata_source":"pith","pith_arxiv_id":"2409.01704","snapshot_observed_at":"2026-07-08T20:35:34.407064Z","title":"General OCR Theory: Towards OCR-2.0 via a Unified End-to-end Model","venue":"cs.CV","work_id":"67b1592f-02b6-4056-b3fe-cc594e484192","year":2024},"citing_paper":{"arxiv_id":"2606.03264","last_updated":"2026-06-02T07:27:03Z","snapshot_observed_at":"2026-08-16T07:12:09.561259Z","submitted_at":"2026-06-02T07:27:03Z","title":"PaddleOCR-VL-1.6: Expanding the Frontier of Document Parsing with Under-Optimized Region Refinement and Progressive Post-Training","version":1},"reference_index":52,"source":"pdf_text","source_observed_at":"2026-06-28T10:39:14.444486Z"},"links":{"cited_paper":"/paper/2409.01704","citing_paper":"/paper/2606.03264"},"observation_digest":"sha256:94046eefd7cc8b00eb1f0806f16661d7637bfcb2d272cfc860d8c3a19bbd1415","observation_id":"58049c56-4317-4cd3-a67c-316ef1a20f74","resolution":{"observed_at":"2026-07-02T02:46:28.805474Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-22T06:32:14.747728+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-22T06:32:14.747728+00:00","source":"crossref"},{"observed_at":"2026-08-22T06:32:06.552537+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-06-28T10:39:14.444486Z","title":"Onechart: Purify the chart structural extraction via one auxiliary token","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2606.03264","last_updated":"2026-06-02T07:27:03Z","snapshot_observed_at":"2026-08-16T07:12:09.561259Z","submitted_at":"2026-06-02T07:27:03Z","title":"PaddleOCR-VL-1.6: Expanding the Frontier of Document Parsing with Under-Optimized Region Refinement and Progressive Post-Training","version":1},"reference_index":53,"source":"pdf_text","source_observed_at":"2026-06-28T10:39:14.444486Z"},"links":{"citing_paper":"/paper/2606.03264"},"observation_digest":"sha256:5a4409491fbf00dde50fb12be6ee8868b988d726aca32356afbdc9d6fc6533ed","observation_id":"fde4ab5a-616f-4c44-bd41-2ae406ee2828","resolution":{"observed_at":"2026-06-28T10:39:14.444486Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2502.13923","last_updated":"2025-02-19T18:00:14Z","snapshot_observed_at":"2026-08-21T16:02:41.546193Z","submitted_at":"2025-02-19T18:00:14Z","title":"Qwen2.5-VL Technical Report","version":1},"cited_work":{"arxiv_id":"2502.13923","doi":"10.48550/arxiv.2502.13923","metadata_source":"pith","pith_arxiv_id":"2502.13923","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Qwen2.5-VL Technical Report","venue":"cs.CV","work_id":"69dffacb-bfe8-442d-be86-48624c60426f","year":2025},"citing_paper":{"arxiv_id":"2606.03264","last_updated":"2026-06-02T07:27:03Z","snapshot_observed_at":"2026-08-16T07:12:09.561259Z","submitted_at":"2026-06-02T07:27:03Z","title":"PaddleOCR-VL-1.6: Expanding the Frontier of Document Parsing with Under-Optimized Region Refinement and Progressive Post-Training","version":1},"reference_index":54,"source":"pdf_text","source_observed_at":"2026-06-28T10:39:14.444486Z"},"links":{"cited_paper":"/paper/2502.13923","citing_paper":"/paper/2606.03264"},"observation_digest":"sha256:5af9d3cc1b92feb0eab9618bba94893d0fec6d01da75bcf921eb0ff2a1e14994","observation_id":"6392d049-253c-49d0-a5fe-ac92bcd22eb5","resolution":{"observed_at":"2026-07-02T02:46:28.814772Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-22T06:32:14.747728+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-22T06:32:14.747728+00:00","source":"crossref"},{"observed_at":"2026-08-08T16:08:16.864468+00:00","source":"crossref_status_cache"},{"observed_at":"2026-08-08T16:08:16.864468+00:00","source":"openalex_status_cache"},{"observed_at":"2026-08-22T06:32:06.552537+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":"2510.12798","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-07-03T05:27:40.003735Z","title":"arXiv preprint arXiv:2510.12798 (2025)","venue":null,"work_id":"18a07d47-4390-4092-9958-8b1fab5c7aad","year":2025},"citing_paper":{"arxiv_id":"2606.03264","last_updated":"2026-06-02T07:27:03Z","snapshot_observed_at":"2026-08-16T07:12:09.561259Z","submitted_at":"2026-06-02T07:27:03Z","title":"PaddleOCR-VL-1.6: Expanding the Frontier of Document Parsing with Under-Optimized Region Refinement and Progressive Post-Training","version":1},"reference_index":55,"source":"pdf_text","source_observed_at":"2026-06-28T10:39:14.444486Z"},"links":{"citing_paper":"/paper/2606.03264"},"observation_digest":"sha256:89daf0bc3ab37febddee77f892db816d3793ad7c9f234dfc5659484cdfe33997","observation_id":"02adb6ae-c580-40ec-9b9a-4fe9ffd27cfd","resolution":{"observed_at":"2026-07-02T02:46:28.838045Z","resolver_source":"arxiv_id","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-22T06:32:14.747728+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-22T06:32:14.747728+00:00","source":"crossref"},{"observed_at":"2026-08-22T06:32:06.552537+00:00","source":"retraction_watch"}],"state":"measured"}}],"paper":{"arxiv_id":"2606.03264","last_updated":"2026-06-02T07:27:03Z","latest_version":1,"primary_category":"cs.CV","snapshot_observed_at":"2026-08-16T07:12:09.561259Z","submitted_at":"2026-06-02T07:27:03Z","title":"PaddleOCR-VL-1.6: Expanding the Frontier of Document Parsing with Under-Optimized Region Refinement and Progressive Post-Training"},"reference_resolution":{"displayed":50,"state_counts":{"malformed_identifier":0,"metadata_mismatch":6,"parse_uncertain":0,"unresolved":18,"verified_exact":26,"verified_fuzzy":0},"total_outbound_references":50},"refusal":"A citation records a reference. It does not transfer a finding from one paper to another.","schema":"pith.paper-citation-record.v1","standing_sources":[{"observed_at":"2026-08-22T06:32:14.747728+00:00","source":"crossref"},{"observed_at":"2026-08-22T06:32:06.552537+00:00","source":"retraction_watch"}],"thesis":"As of 22 August 2026, this Paper Citation Record lists 50 of 50 outbound references and 6 inbound Pith citation observations for arXiv:2606.03264."}