{"as_of":"2026-08-18T10:54:00Z","caps":{"database_statements":6,"inbound":100,"outbound":100},"context_digest":"sha256:3beb2e111617fa716f2c40086a619fb79185a3048e68bbe95f6fcc91172b0b0b","coverage":[{"denominator":18,"lane":"reference_resolution","note":"Typed states for the displayed outbound observations.","records_observed":18,"source":"paper_references, paper_reference_links","source_observed_at":"2026-08-15T20:38:29.748740Z","state":"measured"},{"denominator":18,"lane":"standing_notices","note":"One-hop event checks from named stored sources.","records_observed":18,"source":"scholarly_work_events, retraction_status_cache","source_observed_at":"2026-08-18T06:34:40.430872+00:00","state":"measured"},{"denominator":0,"lane":"inbound_itemization","note":"Pith citing papers itemized under the disclosed page cap.","records_observed":0,"source":"paper_references, paper_reference_links","source_observed_at":null,"state":"measured"},{"denominator":1,"lane":"external_citation_measurements","note":"A source-named dated measurement, never combined with another source.","records_observed":0,"source":"cited_works","source_observed_at":null,"state":"measured"}],"external_citation_measurements":[],"inbound":[],"links":{"evidence":"/evidence","html":"/paper/2505.13535/citation-record","integrity":"/paper/2505.13535/integrity","json":"/paper/2505.13535/citation-record.json","paper":"/paper/2505.13535"},"outbound":[{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-15T20:38:30.584700Z","title":null,"venue":null,"work_id":"dcb4429c-9768-4739-a1c3-2294ac9a7136","year":null},"citing_paper":{"arxiv_id":"2505.13535","last_updated":"2025-05-18T15:49:17Z","snapshot_observed_at":"2026-08-15T20:30:55.368302Z","submitted_at":"2025-05-18T15:49:17Z","title":"Information Extraction from Visually Rich Documents using LLM-based Organization of Documents into Independent Textual Segments","version":1},"reference_index":1,"source":"pdf_text","source_observed_at":"2026-08-15T20:38:29.612834Z"},"links":{"citing_paper":"/paper/2505.13535"},"observation_digest":"sha256:acf185873a68b5c0e9dcfd67a20fee4ba40ea62ce02d3ff242c04ff5da5fae5c","observation_id":"350dea41-d5bf-442e-ba23-2996d430d33f","resolution":{"observed_at":"2026-08-15T20:38:30.590200Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-18T06:34:40.430872+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-18T06:34:40.430872+00:00","source":"crossref"},{"observed_at":"2026-08-18T06:34:34.496301+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-15T20:38:30.358792Z","title":null,"venue":null,"work_id":"d5ffa8f8-0d9c-42ca-b1d2-f4a8860b5592","year":null},"citing_paper":{"arxiv_id":"2505.13535","last_updated":"2025-05-18T15:49:17Z","snapshot_observed_at":"2026-08-15T20:30:55.368302Z","submitted_at":"2025-05-18T15:49:17Z","title":"Information Extraction from Visually Rich Documents using LLM-based Organization of Documents into Independent Textual Segments","version":1},"reference_index":2,"source":"pdf_text","source_observed_at":"2026-08-15T20:38:29.632612Z"},"links":{"citing_paper":"/paper/2505.13535"},"observation_digest":"sha256:8bec14b068a9ae1b76df8df926df846357051617d60ee31a1fdd4e774603c196","observation_id":"2caca5d5-a1a9-4024-85c2-e5545fb3355a","resolution":{"observed_at":"2026-08-15T20:38:30.363294Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-18T06:34:40.430872+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-18T06:34:40.430872+00:00","source":"crossref"},{"observed_at":"2026-08-18T06:34:34.496301+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-15T20:38:30.568340Z","title":null,"venue":null,"work_id":"f6ee7395-6a3d-43ea-a944-b057493fafbe","year":null},"citing_paper":{"arxiv_id":"2505.13535","last_updated":"2025-05-18T15:49:17Z","snapshot_observed_at":"2026-08-15T20:30:55.368302Z","submitted_at":"2025-05-18T15:49:17Z","title":"Information Extraction from Visually Rich Documents using LLM-based Organization of Documents into Independent Textual Segments","version":1},"reference_index":3,"source":"pdf_text","source_observed_at":"2026-08-15T20:38:29.618199Z"},"links":{"citing_paper":"/paper/2505.13535"},"observation_digest":"sha256:58284fd987d30d3ff1fe0e8a00d0fd75dbd2b573f785969aff857311447ed333","observation_id":"4b1e28e9-cc48-4b44-b473-a21bf54aac59","resolution":{"observed_at":"2026-08-15T20:38:30.573577Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-18T06:34:40.430872+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-18T06:34:40.430872+00:00","source":"crossref"},{"observed_at":"2026-08-18T06:34:34.496301+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-15T20:38:30.552268Z","title":"U SE THESE TO HELP CONSTRUCT THE COMPLETE DICTIONARY : {blocks_and_parses} INSTRUCTIONS:","venue":null,"work_id":"897ffdeb-c5b8-4c39-9a68-7233fdee93fc","year":null},"citing_paper":{"arxiv_id":"2505.13535","last_updated":"2025-05-18T15:49:17Z","snapshot_observed_at":"2026-08-15T20:30:55.368302Z","submitted_at":"2025-05-18T15:49:17Z","title":"Information Extraction from Visually Rich Documents using LLM-based Organization of Documents into Independent Textual Segments","version":1},"reference_index":4,"source":"pdf_text","source_observed_at":"2026-08-15T20:38:29.623557Z"},"links":{"citing_paper":"/paper/2505.13535"},"observation_digest":"sha256:6c37f3cc92eabc7f27702ad8ed92c20dfddc7683bb3cbf05285d4000eb500378","observation_id":"5b2cfe11-1094-43e1-84b2-df94d660f656","resolution":{"observed_at":"2026-08-15T20:38:30.557737Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-18T06:34:40.430872+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-18T06:34:40.430872+00:00","source":"crossref"},{"observed_at":"2026-08-18T06:34:34.496301+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-15T20:38:30.487885Z","title":null,"venue":null,"work_id":"07941b19-6659-4d3c-85f6-48b99ffa2be2","year":null},"citing_paper":{"arxiv_id":"2505.13535","last_updated":"2025-05-18T15:49:17Z","snapshot_observed_at":"2026-08-15T20:30:55.368302Z","submitted_at":"2025-05-18T15:49:17Z","title":"Information Extraction from Visually Rich Documents using LLM-based Organization of Documents into Independent Textual Segments","version":1},"reference_index":11,"source":"pdf_text","source_observed_at":"2026-08-15T20:38:29.628301Z"},"links":{"citing_paper":"/paper/2505.13535"},"observation_digest":"sha256:55092c01ad64d5476a70643e84b56c1b004c63fadf68b70cb7778beb69237bc4","observation_id":"80921e7a-af43-4f36-9d57-adc7d09d4dbe","resolution":{"observed_at":"2026-08-15T20:38:30.539898Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-18T06:34:40.430872+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-18T06:34:40.430872+00:00","source":"crossref"},{"observed_at":"2026-08-18T06:34:34.496301+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-15T20:38:30.343045Z","title":null,"venue":null,"work_id":"ff9cdb0d-9dbd-42cd-8fa1-e5db7396d891","year":null},"citing_paper":{"arxiv_id":"2505.13535","last_updated":"2025-05-18T15:49:17Z","snapshot_observed_at":"2026-08-15T20:30:55.368302Z","submitted_at":"2025-05-18T15:49:17Z","title":"Information Extraction from Visually Rich Documents using LLM-based Organization of Documents into Independent Textual Segments","version":1},"reference_index":13,"source":"pdf_text","source_observed_at":"2026-08-15T20:38:29.687997Z"},"links":{"citing_paper":"/paper/2505.13535"},"observation_digest":"sha256:7df791c699413eac0e6b01a1bd57bfa24ab28201f89006034839ed2b0df06628","observation_id":"8a993dcf-e67d-4543-9b31-4e71e5833b4b","resolution":{"observed_at":"2026-08-15T20:38:30.348713Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-18T06:34:40.430872+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-18T06:34:40.430872+00:00","source":"crossref"},{"observed_at":"2026-08-18T06:34:34.496301+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-15T20:38:30.327118Z","title":null,"venue":null,"work_id":"9750b534-22ed-4c0c-b1a0-6f0282bbe7a0","year":null},"citing_paper":{"arxiv_id":"2505.13535","last_updated":"2025-05-18T15:49:17Z","snapshot_observed_at":"2026-08-15T20:30:55.368302Z","submitted_at":"2025-05-18T15:49:17Z","title":"Information Extraction from Visually Rich Documents using LLM-based Organization of Documents into Independent Textual Segments","version":1},"reference_index":14,"source":"pdf_text","source_observed_at":"2026-08-15T20:38:29.729989Z"},"links":{"citing_paper":"/paper/2505.13535"},"observation_digest":"sha256:072c51f815b2741141a6054957dd36670ab2a14038e9eea143f6587db9d0f168","observation_id":"643b3a92-d77a-4640-bfd1-bebafd526e9f","resolution":{"observed_at":"2026-08-15T20:38:30.332060Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-18T06:34:40.430872+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-18T06:34:40.430872+00:00","source":"crossref"},{"observed_at":"2026-08-18T06:34:34.496301+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-15T20:38:30.265663Z","title":"U SE THE REASON FROM PARTIAL PARSES , CHECK IF IT MENTIONS EXACT MATCH","venue":null,"work_id":"4fccf7ca-9b5f-4cc0-bef7-e100430dcf67","year":null},"citing_paper":{"arxiv_id":"2505.13535","last_updated":"2025-05-18T15:49:17Z","snapshot_observed_at":"2026-08-15T20:30:55.368302Z","submitted_at":"2025-05-18T15:49:17Z","title":"Information Extraction from Visually Rich Documents using LLM-based Organization of Documents into Independent Textual Segments","version":1},"reference_index":15,"source":"pdf_text","source_observed_at":"2026-08-15T20:38:29.735184Z"},"links":{"citing_paper":"/paper/2505.13535"},"observation_digest":"sha256:198a19461e186c31ae18c1d15a848fde6be142f91cbae8304227de8e427d6954","observation_id":"c6ff4ed2-2f32-4e20-91ff-34ab0c17b406","resolution":{"observed_at":"2026-08-15T20:38:30.316456Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-18T06:34:40.430872+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-18T06:34:40.430872+00:00","source":"crossref"},{"observed_at":"2026-08-18T06:34:34.496301+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-15T20:38:30.207957Z","title":null,"venue":null,"work_id":"6fa4f49e-b061-43f2-9796-c8314f701528","year":2019},"citing_paper":{"arxiv_id":"2505.13535","last_updated":"2025-05-18T15:49:17Z","snapshot_observed_at":"2026-08-15T20:30:55.368302Z","submitted_at":"2025-05-18T15:49:17Z","title":"Information Extraction from Visually Rich Documents using LLM-based Organization of Documents into Independent Textual Segments","version":1},"reference_index":16,"source":"pdf_text","source_observed_at":"2026-08-15T20:38:29.740017Z"},"links":{"citing_paper":"/paper/2505.13535"},"observation_digest":"sha256:9bb44ec5dab1142330b56d1a6bc075f6c6c4030ebc28159cfc39c70d6ae516ad","observation_id":"4792b0ab-fe58-4165-ad18-1a1b073854c4","resolution":{"observed_at":"2026-08-15T20:38:30.225185Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-18T06:34:40.430872+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-18T06:34:40.430872+00:00","source":"crossref"},{"observed_at":"2026-08-18T06:34:34.496301+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-15T20:38:30.176188Z","title":"It incorporates layout information us- ing cross-attention between bounding boxes and text, and through masked image modeling","venue":null,"work_id":"b02577e4-faeb-48ad-9254-93670ebbb0ab","year":2023},"citing_paper":{"arxiv_id":"2505.13535","last_updated":"2025-05-18T15:49:17Z","snapshot_observed_at":"2026-08-15T20:30:55.368302Z","submitted_at":"2025-05-18T15:49:17Z","title":"Information Extraction from Visually Rich Documents using LLM-based Organization of Documents into Independent Textual Segments","version":1},"reference_index":18,"source":"pdf_text","source_observed_at":"2026-08-15T20:38:29.748740Z"},"links":{"citing_paper":"/paper/2505.13535"},"observation_digest":"sha256:957e10bf7e60f7e2f800f5b041489bd9ee3841154d994066d228d002cfe54458","observation_id":"3bc2beaf-058a-4487-8f85-3db63c344bdb","resolution":{"observed_at":"2026-08-15T20:38:30.181358Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-18T06:34:40.430872+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-18T06:34:40.430872+00:00","source":"crossref"},{"observed_at":"2026-08-18T06:34:34.496301+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-15T20:38:30.191771Z","title":"30 hierarchical entities are annotated manually under top-level entities menu, subtotal and total","venue":null,"work_id":"0c23049a-2c51-4ac8-bdd1-61bc18a1c636","year":2019},"citing_paper":{"arxiv_id":"2505.13535","last_updated":"2025-05-18T15:49:17Z","snapshot_observed_at":"2026-08-15T20:30:55.368302Z","submitted_at":"2025-05-18T15:49:17Z","title":"Information Extraction from Visually Rich Documents using LLM-based Organization of Documents into Independent Textual Segments","version":1},"reference_index":100,"source":"pdf_text","source_observed_at":"2026-08-15T20:38:29.744663Z"},"links":{"citing_paper":"/paper/2505.13535"},"observation_digest":"sha256:f4cc3baa8f63640142335301b3f4cf1b5dc29f875871a64a4ea5df0f60d009e7","observation_id":"7a8e77e2-f7bb-4836-8dd9-99344d27facd","resolution":{"observed_at":"2026-08-15T20:38:30.196743Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-18T06:34:40.430872+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-18T06:34:40.430872+00:00","source":"crossref"},{"observed_at":"2026-08-18T06:34:34.496301+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"1909.04948","last_updated":"2019-10-14T09:55:39Z","snapshot_observed_at":"2026-08-14T01:13:04.875715Z","submitted_at":"2019-09-11T09:51:02Z","title":"BERTgrid: Contextualized Embedding for 2D Document Representation and Understanding","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"1909.04948","snapshot_observed_at":"2026-08-15T20:38:29.507278Z","title":"Association for Computational Linguistics","venue":null,"work_id":null,"year":2013},"citing_paper":{"arxiv_id":"2505.13535","last_updated":"2025-05-18T15:49:17Z","snapshot_observed_at":"2026-08-15T20:30:55.368302Z","submitted_at":"2025-05-18T15:49:17Z","title":"Information Extraction from Visually Rich Documents using LLM-based Organization of Documents into Independent Textual Segments","version":1},"reference_index":2013,"source":"pdf_text","source_observed_at":"2026-08-15T20:38:29.507278Z"},"links":{"cited_paper":"/paper/1909.04948","citing_paper":"/paper/2505.13535"},"observation_digest":"sha256:8feab96455a0a727126b2fe5e143a35267cc44b0869c2d7e52667ffb4907fd96","observation_id":"a762fab5-78b6-4a19-814d-23201f95ab69","resolution":{"observed_at":"2026-08-15T20:38:29.507278Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2106.09685","last_updated":"2021-10-16T18:40:34Z","snapshot_observed_at":"2026-08-17T18:04:53.578114Z","submitted_at":"2021-06-17T17:37:18Z","title":"LoRA: Low-Rank Adaptation of Large Language Models","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2106.09685","snapshot_observed_at":"2026-08-15T20:38:29.528844Z","title":"In 2016 12th IAPR Workshop on Document Analysis Systems (DAS), pages 287–292","venue":null,"work_id":null,"year":2016},"citing_paper":{"arxiv_id":"2505.13535","last_updated":"2025-05-18T15:49:17Z","snapshot_observed_at":"2026-08-15T20:30:55.368302Z","submitted_at":"2025-05-18T15:49:17Z","title":"Information Extraction from Visually Rich Documents using LLM-based Organization of Documents into Independent Textual Segments","version":1},"reference_index":2016,"source":"pdf_text","source_observed_at":"2026-08-15T20:38:29.528844Z"},"links":{"cited_paper":"/paper/2106.09685","citing_paper":"/paper/2505.13535"},"observation_digest":"sha256:7a1ce0af901ae04532521493acd3dfadb75b8643b77c5ffd72827d2bf660209d","observation_id":"aeb51e6c-a349-49b2-98c8-05358ebbf200","resolution":{"observed_at":"2026-08-15T20:38:29.528844Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"1708.07403","last_updated":"2017-08-24T13:40:06Z","snapshot_observed_at":"2026-08-14T20:39:19.356520Z","submitted_at":"2017-08-24T13:40:06Z","title":"CloudScan - A configuration-free invoice analysis system using recurrent neural networks","version":1},"cited_work":{"arxiv_id":"1708.07403","doi":null,"metadata_source":"pith","pith_arxiv_id":"1708.07403","snapshot_observed_at":"2026-08-15T20:38:29.802016Z","title":"CloudScan - A configuration-free invoice analysis system using recurrent neural networks","venue":"cs.CL","work_id":"383ebd8a-75bf-4427-a4e2-8e53b3f6e70a","year":2017},"citing_paper":{"arxiv_id":"2505.13535","last_updated":"2025-05-18T15:49:17Z","snapshot_observed_at":"2026-08-15T20:30:55.368302Z","submitted_at":"2025-05-18T15:49:17Z","title":"Information Extraction from Visually Rich Documents using LLM-based Organization of Documents into Independent Textual Segments","version":1},"reference_index":2017,"source":"pdf_text","source_observed_at":"2026-08-15T20:38:29.565993Z"},"links":{"cited_paper":"/paper/1708.07403","citing_paper":"/paper/2505.13535"},"observation_digest":"sha256:5f5d9086d6dbc539e1d6a451743fc174af5a60137badaa5764256ed4f7153096","observation_id":"94ecd33f-f291-4fcc-9ec3-5de93c4acc9e","resolution":{"observed_at":"2026-08-15T20:38:29.861790Z","resolver_source":"local_arxiv","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-18T06:34:40.430872+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-18T06:34:40.430872+00:00","source":"crossref"},{"observed_at":"2026-08-18T06:34:34.496301+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2005.00642","last_updated":"2021-07-01T08:32:15Z","snapshot_observed_at":"2026-07-06T09:16:59.519524Z","submitted_at":"2020-05-01T22:59:56Z","title":"Spatial Dependency Parsing for Semi-Structured Document Information Extraction","version":3},"cited_work":{"arxiv_id":"2005.00642","doi":null,"metadata_source":"pith","pith_arxiv_id":"2005.00642","snapshot_observed_at":"2026-08-15T20:38:29.938230Z","title":"Spatial Dependency Parsing for Semi-Structured Document Information Extraction","venue":"cs.CL","work_id":"aa0b71b4-7ab6-4d58-af0b-423dc746766c","year":2020},"citing_paper":{"arxiv_id":"2505.13535","last_updated":"2025-05-18T15:49:17Z","snapshot_observed_at":"2026-08-15T20:30:55.368302Z","submitted_at":"2025-05-18T15:49:17Z","title":"Information Extraction from Visually Rich Documents using LLM-based Organization of Documents into Independent Textual Segments","version":1},"reference_index":2019,"source":"pdf_text","source_observed_at":"2026-08-15T20:38:29.535168Z"},"links":{"cited_paper":"/paper/2005.00642","citing_paper":"/paper/2505.13535"},"observation_digest":"sha256:473b2243c3e528fa1ee2446347d7dd73e2e07a4596b138436707401ba6a882f5","observation_id":"17a0415b-86b4-412b-b879-ee142e3c6eb4","resolution":{"observed_at":"2026-08-15T20:38:29.942680Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-18T06:34:40.430872+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-18T06:34:40.430872+00:00","source":"crossref"},{"observed_at":"2026-08-18T06:34:34.496301+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2002.12804","last_updated":"2020-02-28T15:28:49Z","snapshot_observed_at":"2026-08-12T16:37:27.053671Z","submitted_at":"2020-02-28T15:28:49Z","title":"UniLMv2: Pseudo-Masked Language Models for Unified Language Model Pre-Training","version":1},"cited_work":{"arxiv_id":"2002.12804","doi":null,"metadata_source":"pith","pith_arxiv_id":"2002.12804","snapshot_observed_at":"2026-08-15T20:38:30.052406Z","title":"UniLMv2: Pseudo-Masked Language Models for Unified Language Model Pre-Training","venue":"cs.CL","work_id":"796a4915-32c1-4cf4-a903-18b458a7f00c","year":2020},"citing_paper":{"arxiv_id":"2505.13535","last_updated":"2025-05-18T15:49:17Z","snapshot_observed_at":"2026-08-15T20:30:55.368302Z","submitted_at":"2025-05-18T15:49:17Z","title":"Information Extraction from Visually Rich Documents using LLM-based Organization of Documents into Independent Textual Segments","version":1},"reference_index":2020,"source":"pdf_text","source_observed_at":"2026-08-15T20:38:29.453077Z"},"links":{"cited_paper":"/paper/2002.12804","citing_paper":"/paper/2505.13535"},"observation_digest":"sha256:b054a7685c3ef94b23cdc684fb7651a9fbe74c79089dd568a8027f9d08833f71","observation_id":"22497956-9e37-49fe-9d8a-fc320448d193","resolution":{"observed_at":"2026-08-15T20:38:30.157559Z","resolver_source":"local_arxiv","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-18T06:34:40.430872+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-18T06:34:40.430872+00:00","source":"crossref"},{"observed_at":"2026-08-18T06:34:34.496301+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2210.06155","last_updated":"2022-10-14T06:54:17Z","snapshot_observed_at":"2026-08-16T16:24:51.096454Z","submitted_at":"2022-10-12T12:59:24Z","title":"ERNIE-Layout: Layout Knowledge Enhanced Pre-training for Visually-rich Document Understanding","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2210.06155","snapshot_observed_at":"2026-08-15T20:38:29.606582Z","title":"{query_reason}","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2505.13535","last_updated":"2025-05-18T15:49:17Z","snapshot_observed_at":"2026-08-15T20:30:55.368302Z","submitted_at":"2025-05-18T15:49:17Z","title":"Information Extraction from Visually Rich Documents using LLM-based Organization of Documents into Independent Textual Segments","version":1},"reference_index":2022,"source":"pdf_text","source_observed_at":"2026-08-15T20:38:29.606582Z"},"links":{"cited_paper":"/paper/2210.06155","citing_paper":"/paper/2505.13535"},"observation_digest":"sha256:ff9e7d0c9b02780feda4bb289519cfb58674f9b84aa356ac0691158109425133","observation_id":"d6fc3f13-d934-43d7-aa2d-ede36bcd4cf5","resolution":{"observed_at":"2026-08-15T20:38:29.606582Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2401.10825","last_updated":"2024-12-20T15:11:13Z","snapshot_observed_at":"2026-08-17T13:04:24.245421Z","submitted_at":"2024-01-19T17:21:05Z","title":"Recent Advances in Named Entity Recognition: A Comprehensive Survey and Comparative Study","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2401.10825","snapshot_observed_at":"2026-08-15T20:38:29.542231Z","title":"Preprint, arXiv:2401.10825","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2505.13535","last_updated":"2025-05-18T15:49:17Z","snapshot_observed_at":"2026-08-15T20:30:55.368302Z","submitted_at":"2025-05-18T15:49:17Z","title":"Information Extraction from Visually Rich Documents using LLM-based Organization of Documents into Independent Textual Segments","version":1},"reference_index":2024,"source":"pdf_text","source_observed_at":"2026-08-15T20:38:29.542231Z"},"links":{"cited_paper":"/paper/2401.10825","citing_paper":"/paper/2505.13535"},"observation_digest":"sha256:8733a675f784450c527f0dfe2cc21d44050f4d21e60663028c17ec9ab3209748","observation_id":"c75c872d-3ee9-4ca1-8a2c-b299f36b3fd5","resolution":{"observed_at":"2026-08-15T20:38:29.542231Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}}],"paper":{"arxiv_id":"2505.13535","last_updated":"2025-05-18T15:49:17Z","latest_version":1,"primary_category":"cs.IR","snapshot_observed_at":"2026-08-15T20:30:55.368302Z","submitted_at":"2025-05-18T15:49:17Z","title":"Information Extraction from Visually Rich Documents using LLM-based Organization of Documents into Independent Textual Segments"},"reference_resolution":{"displayed":18,"state_counts":{"malformed_identifier":0,"metadata_mismatch":2,"parse_uncertain":0,"unresolved":11,"verified_exact":1,"verified_fuzzy":4},"total_outbound_references":18},"refusal":"A citation records a reference. It does not transfer a finding from one paper to another.","schema":"pith.paper-citation-record.v1","standing_sources":[{"observed_at":"2026-08-18T06:34:40.430872+00:00","source":"crossref"},{"observed_at":"2026-08-18T06:34:34.496301+00:00","source":"retraction_watch"}],"thesis":"As of 18 August 2026, this Paper Citation Record lists 18 of 18 outbound references and 0 inbound Pith citation observations for arXiv:2505.13535."}