{"as_of":"2026-08-13T05:45:00Z","caps":{"database_statements":6,"inbound":100,"outbound":100},"context_digest":"sha256:1857b7812aeb711cd106f7d5ffc3aa6bff6c2467f6545a19e74d90b7a9c00dd7","coverage":[{"denominator":31,"lane":"reference_resolution","note":"Typed states for the displayed outbound observations.","records_observed":31,"source":"paper_references, paper_reference_links","source_observed_at":"2026-08-10T23:39:19.587686Z","state":"measured"},{"denominator":31,"lane":"standing_notices","note":"One-hop event checks from named stored sources.","records_observed":31,"source":"scholarly_work_events, retraction_status_cache","source_observed_at":"2026-08-12T06:34:41.77262+00:00","state":"measured"},{"denominator":0,"lane":"inbound_itemization","note":"Pith citing papers itemized under the disclosed page cap.","records_observed":0,"source":"paper_references, paper_reference_links","source_observed_at":null,"state":"measured"},{"denominator":1,"lane":"external_citation_measurements","note":"A source-named dated measurement, never combined with another source.","records_observed":0,"source":"cited_works","source_observed_at":null,"state":"measured"}],"external_citation_measurements":[],"inbound":[],"links":{"evidence":"/evidence","html":"/paper/2412.20072/citation-record","integrity":"/paper/2412.20072/integrity","json":"/paper/2412.20072/citation-record.json","paper":"/paper/2412.20072"},"outbound":[{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-10T23:39:19.496871Z","title":"Chain-of-thought prompting elicits reasoning in large language models,","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2412.20072","last_updated":"2024-12-31T03:11:03Z","snapshot_observed_at":"2026-08-12T19:56:06.065692Z","submitted_at":"2024-12-28T07:54:14Z","title":"Extract Information from Hybrid Long Documents Leveraging LLMs: A Framework and Dataset","version":2},"reference_index":1,"source":"pdf_text","source_observed_at":"2026-08-10T23:39:19.496871Z"},"links":{"citing_paper":"/paper/2412.20072"},"observation_digest":"sha256:e388ed45ebf9cb392637e3f24a6883ca1300ae942de36147e05a05be3a4e822b","observation_id":"9779d5ef-ad52-4866-8ce8-7dbd49ba260c","resolution":{"observed_at":"2026-08-10T23:39:19.496871Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-10T23:39:19.932810Z","title":"Text2analysis: A benchmark of table question answering with advanced data analysis and unclear queries,","venue":null,"work_id":"dc24998e-90d9-4d34-86da-f92ffbeafa95","year":2024},"citing_paper":{"arxiv_id":"2412.20072","last_updated":"2024-12-31T03:11:03Z","snapshot_observed_at":"2026-08-12T19:56:06.065692Z","submitted_at":"2024-12-28T07:54:14Z","title":"Extract Information from Hybrid Long Documents Leveraging LLMs: A Framework and Dataset","version":2},"reference_index":2,"source":"pdf_text","source_observed_at":"2026-08-10T23:39:19.501136Z"},"links":{"citing_paper":"/paper/2412.20072"},"observation_digest":"sha256:0b078c857f5fdf650497d9b26289d9e5894f9541d9749caffc676486450b92df","observation_id":"8d74f39f-9ed1-4018-ab58-42dd22aaf7da","resolution":{"observed_at":"2026-08-10T23:39:19.935280Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-12T06:34:41.77262+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-12T06:34:41.77262+00:00","source":"crossref"},{"observed_at":"2026-08-12T06:34:36.333875+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2403.10249","last_updated":"2024-03-15T12:37:12Z","snapshot_observed_at":"2026-08-13T00:53:30.393145Z","submitted_at":"2024-03-15T12:37:12Z","title":"A Survey on Game Playing Agents and Large Models: Methods, Applications, and Challenges","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2403.10249","snapshot_observed_at":"2026-08-10T23:39:19.505234Z","title":"A survey on game playing agents and large models: Methods, applications, and challenges,","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2412.20072","last_updated":"2024-12-31T03:11:03Z","snapshot_observed_at":"2026-08-12T19:56:06.065692Z","submitted_at":"2024-12-28T07:54:14Z","title":"Extract Information from Hybrid Long Documents Leveraging LLMs: A Framework and Dataset","version":2},"reference_index":3,"source":"pdf_text","source_observed_at":"2026-08-10T23:39:19.505234Z"},"links":{"cited_paper":"/paper/2403.10249","citing_paper":"/paper/2412.20072"},"observation_digest":"sha256:5098d362e4b7baaf00a6f2672ace05373002c12c9a71538543a138538813e9bf","observation_id":"64866c6f-6117-4ae2-95ef-e9b628981719","resolution":{"observed_at":"2026-08-10T23:39:19.505234Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-10T23:39:19.925369Z","title":"Strago: Harnessing strategic guidance for prompt opti- mization,","venue":null,"work_id":"0d3d8050-d224-4dcf-9036-dd5a3b70fde0","year":2024},"citing_paper":{"arxiv_id":"2412.20072","last_updated":"2024-12-31T03:11:03Z","snapshot_observed_at":"2026-08-12T19:56:06.065692Z","submitted_at":"2024-12-28T07:54:14Z","title":"Extract Information from Hybrid Long Documents Leveraging LLMs: A Framework and Dataset","version":2},"reference_index":4,"source":"pdf_text","source_observed_at":"2026-08-10T23:39:19.508776Z"},"links":{"citing_paper":"/paper/2412.20072"},"observation_digest":"sha256:952aaccc7fc5b3403cf795d40d11151625b7a0c08c2068c974bcb05271630bc5","observation_id":"b43871c2-5fed-418a-a281-ab1f6b683d38","resolution":{"observed_at":"2026-08-10T23:39:19.928036Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-12T06:34:41.77262+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-12T06:34:41.77262+00:00","source":"crossref"},{"observed_at":"2026-08-12T06:34:36.333875+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2305.16344","last_updated":"2024-03-07T13:44:27Z","snapshot_observed_at":"2026-08-03T12:56:22.368722Z","submitted_at":"2023-05-24T10:35:58Z","title":"Enabling and Analyzing How to Efficiently Extract Information from Hybrid Long Documents with LLMs","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2305.16344","snapshot_observed_at":"2026-08-10T23:39:19.511719Z","title":"Enabling and analyzing how to efficiently extract information from hybrid long documents with llms,","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2412.20072","last_updated":"2024-12-31T03:11:03Z","snapshot_observed_at":"2026-08-12T19:56:06.065692Z","submitted_at":"2024-12-28T07:54:14Z","title":"Extract Information from Hybrid Long Documents Leveraging LLMs: A Framework and Dataset","version":2},"reference_index":5,"source":"pdf_text","source_observed_at":"2026-08-10T23:39:19.511719Z"},"links":{"cited_paper":"/paper/2305.16344","citing_paper":"/paper/2412.20072"},"observation_digest":"sha256:2e92b83bfdeda067d4273aca787aabefc635a53adf1bea44d8dbfdf97639584e","observation_id":"f4d08ce1-3f90-4620-acbf-5438e6b0e899","resolution":{"observed_at":"2026-08-10T23:39:19.511719Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-10T23:39:19.514811Z","title":"Large lan- guage models are zero-shot reasoners,","venue":null,"work_id":null,"year":2022},"citing_paper":{"arxiv_id":"2412.20072","last_updated":"2024-12-31T03:11:03Z","snapshot_observed_at":"2026-08-12T19:56:06.065692Z","submitted_at":"2024-12-28T07:54:14Z","title":"Extract Information from Hybrid Long Documents Leveraging LLMs: A Framework and Dataset","version":2},"reference_index":6,"source":"pdf_text","source_observed_at":"2026-08-10T23:39:19.514811Z"},"links":{"citing_paper":"/paper/2412.20072"},"observation_digest":"sha256:f8eeaf4e2ed733ef3c096a94ab1960284ac2eb725c3118295b6ed9698babfcc1","observation_id":"7010662f-5071-41f6-bdb9-02b0bbb01a9e","resolution":{"observed_at":"2026-08-10T23:39:19.514811Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2210.06710","last_updated":"2023-01-23T18:51:00Z","snapshot_observed_at":"2026-08-09T01:12:20.344941Z","submitted_at":"2022-10-13T04:08:24Z","title":"Large Language Models are few(1)-shot Table Reasoners","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2210.06710","snapshot_observed_at":"2026-08-10T23:39:19.517744Z","title":"Large language models are few (1)-shot table reasoners,","venue":null,"work_id":null,"year":2022},"citing_paper":{"arxiv_id":"2412.20072","last_updated":"2024-12-31T03:11:03Z","snapshot_observed_at":"2026-08-12T19:56:06.065692Z","submitted_at":"2024-12-28T07:54:14Z","title":"Extract Information from Hybrid Long Documents Leveraging LLMs: A Framework and Dataset","version":2},"reference_index":7,"source":"pdf_text","source_observed_at":"2026-08-10T23:39:19.517744Z"},"links":{"cited_paper":"/paper/2210.06710","citing_paper":"/paper/2412.20072"},"observation_digest":"sha256:036ef09936626a77289a9ff85778f51d513def2c9b435c9baf61aa39b4c4da84","observation_id":"af1f6ab8-ac54-4f72-af87-61df22330c50","resolution":{"observed_at":"2026-08-10T23:39:19.517744Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2301.13808","last_updated":"2023-04-27T11:24:10Z","snapshot_observed_at":"2026-08-11T03:08:43.900637Z","submitted_at":"2023-01-31T17:51:45Z","title":"Large Language Models are Versatile Decomposers: Decompose Evidence and Questions for Table-based Reasoning","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2301.13808","snapshot_observed_at":"2026-08-10T23:39:19.520599Z","title":"Large language models are versatile decomposers: Decompose evidence and questions for table-based reasoning,","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2412.20072","last_updated":"2024-12-31T03:11:03Z","snapshot_observed_at":"2026-08-12T19:56:06.065692Z","submitted_at":"2024-12-28T07:54:14Z","title":"Extract Information from Hybrid Long Documents Leveraging LLMs: A Framework and Dataset","version":2},"reference_index":8,"source":"pdf_text","source_observed_at":"2026-08-10T23:39:19.520599Z"},"links":{"cited_paper":"/paper/2301.13808","citing_paper":"/paper/2412.20072"},"observation_digest":"sha256:238322ff2f200c51fa3c32117129eb28b54bfbeee3e0194625074953a83f34ab","observation_id":"06354993-17cc-4009-a78a-792fdd2b3484","resolution":{"observed_at":"2026-08-10T23:39:19.520599Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-10T23:39:19.912543Z","title":"Cradle: Empowering foundation agents towards general computer control,","venue":null,"work_id":"a9a02dfe-7d3b-48bc-92ec-eb448a4e0d31","year":2024},"citing_paper":{"arxiv_id":"2412.20072","last_updated":"2024-12-31T03:11:03Z","snapshot_observed_at":"2026-08-12T19:56:06.065692Z","submitted_at":"2024-12-28T07:54:14Z","title":"Extract Information from Hybrid Long Documents Leveraging LLMs: A Framework and Dataset","version":2},"reference_index":9,"source":"pdf_text","source_observed_at":"2026-08-10T23:39:19.524457Z"},"links":{"citing_paper":"/paper/2412.20072"},"observation_digest":"sha256:deef3fe684cda5ff62745fb39df74ec4afb069339223ed5dc5fe052670ad0f5e","observation_id":"49358fae-1e7b-43d4-a6a7-2dd3b0960b96","resolution":{"observed_at":"2026-08-10T23:39:19.915922Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-12T06:34:41.77262+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-12T06:34:41.77262+00:00","source":"crossref"},{"observed_at":"2026-08-12T06:34:36.333875+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-10T23:39:19.904582Z","title":"Embedding-based product retrieval in taobao search,","venue":null,"work_id":"6788b4be-b820-410c-8bdf-7571a6f7a9dd","year":2021},"citing_paper":{"arxiv_id":"2412.20072","last_updated":"2024-12-31T03:11:03Z","snapshot_observed_at":"2026-08-12T19:56:06.065692Z","submitted_at":"2024-12-28T07:54:14Z","title":"Extract Information from Hybrid Long Documents Leveraging LLMs: A Framework and Dataset","version":2},"reference_index":10,"source":"pdf_text","source_observed_at":"2026-08-10T23:39:19.527451Z"},"links":{"citing_paper":"/paper/2412.20072"},"observation_digest":"sha256:b6a88ef90d0f38a282d8dcf34725ce74392b5d854ad4f7e2f05fcfb4f1ba7cb1","observation_id":"7086dcc1-840f-4dd1-b661-2d1eac3dd82b","resolution":{"observed_at":"2026-08-10T23:39:19.907628Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-12T06:34:41.77262+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-12T06:34:41.77262+00:00","source":"crossref"},{"observed_at":"2026-08-12T06:34:36.333875+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-10T23:39:19.530913Z","title":"Can large lan- guage models recall reference location like humans?","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2412.20072","last_updated":"2024-12-31T03:11:03Z","snapshot_observed_at":"2026-08-12T19:56:06.065692Z","submitted_at":"2024-12-28T07:54:14Z","title":"Extract Information from Hybrid Long Documents Leveraging LLMs: A Framework and Dataset","version":2},"reference_index":11,"source":"pdf_text","source_observed_at":"2026-08-10T23:39:19.530913Z"},"links":{"citing_paper":"/paper/2412.20072"},"observation_digest":"sha256:122330f0c7d36fdc8936e3d751a6106943e6c6e32b450590f7c36b969e428974","observation_id":"c5670a7c-f01d-4e0b-b260-e99e90a2a40a","resolution":{"observed_at":"2026-08-10T23:39:19.530913Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2410.03450","last_updated":"2025-05-22T14:19:41Z","snapshot_observed_at":"2026-08-12T22:30:43.542721Z","submitted_at":"2024-10-04T14:10:39Z","title":"MLLM as Retriever: Interactively Learning Multimodal Retrieval for Embodied Agents","version":2},"cited_work":{"arxiv_id":"2410.03450","doi":null,"metadata_source":"pith","pith_arxiv_id":"2410.03450","snapshot_observed_at":"2026-08-10T23:39:19.731607Z","title":"MLLM as Retriever: Interactively Learning Multimodal Retrieval for Embodied Agents","venue":"cs.LG","work_id":"61d007be-d622-4deb-9462-187738f5823f","year":2024},"citing_paper":{"arxiv_id":"2412.20072","last_updated":"2024-12-31T03:11:03Z","snapshot_observed_at":"2026-08-12T19:56:06.065692Z","submitted_at":"2024-12-28T07:54:14Z","title":"Extract Information from Hybrid Long Documents Leveraging LLMs: A Framework and Dataset","version":2},"reference_index":12,"source":"pdf_text","source_observed_at":"2026-08-10T23:39:19.533860Z"},"links":{"cited_paper":"/paper/2410.03450","citing_paper":"/paper/2412.20072"},"observation_digest":"sha256:5463b53fde15dd9746342406cb4d5fc3639031623149a27d555232b7b63dc8dc","observation_id":"af5ff846-a45c-492f-8f45-a67973edf32a","resolution":{"observed_at":"2026-08-10T23:39:19.735906Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-12T06:34:41.77262+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-12T06:34:41.77262+00:00","source":"crossref"},{"observed_at":"2026-08-12T06:34:36.333875+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"1908.10084","last_updated":"2019-08-27T08:50:17Z","snapshot_observed_at":"2026-07-06T08:17:05.681370Z","submitted_at":"2019-08-27T08:50:17Z","title":"Sentence-BERT: Sentence Embeddings using Siamese BERT-Networks","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"1908.10084","snapshot_observed_at":"2026-08-10T23:39:19.536983Z","title":"Sentence-bert: Sentence embeddings using siamese bert-networks,","venue":null,"work_id":null,"year":1908},"citing_paper":{"arxiv_id":"2412.20072","last_updated":"2024-12-31T03:11:03Z","snapshot_observed_at":"2026-08-12T19:56:06.065692Z","submitted_at":"2024-12-28T07:54:14Z","title":"Extract Information from Hybrid Long Documents Leveraging LLMs: A Framework and Dataset","version":2},"reference_index":13,"source":"pdf_text","source_observed_at":"2026-08-10T23:39:19.536983Z"},"links":{"cited_paper":"/paper/1908.10084","citing_paper":"/paper/2412.20072"},"observation_digest":"sha256:1a62e8c97a1070f92c3fa5fd0e2099745c377033245263226e14319f41f6566c","observation_id":"c8dd9f86-a03b-4529-a68f-6f315845a302","resolution":{"observed_at":"2026-08-10T23:39:19.536983Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2011.03228","last_updated":"2020-11-06T08:22:12Z","snapshot_observed_at":"2026-07-06T10:12:12.349833Z","submitted_at":"2020-11-06T08:22:12Z","title":"From Dataset Recycling to Multi-Property Extraction and Beyond","version":1},"cited_work":{"arxiv_id":"2011.03228","doi":null,"metadata_source":"pith","pith_arxiv_id":"2011.03228","snapshot_observed_at":"2026-08-10T23:39:19.710393Z","title":"From Dataset Recycling to Multi-Property Extraction and Beyond","venue":"cs.CL","work_id":"8d056081-98c5-441c-ba41-8eaad8be16bc","year":2020},"citing_paper":{"arxiv_id":"2412.20072","last_updated":"2024-12-31T03:11:03Z","snapshot_observed_at":"2026-08-12T19:56:06.065692Z","submitted_at":"2024-12-28T07:54:14Z","title":"Extract Information from Hybrid Long Documents Leveraging LLMs: A Framework and Dataset","version":2},"reference_index":14,"source":"pdf_text","source_observed_at":"2026-08-10T23:39:19.540661Z"},"links":{"cited_paper":"/paper/2011.03228","citing_paper":"/paper/2412.20072"},"observation_digest":"sha256:211665b3fb7b1701323c7194275aa22edfa6b1ae5fa9c03f3683592a7ddc64b3","observation_id":"76b58022-959f-4cd8-8149-fc414c14d19f","resolution":{"observed_at":"2026-08-10T23:39:19.714894Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-12T06:34:41.77262+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-12T06:34:41.77262+00:00","source":"crossref"},{"observed_at":"2026-08-12T06:34:36.333875+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2302.04914","last_updated":"2024-06-12T14:25:15Z","snapshot_observed_at":"2026-07-06T14:50:14.766783Z","submitted_at":"2023-02-09T19:56:37Z","title":"Flexible, Model-Agnostic Method for Materials Data Extraction from Text Using General Purpose Language Models","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2302.04914","snapshot_observed_at":"2026-08-10T23:39:19.544314Z","title":"Flexible, model-agnostic method for materials data extraction from text using general purpose language models,","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2412.20072","last_updated":"2024-12-31T03:11:03Z","snapshot_observed_at":"2026-08-12T19:56:06.065692Z","submitted_at":"2024-12-28T07:54:14Z","title":"Extract Information from Hybrid Long Documents Leveraging LLMs: A Framework and Dataset","version":2},"reference_index":15,"source":"pdf_text","source_observed_at":"2026-08-10T23:39:19.544314Z"},"links":{"cited_paper":"/paper/2302.04914","citing_paper":"/paper/2412.20072"},"observation_digest":"sha256:ee8a2170717aaed6f8a4bd88e28196184fa4b5f7f03cf1bf688ffdb2898b77d5","observation_id":"60a231c8-c45d-4b3f-b519-2be1839d05d2","resolution":{"observed_at":"2026-08-10T23:39:19.544314Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-10T23:39:19.895145Z","title":"A hybrid ai tool to extract key performance indicators from financial reports for benchmarking,","venue":null,"work_id":"35865790-28d4-4d21-afa1-c22327892450","year":2019},"citing_paper":{"arxiv_id":"2412.20072","last_updated":"2024-12-31T03:11:03Z","snapshot_observed_at":"2026-08-12T19:56:06.065692Z","submitted_at":"2024-12-28T07:54:14Z","title":"Extract Information from Hybrid Long Documents Leveraging LLMs: A Framework and Dataset","version":2},"reference_index":16,"source":"pdf_text","source_observed_at":"2026-08-10T23:39:19.547077Z"},"links":{"citing_paper":"/paper/2412.20072"},"observation_digest":"sha256:6fe4747da181e9a486fdfb5bbbf363d8435ad535d44968668d173f9e67122abb","observation_id":"50a58a9f-4aac-46fc-bc3f-ca57ef03fe74","resolution":{"observed_at":"2026-08-10T23:39:19.898577Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-12T06:34:41.77262+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-12T06:34:41.77262+00:00","source":"crossref"},{"observed_at":"2026-08-12T06:34:36.333875+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-10T23:39:19.886161Z","title":"Spot: A tool for identifying operating segments in financial tables,","venue":null,"work_id":"eca05f9d-08e2-4295-ba2d-01ab38c88d6a","year":2020},"citing_paper":{"arxiv_id":"2412.20072","last_updated":"2024-12-31T03:11:03Z","snapshot_observed_at":"2026-08-12T19:56:06.065692Z","submitted_at":"2024-12-28T07:54:14Z","title":"Extract Information from Hybrid Long Documents Leveraging LLMs: A Framework and Dataset","version":2},"reference_index":17,"source":"pdf_text","source_observed_at":"2026-08-10T23:39:19.550044Z"},"links":{"citing_paper":"/paper/2412.20072"},"observation_digest":"sha256:9d7f0bbefa945fe01fb320612ff85d9e5134b09a29147b836e88c4873a5e06ab","observation_id":"77b16cd9-1924-47d9-b27a-ba540bcae181","resolution":{"observed_at":"2026-08-10T23:39:19.889139Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-12T06:34:41.77262+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-12T06:34:41.77262+00:00","source":"crossref"},{"observed_at":"2026-08-12T06:34:36.333875+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-10T23:39:19.875241Z","title":"Exploring word representations on time expression recognition,","venue":null,"work_id":"6a5f5298-1304-4f01-a772-9b068bd054e3","year":2019},"citing_paper":{"arxiv_id":"2412.20072","last_updated":"2024-12-31T03:11:03Z","snapshot_observed_at":"2026-08-12T19:56:06.065692Z","submitted_at":"2024-12-28T07:54:14Z","title":"Extract Information from Hybrid Long Documents Leveraging LLMs: A Framework and Dataset","version":2},"reference_index":18,"source":"pdf_text","source_observed_at":"2026-08-10T23:39:19.553370Z"},"links":{"citing_paper":"/paper/2412.20072"},"observation_digest":"sha256:df8b3b083f3c59139f5122ca6af9b88d33d671a0d5e951030965bd36aa36c903","observation_id":"29bb98af-7c47-4451-b7e3-7ac145cc55e2","resolution":{"observed_at":"2026-08-10T23:39:19.878256Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-12T06:34:41.77262+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-12T06:34:41.77262+00:00","source":"crossref"},{"observed_at":"2026-08-12T06:34:36.333875+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-10T23:39:19.866374Z","title":"Kpi-bert: A joint named entity recognition and rela- tion extraction model for financial reports,","venue":null,"work_id":"4fccd632-3a4d-4099-af39-4627f06aeb02","year":2022},"citing_paper":{"arxiv_id":"2412.20072","last_updated":"2024-12-31T03:11:03Z","snapshot_observed_at":"2026-08-12T19:56:06.065692Z","submitted_at":"2024-12-28T07:54:14Z","title":"Extract Information from Hybrid Long Documents Leveraging LLMs: A Framework and Dataset","version":2},"reference_index":19,"source":"pdf_text","source_observed_at":"2026-08-10T23:39:19.555628Z"},"links":{"citing_paper":"/paper/2412.20072"},"observation_digest":"sha256:ca5f1f8c9f869b1ec13abe713419cf06eb18a483913d70a72568aaf0397034e5","observation_id":"aa8ebf67-0df3-4e3b-a29d-95afb2960322","resolution":{"observed_at":"2026-08-10T23:39:19.869566Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-12T06:34:41.77262+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-12T06:34:41.77262+00:00","source":"crossref"},{"observed_at":"2026-08-12T06:34:36.333875+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2109.00122","last_updated":"2022-05-07T07:52:39Z","snapshot_observed_at":"2026-08-11T03:07:17.578573Z","submitted_at":"2021-09-01T00:08:14Z","title":"FinQA: A Dataset of Numerical Reasoning over Financial Data","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2109.00122","snapshot_observed_at":"2026-08-10T23:39:19.557580Z","title":"Finqa: A dataset of numerical reasoning over financial data,","venue":null,"work_id":null,"year":2021},"citing_paper":{"arxiv_id":"2412.20072","last_updated":"2024-12-31T03:11:03Z","snapshot_observed_at":"2026-08-12T19:56:06.065692Z","submitted_at":"2024-12-28T07:54:14Z","title":"Extract Information from Hybrid Long Documents Leveraging LLMs: A Framework and Dataset","version":2},"reference_index":20,"source":"pdf_text","source_observed_at":"2026-08-10T23:39:19.557580Z"},"links":{"cited_paper":"/paper/2109.00122","citing_paper":"/paper/2412.20072"},"observation_digest":"sha256:ca224545fb4635cdc49ada8bbfffa462c9f7e02728dbcfa71242af66cdb023b5","observation_id":"ce94a09e-b067-4171-93f0-6852d8c8870b","resolution":{"observed_at":"2026-08-10T23:39:19.557580Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-10T23:39:19.856151Z","title":"Tat-qa: A question answering benchmark on a hybrid of tabular and textual content in finance,","venue":null,"work_id":"f508877b-275b-4e68-9acb-ffdd7c586802","year":2021},"citing_paper":{"arxiv_id":"2412.20072","last_updated":"2024-12-31T03:11:03Z","snapshot_observed_at":"2026-08-12T19:56:06.065692Z","submitted_at":"2024-12-28T07:54:14Z","title":"Extract Information from Hybrid Long Documents Leveraging LLMs: A Framework and Dataset","version":2},"reference_index":21,"source":"pdf_text","source_observed_at":"2026-08-10T23:39:19.560010Z"},"links":{"citing_paper":"/paper/2412.20072"},"observation_digest":"sha256:7196056e08d33fd679db084bd195f1ebe394302cc58829f5cf28df0e95d2a170","observation_id":"f9e22774-70c4-49bc-9842-d416c9122ce0","resolution":{"observed_at":"2026-08-10T23:39:19.860009Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-12T06:34:41.77262+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-12T06:34:41.77262+00:00","source":"crossref"},{"observed_at":"2026-08-12T06:34:36.333875+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-10T23:39:19.845982Z","title":"MultiHiertt: Numerical reasoning over multi hierarchical tabular and textual data,","venue":null,"work_id":"57149fd1-ea70-4c4d-8fc1-1fa19cb51910","year":2022},"citing_paper":{"arxiv_id":"2412.20072","last_updated":"2024-12-31T03:11:03Z","snapshot_observed_at":"2026-08-12T19:56:06.065692Z","submitted_at":"2024-12-28T07:54:14Z","title":"Extract Information from Hybrid Long Documents Leveraging LLMs: A Framework and Dataset","version":2},"reference_index":22,"source":"pdf_text","source_observed_at":"2026-08-10T23:39:19.562114Z"},"links":{"citing_paper":"/paper/2412.20072"},"observation_digest":"sha256:ca8b2c534f0051add61f6779391b0242355b1c6003e372800f785f9b5b3c1d7f","observation_id":"44ffdda4-0ec8-42f6-8702-54aa67fcf80d","resolution":{"observed_at":"2026-08-10T23:39:19.849888Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-12T06:34:41.77262+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-12T06:34:41.77262+00:00","source":"crossref"},{"observed_at":"2026-08-12T06:34:36.333875+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2304.13343","last_updated":"2025-03-18T02:16:56Z","snapshot_observed_at":"2026-08-01T15:59:14.413948Z","submitted_at":"2023-04-26T07:25:31Z","title":"SCM: Enhancing Large Language Model with Self-Controlled Memory Framework","version":4},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2304.13343","snapshot_observed_at":"2026-08-10T23:39:19.564217Z","title":"Unleashing infinite-length input capacity for large-scale lan- guage models with self-controlled memory system,","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2412.20072","last_updated":"2024-12-31T03:11:03Z","snapshot_observed_at":"2026-08-12T19:56:06.065692Z","submitted_at":"2024-12-28T07:54:14Z","title":"Extract Information from Hybrid Long Documents Leveraging LLMs: A Framework and Dataset","version":2},"reference_index":23,"source":"pdf_text","source_observed_at":"2026-08-10T23:39:19.564217Z"},"links":{"cited_paper":"/paper/2304.13343","citing_paper":"/paper/2412.20072"},"observation_digest":"sha256:3e82520b45e102ca4552fb970ca844ada47c5a6610f64be916db9c12789a6163","observation_id":"a5bb3f55-e65c-4042-bab2-6692b8d22e06","resolution":{"observed_at":"2026-08-10T23:39:19.564217Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2304.11633","last_updated":"2023-04-23T12:33:18Z","snapshot_observed_at":"2026-08-06T18:18:16.581812Z","submitted_at":"2023-04-23T12:33:18Z","title":"Evaluating ChatGPT's Information Extraction Capabilities: An Assessment of Performance, Explainability, Calibration, and Faithfulness","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2304.11633","snapshot_observed_at":"2026-08-10T23:39:19.566466Z","title":"Eval- uating chatgpt’s information extraction capabilities: An assessment of performance, explainability, calibration, and faithfulness,","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2412.20072","last_updated":"2024-12-31T03:11:03Z","snapshot_observed_at":"2026-08-12T19:56:06.065692Z","submitted_at":"2024-12-28T07:54:14Z","title":"Extract Information from Hybrid Long Documents Leveraging LLMs: A Framework and Dataset","version":2},"reference_index":24,"source":"pdf_text","source_observed_at":"2026-08-10T23:39:19.566466Z"},"links":{"cited_paper":"/paper/2304.11633","citing_paper":"/paper/2412.20072"},"observation_digest":"sha256:d8c9636f233211f99e1d4fafb60902557e894639c8299df94b52ed78fa0d9da3","observation_id":"9bbbebae-fa20-4561-9893-738aa7e5bd26","resolution":{"observed_at":"2026-08-10T23:39:19.566466Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2302.10205","last_updated":"2024-05-27T07:08:31Z","snapshot_observed_at":"2026-08-11T07:39:47.211296Z","submitted_at":"2023-02-20T12:57:12Z","title":"ChatIE: Zero-Shot Information Extraction via Chatting with ChatGPT","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2302.10205","snapshot_observed_at":"2026-08-10T23:39:19.568872Z","title":"Zero-shot information extraction via chatting with chatgpt,","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2412.20072","last_updated":"2024-12-31T03:11:03Z","snapshot_observed_at":"2026-08-12T19:56:06.065692Z","submitted_at":"2024-12-28T07:54:14Z","title":"Extract Information from Hybrid Long Documents Leveraging LLMs: A Framework and Dataset","version":2},"reference_index":25,"source":"pdf_text","source_observed_at":"2026-08-10T23:39:19.568872Z"},"links":{"cited_paper":"/paper/2302.10205","citing_paper":"/paper/2412.20072"},"observation_digest":"sha256:dc7c32a8763056c21aac80e51c4920b67d77fcfa6e06c3ccfc611acd21b4a400","observation_id":"9ad11bb1-754f-4f02-987e-42ee6499f746","resolution":{"observed_at":"2026-08-10T23:39:19.568872Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2109.08079","last_updated":"2023-06-08T18:33:01Z","snapshot_observed_at":"2026-08-12T16:26:25.821774Z","submitted_at":"2021-09-16T16:10:05Z","title":"Context-NER : Contextual Phrase Generation at Scale","version":4},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2109.08079","snapshot_observed_at":"2026-08-10T23:39:19.571583Z","title":"Context-ner: Contextual phrase generation at scale,","venue":null,"work_id":null,"year":2021},"citing_paper":{"arxiv_id":"2412.20072","last_updated":"2024-12-31T03:11:03Z","snapshot_observed_at":"2026-08-12T19:56:06.065692Z","submitted_at":"2024-12-28T07:54:14Z","title":"Extract Information from Hybrid Long Documents Leveraging LLMs: A Framework and Dataset","version":2},"reference_index":26,"source":"pdf_text","source_observed_at":"2026-08-10T23:39:19.571583Z"},"links":{"cited_paper":"/paper/2109.08079","citing_paper":"/paper/2412.20072"},"observation_digest":"sha256:d88098707f76c647c74ac9ee4bfb33e1d74bee0daaf469df5e8ffbede785c116","observation_id":"9eadba57-a521-4947-ae24-d3737dd05c92","resolution":{"observed_at":"2026-08-10T23:39:19.571583Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2304.10428","last_updated":"2023-10-07T14:25:28Z","snapshot_observed_at":"2026-07-06T15:17:59.640907Z","submitted_at":"2023-04-20T16:17:26Z","title":"GPT-NER: Named Entity Recognition via Large Language Models","version":4},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2304.10428","snapshot_observed_at":"2026-08-10T23:39:19.574325Z","title":"Gpt-ner: Named entity recognition via large language models,","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2412.20072","last_updated":"2024-12-31T03:11:03Z","snapshot_observed_at":"2026-08-12T19:56:06.065692Z","submitted_at":"2024-12-28T07:54:14Z","title":"Extract Information from Hybrid Long Documents Leveraging LLMs: A Framework and Dataset","version":2},"reference_index":27,"source":"pdf_text","source_observed_at":"2026-08-10T23:39:19.574325Z"},"links":{"cited_paper":"/paper/2304.10428","citing_paper":"/paper/2412.20072"},"observation_digest":"sha256:2ca449b0bcb36be3c0e86c135cd3c6454a2a6fe2b36173cc6d4ccd682aed4ee8","observation_id":"82d2636a-615b-4296-99be-38b32b328ec0","resolution":{"observed_at":"2026-08-10T23:39:19.574325Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2305.02105","last_updated":"2023-12-09T02:05:05Z","snapshot_observed_at":"2026-08-09T16:17:13.712329Z","submitted_at":"2023-05-03T13:28:08Z","title":"GPT-RE: In-context Learning for Relation Extraction using Large Language Models","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2305.02105","snapshot_observed_at":"2026-08-10T23:39:19.577042Z","title":"Gpt-re: In-context learning for relation extraction using large language models,","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2412.20072","last_updated":"2024-12-31T03:11:03Z","snapshot_observed_at":"2026-08-12T19:56:06.065692Z","submitted_at":"2024-12-28T07:54:14Z","title":"Extract Information from Hybrid Long Documents Leveraging LLMs: A Framework and Dataset","version":2},"reference_index":28,"source":"pdf_text","source_observed_at":"2026-08-10T23:39:19.577042Z"},"links":{"cited_paper":"/paper/2305.02105","citing_paper":"/paper/2412.20072"},"observation_digest":"sha256:75f703976fb3ff924035eae2dace32205f1ba36c931b633f8d9212b0ece837b4","observation_id":"28bc642d-bf07-49db-ace0-6dda9e42422f","resolution":{"observed_at":"2026-08-10T23:39:19.577042Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2305.01555","last_updated":"2023-06-09T15:59:18Z","snapshot_observed_at":"2026-07-06T15:22:25.008777Z","submitted_at":"2023-05-02T15:55:41Z","title":"How to Unleash the Power of Large Language Models for Few-shot Relation Extraction?","version":4},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2305.01555","snapshot_observed_at":"2026-08-10T23:39:19.580806Z","title":"How to unleash the power of large language models for few-shot relation extraction?","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2412.20072","last_updated":"2024-12-31T03:11:03Z","snapshot_observed_at":"2026-08-12T19:56:06.065692Z","submitted_at":"2024-12-28T07:54:14Z","title":"Extract Information from Hybrid Long Documents Leveraging LLMs: A Framework and Dataset","version":2},"reference_index":29,"source":"pdf_text","source_observed_at":"2026-08-10T23:39:19.580806Z"},"links":{"cited_paper":"/paper/2305.01555","citing_paper":"/paper/2412.20072"},"observation_digest":"sha256:422b16a444eb7ca12b1687d1058ea6cf010196d9aa0c11f048cd187108dfc2fb","observation_id":"698b14ef-5203-40ab-aed5-0a0c56c32e50","resolution":{"observed_at":"2026-08-10T23:39:19.580806Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2305.03513","last_updated":"2023-09-19T14:26:17Z","snapshot_observed_at":"2026-08-11T11:47:25.195630Z","submitted_at":"2023-05-03T19:57:43Z","title":"ChatGraph: Interpretable Text Classification by Converting ChatGPT Knowledge to Graphs","version":2},"cited_work":{"arxiv_id":"2305.03513","doi":null,"metadata_source":"pith","pith_arxiv_id":"2305.03513","snapshot_observed_at":"2026-08-10T23:39:19.622619Z","title":"ChatGraph: Interpretable Text Classification by Converting ChatGPT Knowledge to Graphs","venue":"cs.CL","work_id":"315ab1bd-8a3a-48bb-884d-08173bfd4fe9","year":2023},"citing_paper":{"arxiv_id":"2412.20072","last_updated":"2024-12-31T03:11:03Z","snapshot_observed_at":"2026-08-12T19:56:06.065692Z","submitted_at":"2024-12-28T07:54:14Z","title":"Extract Information from Hybrid Long Documents Leveraging LLMs: A Framework and Dataset","version":2},"reference_index":30,"source":"pdf_text","source_observed_at":"2026-08-10T23:39:19.584869Z"},"links":{"cited_paper":"/paper/2305.03513","citing_paper":"/paper/2412.20072"},"observation_digest":"sha256:004929e1a5f31300f91b818c663e1e8ae1621ce1ad05aacd69663a9696cca69d","observation_id":"3559fbe1-1a83-4bda-905d-55790409b727","resolution":{"observed_at":"2026-08-10T23:39:19.630303Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-12T06:34:41.77262+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-12T06:34:41.77262+00:00","source":"crossref"},{"observed_at":"2026-08-12T06:34:36.333875+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2304.09433","last_updated":"2025-03-07T17:33:50Z","snapshot_observed_at":"2026-08-10T16:50:54.174900Z","submitted_at":"2023-04-19T06:00:26Z","title":"Language Models Enable Simple Systems for Generating Structured Views of Heterogeneous Data Lakes","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2304.09433","snapshot_observed_at":"2026-08-10T23:39:19.587686Z","title":"Language models enable simple systems for gen- erating structured views of heterogeneous data lakes,","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2412.20072","last_updated":"2024-12-31T03:11:03Z","snapshot_observed_at":"2026-08-12T19:56:06.065692Z","submitted_at":"2024-12-28T07:54:14Z","title":"Extract Information from Hybrid Long Documents Leveraging LLMs: A Framework and Dataset","version":2},"reference_index":31,"source":"pdf_text","source_observed_at":"2026-08-10T23:39:19.587686Z"},"links":{"cited_paper":"/paper/2304.09433","citing_paper":"/paper/2412.20072"},"observation_digest":"sha256:bb2f74409911eaf31d20e5b80117334e2148b166a912a0918ea388d864a51a1f","observation_id":"2da0e814-ec99-4e43-85c2-eb1d9e252308","resolution":{"observed_at":"2026-08-10T23:39:19.587686Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}}],"paper":{"arxiv_id":"2412.20072","last_updated":"2024-12-31T03:11:03Z","latest_version":2,"primary_category":"cs.CL","snapshot_observed_at":"2026-08-12T19:56:06.065692Z","submitted_at":"2024-12-28T07:54:14Z","title":"Extract Information from Hybrid Long Documents Leveraging LLMs: A Framework and Dataset"},"reference_resolution":{"displayed":31,"state_counts":{"malformed_identifier":0,"metadata_mismatch":0,"parse_uncertain":0,"unresolved":18,"verified_exact":3,"verified_fuzzy":10},"total_outbound_references":31},"refusal":"A citation records a reference. It does not transfer a finding from one paper to another.","schema":"pith.paper-citation-record.v1","standing_sources":[{"observed_at":"2026-08-12T06:34:41.77262+00:00","source":"crossref"},{"observed_at":"2026-08-12T06:34:36.333875+00:00","source":"retraction_watch"}],"thesis":"As of 13 August 2026, this Paper Citation Record lists 31 of 31 outbound references and 0 inbound Pith citation observations for arXiv:2412.20072."}