{"as_of":"2026-08-08T11:03:00Z","caps":{"database_statements":6,"inbound":100,"outbound":100},"context_digest":"sha256:8448ae265e03b05dbb2dfb580df2c727b66be0aa249e11d514970e36ab4430fd","coverage":[{"denominator":0,"lane":"reference_resolution","note":"Typed states for the displayed outbound observations.","records_observed":0,"source":"paper_references, paper_reference_links","source_observed_at":null,"state":"measured"},{"denominator":5,"lane":"standing_notices","note":"One-hop event checks from named stored sources.","records_observed":5,"source":"scholarly_work_events, retraction_status_cache","source_observed_at":"2026-08-08T06:32:00.761636+00:00","state":"measured"},{"denominator":5,"lane":"inbound_itemization","note":"Pith citing papers itemized under the disclosed page cap.","records_observed":5,"source":"paper_references, paper_reference_links","source_observed_at":"2026-08-07T14:47:15.798080Z","state":"measured"},{"denominator":1,"lane":"external_citation_measurements","note":"A source-named dated measurement, never combined with another source.","records_observed":0,"source":"arxiv_reference","source_observed_at":"2026-06-29T19:13:53.117583Z","state":"measured"}],"external_citation_measurements":[],"inbound":[{"citation":{"cited_paper":{"arxiv_id":"2306.08952","last_updated":"2023-06-27T05:39:25Z","snapshot_observed_at":"2026-07-06T15:42:52.549767Z","submitted_at":"2023-06-15T08:44:41Z","title":"Towards Benchmarking and Improving the Temporal Reasoning Capability of Large Language Models","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2306.08952","snapshot_observed_at":"2026-08-07T14:47:15.798080Z","title":"Towards benchmarking and improving the temporal reasoning capability of large language models.arXiv preprint arXiv:2306.08952, 2023","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2505.17572","last_updated":"2025-05-23T07:30:57Z","snapshot_observed_at":"2026-08-07T21:56:03.875890Z","submitted_at":"2025-05-23T07:30:57Z","title":"USTBench: Benchmarking and Dissecting Spatiotemporal Reasoning of LLMs as Urban Agents","version":1},"reference_index":42,"source":"pdf_text","source_observed_at":"2026-08-07T14:47:15.798080Z"},"links":{"cited_paper":"/paper/2306.08952","citing_paper":"/paper/2505.17572"},"observation_digest":"sha256:6a0be7329ae169c94ea66fd1aa670ee7942540f6c0b86274dea3b996550f204d","observation_id":"c251b9f7-f9ca-43fe-b35f-5a38e669c87a","resolution":{"observed_at":"2026-08-07T14:47:15.798080Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2306.08952","last_updated":"2023-06-27T05:39:25Z","snapshot_observed_at":"2026-07-06T15:42:52.549767Z","submitted_at":"2023-06-15T08:44:41Z","title":"Towards Benchmarking and Improving the Temporal Reasoning Capability of Large Language Models","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2306.08952","snapshot_observed_at":"2026-08-07T14:17:09.966209Z","title":null,"venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2505.19533","last_updated":"2025-05-26T05:39:57Z","snapshot_observed_at":"2026-08-08T00:00:45.680749Z","submitted_at":"2025-05-26T05:39:57Z","title":"ExAnte: A Benchmark for Ex-Ante Inference in Large Language Models","version":1},"reference_index":33,"source":"pdf_text","source_observed_at":"2026-08-07T14:17:09.966209Z"},"links":{"cited_paper":"/paper/2306.08952","citing_paper":"/paper/2505.19533"},"observation_digest":"sha256:ae6b9a5c98c441ee2bb70874679a6bbc1790cc777958d510d6e7b07c3c99b3bc","observation_id":"c5f5e0c8-c4f8-4afa-aa57-d465439c2fc3","resolution":{"observed_at":"2026-08-07T14:17:09.966209Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2306.08952","last_updated":"2023-06-27T05:39:25Z","snapshot_observed_at":"2026-07-06T15:42:52.549767Z","submitted_at":"2023-06-15T08:44:41Z","title":"Towards Benchmarking and Improving the Temporal Reasoning Capability of Large Language Models","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2306.08952","snapshot_observed_at":"2026-08-04T07:01:42.227326Z","title":"T., and Bing, L.Towards benchmarking and improving the temporal reasoning capability of large language models.arXiv preprint arXiv:2306.08952(2023)","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2510.27544","last_updated":"2026-06-06T10:28:21Z","snapshot_observed_at":"2026-08-07T09:42:51.433478Z","submitted_at":"2025-10-31T15:17:55Z","title":"TempoBench: Evaluating Temporal Causal Reasoning in Large Language Models","version":2},"reference_index":44,"source":"pdf_text","source_observed_at":"2026-08-04T07:01:42.227326Z"},"links":{"cited_paper":"/paper/2306.08952","citing_paper":"/paper/2510.27544"},"observation_digest":"sha256:8783472360077cc768bb2d5dcf39e3551b11b17598c5d1581e7b0e540a147549","observation_id":"325538cb-e4b1-437b-94f5-1592e4f6cf6d","resolution":{"observed_at":"2026-08-04T07:01:42.227326Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2306.08952","last_updated":"2023-06-27T05:39:25Z","snapshot_observed_at":"2026-07-06T15:42:52.549767Z","submitted_at":"2023-06-15T08:44:41Z","title":"Towards Benchmarking and Improving the Temporal Reasoning Capability of Large Language Models","version":2},"cited_work":{"arxiv_id":"2306.08952","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2306.08952","snapshot_observed_at":"2026-06-29T19:13:53.117583Z","title":"Towards benchmarking and improving the temporal reasoning capability of large language models","venue":null,"work_id":"60925d02-5597-445d-bcfd-2746efda6a5b","year":2023},"citing_paper":{"arxiv_id":"2605.04243","last_updated":"2026-05-05T19:30:06Z","snapshot_observed_at":"2026-07-06T23:17:04.200299Z","submitted_at":"2026-05-05T19:30:06Z","title":"Temporal Reasoning Is Not the Bottleneck: A Probabilistic Inconsistency Framework for Neuro-Symbolic QA","version":1},"reference_index":80,"source":"pdf_text","source_observed_at":"2026-05-08T17:45:44.270122Z"},"links":{"cited_paper":"/paper/2306.08952","citing_paper":"/paper/2605.04243"},"observation_digest":"sha256:e33297bfbb20b5d90209c18ee3019d7ed39e52280861502c3cf8a80931173a85","observation_id":"216b4c58-b3ed-456f-b19b-bcadb3381aac","resolution":{"observed_at":"2026-05-11T17:16:07.890653Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2306.08952","last_updated":"2023-06-27T05:39:25Z","snapshot_observed_at":"2026-07-06T15:42:52.549767Z","submitted_at":"2023-06-15T08:44:41Z","title":"Towards Benchmarking and Improving the Temporal Reasoning Capability of Large Language Models","version":2},"cited_work":{"arxiv_id":"2306.08952","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2306.08952","snapshot_observed_at":"2026-06-29T19:13:53.117583Z","title":"Towards benchmarking and improving the temporal reasoning capability of large language models","venue":null,"work_id":"60925d02-5597-445d-bcfd-2746efda6a5b","year":2023},"citing_paper":{"arxiv_id":"2606.27881","last_updated":"2026-06-26T09:23:31Z","snapshot_observed_at":"2026-07-07T00:02:04.432452Z","submitted_at":"2026-06-26T09:23:31Z","title":"A Study of Temporal Fusion Strategies for Named Entity Recognition in Historical Texts","version":1},"reference_index":35,"source":"pdf_text","source_observed_at":"2026-06-29T04:52:30.770882Z"},"links":{"cited_paper":"/paper/2306.08952","citing_paper":"/paper/2606.27881"},"observation_digest":"sha256:f5e32a04028b973c64d65058e673ab7c1fef668983355b11d6d4cb9d161837bf","observation_id":"6c8d4a49-3c3d-47e3-b4dd-977a673d21b0","resolution":{"observed_at":"2026-06-29T19:13:53.119924Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}}],"links":{"evidence":"/evidence","html":"/paper/2306.08952/citation-record","integrity":"/paper/2306.08952/integrity","json":"/paper/2306.08952/citation-record.json","paper":"/paper/2306.08952"},"outbound":[],"paper":{"arxiv_id":"2306.08952","last_updated":"2023-06-27T05:39:25Z","latest_version":2,"primary_category":"cs.CL","snapshot_observed_at":"2026-07-06T15:42:52.549767Z","submitted_at":"2023-06-15T08:44:41Z","title":"Towards Benchmarking and Improving the Temporal Reasoning Capability of Large Language Models"},"reference_resolution":{"displayed":0,"state_counts":{"malformed_identifier":0,"metadata_mismatch":0,"parse_uncertain":0,"unresolved":0,"verified_exact":0,"verified_fuzzy":0},"total_outbound_references":0},"refusal":"A citation records a reference. It does not transfer a finding from one paper to another.","schema":"pith.paper-citation-record.v1","standing_sources":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"thesis":"As of 8 August 2026, this Paper Citation Record lists 0 of 0 outbound references and 5 inbound Pith citation observations for arXiv:2306.08952."}