{"as_of":"2026-08-05T20:15:00Z","caps":{"database_statements":6,"inbound":100,"outbound":100},"context_digest":"sha256:6622eecb7687df6bcd4d4dfa5d92120bdbcbff01f92e06c7d8f2121600ce03bf","coverage":[{"denominator":0,"lane":"reference_resolution","note":"Typed states for the displayed outbound observations.","records_observed":0,"source":"paper_references, paper_reference_links","source_observed_at":null,"state":"measured"},{"denominator":20,"lane":"standing_notices","note":"One-hop event checks from named stored sources.","records_observed":20,"source":"scholarly_work_events, retraction_status_cache","source_observed_at":"2026-08-05T06:32:48.257954+00:00","state":"measured"},{"denominator":20,"lane":"inbound_itemization","note":"Pith citing papers itemized under the disclosed page cap.","records_observed":20,"source":"paper_references, paper_reference_links","source_observed_at":"2026-08-03T09:30:02.970530Z","state":"measured"},{"denominator":1,"lane":"external_citation_measurements","note":"A source-named dated measurement, never combined with another source.","records_observed":1,"source":"arxiv_reference","source_observed_at":"2026-08-05T02:28:24.338817Z","state":"measured"}],"external_citation_measurements":[{"count":1,"observed_at":"2026-08-05T02:28:24.338817Z","source":"arxiv_reference"}],"inbound":[{"citation":{"cited_paper":{"arxiv_id":"2508.11987","last_updated":"2025-09-05T09:15:55Z","snapshot_observed_at":"2026-08-05T19:39:12.827378Z","submitted_at":"2025-08-16T08:54:08Z","title":"FutureX: An Advanced Live Benchmark for LLM Agents in Future Prediction","version":3},"cited_work":{"arxiv_id":"2508.11987","doi":"10.48550/arxiv.2508.11987","metadata_source":"arxiv_reference","pith_arxiv_id":"2508.11987","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Futurex: An advanced live benchmark for llm agents in future prediction","venue":"ArXiv.org","work_id":"ac497801-1ebb-45ae-8286-d566aec17d80","year":2025},"citing_paper":{"arxiv_id":"2512.00417","last_updated":"2026-05-15T16:13:36Z","snapshot_observed_at":"2026-07-06T22:37:16.933664Z","submitted_at":"2025-11-29T09:52:34Z","title":"CryptoBench: A Dynamic Benchmark for Expert-Level Evaluation of LLM Agents in Cryptocurrency","version":5},"reference_index":10,"source":"pdf_text","source_observed_at":"2026-05-21T19:04:07.564999Z"},"links":{"cited_paper":"/paper/2508.11987","citing_paper":"/paper/2512.00417"},"observation_digest":"sha256:5f4d0885cfa05baee270f32950fdd6122de97ab44678585298e6e9aac1eecfba","observation_id":"2c6923c7-26be-4d64-987e-d9bbd48a801e","resolution":{"observed_at":"2026-05-21T19:04:19.002830Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2508.11987","last_updated":"2025-09-05T09:15:55Z","snapshot_observed_at":"2026-08-05T19:39:12.827378Z","submitted_at":"2025-08-16T08:54:08Z","title":"FutureX: An Advanced Live Benchmark for LLM Agents in Future Prediction","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2508.11987","snapshot_observed_at":"2026-08-03T09:30:02.970530Z","title":"Futurex: An advanced live benchmark for LLM agents in future prediction","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2601.13836","last_updated":"2026-06-17T10:28:13Z","snapshot_observed_at":"2026-08-03T17:07:22.325179Z","submitted_at":"2026-01-20T10:47:20Z","title":"FutureOmni: Evaluating Future Forecasting from Omni-Modal Context for Multimodal LLMs","version":2},"reference_index":48,"source":"pdf_text","source_observed_at":"2026-08-03T09:30:02.970530Z"},"links":{"cited_paper":"/paper/2508.11987","citing_paper":"/paper/2601.13836"},"observation_digest":"sha256:0ffdf7f22822724fb68f438f2c8c64cd955106ffda2249bcdf412b7ec17eb7d1","observation_id":"9423e84a-2c86-49ad-8a40-bf144231d01e","resolution":{"observed_at":"2026-08-03T09:30:02.970530Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2508.11987","last_updated":"2025-09-05T09:15:55Z","snapshot_observed_at":"2026-08-05T19:39:12.827378Z","submitted_at":"2025-08-16T08:54:08Z","title":"FutureX: An Advanced Live Benchmark for LLM Agents in Future Prediction","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2508.11987","snapshot_observed_at":"2026-08-03T02:44:07.502114Z","title":"Futurex: An advanced live benchmark for llm agents in future prediction.arXiv preprint arXiv:2508.11987, 2025","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2602.18481","last_updated":"2026-05-27T14:41:33Z","snapshot_observed_at":"2026-08-03T07:13:36.244963Z","submitted_at":"2026-02-10T14:29:33Z","title":"AlphaForgeBench: Benchmarking End-to-End Trading Strategy Design with Large Language Models","version":2},"reference_index":19,"source":"pdf_text","source_observed_at":"2026-08-03T02:44:07.502114Z"},"links":{"cited_paper":"/paper/2508.11987","citing_paper":"/paper/2602.18481"},"observation_digest":"sha256:1ff5aea41ba52453c6d0a05eb107d6c3ddd71fee7a61e83396a08bf3b4e947d3","observation_id":"196c64bf-89fb-413e-b28c-2b56b81bc7d9","resolution":{"observed_at":"2026-08-03T02:44:07.502114Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2508.11987","last_updated":"2025-09-05T09:15:55Z","snapshot_observed_at":"2026-08-05T19:39:12.827378Z","submitted_at":"2025-08-16T08:54:08Z","title":"FutureX: An Advanced Live Benchmark for LLM Agents in Future Prediction","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2508.11987","snapshot_observed_at":"2026-08-03T02:57:48.014062Z","title":null,"venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2603.03322","last_updated":"2026-07-31T08:12:53Z","snapshot_observed_at":"2026-08-05T20:15:18.195512Z","submitted_at":"2026-02-10T05:47:22Z","title":"Can Large Language Models Derive New Knowledge? A Dynamic Benchmark for Biological Knowledge Discovery","version":2},"reference_index":34,"source":"pdf_text","source_observed_at":"2026-08-03T02:57:48.014062Z"},"links":{"cited_paper":"/paper/2508.11987","citing_paper":"/paper/2603.03322"},"observation_digest":"sha256:14955f739bdf8b7fddf89471551d4b9f8c65ae4cafa8a9da8e7ceb0e625ebab7","observation_id":"4086e844-9e3c-4fcb-a18a-802f58fa529d","resolution":{"observed_at":"2026-08-03T02:57:48.014062Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2508.11987","last_updated":"2025-09-05T09:15:55Z","snapshot_observed_at":"2026-08-05T19:39:12.827378Z","submitted_at":"2025-08-16T08:54:08Z","title":"FutureX: An Advanced Live Benchmark for LLM Agents in Future Prediction","version":3},"cited_work":{"arxiv_id":"2508.11987","doi":"10.48550/arxiv.2508.11987","metadata_source":"arxiv_reference","pith_arxiv_id":"2508.11987","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Futurex: An advanced live benchmark for llm agents in future prediction","venue":"ArXiv.org","work_id":"ac497801-1ebb-45ae-8286-d566aec17d80","year":2025},"citing_paper":{"arxiv_id":"2604.14199","last_updated":"2026-04-03T06:25:21Z","snapshot_observed_at":"2026-08-02T23:48:38.385069Z","submitted_at":"2026-04-03T06:25:21Z","title":"PolyBench: Benchmarking LLM Forecasting and Trading Capabilities on Live Prediction Market Data","version":1},"reference_index":24,"source":"pdf_text","source_observed_at":"2026-05-13T19:17:22.608745Z"},"links":{"cited_paper":"/paper/2508.11987","citing_paper":"/paper/2604.14199"},"observation_digest":"sha256:b07e6b44f6d875892efbbf3780ab14af618e663f602cfa45061ff3eacb83e5f9","observation_id":"1ed7b871-1ccc-472f-a463-cf997e854763","resolution":{"observed_at":"2026-05-13T19:18:09.288034Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2508.11987","last_updated":"2025-09-05T09:15:55Z","snapshot_observed_at":"2026-08-05T19:39:12.827378Z","submitted_at":"2025-08-16T08:54:08Z","title":"FutureX: An Advanced Live Benchmark for LLM Agents in Future Prediction","version":3},"cited_work":{"arxiv_id":"2508.11987","doi":"10.48550/arxiv.2508.11987","metadata_source":"arxiv_reference","pith_arxiv_id":"2508.11987","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Futurex: An advanced live benchmark for llm agents in future prediction","venue":"ArXiv.org","work_id":"ac497801-1ebb-45ae-8286-d566aec17d80","year":2025},"citing_paper":{"arxiv_id":"2604.15719","last_updated":"2026-05-08T15:22:29Z","snapshot_observed_at":"2026-08-02T09:43:01.889129Z","submitted_at":"2026-04-17T05:43:07Z","title":"Harnessing Pre-Resolution Signals for Future Prediction Agents","version":2},"reference_index":13,"source":"pdf_text","source_observed_at":"2026-05-10T09:01:03.441166Z"},"links":{"cited_paper":"/paper/2508.11987","citing_paper":"/paper/2604.15719"},"observation_digest":"sha256:253b1e256260123c242bb639ada5cda891782b5ff5a0f10d9150fcc9e1bc4d44","observation_id":"48732151-1919-4c5f-8780-dd801c8c8892","resolution":{"observed_at":"2026-05-10T09:03:25.199111Z","resolver_source":"arxiv_id","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2508.11987","last_updated":"2025-09-05T09:15:55Z","snapshot_observed_at":"2026-08-05T19:39:12.827378Z","submitted_at":"2025-08-16T08:54:08Z","title":"FutureX: An Advanced Live Benchmark for LLM Agents in Future Prediction","version":3},"cited_work":{"arxiv_id":"2508.11987","doi":"10.48550/arxiv.2508.11987","metadata_source":"arxiv_reference","pith_arxiv_id":"2508.11987","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Futurex: An advanced live benchmark for llm agents in future prediction","venue":"ArXiv.org","work_id":"ac497801-1ebb-45ae-8286-d566aec17d80","year":2025},"citing_paper":{"arxiv_id":"2604.15719","last_updated":"2026-05-08T15:22:29Z","snapshot_observed_at":"2026-08-02T09:43:01.889129Z","submitted_at":"2026-04-17T05:43:07Z","title":"Harnessing Pre-Resolution Signals for Future Prediction Agents","version":3},"reference_index":13,"source":"pdf_text","source_observed_at":"2026-05-11T02:01:05.101225Z"},"links":{"cited_paper":"/paper/2508.11987","citing_paper":"/paper/2604.15719"},"observation_digest":"sha256:60ba040bb552e3bf8701e1fa0f93d37d6723dc6adeaca58e17d92fcd08640db5","observation_id":"556baddc-4576-4710-9977-bf4e6dec3c1d","resolution":{"observed_at":"2026-05-11T04:00:56.119896Z","resolver_source":"arxiv_id","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2508.11987","last_updated":"2025-09-05T09:15:55Z","snapshot_observed_at":"2026-08-05T19:39:12.827378Z","submitted_at":"2025-08-16T08:54:08Z","title":"FutureX: An Advanced Live Benchmark for LLM Agents in Future Prediction","version":3},"cited_work":{"arxiv_id":"2508.11987","doi":"10.48550/arxiv.2508.11987","metadata_source":"arxiv_reference","pith_arxiv_id":"2508.11987","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Futurex: An advanced live benchmark for llm agents in future prediction","venue":"ArXiv.org","work_id":"ac497801-1ebb-45ae-8286-d566aec17d80","year":2025},"citing_paper":{"arxiv_id":"2604.18576","last_updated":"2026-07-12T18:34:47Z","snapshot_observed_at":"2026-07-16T23:18:42.403825Z","submitted_at":"2026-04-20T17:57:51Z","title":"Agentic Forecasting using Sequential Bayesian Updating of Linguistic Beliefs","version":3},"reference_index":49,"source":"arxiv_source","source_observed_at":"2026-05-10T04:14:17.395227Z"},"links":{"cited_paper":"/paper/2508.11987","citing_paper":"/paper/2604.18576"},"observation_digest":"sha256:4d7da41b0f9a5a13a0663e5b537f1dd807beb3e11fd05535f8549807c71e8d1d","observation_id":"d91d3864-7724-4c92-a51f-c4df51d6db85","resolution":{"observed_at":"2026-05-11T12:06:01.774340Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2508.11987","last_updated":"2025-09-05T09:15:55Z","snapshot_observed_at":"2026-08-05T19:39:12.827378Z","submitted_at":"2025-08-16T08:54:08Z","title":"FutureX: An Advanced Live Benchmark for LLM Agents in Future Prediction","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2508.11987","snapshot_observed_at":"2026-07-14T19:33:20.775008Z","title":"FutureX : An advanced live benchmark for LLM agents in future prediction","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2604.18576","last_updated":"2026-07-12T18:34:47Z","snapshot_observed_at":"2026-07-16T23:18:42.403825Z","submitted_at":"2026-04-20T17:57:51Z","title":"Agentic Forecasting using Sequential Bayesian Updating of Linguistic Beliefs","version":4},"reference_index":49,"source":"arxiv_source","source_observed_at":"2026-07-14T19:33:20.775008Z"},"links":{"cited_paper":"/paper/2508.11987","citing_paper":"/paper/2604.18576"},"observation_digest":"sha256:7b446b837654cc83506666b212f24f23601caee025e582790d2fddc8ce53c15d","observation_id":"62d6fa52-153a-4b9f-a7e6-ed961892fe05","resolution":{"observed_at":"2026-07-14T19:33:20.775008Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2508.11987","last_updated":"2025-09-05T09:15:55Z","snapshot_observed_at":"2026-08-05T19:39:12.827378Z","submitted_at":"2025-08-16T08:54:08Z","title":"FutureX: An Advanced Live Benchmark for LLM Agents in Future Prediction","version":3},"cited_work":{"arxiv_id":"2508.11987","doi":"10.48550/arxiv.2508.11987","metadata_source":"arxiv_reference","pith_arxiv_id":"2508.11987","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Futurex: An advanced live benchmark for llm agents in future prediction","venue":"ArXiv.org","work_id":"ac497801-1ebb-45ae-8286-d566aec17d80","year":2025},"citing_paper":{"arxiv_id":"2604.23072","last_updated":"2026-04-24T23:56:53Z","snapshot_observed_at":"2026-07-06T23:09:19.041332Z","submitted_at":"2026-04-24T23:56:53Z","title":"Analytica: Soft Propositional Reasoning for Robust and Scalable LLM-Driven Analysis","version":1},"reference_index":5,"source":"pdf_text","source_observed_at":"2026-05-08T11:35:36.472216Z"},"links":{"cited_paper":"/paper/2508.11987","citing_paper":"/paper/2604.23072"},"observation_digest":"sha256:298d98abd7b9ac7e380ba00d40e64f8265074602e615f0338818eefebbc0bd4d","observation_id":"7f40aaea-1d2a-4ee3-9c25-17e7c5d5a533","resolution":{"observed_at":"2026-05-11T19:36:13.304505Z","resolver_source":"arxiv_id","status":"malformed_identifier"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2508.11987","last_updated":"2025-09-05T09:15:55Z","snapshot_observed_at":"2026-08-05T19:39:12.827378Z","submitted_at":"2025-08-16T08:54:08Z","title":"FutureX: An Advanced Live Benchmark for LLM Agents in Future Prediction","version":3},"cited_work":{"arxiv_id":"2508.11987","doi":"10.48550/arxiv.2508.11987","metadata_source":"arxiv_reference","pith_arxiv_id":"2508.11987","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Futurex: An advanced live benchmark for llm agents in future prediction","venue":"ArXiv.org","work_id":"ac497801-1ebb-45ae-8286-d566aec17d80","year":2025},"citing_paper":{"arxiv_id":"2604.26235","last_updated":"2026-04-29T02:32:14Z","snapshot_observed_at":"2026-08-02T11:16:41.718755Z","submitted_at":"2026-04-29T02:32:14Z","title":"LATTICE: Evaluating Decision Support Utility of Crypto Agents","version":1},"reference_index":5,"source":"pdf_text","source_observed_at":"2026-05-07T13:30:46.523784Z"},"links":{"cited_paper":"/paper/2508.11987","citing_paper":"/paper/2604.26235"},"observation_digest":"sha256:aaabe72c49f9cec9ee78d87c806250c6d1f8af106c4fea79e5f120cf3b5fcaea","observation_id":"955955d4-8963-4547-b622-7ae398e2a2f3","resolution":{"observed_at":"2026-05-12T08:56:25.143219Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2508.11987","last_updated":"2025-09-05T09:15:55Z","snapshot_observed_at":"2026-08-05T19:39:12.827378Z","submitted_at":"2025-08-16T08:54:08Z","title":"FutureX: An Advanced Live Benchmark for LLM Agents in Future Prediction","version":3},"cited_work":{"arxiv_id":"2508.11987","doi":"10.48550/arxiv.2508.11987","metadata_source":"arxiv_reference","pith_arxiv_id":"2508.11987","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Futurex: An advanced live benchmark for llm agents in future prediction","venue":"ArXiv.org","work_id":"ac497801-1ebb-45ae-8286-d566aec17d80","year":2025},"citing_paper":{"arxiv_id":"2604.27865","last_updated":"2026-04-30T13:47:22Z","snapshot_observed_at":"2026-08-01T02:35:05.866123Z","submitted_at":"2026-04-30T13:47:22Z","title":"KellyBench: A Benchmark for Long-Horizon Sequential Decision Making","version":1},"reference_index":4,"source":"pdf_text","source_observed_at":"2026-05-07T05:45:38.267125Z"},"links":{"cited_paper":"/paper/2508.11987","citing_paper":"/paper/2604.27865"},"observation_digest":"sha256:65c2089beb7ca327149ec1f4b685acb9233bc5ba49741e7cb522c596c2ac2339","observation_id":"8db14471-a901-48f8-9ca3-f91bbcee7cea","resolution":{"observed_at":"2026-05-09T05:05:11.975393Z","resolver_source":"arxiv_id","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2508.11987","last_updated":"2025-09-05T09:15:55Z","snapshot_observed_at":"2026-08-05T19:39:12.827378Z","submitted_at":"2025-08-16T08:54:08Z","title":"FutureX: An Advanced Live Benchmark for LLM Agents in Future Prediction","version":3},"cited_work":{"arxiv_id":"2508.11987","doi":"10.48550/arxiv.2508.11987","metadata_source":"arxiv_reference","pith_arxiv_id":"2508.11987","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Futurex: An advanced live benchmark for llm agents in future prediction","venue":"ArXiv.org","work_id":"ac497801-1ebb-45ae-8286-d566aec17d80","year":2025},"citing_paper":{"arxiv_id":"2605.03762","last_updated":"2026-05-05T13:50:50Z","snapshot_observed_at":"2026-07-06T23:16:40.201701Z","submitted_at":"2026-05-05T13:50:50Z","title":"OracleProto: A Reproducible Framework for Benchmarking LLM Native Forecasting via Knowledge Cutoff and Temporal Masking","version":1},"reference_index":14,"source":"pdf_text","source_observed_at":"2026-05-07T16:29:51.187292Z"},"links":{"cited_paper":"/paper/2508.11987","citing_paper":"/paper/2605.03762"},"observation_digest":"sha256:cdeba4c1811362b6f7cd6dbee5f7ab7ede5d70c816724d7de06a0c50e40bdb7f","observation_id":"9c033e61-f46c-49ea-9e12-1086b2fcf378","resolution":{"observed_at":"2026-05-11T23:41:16.903338Z","resolver_source":"arxiv_id","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2508.11987","last_updated":"2025-09-05T09:15:55Z","snapshot_observed_at":"2026-08-05T19:39:12.827378Z","submitted_at":"2025-08-16T08:54:08Z","title":"FutureX: An Advanced Live Benchmark for LLM Agents in Future Prediction","version":3},"cited_work":{"arxiv_id":"2508.11987","doi":"10.48550/arxiv.2508.11987","metadata_source":"arxiv_reference","pith_arxiv_id":"2508.11987","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Futurex: An advanced live benchmark for llm agents in future prediction","venue":"ArXiv.org","work_id":"ac497801-1ebb-45ae-8286-d566aec17d80","year":2025},"citing_paper":{"arxiv_id":"2605.22681","last_updated":"2026-07-18T20:33:32Z","snapshot_observed_at":"2026-08-02T13:31:38.903992Z","submitted_at":"2026-05-21T16:23:36Z","title":"Scientific reasoning does not reliably translate into scientific forecasting in frontier AI","version":1},"reference_index":22,"source":"pdf_text","source_observed_at":"2026-05-22T05:18:15.358571Z"},"links":{"cited_paper":"/paper/2508.11987","citing_paper":"/paper/2605.22681"},"observation_digest":"sha256:9b5eac37d25b254fedc984d2635cf31ea0aaceb0e9cf5e6cce5c0a84a2ffc3b0","observation_id":"c05e631b-632c-4b40-a531-c32f67559898","resolution":{"observed_at":"2026-05-22T05:21:07.417427Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2508.11987","last_updated":"2025-09-05T09:15:55Z","snapshot_observed_at":"2026-08-05T19:39:12.827378Z","submitted_at":"2025-08-16T08:54:08Z","title":"FutureX: An Advanced Live Benchmark for LLM Agents in Future Prediction","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2508.11987","snapshot_observed_at":"2026-08-02T13:31:44.585483Z","title":"Futurex: An advanced live benchmark for llm agents in future prediction.arXiv preprint arXiv:2508.11987, 2025","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2605.22681","last_updated":"2026-07-18T20:33:32Z","snapshot_observed_at":"2026-08-02T13:31:38.903992Z","submitted_at":"2026-05-21T16:23:36Z","title":"Scientific reasoning does not reliably translate into scientific forecasting in frontier AI","version":2},"reference_index":22,"source":"pdf_text","source_observed_at":"2026-08-02T13:31:44.585483Z"},"links":{"cited_paper":"/paper/2508.11987","citing_paper":"/paper/2605.22681"},"observation_digest":"sha256:5763e2988436c6cd00c490678bb1e00c7f6c0dcaeb2ad6f1aec43d703ac7f79e","observation_id":"072fb71f-1a39-485a-b248-f830b7ee8264","resolution":{"observed_at":"2026-08-02T13:31:44.585483Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2508.11987","last_updated":"2025-09-05T09:15:55Z","snapshot_observed_at":"2026-08-05T19:39:12.827378Z","submitted_at":"2025-08-16T08:54:08Z","title":"FutureX: An Advanced Live Benchmark for LLM Agents in Future Prediction","version":3},"cited_work":{"arxiv_id":"2508.11987","doi":"10.48550/arxiv.2508.11987","metadata_source":"arxiv_reference","pith_arxiv_id":"2508.11987","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Futurex: An advanced live benchmark for llm agents in future prediction","venue":"ArXiv.org","work_id":"ac497801-1ebb-45ae-8286-d566aec17d80","year":2025},"citing_paper":{"arxiv_id":"2606.00644","last_updated":"2026-05-30T09:41:26Z","snapshot_observed_at":"2026-08-04T08:39:53.308041Z","submitted_at":"2026-05-30T09:41:26Z","title":"ForeSci: Evaluating LLM Agents for Forward-Looking AI Research Judgment","version":1},"reference_index":32,"source":"arxiv_source","source_observed_at":"2026-06-28T18:46:16.099080Z"},"links":{"cited_paper":"/paper/2508.11987","citing_paper":"/paper/2606.00644"},"observation_digest":"sha256:a26eb73fe9db4c62e7966b2dc9662df9bb65eb27b3332f92a3df3f8505ac24a7","observation_id":"678056fa-82c0-40f5-a614-ed54691b3b26","resolution":{"observed_at":"2026-06-28T19:52:35.956187Z","resolver_source":"arxiv_id","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2508.11987","last_updated":"2025-09-05T09:15:55Z","snapshot_observed_at":"2026-08-05T19:39:12.827378Z","submitted_at":"2025-08-16T08:54:08Z","title":"FutureX: An Advanced Live Benchmark for LLM Agents in Future Prediction","version":3},"cited_work":{"arxiv_id":"2508.11987","doi":"10.48550/arxiv.2508.11987","metadata_source":"arxiv_reference","pith_arxiv_id":"2508.11987","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Futurex: An advanced live benchmark for llm agents in future prediction","venue":"ArXiv.org","work_id":"ac497801-1ebb-45ae-8286-d566aec17d80","year":2025},"citing_paper":{"arxiv_id":"2606.11816","last_updated":"2026-06-10T08:50:29Z","snapshot_observed_at":"2026-07-06T23:50:49.040880Z","submitted_at":"2026-06-10T08:50:29Z","title":"WorldReasoner: Evaluating Whether Language Model Agents Forecast Events with Valid Reasoning","version":1},"reference_index":45,"source":"arxiv_source","source_observed_at":"2026-06-27T09:47:04.122464Z"},"links":{"cited_paper":"/paper/2508.11987","citing_paper":"/paper/2606.11816"},"observation_digest":"sha256:cc9f88b5a04814a4630375210de78c5479b90c00aabf667650efea0fba0553a5","observation_id":"8c22fc98-30b1-46c3-bb82-0facb0491dc9","resolution":{"observed_at":"2026-07-03T10:58:02.668078Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2508.11987","last_updated":"2025-09-05T09:15:55Z","snapshot_observed_at":"2026-08-05T19:39:12.827378Z","submitted_at":"2025-08-16T08:54:08Z","title":"FutureX: An Advanced Live Benchmark for LLM Agents in Future Prediction","version":3},"cited_work":{"arxiv_id":"2508.11987","doi":"10.48550/arxiv.2508.11987","metadata_source":"arxiv_reference","pith_arxiv_id":"2508.11987","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Futurex: An advanced live benchmark for llm agents in future prediction","venue":"ArXiv.org","work_id":"ac497801-1ebb-45ae-8286-d566aec17d80","year":2025},"citing_paper":{"arxiv_id":"2606.21013","last_updated":"2026-06-19T00:55:00Z","snapshot_observed_at":"2026-08-02T21:12:52.966362Z","submitted_at":"2026-06-19T00:55:00Z","title":"Agentic Time Machine as an Infrastructure for Future-Event Forecasting","version":1},"reference_index":27,"source":"pdf_text","source_observed_at":"2026-06-26T14:42:29.128094Z"},"links":{"cited_paper":"/paper/2508.11987","citing_paper":"/paper/2606.21013"},"observation_digest":"sha256:a0a5a286961a7a53f3fa97fe001e5695a2e1b8435183c040625d05ba0b329dc8","observation_id":"0e6c6222-9efe-4c32-b7eb-4f429394aa3d","resolution":{"observed_at":"2026-07-04T06:09:38.138934Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2508.11987","last_updated":"2025-09-05T09:15:55Z","snapshot_observed_at":"2026-08-05T19:39:12.827378Z","submitted_at":"2025-08-16T08:54:08Z","title":"FutureX: An Advanced Live Benchmark for LLM Agents in Future Prediction","version":3},"cited_work":{"arxiv_id":"2508.11987","doi":"10.48550/arxiv.2508.11987","metadata_source":"arxiv_reference","pith_arxiv_id":"2508.11987","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Futurex: An advanced live benchmark for llm agents in future prediction","venue":"ArXiv.org","work_id":"ac497801-1ebb-45ae-8286-d566aec17d80","year":2025},"citing_paper":{"arxiv_id":"2607.01661","last_updated":"2026-07-02T03:41:25Z","snapshot_observed_at":"2026-07-07T00:07:13.555571Z","submitted_at":"2026-07-02T03:41:25Z","title":"Diverse Evidence, Better Forecasts: Multi-Agent Deliberation Under Information Asymmetry","version":1},"reference_index":22,"source":"arxiv_source","source_observed_at":"2026-07-03T14:40:49.038578Z"},"links":{"cited_paper":"/paper/2508.11987","citing_paper":"/paper/2607.01661"},"observation_digest":"sha256:3d5804ee5c36ef460ac6c7264b467e8871b5da59f646b6f954e8c57e4cb61ffa","observation_id":"90fd05a4-882d-4a58-b040-a52f4bdb448e","resolution":{"observed_at":"2026-07-03T14:48:32.575741Z","resolver_source":"arxiv_id","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2508.11987","last_updated":"2025-09-05T09:15:55Z","snapshot_observed_at":"2026-08-05T19:39:12.827378Z","submitted_at":"2025-08-16T08:54:08Z","title":"FutureX: An Advanced Live Benchmark for LLM Agents in Future Prediction","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2508.11987","snapshot_observed_at":"2026-07-11T19:31:58.491532Z","title":"Futurex: An advanced live benchmark for LLM agents in future prediction.arXiv preprint arXiv:2508.11987, 2025","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2607.04389","last_updated":"2026-07-05T16:34:54Z","snapshot_observed_at":"2026-07-11T19:31:57.214300Z","submitted_at":"2026-07-05T16:34:54Z","title":"Decentralized Aggregation of LLM Predictions via Wagering Mechanisms","version":1},"reference_index":37,"source":"pdf_text","source_observed_at":"2026-07-11T19:31:58.491532Z"},"links":{"cited_paper":"/paper/2508.11987","citing_paper":"/paper/2607.04389"},"observation_digest":"sha256:1dbdb417cd1f7e31f628838e5bd96f25aff0dd189aeaaa9700ece48243c9cb96","observation_id":"498f3c86-cc5f-413e-9032-7a07a0b7f324","resolution":{"observed_at":"2026-07-11T19:31:58.491532Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}}],"links":{"evidence":"/evidence","html":"/paper/2508.11987/citation-record","integrity":"/paper/2508.11987/integrity","json":"/paper/2508.11987/citation-record.json","paper":"/paper/2508.11987"},"outbound":[],"paper":{"arxiv_id":"2508.11987","last_updated":"2025-09-05T09:15:55Z","latest_version":3,"primary_category":"cs.AI","snapshot_observed_at":"2026-08-05T19:39:12.827378Z","submitted_at":"2025-08-16T08:54:08Z","title":"FutureX: An Advanced Live Benchmark for LLM Agents in Future Prediction"},"reference_resolution":{"displayed":0,"state_counts":{"malformed_identifier":0,"metadata_mismatch":0,"parse_uncertain":0,"unresolved":0,"verified_exact":0,"verified_fuzzy":0},"total_outbound_references":0},"refusal":"A citation records a reference. It does not transfer a finding from one paper to another.","schema":"pith.paper-citation-record.v1","standing_sources":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"thesis":"As of 5 August 2026, this Paper Citation Record lists 0 of 0 outbound references and 20 inbound Pith citation observations for arXiv:2508.11987."}