{"as_of":"2026-08-01T02:25:00Z","caps":{"database_statements":6,"inbound":100,"outbound":100},"context_digest":"sha256:b94a01f30e467bf82b0140e4d42276c5d9572977ff3e3cff2b88788ecf260278","coverage":[{"denominator":0,"lane":"reference_resolution","note":"Typed states for the displayed outbound observations.","records_observed":0,"source":"paper_references, paper_reference_links","source_observed_at":null,"state":"measured"},{"denominator":20,"lane":"standing_notices","note":"One-hop event checks from named stored sources.","records_observed":20,"source":"scholarly_work_events, retraction_status_cache","source_observed_at":"2026-07-31T06:34:12.847434+00:00","state":"measured"},{"denominator":20,"lane":"inbound_itemization","note":"Pith citing papers itemized under the disclosed page cap.","records_observed":20,"source":"paper_references, paper_reference_links","source_observed_at":"2026-08-01T00:52:32.240795Z","state":"measured"},{"denominator":1,"lane":"external_citation_measurements","note":"A source-named dated measurement, never combined with another source.","records_observed":0,"source":"arxiv_reference","source_observed_at":"2026-07-10T12:15:01.137692Z","state":"measured"}],"external_citation_measurements":[],"inbound":[{"citation":{"cited_paper":{"arxiv_id":"2504.16074","last_updated":"2025-05-18T14:13:34Z","snapshot_observed_at":"2026-07-06T21:13:09.093282Z","submitted_at":"2025-04-22T17:53:29Z","title":"PHYBench: Holistic Evaluation of Physical Perception and Reasoning in Large Language Models","version":2},"cited_work":{"arxiv_id":"2504.16074","doi":"10.48550/arxiv.2504.16074","metadata_source":"arxiv_reference","pith_arxiv_id":"2504.16074","snapshot_observed_at":"2026-07-10T12:15:01.137692Z","title":"Phybench: Holistic evaluation of physical perception and reasoning in large language models","venue":null,"work_id":"d2ae619c-0cbb-45af-8aa0-0a52a9b8dec1","year":2025},"citing_paper":{"arxiv_id":"2509.26574","last_updated":"2026-05-08T21:22:01Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2025-09-30T17:34:03Z","title":"Probing the Critical Point (CritPt) of AI Reasoning: a Frontier Physics Research Benchmark","version":4},"reference_index":41,"source":"pdf_text","source_observed_at":"2026-05-18T11:52:10.205796Z"},"links":{"cited_paper":"/paper/2504.16074","citing_paper":"/paper/2509.26574"},"observation_digest":"sha256:06fb9fbb335f6ba908c799d6ec23ab52626ee2dc4818fb1f63c17f3162129e04","observation_id":"e0b744c6-cc25-48c7-86ca-448723632ff8","resolution":{"observed_at":"2026-05-18T11:52:35.306254Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-07-31T06:34:12.847434+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-07-31T06:34:12.847434+00:00","source":"crossref"},{"observed_at":"2026-07-31T06:34:08.642788+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2504.16074","last_updated":"2025-05-18T14:13:34Z","snapshot_observed_at":"2026-07-06T21:13:09.093282Z","submitted_at":"2025-04-22T17:53:29Z","title":"PHYBench: Holistic Evaluation of Physical Perception and Reasoning in Large Language Models","version":2},"cited_work":{"arxiv_id":"2504.16074","doi":"10.48550/arxiv.2504.16074","metadata_source":"arxiv_reference","pith_arxiv_id":"2504.16074","snapshot_observed_at":"2026-07-10T12:15:01.137692Z","title":"Phybench: Holistic evaluation of physical perception and reasoning in large language models","venue":null,"work_id":"d2ae619c-0cbb-45af-8aa0-0a52a9b8dec1","year":2025},"citing_paper":{"arxiv_id":"2512.15745","last_updated":"2025-12-24T03:46:46Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2025-12-10T09:26:18Z","title":"LLaDA2.0: Scaling Up Diffusion Language Models to 100B","version":2},"reference_index":26,"source":"pdf_text","source_observed_at":"2026-05-14T18:53:20.911374Z"},"links":{"cited_paper":"/paper/2504.16074","citing_paper":"/paper/2512.15745"},"observation_digest":"sha256:6b2b2212ea67e3327c38ec9d624fc67b982fb4da0ca5ceec31c39fde2b864b7a","observation_id":"cbac3c2e-1c69-46f2-8286-5d56524584fb","resolution":{"observed_at":"2026-05-14T18:53:21.116805Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-07-31T06:34:12.847434+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-07-31T06:34:12.847434+00:00","source":"crossref"},{"observed_at":"2026-07-31T06:34:08.642788+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2504.16074","last_updated":"2025-05-18T14:13:34Z","snapshot_observed_at":"2026-07-06T21:13:09.093282Z","submitted_at":"2025-04-22T17:53:29Z","title":"PHYBench: Holistic Evaluation of Physical Perception and Reasoning in Large Language Models","version":2},"cited_work":{"arxiv_id":"2504.16074","doi":"10.48550/arxiv.2504.16074","metadata_source":"arxiv_reference","pith_arxiv_id":"2504.16074","snapshot_observed_at":"2026-07-10T12:15:01.137692Z","title":"Phybench: Holistic evaluation of physical perception and reasoning in large language models","venue":null,"work_id":"d2ae619c-0cbb-45af-8aa0-0a52a9b8dec1","year":2025},"citing_paper":{"arxiv_id":"2601.08584","last_updated":"2026-01-13T14:06:03Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2026-01-13T14:06:03Z","title":"Ministral 3","version":1},"reference_index":16,"source":"pdf_text","source_observed_at":"2026-05-14T19:12:24.627033Z"},"links":{"cited_paper":"/paper/2504.16074","citing_paper":"/paper/2601.08584"},"observation_digest":"sha256:b0b218554f7dee30bddf2847e1f9a7a632f8c35629917ee0a1c702a8fd2a5bb7","observation_id":"d603fdcf-2b82-42da-b44d-d8cc4cf1ac57","resolution":{"observed_at":"2026-05-14T19:12:24.750113Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-07-31T06:34:12.847434+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-07-31T06:34:12.847434+00:00","source":"crossref"},{"observed_at":"2026-07-31T06:34:08.642788+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2504.16074","last_updated":"2025-05-18T14:13:34Z","snapshot_observed_at":"2026-07-06T21:13:09.093282Z","submitted_at":"2025-04-22T17:53:29Z","title":"PHYBench: Holistic Evaluation of Physical Perception and Reasoning in Large Language Models","version":2},"cited_work":{"arxiv_id":"2504.16074","doi":"10.48550/arxiv.2504.16074","metadata_source":"arxiv_reference","pith_arxiv_id":"2504.16074","snapshot_observed_at":"2026-07-10T12:15:01.137692Z","title":"Phybench: Holistic evaluation of physical perception and reasoning in large language models","venue":null,"work_id":"d2ae619c-0cbb-45af-8aa0-0a52a9b8dec1","year":2025},"citing_paper":{"arxiv_id":"2602.01203","last_updated":"2026-05-27T09:56:04Z","snapshot_observed_at":"2026-07-06T22:44:00.266269Z","submitted_at":"2026-02-01T12:45:39Z","title":"Attention Sink Forges Native MoE in Attention Layers: Sink-Aware Training to Address Head Collapse","version":2},"reference_index":27,"source":"pdf_text","source_observed_at":"2026-05-16T08:47:29.236561Z"},"links":{"cited_paper":"/paper/2504.16074","citing_paper":"/paper/2602.01203"},"observation_digest":"sha256:81af43d8b3a70c9b2b000f19c72dda83ec7ab9736c7cf31772bc9f587fbaff4f","observation_id":"b7c46987-ace9-4a1a-9360-c9abda75a829","resolution":{"observed_at":"2026-05-16T08:47:37.220461Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-07-31T06:34:12.847434+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-07-31T06:34:12.847434+00:00","source":"crossref"},{"observed_at":"2026-07-31T06:34:08.642788+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2504.16074","last_updated":"2025-05-18T14:13:34Z","snapshot_observed_at":"2026-07-06T21:13:09.093282Z","submitted_at":"2025-04-22T17:53:29Z","title":"PHYBench: Holistic Evaluation of Physical Perception and Reasoning in Large Language Models","version":2},"cited_work":{"arxiv_id":"2504.16074","doi":"10.48550/arxiv.2504.16074","metadata_source":"arxiv_reference","pith_arxiv_id":"2504.16074","snapshot_observed_at":"2026-07-10T12:15:01.137692Z","title":"Phybench: Holistic evaluation of physical perception and reasoning in large language models","venue":null,"work_id":"d2ae619c-0cbb-45af-8aa0-0a52a9b8dec1","year":2025},"citing_paper":{"arxiv_id":"2602.07064","last_updated":"2026-04-07T13:49:35Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2026-02-05T14:04:51Z","title":"OmniFysics: Towards Physical Intelligence Evolution via Omni-Modal Signal Processing and Network Optimization","version":2},"reference_index":15,"source":"pdf_text","source_observed_at":"2026-05-16T07:09:46.254851Z"},"links":{"cited_paper":"/paper/2504.16074","citing_paper":"/paper/2602.07064"},"observation_digest":"sha256:36954a86354796b4c6789c3b5220f5233084e0fd6310733e1234c01ba2d4748e","observation_id":"d9d0573e-e1ac-475f-a076-319faa5b08f4","resolution":{"observed_at":"2026-05-16T07:10:43.156734Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-07-31T06:34:12.847434+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-07-31T06:34:12.847434+00:00","source":"crossref"},{"observed_at":"2026-07-31T06:34:08.642788+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2504.16074","last_updated":"2025-05-18T14:13:34Z","snapshot_observed_at":"2026-07-06T21:13:09.093282Z","submitted_at":"2025-04-22T17:53:29Z","title":"PHYBench: Holistic Evaluation of Physical Perception and Reasoning in Large Language Models","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2504.16074","snapshot_observed_at":"2026-07-15T13:27:51.848177Z","title":"Phybench: Holis- tic evaluation of physical perception and reasoning in large language models.arXiv preprint arXiv:2504.16074,","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2603.07109","last_updated":"2026-05-30T07:36:55Z","snapshot_observed_at":"2026-07-15T13:27:48.715737Z","submitted_at":"2026-03-07T08:40:09Z","title":"Vision Language Models Cannot Reason About Physical Transformation","version":2},"reference_index":22,"source":"pdf_text","source_observed_at":"2026-07-15T13:27:51.848177Z"},"links":{"cited_paper":"/paper/2504.16074","citing_paper":"/paper/2603.07109"},"observation_digest":"sha256:b3b7f218923db005638bc8f1619afbbf4f277c632b0afac0b080495be082444b","observation_id":"af40aa66-2d6c-4ea5-aac5-f4df7fab6785","resolution":{"observed_at":"2026-07-15T13:27:51.848177Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2504.16074","last_updated":"2025-05-18T14:13:34Z","snapshot_observed_at":"2026-07-06T21:13:09.093282Z","submitted_at":"2025-04-22T17:53:29Z","title":"PHYBench: Holistic Evaluation of Physical Perception and Reasoning in Large Language Models","version":2},"cited_work":{"arxiv_id":"2504.16074","doi":"10.48550/arxiv.2504.16074","metadata_source":"arxiv_reference","pith_arxiv_id":"2504.16074","snapshot_observed_at":"2026-07-10T12:15:01.137692Z","title":"Phybench: Holistic evaluation of physical perception and reasoning in large language models","venue":null,"work_id":"d2ae619c-0cbb-45af-8aa0-0a52a9b8dec1","year":2025},"citing_paper":{"arxiv_id":"2603.20633","last_updated":"2026-04-17T08:57:17Z","snapshot_observed_at":"2026-07-06T22:49:57.103822Z","submitted_at":"2026-03-21T04:03:45Z","title":"Seed1.8 Model Card: Towards Generalized Real-World Agency","version":3},"reference_index":55,"source":"pdf_text","source_observed_at":"2026-05-15T07:44:02.827006Z"},"links":{"cited_paper":"/paper/2504.16074","citing_paper":"/paper/2603.20633"},"observation_digest":"sha256:948f2d5bf95472747f98681eb8faa85f6a6a148dc31615799a9c8b991f0fa086","observation_id":"969d8e6d-5cbb-49a2-a9a8-34101c6babf9","resolution":{"observed_at":"2026-05-15T07:45:14.405531Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-07-31T06:34:12.847434+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-07-31T06:34:12.847434+00:00","source":"crossref"},{"observed_at":"2026-07-31T06:34:08.642788+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2504.16074","last_updated":"2025-05-18T14:13:34Z","snapshot_observed_at":"2026-07-06T21:13:09.093282Z","submitted_at":"2025-04-22T17:53:29Z","title":"PHYBench: Holistic Evaluation of Physical Perception and Reasoning in Large Language Models","version":2},"cited_work":{"arxiv_id":"2504.16074","doi":"10.48550/arxiv.2504.16074","metadata_source":"arxiv_reference","pith_arxiv_id":"2504.16074","snapshot_observed_at":"2026-07-10T12:15:01.137692Z","title":"Phybench: Holistic evaluation of physical perception and reasoning in large language models","venue":null,"work_id":"d2ae619c-0cbb-45af-8aa0-0a52a9b8dec1","year":2025},"citing_paper":{"arxiv_id":"2604.02934","last_updated":"2026-04-03T10:05:30Z","snapshot_observed_at":"2026-07-06T22:52:10.923214Z","submitted_at":"2026-04-03T10:05:30Z","title":"PolyReal: A Benchmark for Real-World Polymer Science Workflows","version":1},"reference_index":40,"source":"pdf_text","source_observed_at":"2026-05-13T19:55:17.366868Z"},"links":{"cited_paper":"/paper/2504.16074","citing_paper":"/paper/2604.02934"},"observation_digest":"sha256:ed293ea6b12a9355e50309926e2c885678523ad936236b7b4dcdd7c78638ec8d","observation_id":"c20b9fcf-9eaf-4c23-ac30-c378aaf4a9db","resolution":{"observed_at":"2026-05-13T19:58:12.284099Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-07-31T06:34:12.847434+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-07-31T06:34:12.847434+00:00","source":"crossref"},{"observed_at":"2026-07-31T06:34:08.642788+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2504.16074","last_updated":"2025-05-18T14:13:34Z","snapshot_observed_at":"2026-07-06T21:13:09.093282Z","submitted_at":"2025-04-22T17:53:29Z","title":"PHYBench: Holistic Evaluation of Physical Perception and Reasoning in Large Language Models","version":2},"cited_work":{"arxiv_id":"2504.16074","doi":"10.48550/arxiv.2504.16074","metadata_source":"arxiv_reference","pith_arxiv_id":"2504.16074","snapshot_observed_at":"2026-07-10T12:15:01.137692Z","title":"Phybench: Holistic evaluation of physical perception and reasoning in large language models","venue":null,"work_id":"d2ae619c-0cbb-45af-8aa0-0a52a9b8dec1","year":2025},"citing_paper":{"arxiv_id":"2604.15411","last_updated":"2026-04-16T16:22:04Z","snapshot_observed_at":"2026-07-06T23:02:58.361418Z","submitted_at":"2026-04-16T16:22:04Z","title":"PRL-Bench: A Comprehensive Benchmark Evaluating LLMs' Capabilities in Frontier Physics Research","version":1},"reference_index":8,"source":"pdf_text","source_observed_at":"2026-05-10T11:10:21.639856Z"},"links":{"cited_paper":"/paper/2504.16074","citing_paper":"/paper/2604.15411"},"observation_digest":"sha256:0638c3279347451982a736614c379d78c1f704fbd231a3958792b8116c90dd06","observation_id":"aaaff15b-3695-4c8d-a275-dc55b59c8288","resolution":{"observed_at":"2026-05-10T11:20:11.381486Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-07-31T06:34:12.847434+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-07-31T06:34:12.847434+00:00","source":"crossref"},{"observed_at":"2026-07-31T06:34:08.642788+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2504.16074","last_updated":"2025-05-18T14:13:34Z","snapshot_observed_at":"2026-07-06T21:13:09.093282Z","submitted_at":"2025-04-22T17:53:29Z","title":"PHYBench: Holistic Evaluation of Physical Perception and Reasoning in Large Language Models","version":2},"cited_work":{"arxiv_id":"2504.16074","doi":"10.48550/arxiv.2504.16074","metadata_source":"arxiv_reference","pith_arxiv_id":"2504.16074","snapshot_observed_at":"2026-07-10T12:15:01.137692Z","title":"Phybench: Holistic evaluation of physical perception and reasoning in large language models","venue":null,"work_id":"d2ae619c-0cbb-45af-8aa0-0a52a9b8dec1","year":2025},"citing_paper":{"arxiv_id":"2604.23580","last_updated":"2026-04-26T07:37:20Z","snapshot_observed_at":"2026-07-30T20:45:12.910299Z","submitted_at":"2026-04-26T07:37:20Z","title":"PhysCodeBench: Benchmarking Physics-Aware Symbolic Simulation of 3D Scenes via Self-Corrective Multi-Agent Refinement","version":1},"reference_index":22,"source":"pdf_text","source_observed_at":"2026-05-08T06:13:18.967603Z"},"links":{"cited_paper":"/paper/2504.16074","citing_paper":"/paper/2604.23580"},"observation_digest":"sha256:4472fd8b2730d14fac57c4fb5a807eb36b46522fe5b325cff07f12202d423ea9","observation_id":"f727223d-5410-41c5-ab9f-d9748d93aea9","resolution":{"observed_at":"2026-05-11T21:16:13.110357Z","resolver_source":"arxiv_id","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-07-31T06:34:12.847434+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-07-31T06:34:12.847434+00:00","source":"crossref"},{"observed_at":"2026-07-31T06:34:08.642788+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2504.16074","last_updated":"2025-05-18T14:13:34Z","snapshot_observed_at":"2026-07-06T21:13:09.093282Z","submitted_at":"2025-04-22T17:53:29Z","title":"PHYBench: Holistic Evaluation of Physical Perception and Reasoning in Large Language Models","version":2},"cited_work":{"arxiv_id":"2504.16074","doi":"10.48550/arxiv.2504.16074","metadata_source":"arxiv_reference","pith_arxiv_id":"2504.16074","snapshot_observed_at":"2026-07-10T12:15:01.137692Z","title":"Phybench: Holistic evaluation of physical perception and reasoning in large language models","venue":null,"work_id":"d2ae619c-0cbb-45af-8aa0-0a52a9b8dec1","year":2025},"citing_paper":{"arxiv_id":"2604.24443","last_updated":"2026-04-27T13:10:52Z","snapshot_observed_at":"2026-07-06T23:10:29.508602Z","submitted_at":"2026-04-27T13:10:52Z","title":"PhysNote: Self-Knowledge Notes for Evolvable Physical Reasoning in Vision-Language Model","version":1},"reference_index":10,"source":"pdf_text","source_observed_at":"2026-05-08T03:43:18.640857Z"},"links":{"cited_paper":"/paper/2504.16074","citing_paper":"/paper/2604.24443"},"observation_digest":"sha256:bdbd7a560d78782b2e75a181e3b6f98e0a8fa438917a8b1211cf0157baf53e7d","observation_id":"79a6c43f-2664-4ec9-b725-153a6c4c7309","resolution":{"observed_at":"2026-05-11T21:56:26.628602Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-07-31T06:34:12.847434+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-07-31T06:34:12.847434+00:00","source":"crossref"},{"observed_at":"2026-07-31T06:34:08.642788+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2504.16074","last_updated":"2025-05-18T14:13:34Z","snapshot_observed_at":"2026-07-06T21:13:09.093282Z","submitted_at":"2025-04-22T17:53:29Z","title":"PHYBench: Holistic Evaluation of Physical Perception and Reasoning in Large Language Models","version":2},"cited_work":{"arxiv_id":"2504.16074","doi":"10.48550/arxiv.2504.16074","metadata_source":"arxiv_reference","pith_arxiv_id":"2504.16074","snapshot_observed_at":"2026-07-10T12:15:01.137692Z","title":"Phybench: Holistic evaluation of physical perception and reasoning in large language models","venue":null,"work_id":"d2ae619c-0cbb-45af-8aa0-0a52a9b8dec1","year":2025},"citing_paper":{"arxiv_id":"2604.27351","last_updated":"2026-04-30T03:02:27Z","snapshot_observed_at":"2026-07-06T23:12:52.141592Z","submitted_at":"2026-04-30T03:02:27Z","title":"Heterogeneous Scientific Foundation Model Collaboration","version":1},"reference_index":39,"source":"pdf_text","source_observed_at":"2026-05-07T08:50:05.980191Z"},"links":{"cited_paper":"/paper/2504.16074","citing_paper":"/paper/2604.27351"},"observation_digest":"sha256:95ef6ec6f490edaf846de433ea52f9d2561a1f0c9ade5b426509033bc2795481","observation_id":"5d636051-6d1a-4be7-abd3-9f50fa66a361","resolution":{"observed_at":"2026-05-09T04:30:10.962510Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-07-31T06:34:12.847434+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-07-31T06:34:12.847434+00:00","source":"crossref"},{"observed_at":"2026-07-31T06:34:08.642788+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2504.16074","last_updated":"2025-05-18T14:13:34Z","snapshot_observed_at":"2026-07-06T21:13:09.093282Z","submitted_at":"2025-04-22T17:53:29Z","title":"PHYBench: Holistic Evaluation of Physical Perception and Reasoning in Large Language Models","version":2},"cited_work":{"arxiv_id":"2504.16074","doi":"10.48550/arxiv.2504.16074","metadata_source":"arxiv_reference","pith_arxiv_id":"2504.16074","snapshot_observed_at":"2026-07-10T12:15:01.137692Z","title":"Phybench: Holistic evaluation of physical perception and reasoning in large language models","venue":null,"work_id":"d2ae619c-0cbb-45af-8aa0-0a52a9b8dec1","year":2025},"citing_paper":{"arxiv_id":"2605.09636","last_updated":"2026-05-10T16:25:43Z","snapshot_observed_at":"2026-07-30T14:07:22.496707Z","submitted_at":"2026-05-10T16:25:43Z","title":"PDEAgent-Bench: A Multi-Metric, Multi-Library Benchmark for PDE Solver Generation","version":1},"reference_index":37,"source":"pdf_text","source_observed_at":"2026-05-12T02:36:40.696567Z"},"links":{"cited_paper":"/paper/2504.16074","citing_paper":"/paper/2605.09636"},"observation_digest":"sha256:cf418ee25ca63be28c6b1ae53699bbf04973e3558539441f8718458b9c11be2c","observation_id":"c55c8f39-8847-494f-ac22-95dec3f160dd","resolution":{"observed_at":"2026-05-12T07:31:26.675661Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-07-31T06:34:12.847434+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-07-31T06:34:12.847434+00:00","source":"crossref"},{"observed_at":"2026-07-31T06:34:08.642788+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2504.16074","last_updated":"2025-05-18T14:13:34Z","snapshot_observed_at":"2026-07-06T21:13:09.093282Z","submitted_at":"2025-04-22T17:53:29Z","title":"PHYBench: Holistic Evaluation of Physical Perception and Reasoning in Large Language Models","version":2},"cited_work":{"arxiv_id":"2504.16074","doi":"10.48550/arxiv.2504.16074","metadata_source":"arxiv_reference","pith_arxiv_id":"2504.16074","snapshot_observed_at":"2026-07-10T12:15:01.137692Z","title":"Phybench: Holistic evaluation of physical perception and reasoning in large language models","venue":null,"work_id":"d2ae619c-0cbb-45af-8aa0-0a52a9b8dec1","year":2025},"citing_paper":{"arxiv_id":"2606.01538","last_updated":"2026-06-11T05:58:14Z","snapshot_observed_at":"2026-07-06T23:42:07.508474Z","submitted_at":"2026-06-01T01:36:44Z","title":"MPMWorlds: Material-Point-Method Simulations for Inferring and Extrapolating Physical Dynamics","version":2},"reference_index":11,"source":"pdf_text","source_observed_at":"2026-06-28T12:19:27.221596Z"},"links":{"cited_paper":"/paper/2504.16074","citing_paper":"/paper/2606.01538"},"observation_digest":"sha256:28d56b02fbba794323349a02ee9c99a497341389277384fcf21da3d696ec7e86","observation_id":"668e0963-bcde-40db-8433-fc7ccdcea260","resolution":{"observed_at":"2026-07-02T01:16:24.687383Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-07-31T06:34:12.847434+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-07-31T06:34:12.847434+00:00","source":"crossref"},{"observed_at":"2026-07-31T06:34:08.642788+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2504.16074","last_updated":"2025-05-18T14:13:34Z","snapshot_observed_at":"2026-07-06T21:13:09.093282Z","submitted_at":"2025-04-22T17:53:29Z","title":"PHYBench: Holistic Evaluation of Physical Perception and Reasoning in Large Language Models","version":2},"cited_work":{"arxiv_id":"2504.16074","doi":"10.48550/arxiv.2504.16074","metadata_source":"arxiv_reference","pith_arxiv_id":"2504.16074","snapshot_observed_at":"2026-07-10T12:15:01.137692Z","title":"Phybench: Holistic evaluation of physical perception and reasoning in large language models","venue":null,"work_id":"d2ae619c-0cbb-45af-8aa0-0a52a9b8dec1","year":2025},"citing_paper":{"arxiv_id":"2606.07962","last_updated":"2026-06-06T03:40:47Z","snapshot_observed_at":"2026-07-31T20:27:39.414307Z","submitted_at":"2026-06-06T03:40:47Z","title":"ChronoPhyBench: Do MLLMs Truly Understand the World or Merely Exploit Language Priors?","version":1},"reference_index":42,"source":"pdf_text","source_observed_at":"2026-06-27T20:23:18.667677Z"},"links":{"cited_paper":"/paper/2504.16074","citing_paper":"/paper/2606.07962"},"observation_digest":"sha256:67f212b9e621c0af52eedc7104562fd7396016e84aed3c48a88288aeb8adc705","observation_id":"22bbed5b-1b7d-4d4c-87bf-c95da090f955","resolution":{"observed_at":"2026-07-02T20:27:22.619660Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-07-31T06:34:12.847434+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-07-31T06:34:12.847434+00:00","source":"crossref"},{"observed_at":"2026-07-31T06:34:08.642788+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2504.16074","last_updated":"2025-05-18T14:13:34Z","snapshot_observed_at":"2026-07-06T21:13:09.093282Z","submitted_at":"2025-04-22T17:53:29Z","title":"PHYBench: Holistic Evaluation of Physical Perception and Reasoning in Large Language Models","version":2},"cited_work":{"arxiv_id":"2504.16074","doi":"10.48550/arxiv.2504.16074","metadata_source":"arxiv_reference","pith_arxiv_id":"2504.16074","snapshot_observed_at":"2026-07-10T12:15:01.137692Z","title":"Phybench: Holistic evaluation of physical perception and reasoning in large language models","venue":null,"work_id":"d2ae619c-0cbb-45af-8aa0-0a52a9b8dec1","year":2025},"citing_paper":{"arxiv_id":"2606.08034","last_updated":"2026-06-06T07:51:52Z","snapshot_observed_at":"2026-07-06T23:47:33.680573Z","submitted_at":"2026-06-06T07:51:52Z","title":"Sci-Rho: A Multilingual Visually-Grounded Symbolic Benchmark for STEM Problems","version":1},"reference_index":64,"source":"arxiv_source","source_observed_at":"2026-06-27T20:11:02.445626Z"},"links":{"cited_paper":"/paper/2504.16074","citing_paper":"/paper/2606.08034"},"observation_digest":"sha256:e11259259c9dd6f8d8b67e5dbe7833915ec4c48325a717cd205838e69912e2ce","observation_id":"188ce006-ec3f-4cf6-8db8-cef93e7c21fe","resolution":{"observed_at":"2026-07-02T20:47:22.878722Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-07-31T06:34:12.847434+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-07-31T06:34:12.847434+00:00","source":"crossref"},{"observed_at":"2026-07-31T06:34:08.642788+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2504.16074","last_updated":"2025-05-18T14:13:34Z","snapshot_observed_at":"2026-07-06T21:13:09.093282Z","submitted_at":"2025-04-22T17:53:29Z","title":"PHYBench: Holistic Evaluation of Physical Perception and Reasoning in Large Language Models","version":2},"cited_work":{"arxiv_id":"2504.16074","doi":"10.48550/arxiv.2504.16074","metadata_source":"arxiv_reference","pith_arxiv_id":"2504.16074","snapshot_observed_at":"2026-07-10T12:15:01.137692Z","title":"Phybench: Holistic evaluation of physical perception and reasoning in large language models","venue":null,"work_id":"d2ae619c-0cbb-45af-8aa0-0a52a9b8dec1","year":2025},"citing_paper":{"arxiv_id":"2607.00248","last_updated":"2026-06-30T22:57:43Z","snapshot_observed_at":"2026-07-07T00:05:54.401681Z","submitted_at":"2026-06-30T22:57:43Z","title":"Seed2.0 Model Card: Towards Intelligence Frontier for Real-World Complexity","version":1},"reference_index":88,"source":"pdf_text","source_observed_at":"2026-07-02T18:57:46.841456Z"},"links":{"cited_paper":"/paper/2504.16074","citing_paper":"/paper/2607.00248"},"observation_digest":"sha256:b2b0cf9a8328a483d0d5fbb2d9edcec7a8b565defa01ea0e9bb6e4ea01f42d32","observation_id":"323ff441-fde6-4d5e-95e1-41ece48248e6","resolution":{"observed_at":"2026-07-02T19:07:17.387781Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-07-31T06:34:12.847434+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-07-31T06:34:12.847434+00:00","source":"crossref"},{"observed_at":"2026-07-31T06:34:08.642788+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2504.16074","last_updated":"2025-05-18T14:13:34Z","snapshot_observed_at":"2026-07-06T21:13:09.093282Z","submitted_at":"2025-04-22T17:53:29Z","title":"PHYBench: Holistic Evaluation of Physical Perception and Reasoning in Large Language Models","version":2},"cited_work":{"arxiv_id":"2504.16074","doi":"10.48550/arxiv.2504.16074","metadata_source":"arxiv_reference","pith_arxiv_id":"2504.16074","snapshot_observed_at":"2026-07-10T12:15:01.137692Z","title":"Phybench: Holistic evaluation of physical perception and reasoning in large language models","venue":null,"work_id":"d2ae619c-0cbb-45af-8aa0-0a52a9b8dec1","year":2025},"citing_paper":{"arxiv_id":"2607.00276","last_updated":"2026-06-30T23:52:15Z","snapshot_observed_at":"2026-07-07T00:05:59.073754Z","submitted_at":"2026-06-30T23:52:15Z","title":"Testing Frontier Large Language Models' Physics Literacy in Parallel Physical Worlds","version":1},"reference_index":32,"source":"arxiv_source","source_observed_at":"2026-07-02T19:18:43.558804Z"},"links":{"cited_paper":"/paper/2504.16074","citing_paper":"/paper/2607.00276"},"observation_digest":"sha256:0be5d1d8c6969923b8abda088ad663e767f3c5482fd781031320b2e3f661329d","observation_id":"2bb6da02-4f4e-4cd2-bfe9-f1a20a654948","resolution":{"observed_at":"2026-07-02T19:27:18.703824Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-07-31T06:34:12.847434+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-07-31T06:34:12.847434+00:00","source":"crossref"},{"observed_at":"2026-07-31T06:34:08.642788+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2504.16074","last_updated":"2025-05-18T14:13:34Z","snapshot_observed_at":"2026-07-06T21:13:09.093282Z","submitted_at":"2025-04-22T17:53:29Z","title":"PHYBench: Holistic Evaluation of Physical Perception and Reasoning in Large Language Models","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2504.16074","snapshot_observed_at":"2026-07-14T09:17:00.467000Z","title":"Phybench: Holistic evaluation of physical perception and reasoning in large language models,","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2607.10789","last_updated":"2026-07-12T14:42:57Z","snapshot_observed_at":"2026-07-31T04:12:29.287003Z","submitted_at":"2026-07-12T14:42:57Z","title":"Imaging-101: Benchmarking LLM Coding Agents on Scientific Computational Imaging","version":1},"reference_index":39,"source":"pdf_text","source_observed_at":"2026-07-14T09:17:00.467000Z"},"links":{"cited_paper":"/paper/2504.16074","citing_paper":"/paper/2607.10789"},"observation_digest":"sha256:f0c6f69054f267ced91efe3fad34611be78e6b118df997a52ac1590238cc444e","observation_id":"1d9d6a85-fc24-471f-84df-0d55338e06e6","resolution":{"observed_at":"2026-07-14T09:17:00.467000Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2504.16074","last_updated":"2025-05-18T14:13:34Z","snapshot_observed_at":"2026-07-06T21:13:09.093282Z","submitted_at":"2025-04-22T17:53:29Z","title":"PHYBench: Holistic Evaluation of Physical Perception and Reasoning in Large Language Models","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2504.16074","snapshot_observed_at":"2026-08-01T00:52:32.240795Z","title":null,"venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2607.27822","last_updated":"2026-07-30T08:02:33Z","snapshot_observed_at":"2026-08-01T02:04:46.606254Z","submitted_at":"2026-07-30T08:02:33Z","title":"CLVisc Agent for autonomous relativistic hydrodynamics studies","version":1},"reference_index":24,"source":"pdf_text","source_observed_at":"2026-08-01T00:52:32.240795Z"},"links":{"cited_paper":"/paper/2504.16074","citing_paper":"/paper/2607.27822"},"observation_digest":"sha256:18fbae155e5c7803b4425e4df91bd485b9f69ea190aac101d22d4b2f8b3f0828","observation_id":"a9f75970-1fd3-4dcf-9cb4-c4e41ddf764a","resolution":{"observed_at":"2026-08-01T00:52:32.240795Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}}],"links":{"evidence":"/evidence","html":"/paper/2504.16074/citation-record","integrity":"/paper/2504.16074/integrity","json":"/paper/2504.16074/citation-record.json","paper":"/paper/2504.16074"},"outbound":[],"paper":{"arxiv_id":"2504.16074","last_updated":"2025-05-18T14:13:34Z","latest_version":2,"primary_category":"cs.CL","snapshot_observed_at":"2026-07-06T21:13:09.093282Z","submitted_at":"2025-04-22T17:53:29Z","title":"PHYBench: Holistic Evaluation of Physical Perception and Reasoning in Large Language Models"},"reference_resolution":{"displayed":0,"state_counts":{"malformed_identifier":0,"metadata_mismatch":0,"parse_uncertain":0,"unresolved":0,"verified_exact":0,"verified_fuzzy":0},"total_outbound_references":0},"refusal":"A citation records a reference. It does not transfer a finding from one paper to another.","schema":"pith.paper-citation-record.v1","standing_sources":[{"observed_at":"2026-07-31T06:34:12.847434+00:00","source":"crossref"},{"observed_at":"2026-07-31T06:34:08.642788+00:00","source":"retraction_watch"}],"thesis":"As of 1 August 2026, this Paper Citation Record lists 0 of 0 outbound references and 20 inbound Pith citation observations for arXiv:2504.16074."}