{"as_of":"2026-08-04T18:04:00Z","caps":{"database_statements":6,"inbound":100,"outbound":100},"context_digest":"sha256:5fbaab78ecba1d68a4f0f4bd04b8cdf89613e734b291209516445e04bcc40011","coverage":[{"denominator":0,"lane":"reference_resolution","note":"Typed states for the displayed outbound observations.","records_observed":0,"source":"paper_references, paper_reference_links","source_observed_at":null,"state":"measured"},{"denominator":17,"lane":"standing_notices","note":"One-hop event checks from named stored sources.","records_observed":17,"source":"scholarly_work_events, retraction_status_cache","source_observed_at":"2026-08-04T06:34:03.388597+00:00","state":"measured"},{"denominator":17,"lane":"inbound_itemization","note":"Pith citing papers itemized under the disclosed page cap.","records_observed":17,"source":"paper_references, paper_reference_links","source_observed_at":"2026-08-02T10:10:09.242952Z","state":"measured"},{"denominator":1,"lane":"external_citation_measurements","note":"A source-named dated measurement, never combined with another source.","records_observed":0,"source":"arxiv_reference","source_observed_at":"2026-07-04T13:19:50.687532Z","state":"measured"}],"external_citation_measurements":[],"inbound":[{"citation":{"cited_paper":{"arxiv_id":"2406.11802","last_updated":"2024-09-21T06:53:58Z","snapshot_observed_at":"2026-07-06T18:32:23.999019Z","submitted_at":"2024-06-17T17:49:01Z","title":"PhyBench: A Physical Commonsense Benchmark for Evaluating Text-to-Image Models","version":3},"cited_work":{"arxiv_id":"2406.11802","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2406.11802","snapshot_observed_at":"2026-07-04T13:19:50.687532Z","title":"PhyBench: A physical com- monsense benchmark for evaluating text-to-image models","venue":null,"work_id":"b8a57ef1-ee0d-4bb1-8b6d-b9bacbdc8317","year":2024},"citing_paper":{"arxiv_id":"2412.04300","last_updated":"2025-06-04T03:29:18Z","snapshot_observed_at":"2026-08-02T07:59:46.765003Z","submitted_at":"2024-12-05T16:21:01Z","title":"T2I-FactualBench: Benchmarking the Factuality of Text-to-Image Models with Knowledge-Intensive Concepts","version":3},"reference_index":34,"source":"arxiv_source","source_observed_at":"2026-05-23T08:00:12.781392Z"},"links":{"cited_paper":"/paper/2406.11802","citing_paper":"/paper/2412.04300"},"observation_digest":"sha256:fbd292d4afabd9fc69f2a0202caac32154f41b55349083763adff1a3be9c7c2a","observation_id":"a315203a-d6e4-41da-8ea4-77402fa764f2","resolution":{"observed_at":"2026-05-23T08:02:43.359287Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-04T06:34:03.388597+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-04T06:34:03.388597+00:00","source":"crossref"},{"observed_at":"2026-08-04T06:33:57.428241+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2406.11802","last_updated":"2024-09-21T06:53:58Z","snapshot_observed_at":"2026-07-06T18:32:23.999019Z","submitted_at":"2024-06-17T17:49:01Z","title":"PhyBench: A Physical Commonsense Benchmark for Evaluating Text-to-Image Models","version":3},"cited_work":{"arxiv_id":"2406.11802","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2406.11802","snapshot_observed_at":"2026-07-04T13:19:50.687532Z","title":"PhyBench: A physical com- monsense benchmark for evaluating text-to-image models","venue":null,"work_id":"b8a57ef1-ee0d-4bb1-8b6d-b9bacbdc8317","year":2024},"citing_paper":{"arxiv_id":"2503.07265","last_updated":"2026-06-02T17:11:50Z","snapshot_observed_at":"2026-07-06T20:49:51.505079Z","submitted_at":"2025-03-10T12:47:53Z","title":"WISE: A World Knowledge-Informed Semantic Evaluation for Text-to-Image Generation","version":3},"reference_index":31,"source":"pdf_text","source_observed_at":"2026-05-15T16:24:27.407376Z"},"links":{"cited_paper":"/paper/2406.11802","citing_paper":"/paper/2503.07265"},"observation_digest":"sha256:171b0ade4a869e996e7fc49b79fd90a9d6cf59e4a0250817b4f362aaf9613cec","observation_id":"21426675-dd93-4173-83f4-7238c166c849","resolution":{"observed_at":"2026-05-15T16:24:27.688218Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-04T06:34:03.388597+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-04T06:34:03.388597+00:00","source":"crossref"},{"observed_at":"2026-08-04T06:33:57.428241+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2406.11802","last_updated":"2024-09-21T06:53:58Z","snapshot_observed_at":"2026-07-06T18:32:23.999019Z","submitted_at":"2024-06-17T17:49:01Z","title":"PhyBench: A Physical Commonsense Benchmark for Evaluating Text-to-Image Models","version":3},"cited_work":{"arxiv_id":"2406.11802","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2406.11802","snapshot_observed_at":"2026-07-04T13:19:50.687532Z","title":"PhyBench: A physical com- monsense benchmark for evaluating text-to-image models","venue":null,"work_id":"b8a57ef1-ee0d-4bb1-8b6d-b9bacbdc8317","year":2024},"citing_paper":{"arxiv_id":"2507.01908","last_updated":"2026-05-13T06:28:00Z","snapshot_observed_at":"2026-08-02T08:01:40.249812Z","submitted_at":"2025-07-02T17:22:21Z","title":"Reasoning to Edit: Hypothetical Instruction-Based Image Editing with Visual Reasoning","version":3},"reference_index":12,"source":"pdf_text","source_observed_at":"2026-05-19T05:47:19.552825Z"},"links":{"cited_paper":"/paper/2406.11802","citing_paper":"/paper/2507.01908"},"observation_digest":"sha256:e2e28f2e5b20bb5285f657c1bf41fb1d2472b0e8dabca9e155dfc13c8d619c69","observation_id":"fbfd9d6b-b6df-43b1-9b0f-f07e70a87e77","resolution":{"observed_at":"2026-05-19T05:52:07.976102Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-04T06:34:03.388597+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-04T06:34:03.388597+00:00","source":"crossref"},{"observed_at":"2026-08-04T06:33:57.428241+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2406.11802","last_updated":"2024-09-21T06:53:58Z","snapshot_observed_at":"2026-07-06T18:32:23.999019Z","submitted_at":"2024-06-17T17:49:01Z","title":"PhyBench: A Physical Commonsense Benchmark for Evaluating Text-to-Image Models","version":3},"cited_work":{"arxiv_id":"2406.11802","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2406.11802","snapshot_observed_at":"2026-07-04T13:19:50.687532Z","title":"PhyBench: A physical com- monsense benchmark for evaluating text-to-image models","venue":null,"work_id":"b8a57ef1-ee0d-4bb1-8b6d-b9bacbdc8317","year":2024},"citing_paper":{"arxiv_id":"2511.17171","last_updated":"2026-05-22T12:44:18Z","snapshot_observed_at":"2026-08-02T04:31:50.136077Z","submitted_at":"2025-11-21T11:45:22Z","title":"FireScope: Wildfire Risk Raster Prediction with a Chain-of-Thought Oracle","version":5},"reference_index":38,"source":"pdf_text","source_observed_at":"2026-05-17T20:45:15.034196Z"},"links":{"cited_paper":"/paper/2406.11802","citing_paper":"/paper/2511.17171"},"observation_digest":"sha256:dbf5cd434cf87e16e7855fd5cdc971a089c87f02877a44767583dcbae1c2c827","observation_id":"108227d8-9e13-4fff-91c4-2dc5af2d6da8","resolution":{"observed_at":"2026-05-17T20:50:15.312419Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-04T06:34:03.388597+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-04T06:34:03.388597+00:00","source":"crossref"},{"observed_at":"2026-08-04T06:33:57.428241+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2406.11802","last_updated":"2024-09-21T06:53:58Z","snapshot_observed_at":"2026-07-06T18:32:23.999019Z","submitted_at":"2024-06-17T17:49:01Z","title":"PhyBench: A Physical Commonsense Benchmark for Evaluating Text-to-Image Models","version":3},"cited_work":{"arxiv_id":"2406.11802","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2406.11802","snapshot_observed_at":"2026-07-04T13:19:50.687532Z","title":"PhyBench: A physical com- monsense benchmark for evaluating text-to-image models","venue":null,"work_id":"b8a57ef1-ee0d-4bb1-8b6d-b9bacbdc8317","year":2024},"citing_paper":{"arxiv_id":"2511.17171","last_updated":"2026-05-22T12:44:18Z","snapshot_observed_at":"2026-08-02T04:31:50.136077Z","submitted_at":"2025-11-21T11:45:22Z","title":"FireScope: Wildfire Risk Raster Prediction with a Chain-of-Thought Oracle","version":6},"reference_index":38,"source":"pdf_text","source_observed_at":"2026-05-25T07:17:26.381232Z"},"links":{"cited_paper":"/paper/2406.11802","citing_paper":"/paper/2511.17171"},"observation_digest":"sha256:ff0878bc123b7749f1362d573efd7395aff717c1b2bca557de895bce35c7aede","observation_id":"4c6b8518-da53-4f41-ac35-1aa1dd89489c","resolution":{"observed_at":"2026-05-25T07:20:28.791010Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-04T06:34:03.388597+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-04T06:34:03.388597+00:00","source":"crossref"},{"observed_at":"2026-08-04T06:33:57.428241+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2406.11802","last_updated":"2024-09-21T06:53:58Z","snapshot_observed_at":"2026-07-06T18:32:23.999019Z","submitted_at":"2024-06-17T17:49:01Z","title":"PhyBench: A Physical Commonsense Benchmark for Evaluating Text-to-Image Models","version":3},"cited_work":{"arxiv_id":"2406.11802","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2406.11802","snapshot_observed_at":"2026-07-04T13:19:50.687532Z","title":"PhyBench: A physical com- monsense benchmark for evaluating text-to-image models","venue":null,"work_id":"b8a57ef1-ee0d-4bb1-8b6d-b9bacbdc8317","year":2024},"citing_paper":{"arxiv_id":"2512.13281","last_updated":"2026-05-07T09:33:52Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2025-12-15T12:41:23Z","title":"VideoASMR-Bench: Can AI-Generated ASMR Videos Fool VLMs and Humans?","version":4},"reference_index":33,"source":"pdf_text","source_observed_at":"2026-05-16T21:46:43.353305Z"},"links":{"cited_paper":"/paper/2406.11802","citing_paper":"/paper/2512.13281"},"observation_digest":"sha256:a3590d66d98fb90a264342975ef936796f85803c226c5b466df8df612d4171e0","observation_id":"278de7e3-d307-4fa8-8001-dee5ac860fb8","resolution":{"observed_at":"2026-05-16T21:48:34.541702Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-04T06:34:03.388597+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-04T06:34:03.388597+00:00","source":"crossref"},{"observed_at":"2026-08-04T06:33:57.428241+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2406.11802","last_updated":"2024-09-21T06:53:58Z","snapshot_observed_at":"2026-07-06T18:32:23.999019Z","submitted_at":"2024-06-17T17:49:01Z","title":"PhyBench: A Physical Commonsense Benchmark for Evaluating Text-to-Image Models","version":3},"cited_work":{"arxiv_id":"2406.11802","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2406.11802","snapshot_observed_at":"2026-07-04T13:19:50.687532Z","title":"PhyBench: A physical com- monsense benchmark for evaluating text-to-image models","venue":null,"work_id":"b8a57ef1-ee0d-4bb1-8b6d-b9bacbdc8317","year":2024},"citing_paper":{"arxiv_id":"2601.10348","last_updated":"2026-05-21T06:29:24Z","snapshot_observed_at":"2026-08-03T01:57:39.427752Z","submitted_at":"2026-01-15T12:45:05Z","title":"Training-Trajectory-Aware Token Selection","version":2},"reference_index":14,"source":"pdf_text","source_observed_at":"2026-05-22T11:41:21.275802Z"},"links":{"cited_paper":"/paper/2406.11802","citing_paper":"/paper/2601.10348"},"observation_digest":"sha256:cf7e98bda4079436b7730ef598fae4f48acfe22afaed06e0fcc3df2a4e772e83","observation_id":"937aeaef-f27b-434b-b73a-8464971a4497","resolution":{"observed_at":"2026-05-22T11:41:29.533031Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-04T06:34:03.388597+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-04T06:34:03.388597+00:00","source":"crossref"},{"observed_at":"2026-08-04T06:33:57.428241+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2406.11802","last_updated":"2024-09-21T06:53:58Z","snapshot_observed_at":"2026-07-06T18:32:23.999019Z","submitted_at":"2024-06-17T17:49:01Z","title":"PhyBench: A Physical Commonsense Benchmark for Evaluating Text-to-Image Models","version":3},"cited_work":{"arxiv_id":"2406.11802","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2406.11802","snapshot_observed_at":"2026-07-04T13:19:50.687532Z","title":"PhyBench: A physical com- monsense benchmark for evaluating text-to-image models","venue":null,"work_id":"b8a57ef1-ee0d-4bb1-8b6d-b9bacbdc8317","year":2024},"citing_paper":{"arxiv_id":"2604.09415","last_updated":"2026-04-10T15:27:27Z","snapshot_observed_at":"2026-07-06T22:58:17.167367Z","submitted_at":"2026-04-10T15:27:27Z","title":"PhysInOne: Visual Physics Learning and Reasoning in One Suite","version":1},"reference_index":61,"source":"pdf_text","source_observed_at":"2026-05-10T16:39:48.066744Z"},"links":{"cited_paper":"/paper/2406.11802","citing_paper":"/paper/2604.09415"},"observation_digest":"sha256:8cf78e21a6da9369ce67be9c98888ce2958b24b82b261190807f452b3d491c22","observation_id":"ec3782c7-36a8-49fc-81c7-1c9a134ccfa2","resolution":{"observed_at":"2026-05-11T08:25:59.652608Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-04T06:34:03.388597+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-04T06:34:03.388597+00:00","source":"crossref"},{"observed_at":"2026-08-04T06:33:57.428241+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2406.11802","last_updated":"2024-09-21T06:53:58Z","snapshot_observed_at":"2026-07-06T18:32:23.999019Z","submitted_at":"2024-06-17T17:49:01Z","title":"PhyBench: A Physical Commonsense Benchmark for Evaluating Text-to-Image Models","version":3},"cited_work":{"arxiv_id":"2406.11802","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2406.11802","snapshot_observed_at":"2026-07-04T13:19:50.687532Z","title":"PhyBench: A physical com- monsense benchmark for evaluating text-to-image models","venue":null,"work_id":"b8a57ef1-ee0d-4bb1-8b6d-b9bacbdc8317","year":2024},"citing_paper":{"arxiv_id":"2604.28185","last_updated":"2026-04-30T17:59:02Z","snapshot_observed_at":"2026-07-06T23:13:29.310140Z","submitted_at":"2026-04-30T17:59:02Z","title":"Visual Generation in the New Era: An Evolution from Atomic Mapping to Agentic World Modeling","version":1},"reference_index":53,"source":"pdf_text","source_observed_at":"2026-05-07T06:38:04.459129Z"},"links":{"cited_paper":"/paper/2406.11802","citing_paper":"/paper/2604.28185"},"observation_digest":"sha256:9dc79d9065e7be7bcfdab2ee6e1be35aea26c3c4a7b3ded112bdf31b32969167","observation_id":"92504bca-38a7-4ee8-9596-3c11d3bd1cc6","resolution":{"observed_at":"2026-05-12T10:16:29.045357Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-04T06:34:03.388597+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-04T06:34:03.388597+00:00","source":"crossref"},{"observed_at":"2026-08-04T06:33:57.428241+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2406.11802","last_updated":"2024-09-21T06:53:58Z","snapshot_observed_at":"2026-07-06T18:32:23.999019Z","submitted_at":"2024-06-17T17:49:01Z","title":"PhyBench: A Physical Commonsense Benchmark for Evaluating Text-to-Image Models","version":3},"cited_work":{"arxiv_id":"2406.11802","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2406.11802","snapshot_observed_at":"2026-07-04T13:19:50.687532Z","title":"PhyBench: A physical com- monsense benchmark for evaluating text-to-image models","venue":null,"work_id":"b8a57ef1-ee0d-4bb1-8b6d-b9bacbdc8317","year":2024},"citing_paper":{"arxiv_id":"2605.28091","last_updated":"2026-06-25T07:49:03Z","snapshot_observed_at":"2026-07-06T23:37:45.223562Z","submitted_at":"2026-05-27T07:46:43Z","title":"Qwen-Image-Bench: From Generation to Creation in Text-to-Image Evaluation","version":2},"reference_index":9,"source":"pdf_text","source_observed_at":"2026-06-29T14:05:25.988619Z"},"links":{"cited_paper":"/paper/2406.11802","citing_paper":"/paper/2605.28091"},"observation_digest":"sha256:248c3d2641d664d8f7ad4b68bed96efc1498e1752fd3a13a63cbfb9e71c50d63","observation_id":"8537f005-5b6d-4df0-a3a2-ca7a7302589c","resolution":{"observed_at":"2026-06-29T14:13:30.361005Z","resolver_source":"arxiv_id","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-04T06:34:03.388597+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-04T06:34:03.388597+00:00","source":"crossref"},{"observed_at":"2026-08-04T06:33:57.428241+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2406.11802","last_updated":"2024-09-21T06:53:58Z","snapshot_observed_at":"2026-07-06T18:32:23.999019Z","submitted_at":"2024-06-17T17:49:01Z","title":"PhyBench: A Physical Commonsense Benchmark for Evaluating Text-to-Image Models","version":3},"cited_work":{"arxiv_id":"2406.11802","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2406.11802","snapshot_observed_at":"2026-07-04T13:19:50.687532Z","title":"PhyBench: A physical com- monsense benchmark for evaluating text-to-image models","venue":null,"work_id":"b8a57ef1-ee0d-4bb1-8b6d-b9bacbdc8317","year":2024},"citing_paper":{"arxiv_id":"2606.10620","last_updated":"2026-06-09T09:17:55Z","snapshot_observed_at":"2026-08-02T16:01:56.626147Z","submitted_at":"2026-06-09T09:17:55Z","title":"Can Image Models Imagine Time? ImageTime: A Novel Benchmark for Probing Visual World Modeling Through Spatiotemporal Consistency","version":1},"reference_index":22,"source":"pdf_text","source_observed_at":"2026-06-27T13:34:25.079037Z"},"links":{"cited_paper":"/paper/2406.11802","citing_paper":"/paper/2606.10620"},"observation_digest":"sha256:fb371520ac6f34c3c75766ede1f62cfef6c98b2221860f088f0039339dc9b266","observation_id":"b58dbe3f-270d-49de-b664-4c3ada2daaf3","resolution":{"observed_at":"2026-07-03T04:57:38.079383Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-04T06:34:03.388597+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-04T06:34:03.388597+00:00","source":"crossref"},{"observed_at":"2026-08-04T06:33:57.428241+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2406.11802","last_updated":"2024-09-21T06:53:58Z","snapshot_observed_at":"2026-07-06T18:32:23.999019Z","submitted_at":"2024-06-17T17:49:01Z","title":"PhyBench: A Physical Commonsense Benchmark for Evaluating Text-to-Image Models","version":3},"cited_work":{"arxiv_id":"2406.11802","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2406.11802","snapshot_observed_at":"2026-07-04T13:19:50.687532Z","title":"PhyBench: A physical com- monsense benchmark for evaluating text-to-image models","venue":null,"work_id":"b8a57ef1-ee0d-4bb1-8b6d-b9bacbdc8317","year":2024},"citing_paper":{"arxiv_id":"2606.26738","last_updated":"2026-07-15T11:15:44Z","snapshot_observed_at":"2026-08-02T10:10:07.945306Z","submitted_at":"2026-06-25T08:20:34Z","title":"Do Image Editing Models Understand Lighting?","version":1},"reference_index":31,"source":"pdf_text","source_observed_at":"2026-06-26T05:29:42.024146Z"},"links":{"cited_paper":"/paper/2406.11802","citing_paper":"/paper/2606.26738"},"observation_digest":"sha256:95d60f55956eb03bfd260b60da7ac2684a82e4e245822d1d8a9f221e6f79ac5a","observation_id":"65b97e57-5a21-455f-8824-4887aafdc8b0","resolution":{"observed_at":"2026-07-04T13:09:50.412488Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-04T06:34:03.388597+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-04T06:34:03.388597+00:00","source":"crossref"},{"observed_at":"2026-08-04T06:33:57.428241+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2406.11802","last_updated":"2024-09-21T06:53:58Z","snapshot_observed_at":"2026-07-06T18:32:23.999019Z","submitted_at":"2024-06-17T17:49:01Z","title":"PhyBench: A Physical Commonsense Benchmark for Evaluating Text-to-Image Models","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2406.11802","snapshot_observed_at":"2026-08-02T10:10:09.242952Z","title":"Phybench: A physical commonsense benchmark for evaluating text-to-image models.arXiv preprint arXiv:2406.11802, 2024","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2606.26738","last_updated":"2026-07-15T11:15:44Z","snapshot_observed_at":"2026-08-02T10:10:07.945306Z","submitted_at":"2026-06-25T08:20:34Z","title":"Do Image Editing Models Understand Lighting?","version":2},"reference_index":31,"source":"pdf_text","source_observed_at":"2026-08-02T10:10:09.242952Z"},"links":{"cited_paper":"/paper/2406.11802","citing_paper":"/paper/2606.26738"},"observation_digest":"sha256:5dfdf5fe3bf7f50ed811dfef821f6e41c32b1b99a89f66bb15349b254859eab9","observation_id":"bd01808f-6143-4d2b-b265-3ad691bd2e79","resolution":{"observed_at":"2026-08-02T10:10:09.242952Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2406.11802","last_updated":"2024-09-21T06:53:58Z","snapshot_observed_at":"2026-07-06T18:32:23.999019Z","submitted_at":"2024-06-17T17:49:01Z","title":"PhyBench: A Physical Commonsense Benchmark for Evaluating Text-to-Image Models","version":3},"cited_work":{"arxiv_id":"2406.11802","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2406.11802","snapshot_observed_at":"2026-07-04T13:19:50.687532Z","title":"PhyBench: A physical com- monsense benchmark for evaluating text-to-image models","venue":null,"work_id":"b8a57ef1-ee0d-4bb1-8b6d-b9bacbdc8317","year":2024},"citing_paper":{"arxiv_id":"2606.26907","last_updated":"2026-06-26T14:10:09Z","snapshot_observed_at":"2026-07-07T00:01:10.711037Z","submitted_at":"2026-06-25T11:40:12Z","title":"Qwen-Image-Agent: Bridging the Context Gap in Real-World Image Generation","version":1},"reference_index":9,"source":"pdf_text","source_observed_at":"2026-06-26T05:19:21.433049Z"},"links":{"cited_paper":"/paper/2406.11802","citing_paper":"/paper/2606.26907"},"observation_digest":"sha256:2e88e1c238522265708d08fb1e3fac83e41d31de6df6ebeeba0f9dd1d5bca273","observation_id":"4587c0ef-3d86-4507-bf12-af109b8f1e87","resolution":{"observed_at":"2026-07-04T13:19:50.689189Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-04T06:34:03.388597+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-04T06:34:03.388597+00:00","source":"crossref"},{"observed_at":"2026-08-04T06:33:57.428241+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2406.11802","last_updated":"2024-09-21T06:53:58Z","snapshot_observed_at":"2026-07-06T18:32:23.999019Z","submitted_at":"2024-06-17T17:49:01Z","title":"PhyBench: A Physical Commonsense Benchmark for Evaluating Text-to-Image Models","version":3},"cited_work":{"arxiv_id":"2406.11802","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2406.11802","snapshot_observed_at":"2026-07-04T13:19:50.687532Z","title":"PhyBench: A physical com- monsense benchmark for evaluating text-to-image models","venue":null,"work_id":"b8a57ef1-ee0d-4bb1-8b6d-b9bacbdc8317","year":2024},"citing_paper":{"arxiv_id":"2606.26907","last_updated":"2026-06-26T14:10:09Z","snapshot_observed_at":"2026-07-07T00:01:10.711037Z","submitted_at":"2026-06-25T11:40:12Z","title":"Qwen-Image-Agent: Bridging the Context Gap in Real-World Image Generation","version":2},"reference_index":9,"source":"pdf_text","source_observed_at":"2026-06-29T04:58:40.065005Z"},"links":{"cited_paper":"/paper/2406.11802","citing_paper":"/paper/2606.26907"},"observation_digest":"sha256:29bb2006a412abbccb58ad1b7e3cdd197b8d856f7e73fe6bd553633b27cc6f01","observation_id":"446a4e6c-0292-495f-835a-d03e4870eb93","resolution":{"observed_at":"2026-06-29T19:03:52.063184Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-04T06:34:03.388597+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-04T06:34:03.388597+00:00","source":"crossref"},{"observed_at":"2026-08-04T06:33:57.428241+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2406.11802","last_updated":"2024-09-21T06:53:58Z","snapshot_observed_at":"2026-07-06T18:32:23.999019Z","submitted_at":"2024-06-17T17:49:01Z","title":"PhyBench: A Physical Commonsense Benchmark for Evaluating Text-to-Image Models","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2406.11802","snapshot_observed_at":"2026-07-12T02:16:07.323950Z","title":"Phybench: A physical common- sense benchmark for evaluating text-to-image models,","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2607.03470","last_updated":"2026-07-03T16:32:02Z","snapshot_observed_at":"2026-08-02T07:12:40.660479Z","submitted_at":"2026-07-03T16:32:02Z","title":"PhysMirror: Physics-Aware Mirror Object Generation","version":1},"reference_index":14,"source":"pdf_text","source_observed_at":"2026-07-12T02:16:07.323950Z"},"links":{"cited_paper":"/paper/2406.11802","citing_paper":"/paper/2607.03470"},"observation_digest":"sha256:6bdc98fe41f420f24074b3bce861cc277d84dcefb20bff3bf4f8293de5856163","observation_id":"1731c49e-7f2a-45ef-b72c-f05297751bb9","resolution":{"observed_at":"2026-07-12T02:16:07.323950Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2406.11802","last_updated":"2024-09-21T06:53:58Z","snapshot_observed_at":"2026-07-06T18:32:23.999019Z","submitted_at":"2024-06-17T17:49:01Z","title":"PhyBench: A Physical Commonsense Benchmark for Evaluating Text-to-Image Models","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2406.11802","snapshot_observed_at":"2026-08-01T01:54:37.844943Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2607.25641","last_updated":"2026-07-28T12:27:25Z","snapshot_observed_at":"2026-08-03T00:49:38.920938Z","submitted_at":"2026-07-28T12:27:25Z","title":"OmniPhys: Knowledge-Graph-Driven Benchmarking and Collective Optimization for Physical Commonsense in Text-to-Image Generation","version":1},"reference_index":19,"source":"pdf_text","source_observed_at":"2026-08-01T01:54:37.844943Z"},"links":{"cited_paper":"/paper/2406.11802","citing_paper":"/paper/2607.25641"},"observation_digest":"sha256:0fde2aeb3b400e6aed4c5b5f91719601efbae4ad2d27db30d218adaf5bc4e02b","observation_id":"07be468f-00ae-4b8b-85f2-d78c96263e79","resolution":{"observed_at":"2026-08-01T01:54:37.844943Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}}],"links":{"evidence":"/evidence","html":"/paper/2406.11802/citation-record","integrity":"/paper/2406.11802/integrity","json":"/paper/2406.11802/citation-record.json","paper":"/paper/2406.11802"},"outbound":[],"paper":{"arxiv_id":"2406.11802","last_updated":"2024-09-21T06:53:58Z","latest_version":3,"primary_category":"cs.CV","snapshot_observed_at":"2026-07-06T18:32:23.999019Z","submitted_at":"2024-06-17T17:49:01Z","title":"PhyBench: A Physical Commonsense Benchmark for Evaluating Text-to-Image Models"},"reference_resolution":{"displayed":0,"state_counts":{"malformed_identifier":0,"metadata_mismatch":0,"parse_uncertain":0,"unresolved":0,"verified_exact":0,"verified_fuzzy":0},"total_outbound_references":0},"refusal":"A citation records a reference. It does not transfer a finding from one paper to another.","schema":"pith.paper-citation-record.v1","standing_sources":[{"observed_at":"2026-08-04T06:34:03.388597+00:00","source":"crossref"},{"observed_at":"2026-08-04T06:33:57.428241+00:00","source":"retraction_watch"}],"thesis":"As of 4 August 2026, this Paper Citation Record lists 0 of 0 outbound references and 17 inbound Pith citation observations for arXiv:2406.11802."}