{"as_of":"2026-08-07T19:45:00Z","caps":{"database_statements":6,"inbound":100,"outbound":100},"context_digest":"sha256:26b4b113471beb66da3e02b01afc3c3ee703f5e8d6ff67646b570c63b41f3bc6","coverage":[{"denominator":0,"lane":"reference_resolution","note":"Typed states for the displayed outbound observations.","records_observed":0,"source":"paper_references, paper_reference_links","source_observed_at":null,"state":"measured"},{"denominator":15,"lane":"standing_notices","note":"One-hop event checks from named stored sources.","records_observed":15,"source":"scholarly_work_events, retraction_status_cache","source_observed_at":"2026-08-07T06:34:17.273281+00:00","state":"measured"},{"denominator":15,"lane":"inbound_itemization","note":"Pith citing papers itemized under the disclosed page cap.","records_observed":15,"source":"paper_references, paper_reference_links","source_observed_at":"2026-08-07T14:29:14.789184Z","state":"measured"},{"denominator":1,"lane":"external_citation_measurements","note":"A source-named dated measurement, never combined with another source.","records_observed":0,"source":"arxiv_reference","source_observed_at":"2026-07-04T13:09:50.423322Z","state":"measured"}],"external_citation_measurements":[],"inbound":[{"citation":{"cited_paper":{"arxiv_id":"2406.07546","last_updated":"2024-08-12T19:33:52Z","snapshot_observed_at":"2026-07-06T18:29:08.586253Z","submitted_at":"2024-06-11T17:59:48Z","title":"Commonsense-T2I Challenge: Can Text-to-Image Generation Models Understand Commonsense?","version":2},"cited_work":{"arxiv_id":"2406.07546","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2406.07546","snapshot_observed_at":"2026-07-04T13:09:50.423322Z","title":"Commonsense-t2i challenge: Can text-to-image generation models understand commonsense?","venue":null,"work_id":"4e3b9283-2aa2-4823-b228-dc40ffd87114","year":2024},"citing_paper":{"arxiv_id":"2410.05363","last_updated":"2024-10-07T17:56:04Z","snapshot_observed_at":"2026-07-06T19:29:14.335016Z","submitted_at":"2024-10-07T17:56:04Z","title":"Towards World Simulator: Crafting Physical Commonsense-Based Benchmark for Video Generation","version":1},"reference_index":8,"source":"arxiv_source","source_observed_at":"2026-05-18T14:39:59.870039Z"},"links":{"cited_paper":"/paper/2406.07546","citing_paper":"/paper/2410.05363"},"observation_digest":"sha256:bab0b96026f925a124b752a6d641d779b1154fcdc689b40b3b920e8a84d8da72","observation_id":"54752231-047b-4069-be2e-242b0b1df1de","resolution":{"observed_at":"2026-05-18T14:40:00.094694Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2406.07546","last_updated":"2024-08-12T19:33:52Z","snapshot_observed_at":"2026-07-06T18:29:08.586253Z","submitted_at":"2024-06-11T17:59:48Z","title":"Commonsense-T2I Challenge: Can Text-to-Image Generation Models Understand Commonsense?","version":2},"cited_work":{"arxiv_id":"2406.07546","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2406.07546","snapshot_observed_at":"2026-07-04T13:09:50.423322Z","title":"Commonsense-t2i challenge: Can text-to-image generation models understand commonsense?","venue":null,"work_id":"4e3b9283-2aa2-4823-b228-dc40ffd87114","year":2024},"citing_paper":{"arxiv_id":"2412.04300","last_updated":"2025-06-04T03:29:18Z","snapshot_observed_at":"2026-08-02T07:59:46.765003Z","submitted_at":"2024-12-05T16:21:01Z","title":"T2I-FactualBench: Benchmarking the Factuality of Text-to-Image Models with Knowledge-Intensive Concepts","version":3},"reference_index":10,"source":"arxiv_source","source_observed_at":"2026-05-23T08:00:12.781392Z"},"links":{"cited_paper":"/paper/2406.07546","citing_paper":"/paper/2412.04300"},"observation_digest":"sha256:f4a24f980527d0f526ad27f1e188fb958a6fd02ec29a6be177af42b77f1b3d20","observation_id":"afdcc556-0802-4452-b94d-39fe4267e5f9","resolution":{"observed_at":"2026-05-23T08:02:43.281911Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2406.07546","last_updated":"2024-08-12T19:33:52Z","snapshot_observed_at":"2026-07-06T18:29:08.586253Z","submitted_at":"2024-06-11T17:59:48Z","title":"Commonsense-T2I Challenge: Can Text-to-Image Generation Models Understand Commonsense?","version":2},"cited_work":{"arxiv_id":"2406.07546","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2406.07546","snapshot_observed_at":"2026-07-04T13:09:50.423322Z","title":"Commonsense-t2i challenge: Can text-to-image generation models understand commonsense?","venue":null,"work_id":"4e3b9283-2aa2-4823-b228-dc40ffd87114","year":2024},"citing_paper":{"arxiv_id":"2503.07265","last_updated":"2026-06-02T17:11:50Z","snapshot_observed_at":"2026-08-07T17:17:00.060047Z","submitted_at":"2025-03-10T12:47:53Z","title":"WISE: A World Knowledge-Informed Semantic Evaluation for Text-to-Image Generation","version":3},"reference_index":9,"source":"pdf_text","source_observed_at":"2026-05-15T16:24:27.407376Z"},"links":{"cited_paper":"/paper/2406.07546","citing_paper":"/paper/2503.07265"},"observation_digest":"sha256:22a45f44867873960f23e2639ece5d69700472aa61935e969fe3f276dc38c015","observation_id":"adf6ea4b-8d45-4b8b-85df-de20ad9611e0","resolution":{"observed_at":"2026-05-15T16:24:27.579185Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2406.07546","last_updated":"2024-08-12T19:33:52Z","snapshot_observed_at":"2026-07-06T18:29:08.586253Z","submitted_at":"2024-06-11T17:59:48Z","title":"Commonsense-T2I Challenge: Can Text-to-Image Generation Models Understand Commonsense?","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2406.07546","snapshot_observed_at":"2026-08-07T14:29:14.789184Z","title":"Commonsense-t2i challenge: Can text-to-image generation models understand commonsense? arXiv preprint arXiv:2406.07546, 2024","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2505.18730","last_updated":"2025-05-24T14:56:09Z","snapshot_observed_at":"2026-08-07T17:06:29.055596Z","submitted_at":"2025-05-24T14:56:09Z","title":"Align Beyond Prompts: Evaluating World Knowledge Alignment in Text-to-Image Generation","version":1},"reference_index":8,"source":"pdf_text","source_observed_at":"2026-08-07T14:29:14.789184Z"},"links":{"cited_paper":"/paper/2406.07546","citing_paper":"/paper/2505.18730"},"observation_digest":"sha256:22be70bdf30377e02c9a3929ac5e378e2d2ed801772b5efd83e5f3885dbb5c0f","observation_id":"cc46d3a9-9e75-4f0f-83e2-6b198eacd0d1","resolution":{"observed_at":"2026-08-07T14:29:14.789184Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2406.07546","last_updated":"2024-08-12T19:33:52Z","snapshot_observed_at":"2026-07-06T18:29:08.586253Z","submitted_at":"2024-06-11T17:59:48Z","title":"Commonsense-T2I Challenge: Can Text-to-Image Generation Models Understand Commonsense?","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2406.07546","snapshot_observed_at":"2026-08-07T14:17:32.854514Z","title":"Commonsense-t2i challenge: Can text-to-image generation models understand commonsense? arXiv preprint arXiv:2406.07546, 2024","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2505.19415","last_updated":"2025-05-27T20:10:09Z","snapshot_observed_at":"2026-08-07T14:12:18.054397Z","submitted_at":"2025-05-26T02:07:24Z","title":"MMIG-Bench: Towards Comprehensive and Explainable Evaluation of Multi-Modal Image Generation Models","version":2},"reference_index":8,"source":"pdf_text","source_observed_at":"2026-08-07T14:17:32.854514Z"},"links":{"cited_paper":"/paper/2406.07546","citing_paper":"/paper/2505.19415"},"observation_digest":"sha256:ecf2e75aaebcde31d18efba87c262f064355e3472c9a336eb65e0196675dc19d","observation_id":"9b784e3f-2441-487a-a237-cf01ea15bdca","resolution":{"observed_at":"2026-08-07T14:17:32.854514Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2406.07546","last_updated":"2024-08-12T19:33:52Z","snapshot_observed_at":"2026-07-06T18:29:08.586253Z","submitted_at":"2024-06-11T17:59:48Z","title":"Commonsense-T2I Challenge: Can Text-to-Image Generation Models Understand Commonsense?","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2406.07546","snapshot_observed_at":"2026-08-07T12:48:50.711956Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2505.23493","last_updated":"2025-05-29T14:43:46Z","snapshot_observed_at":"2026-08-07T12:42:36.882299Z","submitted_at":"2025-05-29T14:43:46Z","title":"R2I-Bench: Benchmarking Reasoning-Driven Text-to-Image Generation","version":1},"reference_index":8,"source":"arxiv_source","source_observed_at":"2026-08-07T12:48:50.711956Z"},"links":{"cited_paper":"/paper/2406.07546","citing_paper":"/paper/2505.23493"},"observation_digest":"sha256:0d58a2c40f55eeb6c712c342234e1a03b0cf2754b9617fa8b3adce3e955ca8ad","observation_id":"2a1e1bb1-ba1a-4820-90c6-e681d26cbd59","resolution":{"observed_at":"2026-08-07T12:48:50.711956Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2406.07546","last_updated":"2024-08-12T19:33:52Z","snapshot_observed_at":"2026-07-06T18:29:08.586253Z","submitted_at":"2024-06-11T17:59:48Z","title":"Commonsense-T2I Challenge: Can Text-to-Image Generation Models Understand Commonsense?","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2406.07546","snapshot_observed_at":"2026-08-07T12:21:21.501770Z","title":"Commonsense-t2i challenge: Can text-to-image generation models understand commonsense?arXiv preprint arXiv:2406.07546, 2024","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2505.24870","last_updated":"2025-06-06T14:51:40Z","snapshot_observed_at":"2026-08-07T12:10:22.188975Z","submitted_at":"2025-05-30T17:59:26Z","title":"GenSpace: Benchmarking Spatially-Aware Image Generation","version":2},"reference_index":12,"source":"pdf_text","source_observed_at":"2026-08-07T12:21:21.501770Z"},"links":{"cited_paper":"/paper/2406.07546","citing_paper":"/paper/2505.24870"},"observation_digest":"sha256:381281e232512833e16bfb04556aa6d7d5f3209115f46c157c3187327afaafe3","observation_id":"5042c595-d8b2-4957-8d6e-e989a06c2c7b","resolution":{"observed_at":"2026-08-07T12:21:21.501770Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2406.07546","last_updated":"2024-08-12T19:33:52Z","snapshot_observed_at":"2026-07-06T18:29:08.586253Z","submitted_at":"2024-06-11T17:59:48Z","title":"Commonsense-T2I Challenge: Can Text-to-Image Generation Models Understand Commonsense?","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2406.07546","snapshot_observed_at":"2026-08-07T05:25:44.095229Z","title":"Commonsense-t2i challenge: Can text-to-image generation models understand commonsense? arXiv preprint arXiv:2406.07546, 2024","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2506.07977","last_updated":"2025-06-26T15:47:09Z","snapshot_observed_at":"2026-08-07T05:18:51.807244Z","submitted_at":"2025-06-09T17:50:21Z","title":"OneIG-Bench: Omni-dimensional Nuanced Evaluation for Image Generation","version":3},"reference_index":22,"source":"pdf_text","source_observed_at":"2026-08-07T05:25:44.095229Z"},"links":{"cited_paper":"/paper/2406.07546","citing_paper":"/paper/2506.07977"},"observation_digest":"sha256:26078d6270f6420aa8bb65edcc36cf5b48cc1661ff357e3f0c40c7a72e52b51c","observation_id":"16ad961e-f0ce-402b-b9f1-d84bc5b2ff81","resolution":{"observed_at":"2026-08-07T05:25:44.095229Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2406.07546","last_updated":"2024-08-12T19:33:52Z","snapshot_observed_at":"2026-07-06T18:29:08.586253Z","submitted_at":"2024-06-11T17:59:48Z","title":"Commonsense-T2I Challenge: Can Text-to-Image Generation Models Understand Commonsense?","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2406.07546","snapshot_observed_at":"2026-08-06T20:32:59.658017Z","title":"Commonsense-t2i challenge: Can text-to- image generation models understand commonsense? arXiv preprint arXiv:2406.07546, 2024","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2507.02664","last_updated":"2025-07-07T08:00:38Z","snapshot_observed_at":"2026-08-07T07:10:52.209513Z","submitted_at":"2025-07-03T14:26:31Z","title":"AIGI-Holmes: Towards Explainable and Generalizable AI-Generated Image Detection via Multimodal Large Language Models","version":2},"reference_index":26,"source":"pdf_text","source_observed_at":"2026-08-06T20:32:59.658017Z"},"links":{"cited_paper":"/paper/2406.07546","citing_paper":"/paper/2507.02664"},"observation_digest":"sha256:965136c49ebf5015a1234b66585b7f041844f9b037567f54eaefcad532d9d5c9","observation_id":"ba2461e4-cd1c-4b4a-9292-c133b1004b66","resolution":{"observed_at":"2026-08-06T20:32:59.658017Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2406.07546","last_updated":"2024-08-12T19:33:52Z","snapshot_observed_at":"2026-07-06T18:29:08.586253Z","submitted_at":"2024-06-11T17:59:48Z","title":"Commonsense-T2I Challenge: Can Text-to-Image Generation Models Understand Commonsense?","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2406.07546","snapshot_observed_at":"2026-08-04T18:48:03.655845Z","title":"Commonsense-t2i challenge: Can text-to-image generation models understand commonsense?arXiv preprint arXiv:2406.07546, 2024","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2509.09680","last_updated":"2025-09-11T17:59:59Z","snapshot_observed_at":"2026-08-04T18:48:02.197558Z","submitted_at":"2025-09-11T17:59:59Z","title":"FLUX-Reason-6M & PRISM-Bench: A Million-Scale Text-to-Image Reasoning Dataset and Comprehensive Benchmark","version":1},"reference_index":20,"source":"pdf_text","source_observed_at":"2026-08-04T18:48:03.655845Z"},"links":{"cited_paper":"/paper/2406.07546","citing_paper":"/paper/2509.09680"},"observation_digest":"sha256:3f54511d47ba161ae4274a8a8cdb3ad9e98e92df82404dd29e526aa4dcdabd2f","observation_id":"66c32426-e674-488e-8e1d-aaad6cc0f225","resolution":{"observed_at":"2026-08-04T18:48:03.655845Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2406.07546","last_updated":"2024-08-12T19:33:52Z","snapshot_observed_at":"2026-07-06T18:29:08.586253Z","submitted_at":"2024-06-11T17:59:48Z","title":"Commonsense-T2I Challenge: Can Text-to-Image Generation Models Understand Commonsense?","version":2},"cited_work":{"arxiv_id":"2406.07546","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2406.07546","snapshot_observed_at":"2026-07-04T13:09:50.423322Z","title":"Commonsense-t2i challenge: Can text-to-image generation models understand commonsense?","venue":null,"work_id":"4e3b9283-2aa2-4823-b228-dc40ffd87114","year":2024},"citing_paper":{"arxiv_id":"2606.26738","last_updated":"2026-07-15T11:15:44Z","snapshot_observed_at":"2026-08-02T10:10:07.945306Z","submitted_at":"2026-06-25T08:20:34Z","title":"Do Image Editing Models Understand Lighting?","version":1},"reference_index":10,"source":"pdf_text","source_observed_at":"2026-06-26T05:29:42.024146Z"},"links":{"cited_paper":"/paper/2406.07546","citing_paper":"/paper/2606.26738"},"observation_digest":"sha256:809a88474b24b787ba87d74e1d22d99fae94e722f0530b782747d631bcb1d63c","observation_id":"eb48bd77-e076-41ca-a7fe-926222af3742","resolution":{"observed_at":"2026-07-04T13:09:50.425449Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2406.07546","last_updated":"2024-08-12T19:33:52Z","snapshot_observed_at":"2026-07-06T18:29:08.586253Z","submitted_at":"2024-06-11T17:59:48Z","title":"Commonsense-T2I Challenge: Can Text-to-Image Generation Models Understand Commonsense?","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2406.07546","snapshot_observed_at":"2026-08-02T10:10:09.175457Z","title":"Commonsense-t2i challenge: Can text-to-image generation models understand commonsense?arXiv preprint arXiv:2406.07546, 2024","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2606.26738","last_updated":"2026-07-15T11:15:44Z","snapshot_observed_at":"2026-08-02T10:10:07.945306Z","submitted_at":"2026-06-25T08:20:34Z","title":"Do Image Editing Models Understand Lighting?","version":2},"reference_index":10,"source":"pdf_text","source_observed_at":"2026-08-02T10:10:09.175457Z"},"links":{"cited_paper":"/paper/2406.07546","citing_paper":"/paper/2606.26738"},"observation_digest":"sha256:90d56b25e9b2a83c600ea4e4eb45a4789ffc7f931339db3d017396b7c7546bac","observation_id":"9bee8750-2446-469c-90c2-3cc1218ad00e","resolution":{"observed_at":"2026-08-02T10:10:09.175457Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2406.07546","last_updated":"2024-08-12T19:33:52Z","snapshot_observed_at":"2026-07-06T18:29:08.586253Z","submitted_at":"2024-06-11T17:59:48Z","title":"Commonsense-T2I Challenge: Can Text-to-Image Generation Models Understand Commonsense?","version":2},"cited_work":{"arxiv_id":"2406.07546","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2406.07546","snapshot_observed_at":"2026-07-04T13:09:50.423322Z","title":"Commonsense-t2i challenge: Can text-to-image generation models understand commonsense?","venue":null,"work_id":"4e3b9283-2aa2-4823-b228-dc40ffd87114","year":2024},"citing_paper":{"arxiv_id":"2606.30262","last_updated":"2026-06-29T13:09:39Z","snapshot_observed_at":"2026-08-06T12:26:22.536829Z","submitted_at":"2026-06-29T13:09:39Z","title":"Intermediate Text Representation Guided Text-to-Image Generation for Enhancing One-and-Only Alignment","version":1},"reference_index":13,"source":"pdf_text","source_observed_at":"2026-06-30T06:04:54.816368Z"},"links":{"cited_paper":"/paper/2406.07546","citing_paper":"/paper/2606.30262"},"observation_digest":"sha256:1dd0badc347d0c3f633e0ee35d3fb543f2f61ee98e75cf9a091012c121bc7663","observation_id":"611ce9ff-7846-466d-a374-78755633f1e1","resolution":{"observed_at":"2026-06-30T08:24:27.106668Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2406.07546","last_updated":"2024-08-12T19:33:52Z","snapshot_observed_at":"2026-07-06T18:29:08.586253Z","submitted_at":"2024-06-11T17:59:48Z","title":"Commonsense-T2I Challenge: Can Text-to-Image Generation Models Understand Commonsense?","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2406.07546","snapshot_observed_at":"2026-08-01T01:54:37.826437Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2607.25641","last_updated":"2026-07-28T12:27:25Z","snapshot_observed_at":"2026-08-05T01:41:17.552700Z","submitted_at":"2026-07-28T12:27:25Z","title":"OmniPhys: Knowledge-Graph-Driven Benchmarking and Collective Optimization for Physical Commonsense in Text-to-Image Generation","version":1},"reference_index":10,"source":"pdf_text","source_observed_at":"2026-08-01T01:54:37.826437Z"},"links":{"cited_paper":"/paper/2406.07546","citing_paper":"/paper/2607.25641"},"observation_digest":"sha256:f0a3d173d6e328df3d89c8660b4ad5e263951a06de8f245b2ffa309620b16591","observation_id":"64f394f1-01c4-4210-9b78-a5196d16eef7","resolution":{"observed_at":"2026-08-01T01:54:37.826437Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2406.07546","last_updated":"2024-08-12T19:33:52Z","snapshot_observed_at":"2026-07-06T18:29:08.586253Z","submitted_at":"2024-06-11T17:59:48Z","title":"Commonsense-T2I Challenge: Can Text-to-Image Generation Models Understand Commonsense?","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2406.07546","snapshot_observed_at":"2026-08-07T00:15:38.364743Z","title":"Commonsense-t2i challenge: Can text-to- image generation models understand commonsense?, 2024","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2608.04436","last_updated":"2026-08-05T04:26:19Z","snapshot_observed_at":"2026-08-07T19:19:22.437336Z","submitted_at":"2026-08-05T04:26:19Z","title":"ToolArtist: Tool-Using Unified Multimodal Models for Agentic Image Generation","version":1},"reference_index":11,"source":"pdf_text","source_observed_at":"2026-08-07T00:15:38.364743Z"},"links":{"cited_paper":"/paper/2406.07546","citing_paper":"/paper/2608.04436"},"observation_digest":"sha256:b924a6c7883a08242f5d105763c65454ce8466523a02769c0e16b632c8b2a5c9","observation_id":"b8e6eff9-412a-43a9-9940-376db9b11c24","resolution":{"observed_at":"2026-08-07T00:15:38.364743Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}}],"links":{"evidence":"/evidence","html":"/paper/2406.07546/citation-record","integrity":"/paper/2406.07546/integrity","json":"/paper/2406.07546/citation-record.json","paper":"/paper/2406.07546"},"outbound":[],"paper":{"arxiv_id":"2406.07546","last_updated":"2024-08-12T19:33:52Z","latest_version":2,"primary_category":"cs.CV","snapshot_observed_at":"2026-07-06T18:29:08.586253Z","submitted_at":"2024-06-11T17:59:48Z","title":"Commonsense-T2I Challenge: Can Text-to-Image Generation Models Understand Commonsense?"},"reference_resolution":{"displayed":0,"state_counts":{"malformed_identifier":0,"metadata_mismatch":0,"parse_uncertain":0,"unresolved":0,"verified_exact":0,"verified_fuzzy":0},"total_outbound_references":0},"refusal":"A citation records a reference. It does not transfer a finding from one paper to another.","schema":"pith.paper-citation-record.v1","standing_sources":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"thesis":"As of 7 August 2026, this Paper Citation Record lists 0 of 0 outbound references and 15 inbound Pith citation observations for arXiv:2406.07546."}