{"as_of":"2026-08-02T16:02:00Z","caps":{"database_statements":6,"inbound":100,"outbound":100},"context_digest":"sha256:7f4bcbf58420e12d727de98e5e0dfdac997508c726704c684ac5e95b009e250b","coverage":[{"denominator":64,"lane":"reference_resolution","note":"Typed states for the displayed outbound observations.","records_observed":64,"source":"paper_references, paper_reference_links","source_observed_at":"2026-05-12T05:17:30.010064Z","state":"measured"},{"denominator":65,"lane":"standing_notices","note":"One-hop event checks from named stored sources.","records_observed":65,"source":"scholarly_work_events, retraction_status_cache","source_observed_at":"2026-08-02T06:30:47.504484+00:00","state":"measured"},{"denominator":1,"lane":"inbound_itemization","note":"Pith citing papers itemized under the disclosed page cap.","records_observed":1,"source":"paper_references, paper_reference_links","source_observed_at":"2026-07-31T01:50:30.653972Z","state":"measured"},{"denominator":1,"lane":"external_citation_measurements","note":"A source-named dated measurement, never combined with another source.","records_observed":0,"source":"cited_works","source_observed_at":null,"state":"measured"}],"external_citation_measurements":[],"inbound":[{"citation":{"cited_paper":{"arxiv_id":"2605.10806","last_updated":"2026-05-11T16:30:51Z","snapshot_observed_at":"2026-07-06T23:22:42.512039Z","submitted_at":"2026-05-11T16:30:51Z","title":"PhyGround: Benchmarking Physical Reasoning in Generative World Models","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2605.10806","snapshot_observed_at":"2026-07-31T01:50:30.653972Z","title":"Phyground: Benchmarking physical reasoning in generative world models","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2607.28624","last_updated":"2026-07-30T17:59:46Z","snapshot_observed_at":"2026-08-02T15:25:10.761571Z","submitted_at":"2026-07-30T17:59:46Z","title":"PhiZero: A World Model Built Around Physical Language","version":1},"reference_index":29,"source":"pdf_text","source_observed_at":"2026-07-31T01:50:30.653972Z"},"links":{"cited_paper":"/paper/2605.10806","citing_paper":"/paper/2607.28624"},"observation_digest":"sha256:a0493359f27126a6d78bf4f9955172236889530ead7103adeccdf8890d1da085","observation_id":"ceff8e6a-3366-4083-b35a-93810768717e","resolution":{"observed_at":"2026-07-31T01:50:30.653972Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}}],"links":{"evidence":"/evidence","html":"/paper/2605.10806/citation-record","integrity":"/paper/2605.10806/integrity","json":"/paper/2605.10806/citation-record.json","paper":"/paper/2605.10806"},"outbound":[{"citation":{"cited_paper":{"arxiv_id":"2511.21631","last_updated":"2025-11-27T12:16:54Z","snapshot_observed_at":"2026-07-06T22:37:03.716474Z","submitted_at":"2025-11-26T17:59:08Z","title":"Qwen3-VL Technical Report","version":2},"cited_work":{"arxiv_id":"2511.21631","doi":"10.1016/j.neunet.2025.107777","metadata_source":"pith","pith_arxiv_id":"2511.21631","snapshot_observed_at":"2026-07-11T11:50:26.030339Z","title":"Qwen3-VL Technical Report","venue":"cs.CV","work_id":"1fe243aa-e3c0-4da6-b391-4cbcfc88d5c0","year":2025},"citing_paper":{"arxiv_id":"2605.10806","last_updated":"2026-05-11T16:30:51Z","snapshot_observed_at":"2026-07-06T23:22:42.512039Z","submitted_at":"2026-05-11T16:30:51Z","title":"PhyGround: Benchmarking Physical Reasoning in Generative World Models","version":1},"reference_index":1,"source":"pdf_text","source_observed_at":"2026-05-12T05:17:30.010064Z"},"links":{"cited_paper":"/paper/2511.21631","citing_paper":"/paper/2605.10806"},"observation_digest":"sha256:cb77216e93c2c019e189f617432bd0434bfa3c6e08ae7f08ac6cc74818c7198a","observation_id":"e711beae-e78c-45ea-acdd-00f990d9cac0","resolution":{"observed_at":"2026-05-12T05:21:29.357124Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-02T06:30:47.504484+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-02T06:30:47.504484+00:00","source":"crossref"},{"observed_at":"2026-08-02T06:30:44.765945+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2503.14378","last_updated":"2025-03-18T16:10:24Z","snapshot_observed_at":"2026-07-06T20:54:45.985432Z","submitted_at":"2025-03-18T16:10:24Z","title":"Impossible Videos","version":1},"cited_work":{"arxiv_id":"2503.14378","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2503.14378","snapshot_observed_at":"2026-07-01T10:25:41.408031Z","title":"Impossible videos","venue":null,"work_id":"5b20ddf4-07aa-4fa9-b1b0-e9f77b8d7345","year":2025},"citing_paper":{"arxiv_id":"2605.10806","last_updated":"2026-05-11T16:30:51Z","snapshot_observed_at":"2026-07-06T23:22:42.512039Z","submitted_at":"2026-05-11T16:30:51Z","title":"PhyGround: Benchmarking Physical Reasoning in Generative World Models","version":1},"reference_index":2,"source":"pdf_text","source_observed_at":"2026-05-12T05:17:30.010064Z"},"links":{"cited_paper":"/paper/2503.14378","citing_paper":"/paper/2605.10806"},"observation_digest":"sha256:ad5dc8f97d5eb9c0fe84e58809a3838054d0936bd4ff20b667c1227866c8bb3b","observation_id":"ad7fc4c9-cdd6-467a-886f-318f9fa57741","resolution":{"observed_at":"2026-05-12T05:21:29.118097Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-02T06:30:47.504484+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-02T06:30:47.504484+00:00","source":"crossref"},{"observed_at":"2026-08-02T06:30:44.765945+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2406.03520","last_updated":"2024-10-03T17:24:40Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2024-06-05T17:53:55Z","title":"VideoPhy: Evaluating Physical Commonsense for Video Generation","version":2},"cited_work":{"arxiv_id":"2406.03520","doi":null,"metadata_source":"pith","pith_arxiv_id":"2406.03520","snapshot_observed_at":"2026-07-07T15:43:53.689981Z","title":"VideoPhy: Evaluating Physical Commonsense for Video Generation","venue":"cs.CV","work_id":"27ed795c-abbe-4de1-9a7a-2ecf39c354f3","year":2024},"citing_paper":{"arxiv_id":"2605.10806","last_updated":"2026-05-11T16:30:51Z","snapshot_observed_at":"2026-07-06T23:22:42.512039Z","submitted_at":"2026-05-11T16:30:51Z","title":"PhyGround: Benchmarking Physical Reasoning in Generative World Models","version":1},"reference_index":3,"source":"pdf_text","source_observed_at":"2026-05-12T05:17:30.010064Z"},"links":{"cited_paper":"/paper/2406.03520","citing_paper":"/paper/2605.10806"},"observation_digest":"sha256:2b026ef855455d20452cce12e12389eed7595c0ec7dfe38cd4600a05ad988f94","observation_id":"b5475a3e-9663-440d-8d76-0224bcdfbb6e","resolution":{"observed_at":"2026-05-20T11:34:37.994076Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-02T06:30:47.504484+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-02T06:30:47.504484+00:00","source":"crossref"},{"observed_at":"2026-08-02T06:30:44.765945+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2503.06800","last_updated":"2025-03-09T22:49:12Z","snapshot_observed_at":"2026-07-06T20:49:32.001630Z","submitted_at":"2025-03-09T22:49:12Z","title":"VideoPhy-2: A Challenging Action-Centric Physical Commonsense Evaluation in Video Generation","version":1},"cited_work":{"arxiv_id":"2503.06800","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2503.06800","snapshot_observed_at":"2026-07-04T13:19:51.020711Z","title":"Videophy-2: A challenging action-centric physical commonsense evaluation in video generation","venue":null,"work_id":"4dfe3980-dfd5-4917-95e9-179c29e4eb18","year":2025},"citing_paper":{"arxiv_id":"2605.10806","last_updated":"2026-05-11T16:30:51Z","snapshot_observed_at":"2026-07-06T23:22:42.512039Z","submitted_at":"2026-05-11T16:30:51Z","title":"PhyGround: Benchmarking Physical Reasoning in Generative World Models","version":1},"reference_index":4,"source":"pdf_text","source_observed_at":"2026-05-12T05:17:30.010064Z"},"links":{"cited_paper":"/paper/2503.06800","citing_paper":"/paper/2605.10806"},"observation_digest":"sha256:ad85bbe78c4b33582e50cae4ae5ead334126f3d452ec8a902b6cdfd48ceecab4","observation_id":"42134aa3-572d-4f86-9217-752a2cba9445","resolution":{"observed_at":"2026-05-12T05:21:27.838360Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-02T06:30:47.504484+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-02T06:30:47.504484+00:00","source":"crossref"},{"observed_at":"2026-08-02T06:30:44.765945+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Campbell and Julian C","venue":null,"work_id":"a5b2e73a-66db-4d72-92e2-79f22c39137b","year":1963},"citing_paper":{"arxiv_id":"2605.10806","last_updated":"2026-05-11T16:30:51Z","snapshot_observed_at":"2026-07-06T23:22:42.512039Z","submitted_at":"2026-05-11T16:30:51Z","title":"PhyGround: Benchmarking Physical Reasoning in Generative World Models","version":1},"reference_index":5,"source":"pdf_text","source_observed_at":"2026-05-12T05:17:30.010064Z"},"links":{"citing_paper":"/paper/2605.10806"},"observation_digest":"sha256:d1d48f0d912d46222bade882a69a0d5aed18f6ffab5f88373f256fa2c8b55239","observation_id":"fac1361f-d08a-494d-8aed-075a253046ef","resolution":{"observed_at":"2026-05-12T11:41:33.475040Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-02T06:30:47.504484+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-02T06:30:47.504484+00:00","source":"crossref"},{"observed_at":"2026-08-02T06:30:44.765945+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2412.01800","last_updated":"2024-12-02T18:47:25Z","snapshot_observed_at":"2026-08-02T07:22:42.764229Z","submitted_at":"2024-12-02T18:47:25Z","title":"PhysGame: Uncovering Physical Commonsense Violations in Gameplay Videos","version":1},"cited_work":{"arxiv_id":"2412.01800","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2412.01800","snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Physgame: Uncovering physical commonsense violations in gameplay videos","venue":null,"work_id":"46f4f776-4d46-4f0f-bd49-eec743b4e5ef","year":2024},"citing_paper":{"arxiv_id":"2605.10806","last_updated":"2026-05-11T16:30:51Z","snapshot_observed_at":"2026-07-06T23:22:42.512039Z","submitted_at":"2026-05-11T16:30:51Z","title":"PhyGround: Benchmarking Physical Reasoning in Generative World Models","version":1},"reference_index":6,"source":"pdf_text","source_observed_at":"2026-05-12T05:17:30.010064Z"},"links":{"cited_paper":"/paper/2412.01800","citing_paper":"/paper/2605.10806"},"observation_digest":"sha256:1a4ca4e90587d69b2e447981a52735c307a1a618ba08b6713d4720fb0cd56133","observation_id":"805088b8-aac6-4b94-905e-e19c61f61bce","resolution":{"observed_at":"2026-05-12T05:21:27.757534Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-02T06:30:47.504484+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-02T06:30:47.504484+00:00","source":"crossref"},{"observed_at":"2026-08-02T06:30:44.765945+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Nonnaïveté among amazon mechanical turk workers: Consequences and solutions for behavioral researchers.Behavior research methods, 46(1):112–130","venue":null,"work_id":"5a32035d-c477-4c09-9747-2bf6fe5afc60","year":2014},"citing_paper":{"arxiv_id":"2605.10806","last_updated":"2026-05-11T16:30:51Z","snapshot_observed_at":"2026-07-06T23:22:42.512039Z","submitted_at":"2026-05-11T16:30:51Z","title":"PhyGround: Benchmarking Physical Reasoning in Generative World Models","version":1},"reference_index":7,"source":"pdf_text","source_observed_at":"2026-05-12T05:17:30.010064Z"},"links":{"citing_paper":"/paper/2605.10806"},"observation_digest":"sha256:dc7f7270f7a1cc24dc35ad5630c81f834d22c9570905ea3797294b7ea3643210","observation_id":"837f6680-7e3a-4124-8b7c-7d5666031895","resolution":{"observed_at":"2026-05-12T11:41:33.471531Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-02T06:30:47.504484+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-02T06:30:47.504484+00:00","source":"crossref"},{"observed_at":"2026-08-02T06:30:44.765945+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Understanding world or predicting future? a comprehensive survey of world models.ACM Computing Surveys, 58(3):1–38","venue":null,"work_id":"7c8b4f56-1dbd-411c-9db8-324d72daa029","year":2025},"citing_paper":{"arxiv_id":"2605.10806","last_updated":"2026-05-11T16:30:51Z","snapshot_observed_at":"2026-07-06T23:22:42.512039Z","submitted_at":"2026-05-11T16:30:51Z","title":"PhyGround: Benchmarking Physical Reasoning in Generative World Models","version":1},"reference_index":8,"source":"pdf_text","source_observed_at":"2026-05-12T05:17:30.010064Z"},"links":{"citing_paper":"/paper/2605.10806"},"observation_digest":"sha256:bb2e22204244592eb9f1d3e50f15148e4c36719488d1b9c328349ffdec9dca50","observation_id":"04bbf226-7f1c-4a3f-be32-36f8bb923bb5","resolution":{"observed_at":"2026-05-12T11:41:33.480899Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-02T06:30:47.504484+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-02T06:30:47.504484+00:00","source":"crossref"},{"observed_at":"2026-08-02T06:30:44.765945+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Veo 3.https://deepmind.google/models/veo/","venue":null,"work_id":"2b72fa44-bcf9-45ba-afdc-88af619805d6","year":2026},"citing_paper":{"arxiv_id":"2605.10806","last_updated":"2026-05-11T16:30:51Z","snapshot_observed_at":"2026-07-06T23:22:42.512039Z","submitted_at":"2026-05-11T16:30:51Z","title":"PhyGround: Benchmarking Physical Reasoning in Generative World Models","version":1},"reference_index":9,"source":"pdf_text","source_observed_at":"2026-05-12T05:17:30.010064Z"},"links":{"citing_paper":"/paper/2605.10806"},"observation_digest":"sha256:68f8571f0c4ab2efc7539d901a1b65005065e0e3734629f2d095023a55088b4e","observation_id":"2524dd35-b160-4e86-9027-72919e467c5a","resolution":{"observed_at":"2026-05-12T11:41:33.439820Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-02T06:30:47.504484+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-02T06:30:47.504484+00:00","source":"crossref"},{"observed_at":"2026-08-02T06:30:44.765945+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2507.13428","last_updated":"2026-05-26T08:34:24Z","snapshot_observed_at":"2026-07-06T21:58:56.405149Z","submitted_at":"2025-07-17T17:54:09Z","title":"\"PhyWorldBench\": A Comprehensive Evaluation of Physical Realism in Text-to-Video Models","version":3},"cited_work":{"arxiv_id":"2507.13428","doi":null,"metadata_source":"pith","pith_arxiv_id":"2507.13428","snapshot_observed_at":"2026-07-08T07:14:45.185207Z","title":"InThe Thirteenth International Conference on Learning Representations","venue":"cs.CV","work_id":"7c164e09-5340-4f56-b63d-0e810f24f851","year":2025},"citing_paper":{"arxiv_id":"2605.10806","last_updated":"2026-05-11T16:30:51Z","snapshot_observed_at":"2026-07-06T23:22:42.512039Z","submitted_at":"2026-05-11T16:30:51Z","title":"PhyGround: Benchmarking Physical Reasoning in Generative World Models","version":1},"reference_index":10,"source":"pdf_text","source_observed_at":"2026-05-12T05:17:30.010064Z"},"links":{"cited_paper":"/paper/2507.13428","citing_paper":"/paper/2605.10806"},"observation_digest":"sha256:8c3ff3208d972decce2ccf1e3b5f5650fe4df7eb01a6ded1b6f393bcf57d54a4","observation_id":"603714e6-f940-4f6b-a3ab-832a175364c5","resolution":{"observed_at":"2026-05-27T02:05:04.304229Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-02T06:30:47.504484+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-02T06:30:47.504484+00:00","source":"crossref"},{"observed_at":"2026-08-02T06:30:44.765945+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Cosmos world foundation models for physical ai","venue":null,"work_id":"96c69bdc-2745-4296-9981-af60af76ae3f","year":2025},"citing_paper":{"arxiv_id":"2605.10806","last_updated":"2026-05-11T16:30:51Z","snapshot_observed_at":"2026-07-06T23:22:42.512039Z","submitted_at":"2026-05-11T16:30:51Z","title":"PhyGround: Benchmarking Physical Reasoning in Generative World Models","version":1},"reference_index":11,"source":"pdf_text","source_observed_at":"2026-05-12T05:17:30.010064Z"},"links":{"citing_paper":"/paper/2605.10806"},"observation_digest":"sha256:2f8f8f3f0d4c57eaebe06f3d07e0869141a8fb69ea1eb5c494caec6373ebfd4e","observation_id":"e70236ce-1ae2-490c-a316-b0037ff20c03","resolution":{"observed_at":"2026-05-12T11:41:33.445334Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-02T06:30:47.504484+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-02T06:30:47.504484+00:00","source":"crossref"},{"observed_at":"2026-08-02T06:30:44.765945+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2505.00337","last_updated":"2025-05-01T06:34:55Z","snapshot_observed_at":"2026-07-06T21:17:28.488862Z","submitted_at":"2025-05-01T06:34:55Z","title":"T2VPhysBench: A First-Principles Benchmark for Physical Consistency in Text-to-Video Generation","version":1},"cited_work":{"arxiv_id":"2505.00337","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2505.00337","snapshot_observed_at":"2026-07-04T13:59:51.792842Z","title":"T 2 V P hys B ench: A first-principles benchmark for physical consistency in text-to-video generation","venue":null,"work_id":"8e0819a9-997f-40ad-9d64-220808cd9d4c","year":2025},"citing_paper":{"arxiv_id":"2605.10806","last_updated":"2026-05-11T16:30:51Z","snapshot_observed_at":"2026-07-06T23:22:42.512039Z","submitted_at":"2026-05-11T16:30:51Z","title":"PhyGround: Benchmarking Physical Reasoning in Generative World Models","version":1},"reference_index":12,"source":"pdf_text","source_observed_at":"2026-05-12T05:17:30.010064Z"},"links":{"cited_paper":"/paper/2505.00337","citing_paper":"/paper/2605.10806"},"observation_digest":"sha256:1bd7572caffc1d8fe20ddaaa8f5d9e43556f64cc21f68fcd3e963f1ccadd8c59","observation_id":"884c5c48-e4a7-46fc-bf9a-590b8b266281","resolution":{"observed_at":"2026-05-12T05:21:28.228199Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-02T06:30:47.504484+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-02T06:30:47.504484+00:00","source":"crossref"},{"observed_at":"2026-08-02T06:30:44.765945+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2601.03233","last_updated":"2026-01-06T18:24:41Z","snapshot_observed_at":"2026-07-06T22:40:51.136169Z","submitted_at":"2026-01-06T18:24:41Z","title":"LTX-2: Efficient Joint Audio-Visual Foundation Model","version":1},"cited_work":{"arxiv_id":"2601.03233","doi":null,"metadata_source":"pith","pith_arxiv_id":"2601.03233","snapshot_observed_at":"2026-07-10T01:46:41.047035Z","title":"LTX-2: Efficient Joint Audio-Visual Foundation Model","venue":"cs.CV","work_id":"1334671f-ee98-4c53-a1d4-0805281f8d2b","year":2026},"citing_paper":{"arxiv_id":"2605.10806","last_updated":"2026-05-11T16:30:51Z","snapshot_observed_at":"2026-07-06T23:22:42.512039Z","submitted_at":"2026-05-11T16:30:51Z","title":"PhyGround: Benchmarking Physical Reasoning in Generative World Models","version":1},"reference_index":13,"source":"pdf_text","source_observed_at":"2026-05-12T05:17:30.010064Z"},"links":{"cited_paper":"/paper/2601.03233","citing_paper":"/paper/2605.10806"},"observation_digest":"sha256:5ef79144f9ca5db3b74c43c9036928ac5d2e564a3dfab4ccdeb1bdf50089e45d","observation_id":"01a9ba73-d263-4ccc-8df7-065601161be2","resolution":{"observed_at":"2026-05-12T05:21:28.818608Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-02T06:30:47.504484+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-02T06:30:47.504484+00:00","source":"crossref"},{"observed_at":"2026-08-02T06:30:44.765945+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-07-08T06:24:43.213996Z","title":"Video-bench: Human-aligned video generation benchmark","venue":null,"work_id":"53afc07f-f240-4533-adf6-304047280f56","year":2025},"citing_paper":{"arxiv_id":"2605.10806","last_updated":"2026-05-11T16:30:51Z","snapshot_observed_at":"2026-07-06T23:22:42.512039Z","submitted_at":"2026-05-11T16:30:51Z","title":"PhyGround: Benchmarking Physical Reasoning in Generative World Models","version":1},"reference_index":14,"source":"pdf_text","source_observed_at":"2026-05-12T05:17:30.010064Z"},"links":{"citing_paper":"/paper/2605.10806"},"observation_digest":"sha256:ca7e1cc28b1bc74d5e63ad973c05946b92d627493d558796907ac37b92241915","observation_id":"6fae5349-abaa-4d26-9e63-925f2a51af26","resolution":{"observed_at":"2026-05-12T11:41:33.431298Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-02T06:30:47.504484+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-02T06:30:47.504484+00:00","source":"crossref"},{"observed_at":"2026-08-02T06:30:44.765945+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":"2509.22799","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-07-03T10:27:56.742672Z","title":"Huynh-Thu, Q","venue":null,"work_id":"8b5960e1-cfbc-46a2-b0ae-e98f990faa4c","year":2025},"citing_paper":{"arxiv_id":"2605.10806","last_updated":"2026-05-11T16:30:51Z","snapshot_observed_at":"2026-07-06T23:22:42.512039Z","submitted_at":"2026-05-11T16:30:51Z","title":"PhyGround: Benchmarking Physical Reasoning in Generative World Models","version":1},"reference_index":15,"source":"pdf_text","source_observed_at":"2026-05-12T05:17:30.010064Z"},"links":{"citing_paper":"/paper/2605.10806"},"observation_digest":"sha256:24ad52645eabf6a33c548f5c049d3c98e5036908f684427537c83e8f194ec3d4","observation_id":"a1217bf1-d38d-4464-b3eb-9972861502a0","resolution":{"observed_at":"2026-05-12T05:21:28.907937Z","resolver_source":"arxiv_id","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-02T06:30:47.504484+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-02T06:30:47.504484+00:00","source":"crossref"},{"observed_at":"2026-08-02T06:30:44.765945+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-07-08T03:54:30.825321Z","title":"Videoscore: Building automatic metrics to simulate fine-grained human feedback for video generation","venue":null,"work_id":"a95a4387-dcbf-44b4-b3ed-2111b55ee80f","year":2024},"citing_paper":{"arxiv_id":"2605.10806","last_updated":"2026-05-11T16:30:51Z","snapshot_observed_at":"2026-07-06T23:22:42.512039Z","submitted_at":"2026-05-11T16:30:51Z","title":"PhyGround: Benchmarking Physical Reasoning in Generative World Models","version":1},"reference_index":16,"source":"pdf_text","source_observed_at":"2026-05-12T05:17:30.010064Z"},"links":{"citing_paper":"/paper/2605.10806"},"observation_digest":"sha256:6426bcffb6b97b67ee1b636293ab9fc0217e366557e268390bb6567491843665","observation_id":"1ae48bb9-58da-4c9f-ad68-94609dc4b3a8","resolution":{"observed_at":"2026-05-12T11:41:33.422827Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-02T06:30:47.504484+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-02T06:30:47.504484+00:00","source":"crossref"},{"observed_at":"2026-08-02T06:30:44.765945+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-07-05T01:20:38.029679Z","title":"Image quality metrics: Psnr vs","venue":null,"work_id":"fd5133ac-3e98-4ef0-84cc-81c2b24768cf","year":2010},"citing_paper":{"arxiv_id":"2605.10806","last_updated":"2026-05-11T16:30:51Z","snapshot_observed_at":"2026-07-06T23:22:42.512039Z","submitted_at":"2026-05-11T16:30:51Z","title":"PhyGround: Benchmarking Physical Reasoning in Generative World Models","version":1},"reference_index":17,"source":"pdf_text","source_observed_at":"2026-05-12T05:17:30.010064Z"},"links":{"citing_paper":"/paper/2605.10806"},"observation_digest":"sha256:52355a962266bb1a33e7204025c6942803d1eeddac10a0f5e096032c60f221b2","observation_id":"ea49ae1c-a838-426c-a815-a003f06da398","resolution":{"observed_at":"2026-05-12T11:41:33.414017Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-02T06:30:47.504484+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-02T06:30:47.504484+00:00","source":"crossref"},{"observed_at":"2026-08-02T06:30:44.765945+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":"2512.02942","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-07-04T06:19:37.486447Z","title":"Benchmarking scientific understanding and reasoning for video generation using videoscience-bench","venue":null,"work_id":"d440513c-aa40-425f-8f07-255f2ebb82a1","year":2025},"citing_paper":{"arxiv_id":"2605.10806","last_updated":"2026-05-11T16:30:51Z","snapshot_observed_at":"2026-07-06T23:22:42.512039Z","submitted_at":"2026-05-11T16:30:51Z","title":"PhyGround: Benchmarking Physical Reasoning in Generative World Models","version":1},"reference_index":18,"source":"pdf_text","source_observed_at":"2026-05-12T05:17:30.010064Z"},"links":{"citing_paper":"/paper/2605.10806"},"observation_digest":"sha256:1290b21d595b220e9b62f52a72cb5a9bda61fe00caaf875d633a7b92bdf2974e","observation_id":"4b8ca6e1-6940-4119-b055-265b16f8cd07","resolution":{"observed_at":"2026-05-12T05:21:28.798775Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-02T06:30:47.504484+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-02T06:30:47.504484+00:00","source":"crossref"},{"observed_at":"2026-08-02T06:30:44.765945+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Cosmos-eval: Towards explainable evaluation of physics and semantics in text-to-video models","venue":null,"work_id":"c3d43549-7c80-4345-bc58-30dd26a23557","year":null},"citing_paper":{"arxiv_id":"2605.10806","last_updated":"2026-05-11T16:30:51Z","snapshot_observed_at":"2026-07-06T23:22:42.512039Z","submitted_at":"2026-05-11T16:30:51Z","title":"PhyGround: Benchmarking Physical Reasoning in Generative World Models","version":1},"reference_index":19,"source":"pdf_text","source_observed_at":"2026-05-12T05:17:30.010064Z"},"links":{"citing_paper":"/paper/2605.10806"},"observation_digest":"sha256:27225e91bba8bac2b88e4d7ad74d478c94887ec1017f07fcf80a935ca8cfca1b","observation_id":"458345e8-37d5-4f49-9f43-de0f074214e5","resolution":{"observed_at":"2026-05-12T11:41:33.427301Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-02T06:30:47.504484+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-02T06:30:47.504484+00:00","source":"crossref"},{"observed_at":"2026-08-02T06:30:44.765945+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-07-11T01:07:50.611559Z","title":"Vbench: Comprehensive benchmark suite for video generative models","venue":null,"work_id":"2383dbc4-0e5d-40ed-a102-5aac37332eae","year":2024},"citing_paper":{"arxiv_id":"2605.10806","last_updated":"2026-05-11T16:30:51Z","snapshot_observed_at":"2026-07-06T23:22:42.512039Z","submitted_at":"2026-05-11T16:30:51Z","title":"PhyGround: Benchmarking Physical Reasoning in Generative World Models","version":1},"reference_index":20,"source":"pdf_text","source_observed_at":"2026-05-12T05:17:30.010064Z"},"links":{"citing_paper":"/paper/2605.10806"},"observation_digest":"sha256:c0638ac5278d836ac70fee3bc6c865740b5c856bcebfa029563a68f50d45950f","observation_id":"ac2e9143-d96d-4170-841b-7e3578e7c7d5","resolution":{"observed_at":"2026-05-12T11:41:33.449880Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-02T06:30:47.504484+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-02T06:30:47.504484+00:00","source":"crossref"},{"observed_at":"2026-08-02T06:30:44.765945+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Krosnick","venue":null,"work_id":"74dc44ee-a825-433d-8fa1-909e2888f95b","year":1991},"citing_paper":{"arxiv_id":"2605.10806","last_updated":"2026-05-11T16:30:51Z","snapshot_observed_at":"2026-07-06T23:22:42.512039Z","submitted_at":"2026-05-11T16:30:51Z","title":"PhyGround: Benchmarking Physical Reasoning in Generative World Models","version":1},"reference_index":21,"source":"pdf_text","source_observed_at":"2026-05-12T05:17:30.010064Z"},"links":{"citing_paper":"/paper/2605.10806"},"observation_digest":"sha256:7d29fb549dd1f8e512f05c19af41789e50be3c0d229ffc4e9a3a2a9b3765a891","observation_id":"89149ddf-d2ee-4045-a6a3-7d19d0807453","resolution":{"observed_at":"2026-05-12T11:41:33.453722Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-02T06:30:47.504484+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-02T06:30:47.504484+00:00","source":"crossref"},{"observed_at":"2026-08-02T06:30:44.765945+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2502.20694","last_updated":"2025-02-28T03:58:23Z","snapshot_observed_at":"2026-07-06T20:44:08.651242Z","submitted_at":"2025-02-28T03:58:23Z","title":"WorldModelBench: Judging Video Generation Models As World Models","version":1},"cited_work":{"arxiv_id":"2502.20694","doi":"10.48550/arxiv.2502.20694","metadata_source":"pith","pith_arxiv_id":"2502.20694","snapshot_observed_at":"2026-07-10T12:15:01.137692Z","title":"Li, S., Wu, K., Zhang, C., and Zhu, Y","venue":"cs.CV","work_id":"33a5087c-1633-4d52-b891-f34d10b9e7c8","year":2025},"citing_paper":{"arxiv_id":"2605.10806","last_updated":"2026-05-11T16:30:51Z","snapshot_observed_at":"2026-07-06T23:22:42.512039Z","submitted_at":"2026-05-11T16:30:51Z","title":"PhyGround: Benchmarking Physical Reasoning in Generative World Models","version":1},"reference_index":22,"source":"pdf_text","source_observed_at":"2026-05-12T05:17:30.010064Z"},"links":{"cited_paper":"/paper/2502.20694","citing_paper":"/paper/2605.10806"},"observation_digest":"sha256:1ce4c77a560ce3f1720e5067364b78f8c53879791d2e7631a5629a8b487224a1","observation_id":"1e896341-996e-4c94-8a15-38b4b0eadf65","resolution":{"observed_at":"2026-05-12T05:21:28.258267Z","resolver_source":"arxiv_id","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-02T06:30:47.504484+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-02T06:30:47.504484+00:00","source":"crossref"},{"observed_at":"2026-07-10T21:18:51.214081+00:00","source":"crossref_status_cache"},{"observed_at":"2026-07-10T21:18:51.214081+00:00","source":"openalex_status_cache"},{"observed_at":"2026-08-02T06:30:44.765945+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2501.13918","last_updated":"2025-10-27T08:22:57Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2025-01-23T18:55:41Z","title":"Improving Video Generation with Human Feedback","version":2},"cited_work":{"arxiv_id":"2501.13918","doi":null,"metadata_source":"pith","pith_arxiv_id":"2501.13918","snapshot_observed_at":"2026-07-04T13:19:51.115512Z","title":"Improving Video Generation with Human Feedback","venue":"cs.CV","work_id":"cfe4c01d-1cf7-4a00-ba86-06d583ca2cff","year":2025},"citing_paper":{"arxiv_id":"2605.10806","last_updated":"2026-05-11T16:30:51Z","snapshot_observed_at":"2026-07-06T23:22:42.512039Z","submitted_at":"2026-05-11T16:30:51Z","title":"PhyGround: Benchmarking Physical Reasoning in Generative World Models","version":1},"reference_index":23,"source":"pdf_text","source_observed_at":"2026-05-12T05:17:30.010064Z"},"links":{"cited_paper":"/paper/2501.13918","citing_paper":"/paper/2605.10806"},"observation_digest":"sha256:1e85f155fb3b9d2fcb8f74454b4a8c391ad74ef4ad5db45c6b55965dc841abf4","observation_id":"efb4bf65-7fd6-4b0c-81d7-a8c036db7991","resolution":{"observed_at":"2026-05-13T15:30:03.116566Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-02T06:30:47.504484+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-02T06:30:47.504484+00:00","source":"crossref"},{"observed_at":"2026-08-02T06:30:44.765945+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2507.01255","last_updated":"2025-07-02T00:20:06Z","snapshot_observed_at":"2026-07-06T21:50:49.929005Z","submitted_at":"2025-07-02T00:20:06Z","title":"AIGVE-MACS: Unified Multi-Aspect Commenting and Scoring Model for AI-Generated Video Evaluation","version":1},"cited_work":{"arxiv_id":"2507.01255","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2507.01255","snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Aigve-macs: Unified multi-aspect commenting and scoring model for ai-generated video evaluation","venue":null,"work_id":"b0104470-8991-4cc5-be06-099a977c5288","year":2025},"citing_paper":{"arxiv_id":"2605.10806","last_updated":"2026-05-11T16:30:51Z","snapshot_observed_at":"2026-07-06T23:22:42.512039Z","submitted_at":"2026-05-11T16:30:51Z","title":"PhyGround: Benchmarking Physical Reasoning in Generative World Models","version":1},"reference_index":24,"source":"pdf_text","source_observed_at":"2026-05-12T05:17:30.010064Z"},"links":{"cited_paper":"/paper/2507.01255","citing_paper":"/paper/2605.10806"},"observation_digest":"sha256:a2922eae835bab583a9a9c64a211daf37f47284aeb6179e5de5d05e3b0fb0a7f","observation_id":"89572dfe-b377-4713-9158-53ddd38754ba","resolution":{"observed_at":"2026-05-12T05:21:28.617378Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-02T06:30:47.504484+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-02T06:30:47.504484+00:00","source":"crossref"},{"observed_at":"2026-08-02T06:30:44.765945+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Evalcrafter: Benchmarking and evaluating large video generation models","venue":null,"work_id":"6a9b41db-e552-408a-aa02-522edbaa304c","year":2024},"citing_paper":{"arxiv_id":"2605.10806","last_updated":"2026-05-11T16:30:51Z","snapshot_observed_at":"2026-07-06T23:22:42.512039Z","submitted_at":"2026-05-11T16:30:51Z","title":"PhyGround: Benchmarking Physical Reasoning in Generative World Models","version":1},"reference_index":25,"source":"pdf_text","source_observed_at":"2026-05-12T05:17:30.010064Z"},"links":{"citing_paper":"/paper/2605.10806"},"observation_digest":"sha256:6296c715b04b1e49cb5ab028d9c905f93e82b09c11405f8b2c21136c3f96a34d","observation_id":"cdf27064-5e3f-4410-beb6-6a8dfd50a003","resolution":{"observed_at":"2026-05-12T11:41:33.382340Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-02T06:30:47.504484+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-02T06:30:47.504484+00:00","source":"crossref"},{"observed_at":"2026-08-02T06:30:44.765945+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2410.05363","last_updated":"2024-10-07T17:56:04Z","snapshot_observed_at":"2026-07-06T19:29:14.335016Z","submitted_at":"2024-10-07T17:56:04Z","title":"Towards World Simulator: Crafting Physical Commonsense-Based Benchmark for Video Generation","version":1},"cited_work":{"arxiv_id":"2410.05363","doi":null,"metadata_source":"pith","pith_arxiv_id":"2410.05363","snapshot_observed_at":"2026-07-08T07:14:45.198485Z","title":"Towards World Simulator: Crafting Physical Commonsense-Based Benchmark for Video Generation","venue":"cs.CV","work_id":"644e2886-7387-436d-ac11-d848af5fcc71","year":2024},"citing_paper":{"arxiv_id":"2605.10806","last_updated":"2026-05-11T16:30:51Z","snapshot_observed_at":"2026-07-06T23:22:42.512039Z","submitted_at":"2026-05-11T16:30:51Z","title":"PhyGround: Benchmarking Physical Reasoning in Generative World Models","version":1},"reference_index":26,"source":"pdf_text","source_observed_at":"2026-05-12T05:17:30.010064Z"},"links":{"cited_paper":"/paper/2410.05363","citing_paper":"/paper/2605.10806"},"observation_digest":"sha256:7be763033a24bb958f6fad42eab8ed2c5f5607793696fd6a8aeeeeb73211499e","observation_id":"b432c923-e7b5-489e-8d0e-4f3b68b14468","resolution":{"observed_at":"2026-05-18T14:40:00.221207Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-02T06:30:47.504484+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-02T06:30:47.504484+00:00","source":"crossref"},{"observed_at":"2026-08-02T06:30:44.765945+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":"2510.07550","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-07-03T17:58:47.096428Z","title":"Travl: A recipe for making video-language mod- els better judges of physics implausibility","venue":null,"work_id":"c426a017-7a60-4f6b-90f6-7657c04fb4e7","year":2025},"citing_paper":{"arxiv_id":"2605.10806","last_updated":"2026-05-11T16:30:51Z","snapshot_observed_at":"2026-07-06T23:22:42.512039Z","submitted_at":"2026-05-11T16:30:51Z","title":"PhyGround: Benchmarking Physical Reasoning in Generative World Models","version":1},"reference_index":27,"source":"pdf_text","source_observed_at":"2026-05-12T05:17:30.010064Z"},"links":{"citing_paper":"/paper/2605.10806"},"observation_digest":"sha256:3aa908315714ba1a3cf7f6e25272d752ba517fd93061d2261905252f84e1d519","observation_id":"9924653a-6535-4f4b-8da0-ce2497b56f79","resolution":{"observed_at":"2026-05-12T05:21:28.586302Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-02T06:30:47.504484+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-02T06:30:47.504484+00:00","source":"crossref"},{"observed_at":"2026-08-02T06:30:44.765945+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Do generative video models understand physical principles? InProceedings of the IEEE/CVF Winter Conference on Applications of Computer Vision, pages 948–958","venue":null,"work_id":"20863b0b-1454-4a9f-b6b8-775eb09bd754","year":2026},"citing_paper":{"arxiv_id":"2605.10806","last_updated":"2026-05-11T16:30:51Z","snapshot_observed_at":"2026-07-06T23:22:42.512039Z","submitted_at":"2026-05-11T16:30:51Z","title":"PhyGround: Benchmarking Physical Reasoning in Generative World Models","version":1},"reference_index":28,"source":"pdf_text","source_observed_at":"2026-05-12T05:17:30.010064Z"},"links":{"citing_paper":"/paper/2605.10806"},"observation_digest":"sha256:df7e392664917d9ed43372738efeceadb91508c3107c27fdf25e986fc1a577f6","observation_id":"7931133a-73ab-4a23-b457-f6e7a32c27e1","resolution":{"observed_at":"2026-05-12T11:41:33.396219Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-02T06:30:47.504484+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-02T06:30:47.504484+00:00","source":"crossref"},{"observed_at":"2026-08-02T06:30:44.765945+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":null,"venue":null,"work_id":"6634f8c6-592a-4daf-b1c6-ac5000d5a1e6","year":1962},"citing_paper":{"arxiv_id":"2605.10806","last_updated":"2026-05-11T16:30:51Z","snapshot_observed_at":"2026-07-06T23:22:42.512039Z","submitted_at":"2026-05-11T16:30:51Z","title":"PhyGround: Benchmarking Physical Reasoning in Generative World Models","version":1},"reference_index":29,"source":"pdf_text","source_observed_at":"2026-05-12T05:17:30.010064Z"},"links":{"citing_paper":"/paper/2605.10806"},"observation_digest":"sha256:1e7db38cca0f4f2e55c085368e84b8266e6e9d3500e1a870cb2650d99a558afd","observation_id":"abd03aff-b5ff-4218-a812-ade3627a9d27","resolution":{"observed_at":"2026-05-12T11:41:33.324995Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-02T06:30:47.504484+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-02T06:30:47.504484+00:00","source":"crossref"},{"observed_at":"2026-08-02T06:30:44.765945+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":"2603.24458","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-07-03T00:07:28.009236Z","title":"Omniweaving: Towards unified video generation with free-form composition and reasoning.https://arxiv.org/abs/2603.24458","venue":null,"work_id":"c3fe3b8c-232b-435d-bf50-69ec6af6db0f","year":2026},"citing_paper":{"arxiv_id":"2605.10806","last_updated":"2026-05-11T16:30:51Z","snapshot_observed_at":"2026-07-06T23:22:42.512039Z","submitted_at":"2026-05-11T16:30:51Z","title":"PhyGround: Benchmarking Physical Reasoning in Generative World Models","version":1},"reference_index":30,"source":"pdf_text","source_observed_at":"2026-05-12T05:17:30.010064Z"},"links":{"citing_paper":"/paper/2605.10806"},"observation_digest":"sha256:e266e1e1f92a11d9a5b1ceee299c60662b7bf73558bf07f52d589ccff5a87f09","observation_id":"4b51e551-f5f7-45a6-a2cd-55404c8e421d","resolution":{"observed_at":"2026-05-12T05:21:28.757546Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-02T06:30:47.504484+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-02T06:30:47.504484+00:00","source":"crossref"},{"observed_at":"2026-08-02T06:30:44.765945+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Running experiments on amazon mechanical turk.Judgment and Decision making, 5(5):411–419","venue":null,"work_id":"e610fd36-70fc-4087-b2de-e4604518f859","year":2010},"citing_paper":{"arxiv_id":"2605.10806","last_updated":"2026-05-11T16:30:51Z","snapshot_observed_at":"2026-07-06T23:22:42.512039Z","submitted_at":"2026-05-11T16:30:51Z","title":"PhyGround: Benchmarking Physical Reasoning in Generative World Models","version":1},"reference_index":31,"source":"pdf_text","source_observed_at":"2026-05-12T05:17:30.010064Z"},"links":{"citing_paper":"/paper/2605.10806"},"observation_digest":"sha256:8b6cf34eec23a4cad8d0c03bad9b310c9f9a4c3113072403c54f302fb8195af7","observation_id":"24eedc12-6c68-4e90-a1dc-2e76b2d9e580","resolution":{"observed_at":"2026-05-12T11:41:33.320458Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-02T06:30:47.504484+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-02T06:30:47.504484+00:00","source":"crossref"},{"observed_at":"2026-08-02T06:30:44.765945+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-07-09T07:16:04.687562Z","title":"Learning transferable visual models from natural language supervision","venue":null,"work_id":"ad3e05b3-af3a-4fa2-ab30-c45f9f403277","year":2021},"citing_paper":{"arxiv_id":"2605.10806","last_updated":"2026-05-11T16:30:51Z","snapshot_observed_at":"2026-07-06T23:22:42.512039Z","submitted_at":"2026-05-11T16:30:51Z","title":"PhyGround: Benchmarking Physical Reasoning in Generative World Models","version":1},"reference_index":32,"source":"pdf_text","source_observed_at":"2026-05-12T05:17:30.010064Z"},"links":{"citing_paper":"/paper/2605.10806"},"observation_digest":"sha256:707d9503714df491c045ccecc8ef2ab4f81bba7181f0f4e98bbd51afbd450227","observation_id":"c47d1df4-cfe3-4925-998a-a9a584b4c51f","resolution":{"observed_at":"2026-05-12T11:41:33.370350Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-02T06:30:47.504484+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-02T06:30:47.504484+00:00","source":"crossref"},{"observed_at":"2026-08-02T06:30:44.765945+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Reis and Charles M","venue":null,"work_id":"595a8605-13a1-4f4e-bf0e-764f3dce1148","year":2000},"citing_paper":{"arxiv_id":"2605.10806","last_updated":"2026-05-11T16:30:51Z","snapshot_observed_at":"2026-07-06T23:22:42.512039Z","submitted_at":"2026-05-11T16:30:51Z","title":"PhyGround: Benchmarking Physical Reasoning in Generative World Models","version":1},"reference_index":33,"source":"pdf_text","source_observed_at":"2026-05-12T05:17:30.010064Z"},"links":{"citing_paper":"/paper/2605.10806"},"observation_digest":"sha256:e19e92ba4aad6d41dc2dbaec3b7439a9227502ec784a1a3a3501a3b2fcde93fc","observation_id":"52d2bf37-3b2c-4150-aae3-8fe00f3164a9","resolution":{"observed_at":"2026-05-12T11:41:33.328732Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-02T06:30:47.504484+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-02T06:30:47.504484+00:00","source":"crossref"},{"observed_at":"2026-08-02T06:30:44.765945+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2604.16592","last_updated":"2026-04-17T17:51:46Z","snapshot_observed_at":"2026-07-06T23:03:52.631613Z","submitted_at":"2026-04-17T17:51:46Z","title":"Human Cognition in Machines: A Unified Perspective of World Models","version":1},"cited_work":{"arxiv_id":"2604.16592","doi":null,"metadata_source":"pith","pith_arxiv_id":"2604.16592","snapshot_observed_at":"2026-07-02T07:16:45.104866Z","title":"Human Cognition in Machines: A Unified Perspective of World Models","venue":"cs.RO","work_id":"e9f711d2-6321-4819-9a27-738f10976c6f","year":2026},"citing_paper":{"arxiv_id":"2605.10806","last_updated":"2026-05-11T16:30:51Z","snapshot_observed_at":"2026-07-06T23:22:42.512039Z","submitted_at":"2026-05-11T16:30:51Z","title":"PhyGround: Benchmarking Physical Reasoning in Generative World Models","version":1},"reference_index":34,"source":"pdf_text","source_observed_at":"2026-05-12T05:17:30.010064Z"},"links":{"cited_paper":"/paper/2604.16592","citing_paper":"/paper/2605.10806"},"observation_digest":"sha256:c110451ca30e7f9946d7a5b7ba515b7bf8cd8a87305778ccb381abd3de2772f3","observation_id":"b4dcd527-a9ee-4efe-92da-744280dbdd2d","resolution":{"observed_at":"2026-05-12T05:21:28.570348Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-02T06:30:47.504484+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-02T06:30:47.504484+00:00","source":"crossref"},{"observed_at":"2026-08-02T06:30:44.765945+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Shadish, Thomas D","venue":null,"work_id":"85a27950-7c09-47af-8056-3d8ab87fba6f","year":2002},"citing_paper":{"arxiv_id":"2605.10806","last_updated":"2026-05-11T16:30:51Z","snapshot_observed_at":"2026-07-06T23:22:42.512039Z","submitted_at":"2026-05-11T16:30:51Z","title":"PhyGround: Benchmarking Physical Reasoning in Generative World Models","version":1},"reference_index":35,"source":"pdf_text","source_observed_at":"2026-05-12T05:17:30.010064Z"},"links":{"citing_paper":"/paper/2605.10806"},"observation_digest":"sha256:96254402217e8a047466841e98498e220e4aba2e4c239eb5a556b8449f587134","observation_id":"a453b392-dd78-4048-96f7-ff2938345fba","resolution":{"observed_at":"2026-05-12T11:41:33.351301Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-02T06:30:47.504484+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-02T06:30:47.504484+00:00","source":"crossref"},{"observed_at":"2026-08-02T06:30:44.765945+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Vf-eval: Evaluating multimodal llms for generating feedback on aigc videos","venue":null,"work_id":"e67281cd-6f24-4765-9f5b-e66f2bef0ff9","year":2025},"citing_paper":{"arxiv_id":"2605.10806","last_updated":"2026-05-11T16:30:51Z","snapshot_observed_at":"2026-07-06T23:22:42.512039Z","submitted_at":"2026-05-11T16:30:51Z","title":"PhyGround: Benchmarking Physical Reasoning in Generative World Models","version":1},"reference_index":36,"source":"pdf_text","source_observed_at":"2026-05-12T05:17:30.010064Z"},"links":{"citing_paper":"/paper/2605.10806"},"observation_digest":"sha256:d28c570d45a3c126a480c6b59b03e227882a4e0c4dd50d9e387c3408378d1ede","observation_id":"39fda1bd-4801-4ae8-99d8-fe8a5464b9f2","resolution":{"observed_at":"2026-05-12T11:41:33.337033Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-02T06:30:47.504484+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-02T06:30:47.504484+00:00","source":"crossref"},{"observed_at":"2026-08-02T06:30:44.765945+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2502.04076","last_updated":"2025-02-06T13:41:24Z","snapshot_observed_at":"2026-07-06T20:32:14.203915Z","submitted_at":"2025-02-06T13:41:24Z","title":"Content-Rich AIGC Video Quality Assessment via Intricate Text Alignment and Motion-Aware Consistency","version":1},"cited_work":{"arxiv_id":"2502.04076","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2502.04076","snapshot_observed_at":"2026-07-04T16:39:57.276136Z","title":"Content-rich aigc video quality assessment via intricate text alignment and motion-aware consistency","venue":null,"work_id":"32a17dec-b6a9-453f-8bdd-e5c22ff6fa8d","year":2023},"citing_paper":{"arxiv_id":"2605.10806","last_updated":"2026-05-11T16:30:51Z","snapshot_observed_at":"2026-07-06T23:22:42.512039Z","submitted_at":"2026-05-11T16:30:51Z","title":"PhyGround: Benchmarking Physical Reasoning in Generative World Models","version":1},"reference_index":37,"source":"pdf_text","source_observed_at":"2026-05-12T05:17:30.010064Z"},"links":{"cited_paper":"/paper/2502.04076","citing_paper":"/paper/2605.10806"},"observation_digest":"sha256:c75bed1cfed2dad8f367dd275f06c69523c48b545a65fe4cf5d2d9fa76dfe500","observation_id":"fa1bf4dd-37e6-41c9-81ea-bc7a5c8936cd","resolution":{"observed_at":"2026-05-12T05:21:28.553668Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-02T06:30:47.504484+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-02T06:30:47.504484+00:00","source":"crossref"},{"observed_at":"2026-08-02T06:30:44.765945+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2312.11805","last_updated":"2025-05-09T21:04:06Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2023-12-19T02:39:27Z","title":"Gemini: A Family of Highly Capable Multimodal Models","version":5},"cited_work":{"arxiv_id":"2312.11805","doi":"10.1038/nrn2888","metadata_source":"pith","pith_arxiv_id":"2312.11805","snapshot_observed_at":"2026-07-11T11:50:26.030339Z","title":"Gemini: A Family of Highly Capable Multimodal Models","venue":"cs.CL","work_id":"83f7c85b-3f11-450f-ac0c-64d9745220b2","year":2023},"citing_paper":{"arxiv_id":"2605.10806","last_updated":"2026-05-11T16:30:51Z","snapshot_observed_at":"2026-07-06T23:22:42.512039Z","submitted_at":"2026-05-11T16:30:51Z","title":"PhyGround: Benchmarking Physical Reasoning in Generative World Models","version":1},"reference_index":38,"source":"pdf_text","source_observed_at":"2026-05-12T05:17:30.010064Z"},"links":{"cited_paper":"/paper/2312.11805","citing_paper":"/paper/2605.10806"},"observation_digest":"sha256:47d3c60625790cf6e9ab544d69723ffd2be5039b86060becfe73b4dd362d09ae","observation_id":"f50cbe9e-451c-46dd-b4bf-025d89287f95","resolution":{"observed_at":"2026-05-12T05:21:28.578263Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-02T06:30:47.504484+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-02T06:30:47.504484+00:00","source":"crossref"},{"observed_at":"2026-08-02T06:30:44.765945+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2502.01719","last_updated":"2025-02-07T03:54:34Z","snapshot_observed_at":"2026-08-02T02:04:03.925180Z","submitted_at":"2025-02-03T18:56:33Z","title":"MJ-VIDEO: Fine-Grained Benchmarking and Rewarding Video Preferences in Video Generation","version":3},"cited_work":{"arxiv_id":"2502.01719","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2502.01719","snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Mj-video: Fine-grained benchmarking and rewarding video preferences in video generation","venue":null,"work_id":"24ffa43d-520e-494a-87bd-8c4559ac4136","year":2025},"citing_paper":{"arxiv_id":"2605.10806","last_updated":"2026-05-11T16:30:51Z","snapshot_observed_at":"2026-07-06T23:22:42.512039Z","submitted_at":"2026-05-11T16:30:51Z","title":"PhyGround: Benchmarking Physical Reasoning in Generative World Models","version":1},"reference_index":39,"source":"pdf_text","source_observed_at":"2026-05-12T05:17:30.010064Z"},"links":{"cited_paper":"/paper/2502.01719","citing_paper":"/paper/2605.10806"},"observation_digest":"sha256:8fe24e6fe2803cf319758fa87773cd6d693bd7abbdd940c53a648a84d0221223","observation_id":"28313f88-cbbb-403d-bfea-a7e042f23e29","resolution":{"observed_at":"2026-05-12T05:21:28.398672Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-02T06:30:47.504484+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-02T06:30:47.504484+00:00","source":"crossref"},{"observed_at":"2026-08-02T06:30:44.765945+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-07-08T17:35:10.823088Z","title":"Fvd: A new metric for video generation","venue":null,"work_id":"bff3762d-6c6a-4b69-a92b-f3cdc2df44be","year":2019},"citing_paper":{"arxiv_id":"2605.10806","last_updated":"2026-05-11T16:30:51Z","snapshot_observed_at":"2026-07-06T23:22:42.512039Z","submitted_at":"2026-05-11T16:30:51Z","title":"PhyGround: Benchmarking Physical Reasoning in Generative World Models","version":1},"reference_index":40,"source":"pdf_text","source_observed_at":"2026-05-12T05:17:30.010064Z"},"links":{"citing_paper":"/paper/2605.10806"},"observation_digest":"sha256:4b28794822447cb82728ab9b3744861dd41345bb9884ad2958831a1350ed6c17","observation_id":"1c845899-99ac-46f4-b31d-26d97d78121f","resolution":{"observed_at":"2026-05-12T11:41:33.364802Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-02T06:30:47.504484+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-02T06:30:47.504484+00:00","source":"crossref"},{"observed_at":"2026-08-02T06:30:44.765945+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2503.20314","last_updated":"2025-04-19T02:22:42Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2025-03-26T08:25:43Z","title":"Wan: Open and Advanced Large-Scale Video Generative Models","version":2},"cited_work":{"arxiv_id":"2503.20314","doi":"10.1109/19.492748","metadata_source":"pith","pith_arxiv_id":"2503.20314","snapshot_observed_at":"2026-07-11T11:50:26.030339Z","title":"Wan: Open and Advanced Large-Scale Video Generative Models","venue":"cs.CV","work_id":"ad3ebc3b-4224-46c9-b61d-bcf135da0a7c","year":2025},"citing_paper":{"arxiv_id":"2605.10806","last_updated":"2026-05-11T16:30:51Z","snapshot_observed_at":"2026-07-06T23:22:42.512039Z","submitted_at":"2026-05-11T16:30:51Z","title":"PhyGround: Benchmarking Physical Reasoning in Generative World Models","version":1},"reference_index":41,"source":"pdf_text","source_observed_at":"2026-05-12T05:17:30.010064Z"},"links":{"cited_paper":"/paper/2503.20314","citing_paper":"/paper/2605.10806"},"observation_digest":"sha256:23e85f9672cc9bbd346f477fc9dc6fc9f1b26b8674cad86f581c16968b489f30","observation_id":"9923985f-d3cb-46ba-805f-dae2b5773e0c","resolution":{"observed_at":"2026-05-12T05:21:28.487243Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-02T06:30:47.504484+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-02T06:30:47.504484+00:00","source":"crossref"},{"observed_at":"2026-08-02T06:30:44.765945+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2505.12098","last_updated":"2025-05-17T17:49:26Z","snapshot_observed_at":"2026-08-01T15:10:45.025114Z","submitted_at":"2025-05-17T17:49:26Z","title":"LOVE: Benchmarking and Evaluating Text-to-Video Generation and Video-to-Text Interpretation","version":1},"cited_work":{"arxiv_id":"2505.12098","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2505.12098","snapshot_observed_at":"2026-07-02T13:56:59.272327Z","title":"Love: Benchmarking and evaluating text-to-video generation and video-to- text interpretation","venue":null,"work_id":"65b796ec-b5c4-4663-a486-1ac45eb1e875","year":2025},"citing_paper":{"arxiv_id":"2605.10806","last_updated":"2026-05-11T16:30:51Z","snapshot_observed_at":"2026-07-06T23:22:42.512039Z","submitted_at":"2026-05-11T16:30:51Z","title":"PhyGround: Benchmarking Physical Reasoning in Generative World Models","version":1},"reference_index":42,"source":"pdf_text","source_observed_at":"2026-05-12T05:17:30.010064Z"},"links":{"cited_paper":"/paper/2505.12098","citing_paper":"/paper/2605.10806"},"observation_digest":"sha256:373d544d9dc68743b771dbba5941419bbe6e5dc53d375bbf1d06939f1229e8f6","observation_id":"f03c7216-4506-48bb-a9b4-1fe2e7189b63","resolution":{"observed_at":"2026-05-12T05:21:28.543375Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-02T06:30:47.504484+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-02T06:30:47.504484+00:00","source":"crossref"},{"observed_at":"2026-08-02T06:30:44.765945+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Aigv-assessor: Benchmarking and evaluating the perceptual quality of text-to-video generation with lmm","venue":null,"work_id":"55c8933a-47a8-4210-9804-99693fff58fb","year":2025},"citing_paper":{"arxiv_id":"2605.10806","last_updated":"2026-05-11T16:30:51Z","snapshot_observed_at":"2026-07-06T23:22:42.512039Z","submitted_at":"2026-05-11T16:30:51Z","title":"PhyGround: Benchmarking Physical Reasoning in Generative World Models","version":1},"reference_index":43,"source":"pdf_text","source_observed_at":"2026-05-12T05:17:30.010064Z"},"links":{"citing_paper":"/paper/2605.10806"},"observation_digest":"sha256:087d093ee13ae47c4d63cb2308708d467270336243429de8e583d73110947ff0","observation_id":"25d614c6-d84d-4116-858c-d4707edf8d47","resolution":{"observed_at":"2026-05-12T11:41:33.389562Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-02T06:30:47.504484+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-02T06:30:47.504484+00:00","source":"crossref"},{"observed_at":"2026-08-02T06:30:44.765945+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":"2602.20159","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-07-10T01:46:40.917858Z","title":"A very big video reasoning suite","venue":null,"work_id":"4cb52917-7539-4575-a350-92645383c4ef","year":2026},"citing_paper":{"arxiv_id":"2605.10806","last_updated":"2026-05-11T16:30:51Z","snapshot_observed_at":"2026-07-06T23:22:42.512039Z","submitted_at":"2026-05-11T16:30:51Z","title":"PhyGround: Benchmarking Physical Reasoning in Generative World Models","version":1},"reference_index":44,"source":"pdf_text","source_observed_at":"2026-05-12T05:17:30.010064Z"},"links":{"citing_paper":"/paper/2605.10806"},"observation_digest":"sha256:aa27d96f9f2e2e69a4726197b528c309fc28a29d4e839673cd8d063ed931e85e","observation_id":"70a87992-5f73-4b2b-bd90-1636d5db86d2","resolution":{"observed_at":"2026-05-12T05:21:28.563098Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-02T06:30:47.504484+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-02T06:30:47.504484+00:00","source":"crossref"},{"observed_at":"2026-08-02T06:30:44.765945+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2512.01843","last_updated":"2026-05-18T11:10:13Z","snapshot_observed_at":"2026-07-06T22:37:26.652349Z","submitted_at":"2025-12-01T16:28:13Z","title":"PhyDetEx: Detecting and Explaining the Physical Plausibility of T2V Models","version":3},"cited_work":{"arxiv_id":"2512.01843","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2512.01843","snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Phydetex: Detecting and explaining the physical plausibility of t2v models","venue":null,"work_id":"994ad55a-f0eb-44ed-b5a4-dc6dfac22e29","year":2025},"citing_paper":{"arxiv_id":"2605.10806","last_updated":"2026-05-11T16:30:51Z","snapshot_observed_at":"2026-07-06T23:22:42.512039Z","submitted_at":"2026-05-11T16:30:51Z","title":"PhyGround: Benchmarking Physical Reasoning in Generative World Models","version":1},"reference_index":45,"source":"pdf_text","source_observed_at":"2026-05-12T05:17:30.010064Z"},"links":{"cited_paper":"/paper/2512.01843","citing_paper":"/paper/2605.10806"},"observation_digest":"sha256:7064d36eb79329256e48d10f821e57d24dd9ebbfab72c5fd5256d2e1482ad08b","observation_id":"01aec354-73a1-4843-abe6-8df5100cbf10","resolution":{"observed_at":"2026-05-20T00:05:38.433674Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-02T06:30:47.504484+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-02T06:30:47.504484+00:00","source":"crossref"},{"observed_at":"2026-08-02T06:30:44.765945+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-07-08T03:54:30.812941Z","title":"Visionreward: Fine-grained multi-dimensional human preference learning for image and video generation","venue":null,"work_id":"00f4fdb0-f1bf-459c-877a-6ac01a6ea3ff","year":2026},"citing_paper":{"arxiv_id":"2605.10806","last_updated":"2026-05-11T16:30:51Z","snapshot_observed_at":"2026-07-06T23:22:42.512039Z","submitted_at":"2026-05-11T16:30:51Z","title":"PhyGround: Benchmarking Physical Reasoning in Generative World Models","version":1},"reference_index":46,"source":"pdf_text","source_observed_at":"2026-05-12T05:17:30.010064Z"},"links":{"citing_paper":"/paper/2605.10806"},"observation_digest":"sha256:f56f3875acf6a2c98926e5b67bc07940dda978af9f38b05b78dfb55114752828","observation_id":"029bc0e7-cef6-4d6c-b4ab-5c9bdad21d38","resolution":{"observed_at":"2026-05-12T11:41:33.401111Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-02T06:30:47.504484+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-02T06:30:47.504484+00:00","source":"crossref"},{"observed_at":"2026-08-02T06:30:44.765945+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2504.02918","last_updated":"2026-06-29T01:14:54Z","snapshot_observed_at":"2026-07-06T21:03:55.628788Z","submitted_at":"2025-04-03T15:21:17Z","title":"Evaluating Newtonian Mechanics in Video Generative Models with Real Physical Systems","version":3},"cited_work":{"arxiv_id":"2504.02918","doi":"10.48550/arxiv.2504.02918","metadata_source":"pith","pith_arxiv_id":"2504.02918","snapshot_observed_at":"2026-07-10T12:15:01.137692Z","title":"Zhang, D","venue":"cs.CV","work_id":"a5729cf7-9ce0-49a1-a57d-9886470c40b7","year":2025},"citing_paper":{"arxiv_id":"2605.10806","last_updated":"2026-05-11T16:30:51Z","snapshot_observed_at":"2026-07-06T23:22:42.512039Z","submitted_at":"2026-05-11T16:30:51Z","title":"PhyGround: Benchmarking Physical Reasoning in Generative World Models","version":1},"reference_index":47,"source":"pdf_text","source_observed_at":"2026-05-12T05:17:30.010064Z"},"links":{"cited_paper":"/paper/2504.02918","citing_paper":"/paper/2605.10806"},"observation_digest":"sha256:50dead13b06c3534d059867c857c7e479dd0873f1bb4d678b87dbed5035f7a5f","observation_id":"7536e8b1-71cb-44bb-b736-70e487934b53","resolution":{"observed_at":"2026-06-30T03:17:04.940243Z","resolver_source":"arxiv_id","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-02T06:30:47.504484+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-02T06:30:47.504484+00:00","source":"crossref"},{"observed_at":"2026-07-10T21:18:48.706484+00:00","source":"crossref_status_cache"},{"observed_at":"2026-07-10T21:18:48.706484+00:00","source":"openalex_status_cache"},{"observed_at":"2026-08-02T06:30:44.765945+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":"2603.19607","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-30T10:54:36.504697Z","title":"Zhang, P","venue":null,"work_id":"82f208c5-735f-4958-97f5-738ebffd7e4d","year":2026},"citing_paper":{"arxiv_id":"2605.10806","last_updated":"2026-05-11T16:30:51Z","snapshot_observed_at":"2026-07-06T23:22:42.512039Z","submitted_at":"2026-05-11T16:30:51Z","title":"PhyGround: Benchmarking Physical Reasoning in Generative World Models","version":1},"reference_index":48,"source":"pdf_text","source_observed_at":"2026-05-12T05:17:30.010064Z"},"links":{"citing_paper":"/paper/2605.10806"},"observation_digest":"sha256:df29491a660c31b57b806994b04c834220a19356e36945857f839db837bd8e6f","observation_id":"492143e2-cd04-4275-b062-7558d610b776","resolution":{"observed_at":"2026-05-12T05:21:28.811518Z","resolver_source":"arxiv_id","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-02T06:30:47.504484+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-02T06:30:47.504484+00:00","source":"crossref"},{"observed_at":"2026-08-02T06:30:44.765945+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Vq-insight: Teaching vlms for ai-generated video quality understanding via progressive visual reinforcement learning","venue":null,"work_id":"4b6d5001-67a2-4415-b04c-c39dd2f6339f","year":2026},"citing_paper":{"arxiv_id":"2605.10806","last_updated":"2026-05-11T16:30:51Z","snapshot_observed_at":"2026-07-06T23:22:42.512039Z","submitted_at":"2026-05-11T16:30:51Z","title":"PhyGround: Benchmarking Physical Reasoning in Generative World Models","version":1},"reference_index":49,"source":"pdf_text","source_observed_at":"2026-05-12T05:17:30.010064Z"},"links":{"citing_paper":"/paper/2605.10806"},"observation_digest":"sha256:76e6d3759f1c05131e37e8a5414d65f3cf71901bd9f32fd541691add53ab64b0","observation_id":"47de9737-c504-4582-9d94-b904377ac17d","resolution":{"observed_at":"2026-05-12T11:41:33.461840Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-02T06:30:47.504484+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-02T06:30:47.504484+00:00","source":"crossref"},{"observed_at":"2026-08-02T06:30:44.765945+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Q-bench-video: Benchmark the video quality understanding of lmms","venue":null,"work_id":"8320ca5f-a265-43cc-ae8d-98ec5436a32c","year":2025},"citing_paper":{"arxiv_id":"2605.10806","last_updated":"2026-05-11T16:30:51Z","snapshot_observed_at":"2026-07-06T23:22:42.512039Z","submitted_at":"2026-05-11T16:30:51Z","title":"PhyGround: Benchmarking Physical Reasoning in Generative World Models","version":1},"reference_index":50,"source":"pdf_text","source_observed_at":"2026-05-12T05:17:30.010064Z"},"links":{"citing_paper":"/paper/2605.10806"},"observation_digest":"sha256:6e87ea5454cebe490642afd04b768a21091a595232e71f8cec9c16519a6c2a60","observation_id":"75c319f0-613f-4fe6-ab55-e798c2cbb578","resolution":{"observed_at":"2026-05-12T11:41:33.467477Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-02T06:30:47.504484+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-02T06:30:47.504484+00:00","source":"crossref"},{"observed_at":"2026-08-02T06:30:44.765945+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2503.21755","last_updated":"2025-08-20T15:49:30Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2025-03-27T17:57:01Z","title":"VBench-2.0: Advancing Video Generation Benchmark Suite for Intrinsic Faithfulness","version":2},"cited_work":{"arxiv_id":"2503.21755","doi":"10.48550/arxiv.2503.21755","metadata_source":"pith","pith_arxiv_id":"2503.21755","snapshot_observed_at":"2026-07-10T12:15:01.137692Z","title":"VBench-2.0: Advancing Video Generation Benchmark Suite for Intrinsic Faithfulness","venue":"cs.CV","work_id":"14060202-ac5f-48e9-b91a-24d150775431","year":2025},"citing_paper":{"arxiv_id":"2605.10806","last_updated":"2026-05-11T16:30:51Z","snapshot_observed_at":"2026-07-06T23:22:42.512039Z","submitted_at":"2026-05-11T16:30:51Z","title":"PhyGround: Benchmarking Physical Reasoning in Generative World Models","version":1},"reference_index":51,"source":"pdf_text","source_observed_at":"2026-05-12T05:17:30.010064Z"},"links":{"cited_paper":"/paper/2503.21755","citing_paper":"/paper/2605.10806"},"observation_digest":"sha256:bad37abf130bdfff7253e53244bf9ca8b0055364a5cf12fa95f8f4473e70bdde","observation_id":"876ce071-c6a5-494f-b23c-0f48aed884a6","resolution":{"observed_at":"2026-05-14T18:42:03.548540Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-02T06:30:47.504484+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-02T06:30:47.504484+00:00","source":"crossref"},{"observed_at":"2026-08-02T06:30:44.765945+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"spaghetti breaks into pieces","venue":null,"work_id":"90eb0b1c-600c-4d58-b623-92f8274cd58b","year":null},"citing_paper":{"arxiv_id":"2605.10806","last_updated":"2026-05-11T16:30:51Z","snapshot_observed_at":"2026-07-06T23:22:42.512039Z","submitted_at":"2026-05-11T16:30:51Z","title":"PhyGround: Benchmarking Physical Reasoning in Generative World Models","version":1},"reference_index":52,"source":"pdf_text","source_observed_at":"2026-05-12T05:17:30.010064Z"},"links":{"citing_paper":"/paper/2605.10806"},"observation_digest":"sha256:20870c94229fbcd936b6f3ff5554e6f4af608d3524c8383e856f578689b62f5e","observation_id":"4042779e-fc7a-4f2b-bbd1-0d876bf2575b","resolution":{"observed_at":"2026-05-12T11:41:33.457412Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-02T06:30:47.504484+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-02T06:30:47.504484+00:00","source":"crossref"},{"observed_at":"2026-08-02T06:30:44.765945+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":null,"venue":null,"work_id":"e3faef25-fbee-4c25-a2cf-84b8da6500fb","year":null},"citing_paper":{"arxiv_id":"2605.10806","last_updated":"2026-05-11T16:30:51Z","snapshot_observed_at":"2026-07-06T23:22:42.512039Z","submitted_at":"2026-05-11T16:30:51Z","title":"PhyGround: Benchmarking Physical Reasoning in Generative World Models","version":1},"reference_index":53,"source":"pdf_text","source_observed_at":"2026-05-12T05:17:30.010064Z"},"links":{"citing_paper":"/paper/2605.10806"},"observation_digest":"sha256:5f60ffb14e43f09891fbe74f5116673fccfd1758fdee2bb6f49b792379d3360f","observation_id":"da9a4e9a-5cd8-4148-a69d-11a69e08d042","resolution":{"observed_at":"2026-05-12T11:41:33.307689Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-02T06:30:47.504484+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-02T06:30:47.504484+00:00","source":"crossref"},{"observed_at":"2026-08-02T06:30:44.765945+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"rotary filling machine","venue":null,"work_id":"7a50d709-5ff7-452c-931c-7d7503c30ee2","year":null},"citing_paper":{"arxiv_id":"2605.10806","last_updated":"2026-05-11T16:30:51Z","snapshot_observed_at":"2026-07-06T23:22:42.512039Z","submitted_at":"2026-05-11T16:30:51Z","title":"PhyGround: Benchmarking Physical Reasoning in Generative World Models","version":1},"reference_index":54,"source":"pdf_text","source_observed_at":"2026-05-12T05:17:30.010064Z"},"links":{"citing_paper":"/paper/2605.10806"},"observation_digest":"sha256:73c9c8d2da80c1a6a876e27f0bba1b46e7fed02bdced9298f3d2b3daa687213f","observation_id":"d1a8e8b0-7f7f-454d-8e8e-91f83023362a","resolution":{"observed_at":"2026-05-12T11:41:33.292309Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-02T06:30:47.504484+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-02T06:30:47.504484+00:00","source":"crossref"},{"observed_at":"2026-08-02T06:30:44.765945+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Prompts that are only superficially associated through keywords are excluded","venue":null,"work_id":"9cd9607a-eb4e-45ce-b757-71a672fc6e68","year":null},"citing_paper":{"arxiv_id":"2605.10806","last_updated":"2026-05-11T16:30:51Z","snapshot_observed_at":"2026-07-06T23:22:42.512039Z","submitted_at":"2026-05-11T16:30:51Z","title":"PhyGround: Benchmarking Physical Reasoning in Generative World Models","version":1},"reference_index":55,"source":"pdf_text","source_observed_at":"2026-05-12T05:17:30.010064Z"},"links":{"citing_paper":"/paper/2605.10806"},"observation_digest":"sha256:dc6f029498a35cf9a0d49b6a3dff0f481e901c0354b9f18437ebbeae4ffcd898","observation_id":"b3bdd210-3ca5-40f1-a717-96d039fdb3cc","resolution":{"observed_at":"2026-05-12T11:41:33.295500Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-02T06:30:47.504484+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-02T06:30:47.504484+00:00","source":"crossref"},{"observed_at":"2026-08-02T06:30:44.765945+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"The athlete throws a javelin","venue":null,"work_id":"ef4075ba-0282-4f0e-bfb4-ad345228e79c","year":null},"citing_paper":{"arxiv_id":"2605.10806","last_updated":"2026-05-11T16:30:51Z","snapshot_observed_at":"2026-07-06T23:22:42.512039Z","submitted_at":"2026-05-11T16:30:51Z","title":"PhyGround: Benchmarking Physical Reasoning in Generative World Models","version":1},"reference_index":56,"source":"pdf_text","source_observed_at":"2026-05-12T05:17:30.010064Z"},"links":{"citing_paper":"/paper/2605.10806"},"observation_digest":"sha256:ebae95ae595beceb5d3dfb9095be4265e67ab1a8e5add827c0a53ea9481b8360","observation_id":"b6a112b9-d6a4-41b5-9cbe-8952c45410cc","resolution":{"observed_at":"2026-05-12T11:41:33.314157Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-02T06:30:47.504484+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-02T06:30:47.504484+00:00","source":"crossref"},{"observed_at":"2026-08-02T06:30:44.765945+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"The presentation order of videos is randomized independently for each annotator","venue":null,"work_id":"1dc40aff-8544-4fce-9b7b-7842bcffe461","year":null},"citing_paper":{"arxiv_id":"2605.10806","last_updated":"2026-05-11T16:30:51Z","snapshot_observed_at":"2026-07-06T23:22:42.512039Z","submitted_at":"2026-05-11T16:30:51Z","title":"PhyGround: Benchmarking Physical Reasoning in Generative World Models","version":1},"reference_index":57,"source":"pdf_text","source_observed_at":"2026-05-12T05:17:30.010064Z"},"links":{"citing_paper":"/paper/2605.10806"},"observation_digest":"sha256:b798a268f3f1daccb8b611be008d66a3f43ecb89546fc69f030e447f147d3332","observation_id":"086ad89c-0463-418b-a49c-2ab194d21d0f","resolution":{"observed_at":"2026-05-12T11:41:33.332304Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-02T06:30:47.504484+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-02T06:30:47.504484+00:00","source":"crossref"},{"observed_at":"2026-08-02T06:30:44.765945+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"The task is described without revealing our hypotheses, and model identities are hidden from annotators","venue":null,"work_id":"34c0b2ff-6357-4ed3-9e5c-38b3811d23a0","year":null},"citing_paper":{"arxiv_id":"2605.10806","last_updated":"2026-05-11T16:30:51Z","snapshot_observed_at":"2026-07-06T23:22:42.512039Z","submitted_at":"2026-05-11T16:30:51Z","title":"PhyGround: Benchmarking Physical Reasoning in Generative World Models","version":1},"reference_index":58,"source":"pdf_text","source_observed_at":"2026-05-12T05:17:30.010064Z"},"links":{"citing_paper":"/paper/2605.10806"},"observation_digest":"sha256:80a39bc395ad35c90da778bf250fcbe834e8b9d6a52d0d42bb0ba50b2c3708e9","observation_id":"74c157d8-d976-446e-92c8-673133bc481a","resolution":{"observed_at":"2026-05-12T11:41:33.377798Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-02T06:30:47.504484+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-02T06:30:47.504484+00:00","source":"crossref"},{"observed_at":"2026-08-02T06:30:44.765945+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"The module presents example videos that illustrate different levels of physical realism and explains how to apply the five-point Likert scale","venue":null,"work_id":"b32f7f9f-56ef-4086-8c79-78d5a8049761","year":null},"citing_paper":{"arxiv_id":"2605.10806","last_updated":"2026-05-11T16:30:51Z","snapshot_observed_at":"2026-07-06T23:22:42.512039Z","submitted_at":"2026-05-11T16:30:51Z","title":"PhyGround: Benchmarking Physical Reasoning in Generative World Models","version":1},"reference_index":59,"source":"pdf_text","source_observed_at":"2026-05-12T05:17:30.010064Z"},"links":{"citing_paper":"/paper/2605.10806"},"observation_digest":"sha256:2411618e3b365fcadc84bbc9a0a410926823df357fe37ff16539ea44e9997d8c","observation_id":"a296956f-5f00-4996-9df8-302c71552216","resolution":{"observed_at":"2026-05-12T11:41:33.342583Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-02T06:30:47.504484+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-02T06:30:47.504484+00:00","source":"crossref"},{"observed_at":"2026-08-02T06:30:44.765945+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Annotators evaluate three general dimensions, semantic alignment, physical temporal validity, and object persistence, as well as the applicable physical laws for that video","venue":null,"work_id":"d335b989-aec7-48d9-9732-37bc336c75f3","year":null},"citing_paper":{"arxiv_id":"2605.10806","last_updated":"2026-05-11T16:30:51Z","snapshot_observed_at":"2026-07-06T23:22:42.512039Z","submitted_at":"2026-05-11T16:30:51Z","title":"PhyGround: Benchmarking Physical Reasoning in Generative World Models","version":1},"reference_index":60,"source":"pdf_text","source_observed_at":"2026-05-12T05:17:30.010064Z"},"links":{"citing_paper":"/paper/2605.10806"},"observation_digest":"sha256:c8f4a7f51b5ef3fa8b91f79f076821548635620fb4591e8597482aa051562d6d","observation_id":"a925d8a5-e5f4-444a-9a27-5be0fc119fa8","resolution":{"observed_at":"2026-05-12T11:41:33.355422Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-02T06:30:47.504484+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-02T06:30:47.504484+00:00","source":"crossref"},{"observed_at":"2026-08-02T06:30:44.765945+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"std= 0 means assigning the same score to every dimension of every video and provides no discriminating information whatsoever; std<0.3 is treated as near-constant","venue":null,"work_id":"93c96285-ee2f-43b1-ac94-57b4620f154b","year":null},"citing_paper":{"arxiv_id":"2605.10806","last_updated":"2026-05-11T16:30:51Z","snapshot_observed_at":"2026-07-06T23:22:42.512039Z","submitted_at":"2026-05-11T16:30:51Z","title":"PhyGround: Benchmarking Physical Reasoning in Generative World Models","version":1},"reference_index":61,"source":"pdf_text","source_observed_at":"2026-05-12T05:17:30.010064Z"},"links":{"citing_paper":"/paper/2605.10806"},"observation_digest":"sha256:0c710424b2ba4fb733fe598b5e33e8dc6018dfe4f5007c7c17e4541cb9ab3d38","observation_id":"a2449bc4-6e59-4804-bb03-06e8485e1f4f","resolution":{"observed_at":"2026-05-12T11:41:33.361808Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-02T06:30:47.504484+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-02T06:30:47.504484+00:00","source":"crossref"},{"observed_at":"2026-08-02T06:30:44.765945+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"A 100% copy-paste rate means the annotator does not differentiate evaluation dimensions (e.g., gravity vs","venue":null,"work_id":"692bfee7-6924-40cb-aebf-a06a8ada8d15","year":null},"citing_paper":{"arxiv_id":"2605.10806","last_updated":"2026-05-11T16:30:51Z","snapshot_observed_at":"2026-07-06T23:22:42.512039Z","submitted_at":"2026-05-11T16:30:51Z","title":"PhyGround: Benchmarking Physical Reasoning in Generative World Models","version":1},"reference_index":62,"source":"pdf_text","source_observed_at":"2026-05-12T05:17:30.010064Z"},"links":{"citing_paper":"/paper/2605.10806"},"observation_digest":"sha256:4814bfabfa4fac2184f1ca0ace4d37408465bc4114e156e422d395797c0ae767","observation_id":"fac7664f-e408-42e1-9f53-f27680acdd48","resolution":{"observed_at":"2026-05-12T11:41:33.406474Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-02T06:30:47.504484+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-02T06:30:47.504484+00:00","source":"crossref"},{"observed_at":"2026-08-02T06:30:44.765945+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":null,"venue":null,"work_id":"fef2e394-bb24-41a6-b62a-a54d5f97adea","year":null},"citing_paper":{"arxiv_id":"2605.10806","last_updated":"2026-05-11T16:30:51Z","snapshot_observed_at":"2026-07-06T23:22:42.512039Z","submitted_at":"2026-05-11T16:30:51Z","title":"PhyGround: Benchmarking Physical Reasoning in Generative World Models","version":1},"reference_index":63,"source":"pdf_text","source_observed_at":"2026-05-12T05:17:30.010064Z"},"links":{"citing_paper":"/paper/2605.10806"},"observation_digest":"sha256:81fee815b0219b51b620f4e06a95280fa9dedf0f76d7893877b0203e7c5986ee","observation_id":"2c593aa0-ce88-4673-841a-6713409258f2","resolution":{"observed_at":"2026-05-12T11:41:33.347236Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-02T06:30:47.504484+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-02T06:30:47.504484+00:00","source":"crossref"},{"observed_at":"2026-08-02T06:30:44.765945+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":"1000.0794","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"skip when hard","venue":null,"work_id":"16cc253e-e267-4cbd-9f5d-6ccb195ad79f","year":2037},"citing_paper":{"arxiv_id":"2605.10806","last_updated":"2026-05-11T16:30:51Z","snapshot_observed_at":"2026-07-06T23:22:42.512039Z","submitted_at":"2026-05-11T16:30:51Z","title":"PhyGround: Benchmarking Physical Reasoning in Generative World Models","version":1},"reference_index":64,"source":"pdf_text","source_observed_at":"2026-05-12T05:17:30.010064Z"},"links":{"citing_paper":"/paper/2605.10806"},"observation_digest":"sha256:2b7a1febd0b243604daf45dae2d59465f42e693e7dbd88e02a7e438fe78ea11b","observation_id":"9aa514a1-9e6b-47f4-a857-4a95ae8d313d","resolution":{"observed_at":"2026-05-12T05:21:28.187296Z","resolver_source":"arxiv_id","status":"malformed_identifier"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-02T06:30:47.504484+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-02T06:30:47.504484+00:00","source":"crossref"},{"observed_at":"2026-08-02T06:30:44.765945+00:00","source":"retraction_watch"}],"state":"measured"}}],"paper":{"arxiv_id":"2605.10806","last_updated":"2026-05-11T16:30:51Z","latest_version":1,"primary_category":"cs.CV","snapshot_observed_at":"2026-07-06T23:22:42.512039Z","submitted_at":"2026-05-11T16:30:51Z","title":"PhyGround: Benchmarking Physical Reasoning in Generative World Models"},"reference_resolution":{"displayed":64,"state_counts":{"malformed_identifier":1,"metadata_mismatch":4,"parse_uncertain":0,"unresolved":3,"verified_exact":23,"verified_fuzzy":33},"total_outbound_references":64},"refusal":"A citation records a reference. It does not transfer a finding from one paper to another.","schema":"pith.paper-citation-record.v1","standing_sources":[{"observed_at":"2026-08-02T06:30:47.504484+00:00","source":"crossref"},{"observed_at":"2026-08-02T06:30:44.765945+00:00","source":"retraction_watch"}],"thesis":"As of 2 August 2026, this Paper Citation Record lists 64 of 64 outbound references and 1 inbound Pith citation observation for arXiv:2605.10806."}