{"as_of":"2026-08-05T18:28:00Z","caps":{"database_statements":6,"inbound":100,"outbound":100},"context_digest":"sha256:9ac7ee693b954fcd44e3fc03300e215f1fecd6bc67759d412f5dc92bf5b1b961","coverage":[{"denominator":64,"lane":"reference_resolution","note":"Typed states for the displayed outbound observations.","records_observed":64,"source":"paper_references, paper_reference_links","source_observed_at":"2026-05-21T17:06:34.398973Z","state":"measured"},{"denominator":70,"lane":"standing_notices","note":"One-hop event checks from named stored sources.","records_observed":70,"source":"scholarly_work_events, retraction_status_cache","source_observed_at":"2026-08-05T06:32:48.257954+00:00","state":"measured"},{"denominator":6,"lane":"inbound_itemization","note":"Pith citing papers itemized under the disclosed page cap.","records_observed":6,"source":"paper_references, paper_reference_links","source_observed_at":"2026-08-03T13:31:10.576156Z","state":"measured"},{"denominator":1,"lane":"external_citation_measurements","note":"A source-named dated measurement, never combined with another source.","records_observed":0,"source":"arxiv_reference","source_observed_at":"2026-05-13T17:08:00.948705Z","state":"measured"}],"external_citation_measurements":[],"inbound":[{"citation":{"cited_paper":{"arxiv_id":"2512.23180","last_updated":"2026-05-16T07:53:20Z","snapshot_observed_at":"2026-08-03T04:34:56.955960Z","submitted_at":"2025-12-29T03:40:05Z","title":"GaussianDWM: 3D Gaussian Driving World Model for Unified Scene Understanding and Multi-Modal Generation","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2512.23180","snapshot_observed_at":"2026-08-03T13:31:10.576156Z","title":"GaussianDWM: 3D Gaussian Driving World Model for Unified Scene Understanding and Multi -Modal Generation,","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2512.24075","last_updated":"2026-05-25T12:57:46Z","snapshot_observed_at":"2026-08-05T14:10:28.899165Z","submitted_at":"2025-12-30T08:36:35Z","title":"Evolutionary Physics-Informed Temporal Fusion for Lane-Change Intention Prediction","version":5},"reference_index":21,"source":"pdf_text","source_observed_at":"2026-08-03T13:31:10.576156Z"},"links":{"cited_paper":"/paper/2512.23180","citing_paper":"/paper/2512.24075"},"observation_digest":"sha256:de3786859c55720ae4bac79e71bb06933234b2ed855c67c6523bd2054815ab8b","observation_id":"ea80281b-b8c8-40b7-a35f-c1b13e144ed7","resolution":{"observed_at":"2026-08-03T13:31:10.576156Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2512.23180","last_updated":"2026-05-16T07:53:20Z","snapshot_observed_at":"2026-08-03T04:34:56.955960Z","submitted_at":"2025-12-29T03:40:05Z","title":"GaussianDWM: 3D Gaussian Driving World Model for Unified Scene Understanding and Multi-Modal Generation","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2512.23180","snapshot_observed_at":"2026-07-13T19:53:53.219607Z","title":"Gaussiandwm: 3d gaussian driving world model for unified scene understanding and multi-modal generation.arXiv preprint arXiv:2512.23180,","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2603.23082","last_updated":"2026-06-05T20:06:49Z","snapshot_observed_at":"2026-07-30T13:44:48.923884Z","submitted_at":"2026-03-24T11:25:52Z","title":"Spatial navigation in preclinical Alzheimer's disease: A review","version":2},"reference_index":14,"source":"pdf_text","source_observed_at":"2026-07-13T19:53:53.219607Z"},"links":{"cited_paper":"/paper/2512.23180","citing_paper":"/paper/2603.23082"},"observation_digest":"sha256:92f2987a72df31ccb0186d45cd8d5300313c5aa8f25d12c85185e92923e95608","observation_id":"9b3e77ca-d219-479a-a9ed-f9b9dfc89605","resolution":{"observed_at":"2026-07-13T19:53:53.219607Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2512.23180","last_updated":"2026-05-16T07:53:20Z","snapshot_observed_at":"2026-08-03T04:34:56.955960Z","submitted_at":"2025-12-29T03:40:05Z","title":"GaussianDWM: 3D Gaussian Driving World Model for Unified Scene Understanding and Multi-Modal Generation","version":3},"cited_work":{"arxiv_id":"2512.23180","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2512.23180","snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Gaussiandwm: 3d gaussian driving world model for unified scene understanding and multi-modal generation.arXiv preprint arXiv:2512.23180","venue":null,"work_id":"b62d9196-ae8a-4859-987a-617456f68914","year":2025},"citing_paper":{"arxiv_id":"2604.04055","last_updated":"2026-04-05T10:59:48Z","snapshot_observed_at":"2026-07-06T22:53:06.517678Z","submitted_at":"2026-04-05T10:59:48Z","title":"DINO-VO: Learning Where to Focus for Enhanced State Estimation","version":1},"reference_index":11,"source":"pdf_text","source_observed_at":"2026-05-13T17:05:52.493660Z"},"links":{"cited_paper":"/paper/2512.23180","citing_paper":"/paper/2604.04055"},"observation_digest":"sha256:ca481e176143e4eca7a640389777a8cd38490f22dc73c9cda1717f4ff4983ab6","observation_id":"5472f68e-b5be-45cc-862f-2f8792081ba3","resolution":{"observed_at":"2026-05-20T00:03:01.148630Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2512.23180","last_updated":"2026-05-16T07:53:20Z","snapshot_observed_at":"2026-08-03T04:34:56.955960Z","submitted_at":"2025-12-29T03:40:05Z","title":"GaussianDWM: 3D Gaussian Driving World Model for Unified Scene Understanding and Multi-Modal Generation","version":3},"cited_work":{"arxiv_id":"2512.23180","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2512.23180","snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Gaussiandwm: 3d gaussian driving world model for unified scene understanding and multi-modal generation.arXiv preprint arXiv:2512.23180","venue":null,"work_id":"b62d9196-ae8a-4859-987a-617456f68914","year":2025},"citing_paper":{"arxiv_id":"2604.05908","last_updated":"2026-04-07T14:11:54Z","snapshot_observed_at":"2026-07-06T22:54:30.567988Z","submitted_at":"2026-04-07T14:11:54Z","title":"Appearance Decomposition Gaussian Splatting for Multi-Traversal Reconstruction","version":1},"reference_index":8,"source":"pdf_text","source_observed_at":"2026-05-10T19:11:56.114485Z"},"links":{"cited_paper":"/paper/2512.23180","citing_paper":"/paper/2604.05908"},"observation_digest":"sha256:d4e9af0242fa301ec98fc957be51ef359c856dfc94346e502c8c3a62a3f1f658","observation_id":"5e55bb5d-c3d0-4818-825d-0134a35f8e41","resolution":{"observed_at":"2026-05-20T00:03:01.148630Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2512.23180","last_updated":"2026-05-16T07:53:20Z","snapshot_observed_at":"2026-08-03T04:34:56.955960Z","submitted_at":"2025-12-29T03:40:05Z","title":"GaussianDWM: 3D Gaussian Driving World Model for Unified Scene Understanding and Multi-Modal Generation","version":3},"cited_work":{"arxiv_id":"2512.23180","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2512.23180","snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Gaussiandwm: 3d gaussian driving world model for unified scene understanding and multi-modal generation.arXiv preprint arXiv:2512.23180","venue":null,"work_id":"b62d9196-ae8a-4859-987a-617456f68914","year":2025},"citing_paper":{"arxiv_id":"2604.23902","last_updated":"2026-04-26T22:26:20Z","snapshot_observed_at":"2026-07-06T23:10:02.659384Z","submitted_at":"2026-04-26T22:26:20Z","title":"LLM-Augmented Traffic Signal Control with LSTM-Based Traffic State Prediction and Safety-Constrained Decision Support","version":1},"reference_index":15,"source":"pdf_text","source_observed_at":"2026-05-08T05:59:22.177960Z"},"links":{"cited_paper":"/paper/2512.23180","citing_paper":"/paper/2604.23902"},"observation_digest":"sha256:8049747fcb01375837a83830f5a4dc2b3f89b1b9be7a0eb7c1cf015556fa5aa0","observation_id":"613b953d-2777-404a-9919-556c9747648b","resolution":{"observed_at":"2026-05-20T00:03:01.148630Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2512.23180","last_updated":"2026-05-16T07:53:20Z","snapshot_observed_at":"2026-08-03T04:34:56.955960Z","submitted_at":"2025-12-29T03:40:05Z","title":"GaussianDWM: 3D Gaussian Driving World Model for Unified Scene Understanding and Multi-Modal Generation","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2512.23180","snapshot_observed_at":"2026-07-14T13:19:08.066069Z","title":null,"venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2607.10237","last_updated":"2026-07-11T09:53:52Z","snapshot_observed_at":"2026-08-05T13:03:09.634573Z","submitted_at":"2026-07-11T09:53:52Z","title":"CoSAG: Compact Semantic Anchor Gaussians via Training-Free Rate-Distortion Coding","version":1},"reference_index":45,"source":"arxiv_source","source_observed_at":"2026-07-14T13:19:08.066069Z"},"links":{"cited_paper":"/paper/2512.23180","citing_paper":"/paper/2607.10237"},"observation_digest":"sha256:8d1cf9b76b3871e1cd2743bd36488776cd8612267b4195b74c27df5435e577ad","observation_id":"8256753b-ba6b-4e00-92b4-ce3ec79adf70","resolution":{"observed_at":"2026-07-14T13:19:08.066069Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}}],"links":{"evidence":"/evidence","html":"/paper/2512.23180/citation-record","integrity":"/paper/2512.23180/integrity","json":"/paper/2512.23180/citation-record.json","paper":"/paper/2512.23180"},"outbound":[{"citation":{"cited_paper":{"arxiv_id":"2303.08774","last_updated":"2024-03-04T06:01:33Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2023-03-15T17:15:04Z","title":"GPT-4 Technical Report","version":6},"cited_work":{"arxiv_id":"2303.08774","doi":"10.1002/tea.20265","metadata_source":"pith","pith_arxiv_id":"2303.08774","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"GPT-4 Technical Report","venue":"cs.CL","work_id":"b928e041-6991-4c08-8c81-0359e4097c7b","year":2023},"citing_paper":{"arxiv_id":"2512.23180","last_updated":"2026-05-16T07:53:20Z","snapshot_observed_at":"2026-08-03T04:34:56.955960Z","submitted_at":"2025-12-29T03:40:05Z","title":"GaussianDWM: 3D Gaussian Driving World Model for Unified Scene Understanding and Multi-Modal Generation","version":3},"reference_index":1,"source":"pdf_text","source_observed_at":"2026-05-21T17:06:34.398973Z"},"links":{"cited_paper":"/paper/2303.08774","citing_paper":"/paper/2512.23180"},"observation_digest":"sha256:0ba0c1cba10939a0fab0c74972620e67c0faf13662ce3563635a31acf77fb28a","observation_id":"41be73b1-d803-4c8f-9772-7e0dc7f0f13d","resolution":{"observed_at":"2026-05-21T17:10:25.121871Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2311.15127","last_updated":"2023-11-25T22:28:38Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2023-11-25T22:28:38Z","title":"Stable Video Diffusion: Scaling Latent Video Diffusion Models to Large Datasets","version":1},"cited_work":{"arxiv_id":"2311.15127","doi":"10.48550/arxiv.2311.15127","metadata_source":"pith","pith_arxiv_id":"2311.15127","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Stable Video Diffusion: Scaling Latent Video Diffusion Models to Large Datasets","venue":"cs.CV","work_id":"4f68eada-27e3-437a-a2fe-6e4ca524d0d3","year":2023},"citing_paper":{"arxiv_id":"2512.23180","last_updated":"2026-05-16T07:53:20Z","snapshot_observed_at":"2026-08-03T04:34:56.955960Z","submitted_at":"2025-12-29T03:40:05Z","title":"GaussianDWM: 3D Gaussian Driving World Model for Unified Scene Understanding and Multi-Modal Generation","version":3},"reference_index":2,"source":"pdf_text","source_observed_at":"2026-05-21T17:06:34.398973Z"},"links":{"cited_paper":"/paper/2311.15127","citing_paper":"/paper/2512.23180"},"observation_digest":"sha256:8504d05f8c3cf382d90fa87032de366b93268a56f6599b20e7951a70d4e8fe71","observation_id":"c3e9ddd1-9cf0-4fff-927c-4772dfa87640","resolution":{"observed_at":"2026-05-21T17:10:25.095962Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-07-11T15:53:24.472993+00:00","source":"crossref_status_cache"},{"observed_at":"2026-07-11T15:53:24.472993+00:00","source":"openalex_status_cache"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"nuscenes: A multi- modal dataset for autonomous driving","venue":null,"work_id":"8192dcaf-da9a-4a10-a064-5fc0746066d9","year":2020},"citing_paper":{"arxiv_id":"2512.23180","last_updated":"2026-05-16T07:53:20Z","snapshot_observed_at":"2026-08-03T04:34:56.955960Z","submitted_at":"2025-12-29T03:40:05Z","title":"GaussianDWM: 3D Gaussian Driving World Model for Unified Scene Understanding and Multi-Modal Generation","version":3},"reference_index":3,"source":"pdf_text","source_observed_at":"2026-05-21T17:06:34.398973Z"},"links":{"citing_paper":"/paper/2512.23180"},"observation_digest":"sha256:68bb371ad2582a4ce3e1e6d908ed79e1ea5b83b174336bb2b8c1b1a5eda99e37","observation_id":"5284ef9e-cd56-4670-8711-e36663f8e5d3","resolution":{"observed_at":"2026-05-21T17:10:26.072050Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2311.18561","last_updated":"2025-08-23T01:56:43Z","snapshot_observed_at":"2026-07-06T16:55:01.268118Z","submitted_at":"2023-11-30T13:53:50Z","title":"Periodic Vibration Gaussian: Dynamic Urban Scene Reconstruction and Real-time Rendering","version":3},"cited_work":{"arxiv_id":"2311.18561","doi":"10.48550/arxiv.2311.18561","metadata_source":"arxiv_reference","pith_arxiv_id":"2311.18561","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Periodic vibration gaussian: Dynamic urban scene reconstruction and real-time rendering","venue":"arXiv (Cornell University)","work_id":"b2ef8501-a724-4c23-9a5c-eb58122910f8","year":2023},"citing_paper":{"arxiv_id":"2512.23180","last_updated":"2026-05-16T07:53:20Z","snapshot_observed_at":"2026-08-03T04:34:56.955960Z","submitted_at":"2025-12-29T03:40:05Z","title":"GaussianDWM: 3D Gaussian Driving World Model for Unified Scene Understanding and Multi-Modal Generation","version":3},"reference_index":4,"source":"pdf_text","source_observed_at":"2026-05-21T17:06:34.398973Z"},"links":{"cited_paper":"/paper/2311.18561","citing_paper":"/paper/2512.23180"},"observation_digest":"sha256:76b5638a791e18ef967bce9e7b748a9d551511dbc00a74ce0a0f0249c9cff693","observation_id":"2505e440-b2ce-441b-a6fb-784fb1ed269a","resolution":{"observed_at":"2026-05-21T17:10:25.109113Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2504.08361","last_updated":"2025-04-11T08:51:23Z","snapshot_observed_at":"2026-08-04T08:32:17.676943Z","submitted_at":"2025-04-11T08:51:23Z","title":"SN-LiDAR: Semantic Neural Fields for Novel Space-time View LiDAR Synthesis","version":1},"cited_work":{"arxiv_id":"2504.08361","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2504.08361","snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Sn-lidar: Semantic neural fields for novel space-time view lidar syn- thesis","venue":null,"work_id":"c7f585bb-58a7-41bb-9b8b-db1a07d724c1","year":2025},"citing_paper":{"arxiv_id":"2512.23180","last_updated":"2026-05-16T07:53:20Z","snapshot_observed_at":"2026-08-03T04:34:56.955960Z","submitted_at":"2025-12-29T03:40:05Z","title":"GaussianDWM: 3D Gaussian Driving World Model for Unified Scene Understanding and Multi-Modal Generation","version":3},"reference_index":5,"source":"pdf_text","source_observed_at":"2026-05-21T17:06:34.398973Z"},"links":{"cited_paper":"/paper/2504.08361","citing_paper":"/paper/2512.23180"},"observation_digest":"sha256:eff2fdd831c8d0ea8dcb4c84091a2e345f2435a9b944a3c2fadf11b015be8bd3","observation_id":"c058bf0a-32c7-4f3f-8a6a-ec472da5f164","resolution":{"observed_at":"2026-05-21T17:10:25.131282Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"How far are we to gpt-4v? closing the gap to commercial multimodal models with open-source suites.Science China Information Sciences, 67(12):220101","venue":null,"work_id":"2b877eeb-1eae-493d-a112-aeb4ce5d39c6","year":null},"citing_paper":{"arxiv_id":"2512.23180","last_updated":"2026-05-16T07:53:20Z","snapshot_observed_at":"2026-08-03T04:34:56.955960Z","submitted_at":"2025-12-29T03:40:05Z","title":"GaussianDWM: 3D Gaussian Driving World Model for Unified Scene Understanding and Multi-Modal Generation","version":3},"reference_index":6,"source":"pdf_text","source_observed_at":"2026-05-21T17:06:34.398973Z"},"links":{"citing_paper":"/paper/2512.23180"},"observation_digest":"sha256:97b706f8fadc414ec36fd673d0b4dd0520090adee153d775e575b006eb06881a","observation_id":"7b997547-040e-4f6f-8ae8-bb43eeb053d9","resolution":{"observed_at":"2026-05-21T17:10:25.956530Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Omnire: Omni urban scene reconstruction","venue":null,"work_id":"cb1a25a0-04ee-4160-83c2-67ab0c5a1ae1","year":2025},"citing_paper":{"arxiv_id":"2512.23180","last_updated":"2026-05-16T07:53:20Z","snapshot_observed_at":"2026-08-03T04:34:56.955960Z","submitted_at":"2025-12-29T03:40:05Z","title":"GaussianDWM: 3D Gaussian Driving World Model for Unified Scene Understanding and Multi-Modal Generation","version":3},"reference_index":7,"source":"pdf_text","source_observed_at":"2026-05-21T17:06:34.398973Z"},"links":{"citing_paper":"/paper/2512.23180"},"observation_digest":"sha256:8fcfabf75b9729b9b80755aa2940ba277c7e6440674ca4cbd688db7424bc6363","observation_id":"53a8afd7-33b1-4cff-98aa-c60580fdcf2c","resolution":{"observed_at":"2026-05-21T17:10:25.953122Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2312.09076","last_updated":"2026-07-10T06:33:14Z","snapshot_observed_at":"2026-08-05T14:08:06.248024Z","submitted_at":"2023-12-14T16:11:42Z","title":"ProSGNeRF: Progressive Dynamic Neural Scene Graph with Frequency Modulated Foundation Model in Urban Scenes","version":4},"cited_work":{"arxiv_id":"2312.09076","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2312.09076","snapshot_observed_at":"2026-07-13T02:17:00.989019Z","title":"Prosgnerf: Progressive dynamic neural scene graph with frequency modulated auto-encoder in urban scenes","venue":null,"work_id":"4088f50f-ed41-453f-9253-0a0e9d6b5ee7","year":2023},"citing_paper":{"arxiv_id":"2512.23180","last_updated":"2026-05-16T07:53:20Z","snapshot_observed_at":"2026-08-03T04:34:56.955960Z","submitted_at":"2025-12-29T03:40:05Z","title":"GaussianDWM: 3D Gaussian Driving World Model for Unified Scene Understanding and Multi-Modal Generation","version":3},"reference_index":8,"source":"pdf_text","source_observed_at":"2026-05-21T17:06:34.398973Z"},"links":{"cited_paper":"/paper/2312.09076","citing_paper":"/paper/2512.23180"},"observation_digest":"sha256:b96c8fe44cd8f059733a16f35712ae1bdd68f729c2009fce03e776541c0abac2","observation_id":"4cd816e4-7b20-48c0-ac77-4bc1a69215b4","resolution":{"observed_at":"2026-07-13T02:17:00.989019Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Plgslam: Progressive neural scene represenation with local to global bundle adjustment","venue":null,"work_id":"e90c14be-e0aa-44f4-a231-39f022dfed4c","year":2024},"citing_paper":{"arxiv_id":"2512.23180","last_updated":"2026-05-16T07:53:20Z","snapshot_observed_at":"2026-08-03T04:34:56.955960Z","submitted_at":"2025-12-29T03:40:05Z","title":"GaussianDWM: 3D Gaussian Driving World Model for Unified Scene Understanding and Multi-Modal Generation","version":3},"reference_index":9,"source":"pdf_text","source_observed_at":"2026-05-21T17:06:34.398973Z"},"links":{"citing_paper":"/paper/2512.23180"},"observation_digest":"sha256:ceb7ad2f4941b8544c20e76c882faf11e12a5a827a7e8e67b636d816738e20a4","observation_id":"917f9d43-2d73-4b9a-b2e5-2bffed6dd604","resolution":{"observed_at":"2026-05-21T17:10:25.959863Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":"2512.03422","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-07-09T10:46:11.607069Z","title":"What is the best 3d scene representation for robotics? from geometric to foundation models","venue":null,"work_id":"e2f37cab-d139-438b-9969-f14a0c9fc8f9","year":2025},"citing_paper":{"arxiv_id":"2512.23180","last_updated":"2026-05-16T07:53:20Z","snapshot_observed_at":"2026-08-03T04:34:56.955960Z","submitted_at":"2025-12-29T03:40:05Z","title":"GaussianDWM: 3D Gaussian Driving World Model for Unified Scene Understanding and Multi-Modal Generation","version":3},"reference_index":10,"source":"pdf_text","source_observed_at":"2026-05-21T17:06:34.398973Z"},"links":{"citing_paper":"/paper/2512.23180"},"observation_digest":"sha256:782200f155ffe2e4dfb9afe45dd92a7528f3c9162b4d2c5bd37a6ea361cfb9cb","observation_id":"bc5b648f-dfdd-44ad-a3a4-db41b2c98c47","resolution":{"observed_at":"2026-05-21T17:10:25.100430Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2506.18678","last_updated":"2025-08-19T06:02:44Z","snapshot_observed_at":"2026-07-06T21:46:20.835997Z","submitted_at":"2025-06-23T14:22:29Z","title":"MCN-SLAM: Multi-Agent Collaborative Neural SLAM with Hybrid Implicit Neural Scene Representation","version":2},"cited_work":{"arxiv_id":"2506.18678","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2506.18678","snapshot_observed_at":"2026-06-29T09:13:16.628409Z","title":"Mcn-slam: Multi-agent collaborative neural slam with hybrid implicit neural scene representation","venue":null,"work_id":"2ba917c0-33fd-4fcd-b06a-62edba7df777","year":2025},"citing_paper":{"arxiv_id":"2512.23180","last_updated":"2026-05-16T07:53:20Z","snapshot_observed_at":"2026-08-03T04:34:56.955960Z","submitted_at":"2025-12-29T03:40:05Z","title":"GaussianDWM: 3D Gaussian Driving World Model for Unified Scene Understanding and Multi-Modal Generation","version":3},"reference_index":11,"source":"pdf_text","source_observed_at":"2026-05-21T17:06:34.398973Z"},"links":{"cited_paper":"/paper/2506.18678","citing_paper":"/paper/2512.23180"},"observation_digest":"sha256:b6565378e0ad806fac9627b82d3dce1f588be17552f3a29c13ffe13f2a29eb6a","observation_id":"885d050a-db71-4834-b741-95ff520237b4","resolution":{"observed_at":"2026-05-21T17:10:25.136388Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Mne-slam: Multi-agent neural slam for mobile robots","venue":null,"work_id":"fa153a21-5d2b-4561-bd35-3fb1a3b3b95b","year":2025},"citing_paper":{"arxiv_id":"2512.23180","last_updated":"2026-05-16T07:53:20Z","snapshot_observed_at":"2026-08-03T04:34:56.955960Z","submitted_at":"2025-12-29T03:40:05Z","title":"GaussianDWM: 3D Gaussian Driving World Model for Unified Scene Understanding and Multi-Modal Generation","version":3},"reference_index":12,"source":"pdf_text","source_observed_at":"2026-05-21T17:06:34.398973Z"},"links":{"citing_paper":"/paper/2512.23180"},"observation_digest":"sha256:85933d1f39e27f93ffada14a97e93dc3c65a01984ab96fe8d52453dad08f2b82","observation_id":"781b0892-0ffd-437d-957f-cc935670672f","resolution":{"observed_at":"2026-05-21T17:10:25.950372Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":null,"venue":null,"work_id":"42213c9c-1c6e-415d-8625-b22e7ae682a9","year":2025},"citing_paper":{"arxiv_id":"2512.23180","last_updated":"2026-05-16T07:53:20Z","snapshot_observed_at":"2026-08-03T04:34:56.955960Z","submitted_at":"2025-12-29T03:40:05Z","title":"GaussianDWM: 3D Gaussian Driving World Model for Unified Scene Understanding and Multi-Modal Generation","version":3},"reference_index":13,"source":"pdf_text","source_observed_at":"2026-05-21T17:06:34.398973Z"},"links":{"citing_paper":"/paper/2512.23180"},"observation_digest":"sha256:c330bc4cc1bc98b8f963b294d69bb12822e50b965e29b4e381240410223f6a34","observation_id":"0b48d5e7-a59c-467c-91b8-2b388745e3e6","resolution":{"observed_at":"2026-05-21T17:10:26.068813Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":"2505.18992","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-30T08:44:28.143258Z","title":"Vpgs-slam: V oxel-based progressive 3d gaussian slam in large-scale scenes","venue":null,"work_id":"21be5359-9479-4b51-b594-1bbd64ec35cf","year":2025},"citing_paper":{"arxiv_id":"2512.23180","last_updated":"2026-05-16T07:53:20Z","snapshot_observed_at":"2026-08-03T04:34:56.955960Z","submitted_at":"2025-12-29T03:40:05Z","title":"GaussianDWM: 3D Gaussian Driving World Model for Unified Scene Understanding and Multi-Modal Generation","version":3},"reference_index":14,"source":"pdf_text","source_observed_at":"2026-05-21T17:06:34.398973Z"},"links":{"citing_paper":"/paper/2512.23180"},"observation_digest":"sha256:d9e1a2b85f2af385e1a74efd38f6f5828235c7b06a3386f2bcd096446808f60e","observation_id":"624ad7d5-8587-4c24-8d90-84b11c97f0a9","resolution":{"observed_at":"2026-05-21T17:10:25.192004Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-07-09T09:16:07.203676Z","title":"Scaling recti- fied flow transformers for high-resolution image synthesis","venue":null,"work_id":"a92e35c5-be55-4ec3-80cd-8bf130aa79ac","year":null},"citing_paper":{"arxiv_id":"2512.23180","last_updated":"2026-05-16T07:53:20Z","snapshot_observed_at":"2026-08-03T04:34:56.955960Z","submitted_at":"2025-12-29T03:40:05Z","title":"GaussianDWM: 3D Gaussian Driving World Model for Unified Scene Understanding and Multi-Modal Generation","version":3},"reference_index":15,"source":"pdf_text","source_observed_at":"2026-05-21T17:06:34.398973Z"},"links":{"citing_paper":"/paper/2512.23180"},"observation_digest":"sha256:9a337458e2f6fc878c0433b8c4f34129bd4be16506494d30c666ccb319b6cd1b","observation_id":"af20115f-79e0-430d-a306-b26fc0222edb","resolution":{"observed_at":"2026-05-21T17:10:26.065974Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2310.02601","last_updated":"2024-05-03T04:50:27Z","snapshot_observed_at":"2026-07-06T16:27:31.403609Z","submitted_at":"2023-10-04T06:14:06Z","title":"MagicDrive: Street View Generation with Diverse 3D Geometry Control","version":7},"cited_work":{"arxiv_id":"2310.02601","doi":null,"metadata_source":"pith","pith_arxiv_id":"2310.02601","snapshot_observed_at":"2026-07-08T03:24:28.755891Z","title":"Magicdrive: Street view generation with diverse 3d geometry control","venue":"cs.CV","work_id":"0d62d074-2eec-4693-9b9b-3dd8367909d9","year":2023},"citing_paper":{"arxiv_id":"2512.23180","last_updated":"2026-05-16T07:53:20Z","snapshot_observed_at":"2026-08-03T04:34:56.955960Z","submitted_at":"2025-12-29T03:40:05Z","title":"GaussianDWM: 3D Gaussian Driving World Model for Unified Scene Understanding and Multi-Modal Generation","version":3},"reference_index":16,"source":"pdf_text","source_observed_at":"2026-05-21T17:06:34.398973Z"},"links":{"cited_paper":"/paper/2310.02601","citing_paper":"/paper/2512.23180"},"observation_digest":"sha256:8c844feb01a5abc08ee349254fb2f090c0dff0fb40ab711731c89175cfc58694","observation_id":"5d4f4b16-3e17-41b0-a456-f9dc19645bc1","resolution":{"observed_at":"2026-05-21T17:10:25.187814Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Vista: A generalizable driving world model with high fidelity and versatile controllability.Advances in Neural Information Processing Systems, 37:91560–91596","venue":null,"work_id":"c07468d2-0eb1-4912-a1e4-468777983ac7","year":2024},"citing_paper":{"arxiv_id":"2512.23180","last_updated":"2026-05-16T07:53:20Z","snapshot_observed_at":"2026-08-03T04:34:56.955960Z","submitted_at":"2025-12-29T03:40:05Z","title":"GaussianDWM: 3D Gaussian Driving World Model for Unified Scene Understanding and Multi-Modal Generation","version":3},"reference_index":17,"source":"pdf_text","source_observed_at":"2026-05-21T17:06:34.398973Z"},"links":{"citing_paper":"/paper/2512.23180"},"observation_digest":"sha256:0fb934db5f30d589ff9d7c27393f757bd3a9514e289a28b1f74b58d0c5792228","observation_id":"b4d6c827-e202-40dd-b3f3-189e59bac8ae","resolution":{"observed_at":"2026-05-21T17:10:26.062079Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Mocount: Motion-based repetitive ac- tion counting","venue":null,"work_id":"c2ebc357-13e2-4c57-8b96-586677c52fe5","year":2025},"citing_paper":{"arxiv_id":"2512.23180","last_updated":"2026-05-16T07:53:20Z","snapshot_observed_at":"2026-08-03T04:34:56.955960Z","submitted_at":"2025-12-29T03:40:05Z","title":"GaussianDWM: 3D Gaussian Driving World Model for Unified Scene Understanding and Multi-Modal Generation","version":3},"reference_index":18,"source":"pdf_text","source_observed_at":"2026-05-21T17:06:34.398973Z"},"links":{"citing_paper":"/paper/2512.23180"},"observation_digest":"sha256:82f2ef52978f8f42a54ec8a099291f191a87bcf8d85097c7493a2fe4169b8307","observation_id":"15da5d56-1ed3-45fb-8f84-1e11931f69be","resolution":{"observed_at":"2026-05-21T17:10:26.058765Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Dist-4d: Disentangled spa- tiotemporal diffusion with metric depth for 4d driving scene generation","venue":null,"work_id":"634289bc-f73c-4aff-baa4-6041eab8c535","year":null},"citing_paper":{"arxiv_id":"2512.23180","last_updated":"2026-05-16T07:53:20Z","snapshot_observed_at":"2026-08-03T04:34:56.955960Z","submitted_at":"2025-12-29T03:40:05Z","title":"GaussianDWM: 3D Gaussian Driving World Model for Unified Scene Understanding and Multi-Modal Generation","version":3},"reference_index":19,"source":"pdf_text","source_observed_at":"2026-05-21T17:06:34.398973Z"},"links":{"citing_paper":"/paper/2512.23180"},"observation_digest":"sha256:d59b546eeeb8a8b7aeb8df08b18ebdbbeda0d31ecb017ec382d86ab349476988","observation_id":"8d4dfb09-b038-4e2b-b586-5b735883cde3","resolution":{"observed_at":"2026-05-21T17:10:26.055578Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2503.15208","last_updated":"2025-03-19T13:49:48Z","snapshot_observed_at":"2026-08-01T23:42:29.749765Z","submitted_at":"2025-03-19T13:49:48Z","title":"DiST-4D: Disentangled Spatiotemporal Diffusion with Metric Depth for 4D Driving Scene Generation","version":1},"cited_work":{"arxiv_id":"2503.15208","doi":null,"metadata_source":"pith","pith_arxiv_id":"2503.15208","snapshot_observed_at":"2026-07-08T03:24:28.768344Z","title":"Dist-4d: Disentangled spatiotemporal diffusion with metric depth for 4d driving scene generation","venue":"cs.CV","work_id":"48e2d19a-8aee-4a66-b8d4-1bc3fc521e17","year":2025},"citing_paper":{"arxiv_id":"2512.23180","last_updated":"2026-05-16T07:53:20Z","snapshot_observed_at":"2026-08-03T04:34:56.955960Z","submitted_at":"2025-12-29T03:40:05Z","title":"GaussianDWM: 3D Gaussian Driving World Model for Unified Scene Understanding and Multi-Modal Generation","version":3},"reference_index":20,"source":"pdf_text","source_observed_at":"2026-05-21T17:06:34.398973Z"},"links":{"cited_paper":"/paper/2503.15208","citing_paper":"/paper/2512.23180"},"observation_digest":"sha256:abd28a5245d7a2933c8701cd32e4484315fc4719f430aa034deba9419852acd9","observation_id":"6dbf3dd0-572f-40e1-8dcf-9a7bb3181a74","resolution":{"observed_at":"2026-05-21T17:10:25.148944Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2507.00886","last_updated":"2025-07-01T15:52:59Z","snapshot_observed_at":"2026-08-04T07:18:03.803134Z","submitted_at":"2025-07-01T15:52:59Z","title":"GaussianVLM: Scene-centric 3D Vision-Language Models using Language-aligned Gaussian Splats for Embodied Reasoning and Beyond","version":1},"cited_work":{"arxiv_id":"2507.00886","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2507.00886","snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Gaussianvlm: Scene-centric 3d vision-language models using language- aligned gaussian splats for embodied reasoning and beyond","venue":null,"work_id":"4bd5b345-8b04-4090-b596-414a2afcc851","year":2025},"citing_paper":{"arxiv_id":"2512.23180","last_updated":"2026-05-16T07:53:20Z","snapshot_observed_at":"2026-08-03T04:34:56.955960Z","submitted_at":"2025-12-29T03:40:05Z","title":"GaussianDWM: 3D Gaussian Driving World Model for Unified Scene Understanding and Multi-Modal Generation","version":3},"reference_index":21,"source":"pdf_text","source_observed_at":"2026-05-21T17:06:34.398973Z"},"links":{"cited_paper":"/paper/2507.00886","citing_paper":"/paper/2512.23180"},"observation_digest":"sha256:aba5763478e7a912946f7a4558edb222d42b4377c3d6fa6e35bd4fe01b3c83dd","observation_id":"86b7e462-9ced-4beb-bd69-2583450238f1","resolution":{"observed_at":"2026-05-21T17:10:25.178774Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"3d-llm: In- jecting the 3d world into large language models.Advances 17 in Neural Information Processing Systems, 36:20482–20494","venue":null,"work_id":"8b966b8d-9e9d-4c03-9065-4d6f2dd20e09","year":null},"citing_paper":{"arxiv_id":"2512.23180","last_updated":"2026-05-16T07:53:20Z","snapshot_observed_at":"2026-08-03T04:34:56.955960Z","submitted_at":"2025-12-29T03:40:05Z","title":"GaussianDWM: 3D Gaussian Driving World Model for Unified Scene Understanding and Multi-Modal Generation","version":3},"reference_index":22,"source":"pdf_text","source_observed_at":"2026-05-21T17:06:34.398973Z"},"links":{"citing_paper":"/paper/2512.23180"},"observation_digest":"sha256:e5e884aaba48e382f8bb541bc28053758a19e59d3e71fd5316e606806808c857","observation_id":"912388c2-61f7-4054-bad7-55fe772b783d","resolution":{"observed_at":"2026-05-21T17:10:26.052326Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2309.17080","last_updated":"2023-09-29T09:20:37Z","snapshot_observed_at":"2026-07-06T16:25:21.571679Z","submitted_at":"2023-09-29T09:20:37Z","title":"GAIA-1: A Generative World Model for Autonomous Driving","version":1},"cited_work":{"arxiv_id":"2309.17080","doi":null,"metadata_source":"pith","pith_arxiv_id":"2309.17080","snapshot_observed_at":"2026-07-10T11:47:02.949785Z","title":"GAIA-1: A Generative World Model for Autonomous Driving","venue":"cs.CV","work_id":"313484e6-a442-4522-8e19-d07e502844a8","year":2023},"citing_paper":{"arxiv_id":"2512.23180","last_updated":"2026-05-16T07:53:20Z","snapshot_observed_at":"2026-08-03T04:34:56.955960Z","submitted_at":"2025-12-29T03:40:05Z","title":"GaussianDWM: 3D Gaussian Driving World Model for Unified Scene Understanding and Multi-Modal Generation","version":3},"reference_index":24,"source":"pdf_text","source_observed_at":"2026-05-21T17:06:34.398973Z"},"links":{"cited_paper":"/paper/2309.17080","citing_paper":"/paper/2512.23180"},"observation_digest":"sha256:8b212049fd3172e74f05a293e0c8303160e125aaa808eecc57087d369a0910cd","observation_id":"58fdc6c1-aacf-42d3-8dd1-ac838e2a0feb","resolution":{"observed_at":"2026-05-21T17:10:25.153507Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-07-09T07:46:04.037239Z","title":"3d gaussian splatting for real-time radiance field rendering.ACM Trans","venue":null,"work_id":"f02b3a8e-cb41-4f1f-ae1c-75b487ba20f9","year":null},"citing_paper":{"arxiv_id":"2512.23180","last_updated":"2026-05-16T07:53:20Z","snapshot_observed_at":"2026-08-03T04:34:56.955960Z","submitted_at":"2025-12-29T03:40:05Z","title":"GaussianDWM: 3D Gaussian Driving World Model for Unified Scene Understanding and Multi-Modal Generation","version":3},"reference_index":25,"source":"pdf_text","source_observed_at":"2026-05-21T17:06:34.398973Z"},"links":{"citing_paper":"/paper/2512.23180"},"observation_digest":"sha256:bb270856ff0ea6212283dbfc66536a986a3d48cecd58a9364cd8550c66c6afa7","observation_id":"d92bbf53-ac26-4ddd-ab65-9a832bd97808","resolution":{"observed_at":"2026-05-21T17:10:26.049436Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Segment any- thing","venue":null,"work_id":"84623829-7b68-4163-a3a0-f05bbca1bf39","year":2023},"citing_paper":{"arxiv_id":"2512.23180","last_updated":"2026-05-16T07:53:20Z","snapshot_observed_at":"2026-08-03T04:34:56.955960Z","submitted_at":"2025-12-29T03:40:05Z","title":"GaussianDWM: 3D Gaussian Driving World Model for Unified Scene Understanding and Multi-Modal Generation","version":3},"reference_index":26,"source":"pdf_text","source_observed_at":"2026-05-21T17:06:34.398973Z"},"links":{"citing_paper":"/paper/2512.23180"},"observation_digest":"sha256:7fb4b71130f2500109ca4cd3dec8c37667bb9f74a04f02ca6911a42adb0d41e9","observation_id":"5724e774-4ef0-4ad0-b01a-c8b1da45ae31","resolution":{"observed_at":"2026-05-21T17:10:26.046450Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2509.07996","last_updated":"2026-07-20T17:34:32Z","snapshot_observed_at":"2026-08-05T06:03:57.824747Z","submitted_at":"2025-09-04T17:59:58Z","title":"3D and 4D World Modeling: A Survey","version":4},"cited_work":{"arxiv_id":"2509.07996","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2509.07996","snapshot_observed_at":"2026-07-21T03:22:29.733321Z","title":"3d and 4d world modeling: A survey","venue":null,"work_id":"0197d8fe-e11f-40f4-9254-87e63d981bf9","year":2025},"citing_paper":{"arxiv_id":"2512.23180","last_updated":"2026-05-16T07:53:20Z","snapshot_observed_at":"2026-08-03T04:34:56.955960Z","submitted_at":"2025-12-29T03:40:05Z","title":"GaussianDWM: 3D Gaussian Driving World Model for Unified Scene Understanding and Multi-Modal Generation","version":3},"reference_index":27,"source":"pdf_text","source_observed_at":"2026-05-21T17:06:34.398973Z"},"links":{"cited_paper":"/paper/2509.07996","citing_paper":"/paper/2512.23180"},"observation_digest":"sha256:ea26e81273e7dec8e105bc0f2c6c2fb3c51227ce1edd58a979e249710d6728aa","observation_id":"2d6c893e-77f5-41c8-9696-3fb7812e0860","resolution":{"observed_at":"2026-07-21T03:22:29.733321Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Uniscene: Unified occupancy-centric driving scene generation","venue":null,"work_id":"23a2aa5d-5fc8-4656-972d-5824236c2ad3","year":2025},"citing_paper":{"arxiv_id":"2512.23180","last_updated":"2026-05-16T07:53:20Z","snapshot_observed_at":"2026-08-03T04:34:56.955960Z","submitted_at":"2025-12-29T03:40:05Z","title":"GaussianDWM: 3D Gaussian Driving World Model for Unified Scene Understanding and Multi-Modal Generation","version":3},"reference_index":28,"source":"pdf_text","source_observed_at":"2026-05-21T17:06:34.398973Z"},"links":{"citing_paper":"/paper/2512.23180"},"observation_digest":"sha256:3d910ce2c5dd6791fb42244465ae70c583d4dc3d8b9e0461251833609f7f596f","observation_id":"77071d20-82d4-42c7-bf1a-3120ab6ae60c","resolution":{"observed_at":"2026-05-21T17:10:26.043249Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Human motion instruction tuning","venue":null,"work_id":"9a5f1b37-15f6-4889-bcb8-63f805dac616","year":2025},"citing_paper":{"arxiv_id":"2512.23180","last_updated":"2026-05-16T07:53:20Z","snapshot_observed_at":"2026-08-03T04:34:56.955960Z","submitted_at":"2025-12-29T03:40:05Z","title":"GaussianDWM: 3D Gaussian Driving World Model for Unified Scene Understanding and Multi-Modal Generation","version":3},"reference_index":29,"source":"pdf_text","source_observed_at":"2026-05-21T17:06:34.398973Z"},"links":{"citing_paper":"/paper/2512.23180"},"observation_digest":"sha256:2c474d0dc58e8b30e462910af82a38bfabf3201b8691dc385432a1cacffa6ca3","observation_id":"d5dd6ddc-ea0d-4dd4-9e53-5b3820216edb","resolution":{"observed_at":"2026-05-21T17:10:26.040268Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Bevformer: learning bird’s-eye-view representation from lidar-camera via spatiotemporal transformers.IEEE Transactions on Pat- tern Analysis and Machine Intelligence","venue":null,"work_id":"fb3b7a81-9fe3-49e5-adc0-3d702747f6c3","year":2024},"citing_paper":{"arxiv_id":"2512.23180","last_updated":"2026-05-16T07:53:20Z","snapshot_observed_at":"2026-08-03T04:34:56.955960Z","submitted_at":"2025-12-29T03:40:05Z","title":"GaussianDWM: 3D Gaussian Driving World Model for Unified Scene Understanding and Multi-Modal Generation","version":3},"reference_index":30,"source":"pdf_text","source_observed_at":"2026-05-21T17:06:34.398973Z"},"links":{"citing_paper":"/paper/2512.23180"},"observation_digest":"sha256:e6fc7b639dfdf6e7b4ae7bf4cfe3de178c18225c5f5fa316ea7ad112aa5c93f4","observation_id":"0fb91e00-d989-4072-98b4-440e7d74e00b","resolution":{"observed_at":"2026-05-21T17:10:26.037040Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-07-08T16:45:08.094292Z","title":"Rouge: A package for automatic evaluation of summaries","venue":null,"work_id":"cfa238ae-1503-4956-ab35-8dec3a13b5a0","year":2004},"citing_paper":{"arxiv_id":"2512.23180","last_updated":"2026-05-16T07:53:20Z","snapshot_observed_at":"2026-08-03T04:34:56.955960Z","submitted_at":"2025-12-29T03:40:05Z","title":"GaussianDWM: 3D Gaussian Driving World Model for Unified Scene Understanding and Multi-Modal Generation","version":3},"reference_index":31,"source":"pdf_text","source_observed_at":"2026-05-21T17:06:34.398973Z"},"links":{"citing_paper":"/paper/2512.23180"},"observation_digest":"sha256:7bc55f19ed2e9b190f9e4d44140ed79f78fdda8ccd2714f8f85e1e3bca753c73","observation_id":"65e38c08-4e44-44bd-96a7-3768949905e8","resolution":{"observed_at":"2026-05-21T17:10:26.034077Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Improved baselines with visual instruction tuning","venue":null,"work_id":"696d8adb-61d8-40c9-999c-89ff9b54e76a","year":2024},"citing_paper":{"arxiv_id":"2512.23180","last_updated":"2026-05-16T07:53:20Z","snapshot_observed_at":"2026-08-03T04:34:56.955960Z","submitted_at":"2025-12-29T03:40:05Z","title":"GaussianDWM: 3D Gaussian Driving World Model for Unified Scene Understanding and Multi-Modal Generation","version":3},"reference_index":32,"source":"pdf_text","source_observed_at":"2026-05-21T17:06:34.398973Z"},"links":{"citing_paper":"/paper/2512.23180"},"observation_digest":"sha256:f66a588c40e83bc8163e45a76a5612ad5d1d24e28e100ae8573e49d57b1746b2","observation_id":"a2b336cb-d676-4884-951b-110e4f089cda","resolution":{"observed_at":"2026-05-21T17:10:26.031555Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Petr: Position embedding transformation for multi-view 3d object detection","venue":null,"work_id":"6775fe74-f5fa-4c19-bb3d-0e9c85c74a91","year":2022},"citing_paper":{"arxiv_id":"2512.23180","last_updated":"2026-05-16T07:53:20Z","snapshot_observed_at":"2026-08-03T04:34:56.955960Z","submitted_at":"2025-12-29T03:40:05Z","title":"GaussianDWM: 3D Gaussian Driving World Model for Unified Scene Understanding and Multi-Modal Generation","version":3},"reference_index":33,"source":"pdf_text","source_observed_at":"2026-05-21T17:06:34.398973Z"},"links":{"citing_paper":"/paper/2512.23180"},"observation_digest":"sha256:c81b31aba00e51a22e0ea45dec46d54af2573dfac19da777a3e76c2b72d86886","observation_id":"77d3fde9-1b54-4a03-b293-00e2127e77ed","resolution":{"observed_at":"2026-05-21T17:10:26.028720Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Dreamdrive: Generative 4d scene modeling from street view images","venue":null,"work_id":"54e30a13-6034-44f5-aeeb-e966f70211dc","year":2025},"citing_paper":{"arxiv_id":"2512.23180","last_updated":"2026-05-16T07:53:20Z","snapshot_observed_at":"2026-08-03T04:34:56.955960Z","submitted_at":"2025-12-29T03:40:05Z","title":"GaussianDWM: 3D Gaussian Driving World Model for Unified Scene Understanding and Multi-Modal Generation","version":3},"reference_index":34,"source":"pdf_text","source_observed_at":"2026-05-21T17:06:34.398973Z"},"links":{"citing_paper":"/paper/2512.23180"},"observation_digest":"sha256:eb960f73ca72888ae0fa7162e1892814a6c353a66728f7b1e156cfe341c3bf43","observation_id":"a9a962a8-b05d-43cf-8abb-1bcb2e7dce43","resolution":{"observed_at":"2026-05-21T17:10:26.025930Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Nerf: Representing scenes as neural radiance fields for view syn- thesis.Communications of the ACM, 65(1):99–106","venue":null,"work_id":"3636ffa2-b06c-4e63-8205-b631ee117bfb","year":2021},"citing_paper":{"arxiv_id":"2512.23180","last_updated":"2026-05-16T07:53:20Z","snapshot_observed_at":"2026-08-03T04:34:56.955960Z","submitted_at":"2025-12-29T03:40:05Z","title":"GaussianDWM: 3D Gaussian Driving World Model for Unified Scene Understanding and Multi-Modal Generation","version":3},"reference_index":35,"source":"pdf_text","source_observed_at":"2026-05-21T17:06:34.398973Z"},"links":{"citing_paper":"/paper/2512.23180"},"observation_digest":"sha256:3e597d068c4207a41237a6a23eccb5a1a11c0026a4855b8527628264ba7c2dad","observation_id":"972ac40c-f240-4684-a435-3610640f5757","resolution":{"observed_at":"2026-05-21T17:10:26.022990Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Neural scene graphs for dynamic scenes","venue":null,"work_id":"d117481d-36d0-4e4e-80fa-36c931da3e38","year":2021},"citing_paper":{"arxiv_id":"2512.23180","last_updated":"2026-05-16T07:53:20Z","snapshot_observed_at":"2026-08-03T04:34:56.955960Z","submitted_at":"2025-12-29T03:40:05Z","title":"GaussianDWM: 3D Gaussian Driving World Model for Unified Scene Understanding and Multi-Modal Generation","version":3},"reference_index":36,"source":"pdf_text","source_observed_at":"2026-05-21T17:06:34.398973Z"},"links":{"citing_paper":"/paper/2512.23180"},"observation_digest":"sha256:6f366ef7852304b4f7283fd9398a1a52595a5976be33f2cb464572e7a1efa1e0","observation_id":"cd8d2633-665f-4da9-9a18-65f516259708","resolution":{"observed_at":"2026-05-21T17:10:26.019645Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-07-08T16:45:08.098053Z","title":"Bleu: a method for automatic evaluation of machine translation","venue":null,"work_id":"2670955d-43d2-4a62-9370-91b0ea9b3281","year":null},"citing_paper":{"arxiv_id":"2512.23180","last_updated":"2026-05-16T07:53:20Z","snapshot_observed_at":"2026-08-03T04:34:56.955960Z","submitted_at":"2025-12-29T03:40:05Z","title":"GaussianDWM: 3D Gaussian Driving World Model for Unified Scene Understanding and Multi-Modal Generation","version":3},"reference_index":37,"source":"pdf_text","source_observed_at":"2026-05-21T17:06:34.398973Z"},"links":{"citing_paper":"/paper/2512.23180"},"observation_digest":"sha256:e773aec71ef018c54e903628cd9e2b575b0bf5950800842c027e8ad61c8192f7","observation_id":"204d4794-79fe-49ee-9cf3-26ce48aaea63","resolution":{"observed_at":"2026-05-21T17:10:26.016813Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"A lesson in splats: Teacher-guided diffusion for 3d gaussian splats generation with 2d supervision","venue":null,"work_id":"50e2e6fa-418b-41b4-8325-8c81801663a8","year":2025},"citing_paper":{"arxiv_id":"2512.23180","last_updated":"2026-05-16T07:53:20Z","snapshot_observed_at":"2026-08-03T04:34:56.955960Z","submitted_at":"2025-12-29T03:40:05Z","title":"GaussianDWM: 3D Gaussian Driving World Model for Unified Scene Understanding and Multi-Modal Generation","version":3},"reference_index":38,"source":"pdf_text","source_observed_at":"2026-05-21T17:06:34.398973Z"},"links":{"citing_paper":"/paper/2512.23180"},"observation_digest":"sha256:fe3a028684ff0e62423e313beab1b4f6598f1b514af338cbae640f1fbc340087","observation_id":"a5ecf34f-81cc-4998-ac08-a0498dcbcc04","resolution":{"observed_at":"2026-05-21T17:10:26.013952Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Desire-gs: 4d street gaussians for static-dynamic decomposition and surface reconstruction for urban driving scenes","venue":null,"work_id":"e718d8c1-edcb-4223-a81d-f55916f506c0","year":2025},"citing_paper":{"arxiv_id":"2512.23180","last_updated":"2026-05-16T07:53:20Z","snapshot_observed_at":"2026-08-03T04:34:56.955960Z","submitted_at":"2025-12-29T03:40:05Z","title":"GaussianDWM: 3D Gaussian Driving World Model for Unified Scene Understanding and Multi-Modal Generation","version":3},"reference_index":39,"source":"pdf_text","source_observed_at":"2026-05-21T17:06:34.398973Z"},"links":{"citing_paper":"/paper/2512.23180"},"observation_digest":"sha256:8f5262a600812561160a41c4bf4599eaab0274b00951201eab4d0e0aab8cc358","observation_id":"ec19c342-9138-4638-ae7d-4f5dff1f7e68","resolution":{"observed_at":"2026-05-21T17:10:26.010944Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Langsplat: 3d language gaussian splatting","venue":null,"work_id":"aaa18575-6a23-42a8-ab9a-02c1f6b539a4","year":2024},"citing_paper":{"arxiv_id":"2512.23180","last_updated":"2026-05-16T07:53:20Z","snapshot_observed_at":"2026-08-03T04:34:56.955960Z","submitted_at":"2025-12-29T03:40:05Z","title":"GaussianDWM: 3D Gaussian Driving World Model for Unified Scene Understanding and Multi-Modal Generation","version":3},"reference_index":40,"source":"pdf_text","source_observed_at":"2026-05-21T17:06:34.398973Z"},"links":{"citing_paper":"/paper/2512.23180"},"observation_digest":"sha256:084ac472f53fcd771ad5569998db6db3fa29f54fd0c5b08299dd85e8d6628dbe","observation_id":"43f3beeb-026e-4386-b4f3-88ac326ee877","resolution":{"observed_at":"2026-05-21T17:10:26.007857Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Drivelm: Driving with graph visual question answering","venue":null,"work_id":"48987aea-ccfd-452d-a682-952121fac63b","year":2024},"citing_paper":{"arxiv_id":"2512.23180","last_updated":"2026-05-16T07:53:20Z","snapshot_observed_at":"2026-08-03T04:34:56.955960Z","submitted_at":"2025-12-29T03:40:05Z","title":"GaussianDWM: 3D Gaussian Driving World Model for Unified Scene Understanding and Multi-Modal Generation","version":3},"reference_index":41,"source":"pdf_text","source_observed_at":"2026-05-21T17:06:34.398973Z"},"links":{"citing_paper":"/paper/2512.23180"},"observation_digest":"sha256:42ba395b66cc447a60f3a6be38afdb8743da4fd52dc895f2a13351706871cf21","observation_id":"92c859f4-c4ca-4391-b504-242bf9d40226","resolution":{"observed_at":"2026-05-21T17:10:26.005001Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Driv- ingforward: Feed-forward 3d gaussian splatting for driving scene reconstruction from flexible surround-view input","venue":null,"work_id":"8aff2904-5e01-48e0-b225-7fd8d8740235","year":2025},"citing_paper":{"arxiv_id":"2512.23180","last_updated":"2026-05-16T07:53:20Z","snapshot_observed_at":"2026-08-03T04:34:56.955960Z","submitted_at":"2025-12-29T03:40:05Z","title":"GaussianDWM: 3D Gaussian Driving World Model for Unified Scene Understanding and Multi-Modal Generation","version":3},"reference_index":42,"source":"pdf_text","source_observed_at":"2026-05-21T17:06:34.398973Z"},"links":{"citing_paper":"/paper/2512.23180"},"observation_digest":"sha256:d42235f94566ca0b401929f58731b6020946821257e008c4447400a7e3b9a03d","observation_id":"26ebf199-43b9-4ab6-97de-96e7ec64d2ea","resolution":{"observed_at":"2026-05-21T17:10:26.001855Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Suds: Scalable urban dynamic scenes","venue":null,"work_id":"627ac5c7-7f98-4325-974b-906525d9486b","year":2023},"citing_paper":{"arxiv_id":"2512.23180","last_updated":"2026-05-16T07:53:20Z","snapshot_observed_at":"2026-08-03T04:34:56.955960Z","submitted_at":"2025-12-29T03:40:05Z","title":"GaussianDWM: 3D Gaussian Driving World Model for Unified Scene Understanding and Multi-Modal Generation","version":3},"reference_index":43,"source":"pdf_text","source_observed_at":"2026-05-21T17:06:34.398973Z"},"links":{"citing_paper":"/paper/2512.23180"},"observation_digest":"sha256:451fca0db38aff0f3c2ef9cd92faf75d0719cb79356cd311a46b994c285aedb6","observation_id":"89696230-600f-455b-a254-00e1b1ae629f","resolution":{"observed_at":"2026-05-21T17:10:25.998893Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Cider: Consensus-based image description evalua- tion","venue":null,"work_id":"13f6d756-4dcc-4004-9136-9a16aff1fb8f","year":2015},"citing_paper":{"arxiv_id":"2512.23180","last_updated":"2026-05-16T07:53:20Z","snapshot_observed_at":"2026-08-03T04:34:56.955960Z","submitted_at":"2025-12-29T03:40:05Z","title":"GaussianDWM: 3D Gaussian Driving World Model for Unified Scene Understanding and Multi-Modal Generation","version":3},"reference_index":44,"source":"pdf_text","source_observed_at":"2026-05-21T17:06:34.398973Z"},"links":{"citing_paper":"/paper/2512.23180"},"observation_digest":"sha256:eff21633f276bd60031d18368a08b27caaef845374ace4a893c8322435c8578b","observation_id":"30cf4b65-afde-4d84-8480-50f97039f731","resolution":{"observed_at":"2026-05-21T17:10:25.995959Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Freevs: Generative view synthesis on free driv- ing trajectory","venue":null,"work_id":"76acc07a-e1e5-469b-92cf-f74534cdec89","year":2025},"citing_paper":{"arxiv_id":"2512.23180","last_updated":"2026-05-16T07:53:20Z","snapshot_observed_at":"2026-08-03T04:34:56.955960Z","submitted_at":"2025-12-29T03:40:05Z","title":"GaussianDWM: 3D Gaussian Driving World Model for Unified Scene Understanding and Multi-Modal Generation","version":3},"reference_index":45,"source":"pdf_text","source_observed_at":"2026-05-21T17:06:34.398973Z"},"links":{"citing_paper":"/paper/2512.23180"},"observation_digest":"sha256:6693b944b930217180bfb6160dff1a092db3a9cf7e92f9cec73fd397e08bcd5a","observation_id":"d6d7c433-764e-4e6e-85e6-d8585c2917e0","resolution":{"observed_at":"2026-05-21T17:10:25.993538Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Omnidrive: A holistic vision-language dataset for au- tonomous driving with counterfactual reasoning","venue":null,"work_id":"33a800fa-478f-4ba1-a561-1e74b149bbb7","year":2025},"citing_paper":{"arxiv_id":"2512.23180","last_updated":"2026-05-16T07:53:20Z","snapshot_observed_at":"2026-08-03T04:34:56.955960Z","submitted_at":"2025-12-29T03:40:05Z","title":"GaussianDWM: 3D Gaussian Driving World Model for Unified Scene Understanding and Multi-Modal Generation","version":3},"reference_index":46,"source":"pdf_text","source_observed_at":"2026-05-21T17:06:34.398973Z"},"links":{"citing_paper":"/paper/2512.23180"},"observation_digest":"sha256:14be8f35ee0d7d56e43c85de863999b8a34862ed13c53b222bb39994cee3a1b8","observation_id":"4bad99bb-55ef-4021-8ee7-0dc5ca1df234","resolution":{"observed_at":"2026-05-21T17:10:25.990708Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Driving into the future: Multiview visual forecasting and planning with world model for au- tonomous driving","venue":null,"work_id":"26a66279-0174-402a-a840-618a6453d466","year":2024},"citing_paper":{"arxiv_id":"2512.23180","last_updated":"2026-05-16T07:53:20Z","snapshot_observed_at":"2026-08-03T04:34:56.955960Z","submitted_at":"2025-12-29T03:40:05Z","title":"GaussianDWM: 3D Gaussian Driving World Model for Unified Scene Understanding and Multi-Modal Generation","version":3},"reference_index":47,"source":"pdf_text","source_observed_at":"2026-05-21T17:06:34.398973Z"},"links":{"citing_paper":"/paper/2512.23180"},"observation_digest":"sha256:4e34a2465c7b17a1eb9b2a2cf593b521ddeb6a6731f4942452487fb528ce72eb","observation_id":"570ef81b-118f-4e0b-9621-5568a715d339","resolution":{"observed_at":"2026-05-21T17:10:25.987613Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2507.11001","last_updated":"2025-07-15T05:37:24Z","snapshot_observed_at":"2026-07-06T21:57:22.011378Z","submitted_at":"2025-07-15T05:37:24Z","title":"Learning to Tune Like an Expert: Interpretable and Scene-Aware Navigation via MLLM Reasoning and CVAE-Based Adaptation","version":1},"cited_work":{"arxiv_id":"2507.11001","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2507.11001","snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Learning to tune like an expert: Interpretable and scene- aware navigation via mllm reasoning and cvae-based adapta- tion","venue":null,"work_id":"20169e57-32b7-4cc3-af91-bd0d415c2767","year":2025},"citing_paper":{"arxiv_id":"2512.23180","last_updated":"2026-05-16T07:53:20Z","snapshot_observed_at":"2026-08-03T04:34:56.955960Z","submitted_at":"2025-12-29T03:40:05Z","title":"GaussianDWM: 3D Gaussian Driving World Model for Unified Scene Understanding and Multi-Modal Generation","version":3},"reference_index":48,"source":"pdf_text","source_observed_at":"2026-05-21T17:06:34.398973Z"},"links":{"cited_paper":"/paper/2507.11001","citing_paper":"/paper/2512.23180"},"observation_digest":"sha256:59d7d34aa9e0c3549d177760fffdc00f8e6414cba4a8e435a83d7b7ac024f68f","observation_id":"e2751359-d733-4f8b-8c42-6fffe58c58fd","resolution":{"observed_at":"2026-05-21T17:10:25.117502Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2505.09406","last_updated":"2025-05-14T14:02:49Z","snapshot_observed_at":"2026-07-06T21:23:47.292610Z","submitted_at":"2025-05-14T14:02:49Z","title":"FreeDriveRF: Monocular RGB Dynamic NeRF without Poses for Autonomous Driving via Point-Level Dynamic-Static Decoupling","version":1},"cited_work":{"arxiv_id":"2505.09406","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2505.09406","snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Freedriverf: Monocu- lar rgb dynamic nerf without poses for autonomous driving via point-level dynamic-static decoupling","venue":null,"work_id":"c91d538b-a42d-4c6a-a7ff-fabf2df3f3df","year":2025},"citing_paper":{"arxiv_id":"2512.23180","last_updated":"2026-05-16T07:53:20Z","snapshot_observed_at":"2026-08-03T04:34:56.955960Z","submitted_at":"2025-12-29T03:40:05Z","title":"GaussianDWM: 3D Gaussian Driving World Model for Unified Scene Understanding and Multi-Modal Generation","version":3},"reference_index":49,"source":"pdf_text","source_observed_at":"2026-05-21T17:06:34.398973Z"},"links":{"cited_paper":"/paper/2505.09406","citing_paper":"/paper/2512.23180"},"observation_digest":"sha256:f5c7aa73a339fcb27e1499caf13b7796bcc1358f083b26f1b1fc043a6b60033a","observation_id":"0c78d41a-6594-41db-8417-259c653bfa7d","resolution":{"observed_at":"2026-05-21T17:10:25.113160Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Dynamicrafter: Animating open-domain images with video diffusion priors","venue":null,"work_id":"7aef753c-12a7-4cdd-bb10-567f4d4bd6a0","year":2024},"citing_paper":{"arxiv_id":"2512.23180","last_updated":"2026-05-16T07:53:20Z","snapshot_observed_at":"2026-08-03T04:34:56.955960Z","submitted_at":"2025-12-29T03:40:05Z","title":"GaussianDWM: 3D Gaussian Driving World Model for Unified Scene Understanding and Multi-Modal Generation","version":3},"reference_index":50,"source":"pdf_text","source_observed_at":"2026-05-21T17:06:34.398973Z"},"links":{"citing_paper":"/paper/2512.23180"},"observation_digest":"sha256:c34e6a745cda57e3fd636bff352e864c70bb970eec4b440396bbfb35894acf40","observation_id":"4f9e0473-1be6-4a87-b4c6-c8487cfaf7d7","resolution":{"observed_at":"2026-05-21T17:10:25.984189Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Cape: Camera view position embedding for multi-view 3d object detection","venue":null,"work_id":"a9c89f8a-5200-425e-a44d-17a8202f6519","year":2023},"citing_paper":{"arxiv_id":"2512.23180","last_updated":"2026-05-16T07:53:20Z","snapshot_observed_at":"2026-08-03T04:34:56.955960Z","submitted_at":"2025-12-29T03:40:05Z","title":"GaussianDWM: 3D Gaussian Driving World Model for Unified Scene Understanding and Multi-Modal Generation","version":3},"reference_index":51,"source":"pdf_text","source_observed_at":"2026-05-21T17:06:34.398973Z"},"links":{"citing_paper":"/paper/2512.23180"},"observation_digest":"sha256:59af236eb23b3be7742ce364ceac2d89597d994737ea442869f077950ce0a028","observation_id":"d5a040d2-0be0-445a-ad58-c6785259571a","resolution":{"observed_at":"2026-05-21T17:10:25.981152Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Drivegpt4: Interpretable end-to-end autonomous driving via large language model.IEEE Robotics and Automation Let- ters","venue":null,"work_id":"2d08e603-d87a-4e4c-8e0f-75f644fb62ce","year":2024},"citing_paper":{"arxiv_id":"2512.23180","last_updated":"2026-05-16T07:53:20Z","snapshot_observed_at":"2026-08-03T04:34:56.955960Z","submitted_at":"2025-12-29T03:40:05Z","title":"GaussianDWM: 3D Gaussian Driving World Model for Unified Scene Understanding and Multi-Modal Generation","version":3},"reference_index":52,"source":"pdf_text","source_observed_at":"2026-05-21T17:06:34.398973Z"},"links":{"citing_paper":"/paper/2512.23180"},"observation_digest":"sha256:516d50bd6c0e717e67af847f283730dd805f5c048d26230550912e7954d39505","observation_id":"ac527803-74d2-47a8-ba17-3021452ec1e4","resolution":{"observed_at":"2026-05-21T17:10:25.978254Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Drivingsphere: Building a high-fidelity 4d world for closed- loop simulation","venue":null,"work_id":"9addbdd2-1e35-42b2-ba81-c96cce20192c","year":2025},"citing_paper":{"arxiv_id":"2512.23180","last_updated":"2026-05-16T07:53:20Z","snapshot_observed_at":"2026-08-03T04:34:56.955960Z","submitted_at":"2025-12-29T03:40:05Z","title":"GaussianDWM: 3D Gaussian Driving World Model for Unified Scene Understanding and Multi-Modal Generation","version":3},"reference_index":53,"source":"pdf_text","source_observed_at":"2026-05-21T17:06:34.398973Z"},"links":{"citing_paper":"/paper/2512.23180"},"observation_digest":"sha256:b45f95f6c80ae6f87552798cae6eaa71d494930402eae2fb4f220bcd720c0213","observation_id":"f4a4d8d2-6abd-45f2-a018-5bf6acd2293a","resolution":{"observed_at":"2026-05-21T17:10:25.974931Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Street gaussians: Modeling dynamic urban scenes with gaussian splatting","venue":null,"work_id":"8adaa207-3cf9-4871-873f-edb1833a51a5","year":2024},"citing_paper":{"arxiv_id":"2512.23180","last_updated":"2026-05-16T07:53:20Z","snapshot_observed_at":"2026-08-03T04:34:56.955960Z","submitted_at":"2025-12-29T03:40:05Z","title":"GaussianDWM: 3D Gaussian Driving World Model for Unified Scene Understanding and Multi-Modal Generation","version":3},"reference_index":54,"source":"pdf_text","source_observed_at":"2026-05-21T17:06:34.398973Z"},"links":{"citing_paper":"/paper/2512.23180"},"observation_digest":"sha256:5a8cbec83403ba3c73504adc8c0376e3ba7604f2b6b0508085cb9bf58c5748be","observation_id":"b7edaaa9-10b7-40d2-8efd-cab150dfa4b5","resolution":{"observed_at":"2026-05-21T17:10:25.971991Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2505.09388","last_updated":"2025-05-14T13:41:34Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2025-05-14T13:41:34Z","title":"Qwen3 Technical Report","version":1},"cited_work":{"arxiv_id":"2505.09388","doi":"10.1016/j.aiopen.2022.12","metadata_source":"pith","pith_arxiv_id":"2505.09388","snapshot_observed_at":"2026-07-11T11:50:26.030339Z","title":"Qwen3 Technical Report","venue":"cs.CL","work_id":"25a4e30c-1232-48e7-9925-02fa12ba7c9e","year":2025},"citing_paper":{"arxiv_id":"2512.23180","last_updated":"2026-05-16T07:53:20Z","snapshot_observed_at":"2026-08-03T04:34:56.955960Z","submitted_at":"2025-12-29T03:40:05Z","title":"GaussianDWM: 3D Gaussian Driving World Model for Unified Scene Understanding and Multi-Modal Generation","version":3},"reference_index":55,"source":"pdf_text","source_observed_at":"2026-05-21T17:06:34.398973Z"},"links":{"cited_paper":"/paper/2505.09388","citing_paper":"/paper/2512.23180"},"observation_digest":"sha256:e460332f19c51bb7c92206eb99ef23023c6dcc00fa7efba1502c25bc158cee57","observation_id":"d43fb6d7-b8a2-4d71-a1e9-0a91d583166c","resolution":{"observed_at":"2026-05-21T17:10:25.144067Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2311.02077","last_updated":"2023-11-03T17:59:55Z","snapshot_observed_at":"2026-08-03T23:31:39.294900Z","submitted_at":"2023-11-03T17:59:55Z","title":"EmerNeRF: Emergent Spatial-Temporal Scene Decomposition via Self-Supervision","version":1},"cited_work":{"arxiv_id":"2311.02077","doi":null,"metadata_source":"pith","pith_arxiv_id":"2311.02077","snapshot_observed_at":"2026-07-08T03:24:28.762183Z","title":"Emernerf: Emergent spatial-temporal scene decomposition via self-supervision","venue":"cs.CV","work_id":"a7c91f0a-e32a-47fd-ab9b-aec98b92c370","year":2023},"citing_paper":{"arxiv_id":"2512.23180","last_updated":"2026-05-16T07:53:20Z","snapshot_observed_at":"2026-08-03T04:34:56.955960Z","submitted_at":"2025-12-29T03:40:05Z","title":"GaussianDWM: 3D Gaussian Driving World Model for Unified Scene Understanding and Multi-Modal Generation","version":3},"reference_index":56,"source":"pdf_text","source_observed_at":"2026-05-21T17:06:34.398973Z"},"links":{"cited_paper":"/paper/2311.02077","citing_paper":"/paper/2512.23180"},"observation_digest":"sha256:1a6fd8810bef7b60a6d1d683d9a5c5379c67cf17c272b3f1c7ec2e7caee531de","observation_id":"c6e30895-495d-4f34-a1c1-f4e46c8085d4","resolution":{"observed_at":"2026-05-21T17:10:25.104651Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2501.00602","last_updated":"2024-12-31T18:59:58Z","snapshot_observed_at":"2026-07-06T20:15:17.786867Z","submitted_at":"2024-12-31T18:59:58Z","title":"STORM: Spatio-Temporal Reconstruction Model for Large-Scale Outdoor Scenes","version":1},"cited_work":{"arxiv_id":"2501.00602","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2501.00602","snapshot_observed_at":"2026-07-03T04:37:36.984858Z","title":"Storm: Spatio-temporal re- construction model for large-scale outdoor scenes","venue":null,"work_id":"d97376a5-166b-42a7-b3b0-35bb4de04281","year":2024},"citing_paper":{"arxiv_id":"2512.23180","last_updated":"2026-05-16T07:53:20Z","snapshot_observed_at":"2026-08-03T04:34:56.955960Z","submitted_at":"2025-12-29T03:40:05Z","title":"GaussianDWM: 3D Gaussian Driving World Model for Unified Scene Understanding and Multi-Modal Generation","version":3},"reference_index":57,"source":"pdf_text","source_observed_at":"2026-05-21T17:06:34.398973Z"},"links":{"cited_paper":"/paper/2501.00602","citing_paper":"/paper/2512.23180"},"observation_digest":"sha256:07e34795ccf0bc98c8a8197c644a0430e9b3370fc5bc4f8ad1589f2aae0aa142","observation_id":"4307bcd8-2281-4233-92a0-41ccddbab1e1","resolution":{"observed_at":"2026-05-21T17:10:25.126719Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2309.13101","last_updated":"2023-11-19T14:50:08Z","snapshot_observed_at":"2026-08-04T05:15:02.372597Z","submitted_at":"2023-09-22T16:04:02Z","title":"Deformable 3D Gaussians for High-Fidelity Monocular Dynamic Scene Reconstruction","version":2},"cited_work":{"arxiv_id":"2309.13101","doi":"10.48550/arxiv.2309.13101","metadata_source":"arxiv_reference","pith_arxiv_id":"2309.13101","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Deformable 3d gaussians for high-fidelity monocular dynamic scene reconstruction","venue":"arXiv (Cornell University)","work_id":"bb973e91-b2d8-4ab6-9b92-7c3b13b2f9a4","year":2023},"citing_paper":{"arxiv_id":"2512.23180","last_updated":"2026-05-16T07:53:20Z","snapshot_observed_at":"2026-08-03T04:34:56.955960Z","submitted_at":"2025-12-29T03:40:05Z","title":"GaussianDWM: 3D Gaussian Driving World Model for Unified Scene Understanding and Multi-Modal Generation","version":3},"reference_index":58,"source":"pdf_text","source_observed_at":"2026-05-21T17:06:34.398973Z"},"links":{"cited_paper":"/paper/2309.13101","citing_paper":"/paper/2512.23180"},"observation_digest":"sha256:c9b4d97a03e8e1058c3bedcd4e714ca9f0898ece3b5a6ef5728d594514831b45","observation_id":"4892945c-1197-4045-b17e-0b6982544b6f","resolution":{"observed_at":"2026-05-21T17:10:25.183203Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Visual point cloud forecasting enables scalable autonomous driving","venue":null,"work_id":"4fa7dde4-17cc-4582-a447-60e9c7a3c38d","year":2024},"citing_paper":{"arxiv_id":"2512.23180","last_updated":"2026-05-16T07:53:20Z","snapshot_observed_at":"2026-08-03T04:34:56.955960Z","submitted_at":"2025-12-29T03:40:05Z","title":"GaussianDWM: 3D Gaussian Driving World Model for Unified Scene Understanding and Multi-Modal Generation","version":3},"reference_index":59,"source":"pdf_text","source_observed_at":"2026-05-21T17:06:34.398973Z"},"links":{"citing_paper":"/paper/2512.23180"},"observation_digest":"sha256:2f9815d8c67390f540ea44d369188df226581c103633187d27637c81c9eb4c6a","observation_id":"abc6635c-121c-42e2-b1f7-518a0714f065","resolution":{"observed_at":"2026-05-21T17:10:25.969213Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Drivedreamer-2: Llm-enhanced world models for diverse driving video generation","venue":null,"work_id":"473a7752-b988-4ed1-a698-48a40f8244b4","year":2025},"citing_paper":{"arxiv_id":"2512.23180","last_updated":"2026-05-16T07:53:20Z","snapshot_observed_at":"2026-08-03T04:34:56.955960Z","submitted_at":"2025-12-29T03:40:05Z","title":"GaussianDWM: 3D Gaussian Driving World Model for Unified Scene Understanding and Multi-Modal Generation","version":3},"reference_index":60,"source":"pdf_text","source_observed_at":"2026-05-21T17:06:34.398973Z"},"links":{"citing_paper":"/paper/2512.23180"},"observation_digest":"sha256:319124975a915be6e322fdbd27ad8184efc517966262e722fa3083814d9da069","observation_id":"1448b9dd-b4fa-4bac-a880-84ad3ac2ee60","resolution":{"observed_at":"2026-05-21T17:10:25.966280Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2505.08725","last_updated":"2025-05-13T16:36:51Z","snapshot_observed_at":"2026-07-06T21:23:19.117703Z","submitted_at":"2025-05-13T16:36:51Z","title":"Extending Large Vision-Language Model for Diverse Interactive Tasks in Autonomous Driving","version":1},"cited_work":{"arxiv_id":"2505.08725","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2505.08725","snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Extending large vision-language model for diverse interactive tasks in autonomous driving","venue":null,"work_id":"85272d91-8588-4e97-9460-c30a2adfa147","year":2025},"citing_paper":{"arxiv_id":"2512.23180","last_updated":"2026-05-16T07:53:20Z","snapshot_observed_at":"2026-08-03T04:34:56.955960Z","submitted_at":"2025-12-29T03:40:05Z","title":"GaussianDWM: 3D Gaussian Driving World Model for Unified Scene Understanding and Multi-Modal Generation","version":3},"reference_index":61,"source":"pdf_text","source_observed_at":"2026-05-21T17:06:34.398973Z"},"links":{"cited_paper":"/paper/2505.08725","citing_paper":"/paper/2512.23180"},"observation_digest":"sha256:cdcc39662e8e938b106cae2c2773fe6fa0ddb2cf40ef610a6ee0b4b24f98d334","observation_id":"a561ba29-651a-49e2-bdc8-6f643e12bc5e","resolution":{"observed_at":"2026-05-21T17:10:25.162426Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2403.09631","last_updated":"2024-03-14T17:58:41Z","snapshot_observed_at":"2026-07-06T17:44:45.924645Z","submitted_at":"2024-03-14T17:58:41Z","title":"3D-VLA: A 3D Vision-Language-Action Generative World Model","version":1},"cited_work":{"arxiv_id":"2403.09631","doi":"10.48550/arxiv.2403.09631","metadata_source":"pith","pith_arxiv_id":"2403.09631","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"3D-VLA: A 3D Vision-Language-Action Generative World Model","venue":"cs.CV","work_id":"aebf924c-e761-437e-9cee-f1ccc2e427bd","year":2024},"citing_paper":{"arxiv_id":"2512.23180","last_updated":"2026-05-16T07:53:20Z","snapshot_observed_at":"2026-08-03T04:34:56.955960Z","submitted_at":"2025-12-29T03:40:05Z","title":"GaussianDWM: 3D Gaussian Driving World Model for Unified Scene Understanding and Multi-Modal Generation","version":3},"reference_index":62,"source":"pdf_text","source_observed_at":"2026-05-21T17:06:34.398973Z"},"links":{"cited_paper":"/paper/2403.09631","citing_paper":"/paper/2512.23180"},"observation_digest":"sha256:54cdcbde07e1ef601e72657ba04d215de63d921a4a53484a53c2b31c43ac906d","observation_id":"c0db86ef-5a47-47d0-8270-c69c21a01831","resolution":{"observed_at":"2026-05-21T17:10:25.157527Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Drivinggaussian: Composite gaussian splatting for surrounding dynamic au- tonomous driving scenes","venue":null,"work_id":"44e98d16-a088-4dfe-9c30-1bc575db5b40","year":2024},"citing_paper":{"arxiv_id":"2512.23180","last_updated":"2026-05-16T07:53:20Z","snapshot_observed_at":"2026-08-03T04:34:56.955960Z","submitted_at":"2025-12-29T03:40:05Z","title":"GaussianDWM: 3D Gaussian Driving World Model for Unified Scene Understanding and Multi-Modal Generation","version":3},"reference_index":63,"source":"pdf_text","source_observed_at":"2026-05-21T17:06:34.398973Z"},"links":{"citing_paper":"/paper/2512.23180"},"observation_digest":"sha256:e15c9e34f2fd6ce67844c81340065f197266b9ad45e1b72348721df08f838b9e","observation_id":"df848589-0fbf-4912-8a6c-53760c492a0d","resolution":{"observed_at":"2026-05-21T17:10:25.963205Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2501.14729","last_updated":"2025-08-13T09:10:30Z","snapshot_observed_at":"2026-08-02T16:02:02.988357Z","submitted_at":"2025-01-24T18:59:51Z","title":"HERMES: A Unified Self-Driving World Model for Simultaneous 3D Scene Understanding and Generation","version":3},"cited_work":{"arxiv_id":"2501.14729","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2501.14729","snapshot_observed_at":"2026-07-01T06:55:28.848054Z","title":"Hermes: A unified self-driving world model for simultaneous 3d scene understanding and generation","venue":null,"work_id":"5aa3c3ee-e31f-4748-b510-3cced0c9c63b","year":2025},"citing_paper":{"arxiv_id":"2512.23180","last_updated":"2026-05-16T07:53:20Z","snapshot_observed_at":"2026-08-03T04:34:56.955960Z","submitted_at":"2025-12-29T03:40:05Z","title":"GaussianDWM: 3D Gaussian Driving World Model for Unified Scene Understanding and Multi-Modal Generation","version":3},"reference_index":64,"source":"pdf_text","source_observed_at":"2026-05-21T17:06:34.398973Z"},"links":{"cited_paper":"/paper/2501.14729","citing_paper":"/paper/2512.23180"},"observation_digest":"sha256:b1131de4c1087e9120cbd5530ac2e9a557fa2f93d69a5d0e081d71757a2a5238","observation_id":"76b4ae62-c7e5-42df-959a-a0c1abf1c81a","resolution":{"observed_at":"2026-05-21T17:10:25.174828Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2503.10604","last_updated":"2025-03-13T17:48:41Z","snapshot_observed_at":"2026-08-01T22:59:29.034612Z","submitted_at":"2025-03-13T17:48:41Z","title":"MuDG: Taming Multi-modal Diffusion with Gaussian Splatting for Urban Scene Reconstruction","version":1},"cited_work":{"arxiv_id":"2503.10604","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2503.10604","snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Mudg: Taming multi-modal diffusion with gaussian splatting for urban scene reconstruction","venue":null,"work_id":"9ff45963-09d0-48b7-b382-3f5d5ad780cb","year":2025},"citing_paper":{"arxiv_id":"2512.23180","last_updated":"2026-05-16T07:53:20Z","snapshot_observed_at":"2026-08-03T04:34:56.955960Z","submitted_at":"2025-12-29T03:40:05Z","title":"GaussianDWM: 3D Gaussian Driving World Model for Unified Scene Understanding and Multi-Modal Generation","version":3},"reference_index":65,"source":"pdf_text","source_observed_at":"2026-05-21T17:06:34.398973Z"},"links":{"cited_paper":"/paper/2503.10604","citing_paper":"/paper/2512.23180"},"observation_digest":"sha256:ed1945125a3aa58fca9558f6a2b6af6ca1055edadb9325084ed45f0a112bebd2","observation_id":"916059ab-6930-4e5e-bc2a-be41c95284dd","resolution":{"observed_at":"2026-05-21T17:10:25.170652Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}}],"paper":{"arxiv_id":"2512.23180","last_updated":"2026-05-16T07:53:20Z","latest_version":3,"primary_category":"cs.CV","snapshot_observed_at":"2026-08-03T04:34:56.955960Z","submitted_at":"2025-12-29T03:40:05Z","title":"GaussianDWM: 3D Gaussian Driving World Model for Unified Scene Understanding and Multi-Modal Generation"},"reference_resolution":{"displayed":64,"state_counts":{"malformed_identifier":0,"metadata_mismatch":0,"parse_uncertain":0,"unresolved":1,"verified_exact":23,"verified_fuzzy":40},"total_outbound_references":64},"refusal":"A citation records a reference. It does not transfer a finding from one paper to another.","schema":"pith.paper-citation-record.v1","standing_sources":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"thesis":"As of 5 August 2026, this Paper Citation Record lists 64 of 64 outbound references and 6 inbound Pith citation observations for arXiv:2512.23180."}