{"as_of":"2026-08-09T13:09:00Z","caps":{"database_statements":6,"inbound":100,"outbound":100},"context_digest":"sha256:d61eaa31cf0461d4175698ba9b4c45af5a24ca63f85e2de7b76c53cb42cd8c89","coverage":[{"denominator":55,"lane":"reference_resolution","note":"Typed states for the displayed outbound observations.","records_observed":55,"source":"paper_references, paper_reference_links","source_observed_at":"2026-08-07T04:51:34.522907Z","state":"measured"},{"denominator":64,"lane":"standing_notices","note":"One-hop event checks from named stored sources.","records_observed":64,"source":"scholarly_work_events, retraction_status_cache","source_observed_at":"2026-08-09T06:31:02.800959+00:00","state":"measured"},{"denominator":9,"lane":"inbound_itemization","note":"Pith citing papers itemized under the disclosed page cap.","records_observed":9,"source":"paper_references, paper_reference_links","source_observed_at":"2026-08-04T09:37:24.204341Z","state":"measured"},{"denominator":1,"lane":"external_citation_measurements","note":"A source-named dated measurement, never combined with another source.","records_observed":0,"source":"arxiv_reference","source_observed_at":"2026-05-25T04:26:37.591584Z","state":"measured"}],"external_citation_measurements":[],"inbound":[{"citation":{"cited_paper":{"arxiv_id":"2506.09565","last_updated":"2025-06-13T08:30:38Z","snapshot_observed_at":"2026-08-09T06:12:36.730836Z","submitted_at":"2025-06-11T09:56:39Z","title":"SemanticSplat: Feed-Forward 3D Scene Understanding with Language-Aware Gaussian Fields","version":2},"cited_work":{"arxiv_id":"2506.09565","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2506.09565","snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Semanticsplat: Feed-forward 3d scene understanding with language-aware gaussian fields","venue":null,"work_id":"02679082-1066-4d13-9a3e-d7bd2befbcc0","year":2025},"citing_paper":{"arxiv_id":"2508.09977","last_updated":"2026-04-11T08:26:14Z","snapshot_observed_at":"2026-07-06T22:12:31.023627Z","submitted_at":"2025-08-13T17:44:39Z","title":"A Survey on 3D Gaussian Splatting Applications: Segmentation, Editing, and Generation","version":4},"reference_index":120,"source":"pdf_text","source_observed_at":"2026-05-18T22:39:24.200019Z"},"links":{"cited_paper":"/paper/2506.09565","citing_paper":"/paper/2508.09977"},"observation_digest":"sha256:a66ee5a88ef2147dc9b412767ce4a40865673fee390ffa077d6a224adbc90f0d","observation_id":"19e681b1-f7b5-46dc-911b-e549ea5d0e77","resolution":{"observed_at":"2026-05-18T22:41:53.301673Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2506.09565","last_updated":"2025-06-13T08:30:38Z","snapshot_observed_at":"2026-08-09T06:12:36.730836Z","submitted_at":"2025-06-11T09:56:39Z","title":"SemanticSplat: Feed-Forward 3D Scene Understanding with Language-Aware Gaussian Fields","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2506.09565","snapshot_observed_at":"2026-08-04T09:37:24.204341Z","title":"SemanticSplat: Feed-Forward 3D Scene Understanding with Language-Aware Gaus- sian Fields,","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2510.14546","last_updated":"2026-08-03T08:20:35Z","snapshot_observed_at":"2026-08-08T21:04:22.749605Z","submitted_at":"2025-10-16T10:41:31Z","title":"QuASH: Using Natural-Language Heuristics to Query Visual-Language Robotic Maps","version":2},"reference_index":14,"source":"pdf_text","source_observed_at":"2026-08-04T09:37:24.204341Z"},"links":{"cited_paper":"/paper/2506.09565","citing_paper":"/paper/2510.14546"},"observation_digest":"sha256:a080a595ee1f7e27871dba76f19608b29bc19f212b1e3c4ebb3ee3d99293d7e6","observation_id":"341d7b6d-3202-428f-a1ba-665d191b90b0","resolution":{"observed_at":"2026-08-04T09:37:24.204341Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2506.09565","last_updated":"2025-06-13T08:30:38Z","snapshot_observed_at":"2026-08-09T06:12:36.730836Z","submitted_at":"2025-06-11T09:56:39Z","title":"SemanticSplat: Feed-Forward 3D Scene Understanding with Language-Aware Gaussian Fields","version":2},"cited_work":{"arxiv_id":"2506.09565","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2506.09565","snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Semanticsplat: Feed-forward 3d scene understanding with language-aware gaussian fields","venue":null,"work_id":"02679082-1066-4d13-9a3e-d7bd2befbcc0","year":2025},"citing_paper":{"arxiv_id":"2512.17541","last_updated":"2026-04-05T17:51:43Z","snapshot_observed_at":"2026-07-06T22:39:33.919723Z","submitted_at":"2025-12-19T13:04:13Z","title":"FLEG: Feed-Forward Language Embedded Gaussian Splatting from Any Views via Compact Semantic Representation","version":2},"reference_index":14,"source":"pdf_text","source_observed_at":"2026-05-16T20:46:04.875912Z"},"links":{"cited_paper":"/paper/2506.09565","citing_paper":"/paper/2512.17541"},"observation_digest":"sha256:2e3b12b203d475f5bba807d0aeb293fff3225605e806fe17e3e631f789c86f36","observation_id":"eb902cba-907f-4175-a9d3-173b5b4e5133","resolution":{"observed_at":"2026-05-16T20:48:32.612170Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2506.09565","last_updated":"2025-06-13T08:30:38Z","snapshot_observed_at":"2026-08-09T06:12:36.730836Z","submitted_at":"2025-06-11T09:56:39Z","title":"SemanticSplat: Feed-Forward 3D Scene Understanding with Language-Aware Gaussian Fields","version":2},"cited_work":{"arxiv_id":"2506.09565","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2506.09565","snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Semanticsplat: Feed-forward 3d scene understanding with language-aware gaussian fields","venue":null,"work_id":"02679082-1066-4d13-9a3e-d7bd2befbcc0","year":2025},"citing_paper":{"arxiv_id":"2603.08096","last_updated":"2026-04-20T17:13:59Z","snapshot_observed_at":"2026-07-06T22:48:21.803560Z","submitted_at":"2026-03-09T08:37:05Z","title":"TrianguLang: Geometry-Aware Semantic Consensus for Pose-Free 3D Localization","version":3},"reference_index":21,"source":"pdf_text","source_observed_at":"2026-05-15T15:12:23.459579Z"},"links":{"cited_paper":"/paper/2506.09565","citing_paper":"/paper/2603.08096"},"observation_digest":"sha256:bebd426d775d8432490221a72fc9450ae8f6c2477cbf7753cd87bd4680472c0f","observation_id":"ee158deb-ee13-4206-b957-c4943bab18f1","resolution":{"observed_at":"2026-05-15T15:16:09.975295Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2506.09565","last_updated":"2025-06-13T08:30:38Z","snapshot_observed_at":"2026-08-09T06:12:36.730836Z","submitted_at":"2025-06-11T09:56:39Z","title":"SemanticSplat: Feed-Forward 3D Scene Understanding with Language-Aware Gaussian Fields","version":2},"cited_work":{"arxiv_id":"2506.09565","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2506.09565","snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Semanticsplat: Feed-forward 3d scene understanding with language-aware gaussian fields","venue":null,"work_id":"02679082-1066-4d13-9a3e-d7bd2befbcc0","year":2025},"citing_paper":{"arxiv_id":"2604.09862","last_updated":"2026-04-10T19:45:24Z","snapshot_observed_at":"2026-08-03T00:48:42.155502Z","submitted_at":"2026-04-10T19:45:24Z","title":"FF3R: Feedforward Feature 3D Reconstruction from Unconstrained views","version":1},"reference_index":17,"source":"pdf_text","source_observed_at":"2026-05-10T17:47:52.901901Z"},"links":{"cited_paper":"/paper/2506.09565","citing_paper":"/paper/2604.09862"},"observation_digest":"sha256:2314206c50f43db678dcbbe62b37d86f6e86102b860b36cd43300ace05df4e99","observation_id":"0533a255-e42e-4f46-a915-edbcb158385f","resolution":{"observed_at":"2026-05-11T06:05:59.628613Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2506.09565","last_updated":"2025-06-13T08:30:38Z","snapshot_observed_at":"2026-08-09T06:12:36.730836Z","submitted_at":"2025-06-11T09:56:39Z","title":"SemanticSplat: Feed-Forward 3D Scene Understanding with Language-Aware Gaussian Fields","version":2},"cited_work":{"arxiv_id":"2506.09565","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2506.09565","snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Semanticsplat: Feed-forward 3d scene understanding with language-aware gaussian fields","venue":null,"work_id":"02679082-1066-4d13-9a3e-d7bd2befbcc0","year":2025},"citing_paper":{"arxiv_id":"2604.10573","last_updated":"2026-04-12T10:36:18Z","snapshot_observed_at":"2026-08-02T12:13:25.467849Z","submitted_at":"2026-04-12T10:36:18Z","title":"Learning 3D Representations for Spatial Intelligence from Unposed Multi-View Images","version":1},"reference_index":35,"source":"pdf_text","source_observed_at":"2026-05-10T15:44:39.322417Z"},"links":{"cited_paper":"/paper/2506.09565","citing_paper":"/paper/2604.10573"},"observation_digest":"sha256:bdba8d6ab47f75af2d8285beddadea04d94f772cb27cdc1b83f5a814f4c238af","observation_id":"d40c7e1c-b85c-46eb-9bc5-586eda8844d1","resolution":{"observed_at":"2026-05-11T09:56:03.517065Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2506.09565","last_updated":"2025-06-13T08:30:38Z","snapshot_observed_at":"2026-08-09T06:12:36.730836Z","submitted_at":"2025-06-11T09:56:39Z","title":"SemanticSplat: Feed-Forward 3D Scene Understanding with Language-Aware Gaussian Fields","version":2},"cited_work":{"arxiv_id":"2506.09565","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2506.09565","snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Semanticsplat: Feed-forward 3d scene understanding with language-aware gaussian fields","venue":null,"work_id":"02679082-1066-4d13-9a3e-d7bd2befbcc0","year":2025},"citing_paper":{"arxiv_id":"2604.11401","last_updated":"2026-04-13T12:42:28Z","snapshot_observed_at":"2026-08-03T03:37:10.899621Z","submitted_at":"2026-04-13T12:42:28Z","title":"GS4City: Hierarchical Semantic Gaussian Splatting via City-Model Priors","version":1},"reference_index":18,"source":"pdf_text","source_observed_at":"2026-05-10T15:38:21.644980Z"},"links":{"cited_paper":"/paper/2506.09565","citing_paper":"/paper/2604.11401"},"observation_digest":"sha256:b5550f8dcda6b9e9377baff3a8804460c150f5bb603599235f7b6204ba8b075c","observation_id":"b11c7360-13c5-4c2d-b94e-75943fb85a5d","resolution":{"observed_at":"2026-05-11T10:06:04.848915Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2506.09565","last_updated":"2025-06-13T08:30:38Z","snapshot_observed_at":"2026-08-09T06:12:36.730836Z","submitted_at":"2025-06-11T09:56:39Z","title":"SemanticSplat: Feed-Forward 3D Scene Understanding with Language-Aware Gaussian Fields","version":2},"cited_work":{"arxiv_id":"2506.09565","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2506.09565","snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Semanticsplat: Feed-forward 3d scene understanding with language-aware gaussian fields","venue":null,"work_id":"02679082-1066-4d13-9a3e-d7bd2befbcc0","year":2025},"citing_paper":{"arxiv_id":"2605.23287","last_updated":"2026-05-22T06:59:00Z","snapshot_observed_at":"2026-07-06T23:33:29.550551Z","submitted_at":"2026-05-22T06:59:00Z","title":"LangFlash: Feed-forward 3D Language Gaussian Splatting from Sparse Unposed Images","version":1},"reference_index":21,"source":"pdf_text","source_observed_at":"2026-05-25T04:26:13.921964Z"},"links":{"cited_paper":"/paper/2506.09565","citing_paper":"/paper/2605.23287"},"observation_digest":"sha256:4200c86d05a8e2db5ff10f4c7092d0af566b69eaf1e290abe1b20bf2808ee10c","observation_id":"0ea7b131-3685-459b-a895-d88220c875c1","resolution":{"observed_at":"2026-05-25T04:26:37.594976Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2506.09565","last_updated":"2025-06-13T08:30:38Z","snapshot_observed_at":"2026-08-09T06:12:36.730836Z","submitted_at":"2025-06-11T09:56:39Z","title":"SemanticSplat: Feed-Forward 3D Scene Understanding with Language-Aware Gaussian Fields","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2506.09565","snapshot_observed_at":"2026-08-01T23:11:20.341168Z","title":"arXiv preprint arXiv:2506.09565 (2025)","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2607.15536","last_updated":"2026-07-17T01:12:26Z","snapshot_observed_at":"2026-08-08T01:12:13.969621Z","submitted_at":"2026-07-17T01:12:26Z","title":"E3DGS: Unified Geometric-Photometric Equivariance for 3D Gaussian Splatting via Color-as-Geometry Embedding","version":1},"reference_index":22,"source":"pdf_text","source_observed_at":"2026-08-01T23:11:20.341168Z"},"links":{"cited_paper":"/paper/2506.09565","citing_paper":"/paper/2607.15536"},"observation_digest":"sha256:c65b83f366b7c3a7a0b0dda62582700d6e34ad8252ab42db51f22e367db5ca87","observation_id":"ac95f349-bf93-4b7d-9a02-274ac506a627","resolution":{"observed_at":"2026-08-01T23:11:20.341168Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}}],"links":{"evidence":"/evidence","html":"/paper/2506.09565/citation-record","integrity":"/paper/2506.09565/integrity","json":"/paper/2506.09565/citation-record.json","paper":"/paper/2506.09565"},"outbound":[{"citation":{"cited_paper":{"arxiv_id":"2108.07258","last_updated":"2022-07-12T23:45:14Z","snapshot_observed_at":"2026-08-02T09:20:40.804790Z","submitted_at":"2021-08-16T17:50:08Z","title":"On the Opportunities and Risks of Foundation Models","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2108.07258","snapshot_observed_at":"2026-08-07T04:51:28.253320Z","title":"On the opportunities and risks of foundation models.arXiv preprint arXiv:2108.07258, 2021","venue":null,"work_id":null,"year":2021},"citing_paper":{"arxiv_id":"2506.09565","last_updated":"2025-06-13T08:30:38Z","snapshot_observed_at":"2026-08-09T06:12:36.730836Z","submitted_at":"2025-06-11T09:56:39Z","title":"SemanticSplat: Feed-Forward 3D Scene Understanding with Language-Aware Gaussian Fields","version":2},"reference_index":1,"source":"pdf_text","source_observed_at":"2026-08-07T04:51:28.253320Z"},"links":{"cited_paper":"/paper/2108.07258","citing_paper":"/paper/2506.09565"},"observation_digest":"sha256:200921af215ba61ec731691987adb090a6dc44c1e450be1df544ebf8e92c06cd","observation_id":"36f8223e-d108-44aa-9b67-f09620739081","resolution":{"observed_at":"2026-08-07T04:51:28.253320Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T04:51:35.533022Z","title":"Segment any 3d gaussians","venue":null,"work_id":"aed4575c-0cb4-4c6d-a28e-4516f26ba1fb","year":1971},"citing_paper":{"arxiv_id":"2506.09565","last_updated":"2025-06-13T08:30:38Z","snapshot_observed_at":"2026-08-09T06:12:36.730836Z","submitted_at":"2025-06-11T09:56:39Z","title":"SemanticSplat: Feed-Forward 3D Scene Understanding with Language-Aware Gaussian Fields","version":2},"reference_index":2,"source":"pdf_text","source_observed_at":"2026-08-07T04:51:28.316560Z"},"links":{"citing_paper":"/paper/2506.09565"},"observation_digest":"sha256:96dc969505595fb327cc265487d39356352a11e4ba6d4d02e7992e29deb641eb","observation_id":"92873a2f-2106-4410-aeb6-090c09d3d780","resolution":{"observed_at":"2026-08-07T04:51:35.537741Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T04:51:35.516880Z","title":"pixelsplat: 3d gaussian splats from image pairs for scalable generalizable 3d reconstruction","venue":null,"work_id":"e4966c8e-aa5b-4ae7-8a51-c531efe559a4","year":2024},"citing_paper":{"arxiv_id":"2506.09565","last_updated":"2025-06-13T08:30:38Z","snapshot_observed_at":"2026-08-09T06:12:36.730836Z","submitted_at":"2025-06-11T09:56:39Z","title":"SemanticSplat: Feed-Forward 3D Scene Understanding with Language-Aware Gaussian Fields","version":2},"reference_index":3,"source":"pdf_text","source_observed_at":"2026-08-07T04:51:28.386148Z"},"links":{"citing_paper":"/paper/2506.09565"},"observation_digest":"sha256:9ac959ba6b0c4288ea68fa502d83395b7f9d6d75d265aff26a79484f67df5a5c","observation_id":"adf41de5-51d2-48aa-bfd0-982f3de8c157","resolution":{"observed_at":"2026-08-07T04:51:35.521762Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T04:51:35.500526Z","title":"Lara: Efficient large-baseline radiance fields","venue":null,"work_id":"cbe43053-e7fb-4752-afba-e10c9bb6ce4c","year":2024},"citing_paper":{"arxiv_id":"2506.09565","last_updated":"2025-06-13T08:30:38Z","snapshot_observed_at":"2026-08-09T06:12:36.730836Z","submitted_at":"2025-06-11T09:56:39Z","title":"SemanticSplat: Feed-Forward 3D Scene Understanding with Language-Aware Gaussian Fields","version":2},"reference_index":4,"source":"pdf_text","source_observed_at":"2026-08-07T04:51:28.454663Z"},"links":{"citing_paper":"/paper/2506.09565"},"observation_digest":"sha256:e96717716bdd0515df3f6e57593ef34c1af1a43aeb6152f0f5f32e91212b767b","observation_id":"f7785709-3860-4f15-b08f-3baa71b7f97e","resolution":{"observed_at":"2026-08-07T04:51:35.505764Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T04:51:28.579966Z","title":"Feat2gs: Probing visual foundation models with gaussian splatting.arXiv preprint arXiv:2412.09606, 2024","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2506.09565","last_updated":"2025-06-13T08:30:38Z","snapshot_observed_at":"2026-08-09T06:12:36.730836Z","submitted_at":"2025-06-11T09:56:39Z","title":"SemanticSplat: Feed-Forward 3D Scene Understanding with Language-Aware Gaussian Fields","version":2},"reference_index":5,"source":"pdf_text","source_observed_at":"2026-08-07T04:51:28.579966Z"},"links":{"citing_paper":"/paper/2506.09565"},"observation_digest":"sha256:e704c324fb5c41a4997d455c42e6c40bf381014277920e17b4bd382572e8dc1e","observation_id":"84ae9e75-33f5-479b-96b0-c12269dfe54e","resolution":{"observed_at":"2026-08-07T04:51:28.579966Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T04:51:28.666040Z","title":"Mvsplat: Efficient 3d gaussian splatting from sparse multi-view images.ECCV, 2024","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2506.09565","last_updated":"2025-06-13T08:30:38Z","snapshot_observed_at":"2026-08-09T06:12:36.730836Z","submitted_at":"2025-06-11T09:56:39Z","title":"SemanticSplat: Feed-Forward 3D Scene Understanding with Language-Aware Gaussian Fields","version":2},"reference_index":6,"source":"pdf_text","source_observed_at":"2026-08-07T04:51:28.666040Z"},"links":{"citing_paper":"/paper/2506.09565"},"observation_digest":"sha256:c08f10b9abd9a43a4e8ff89dd357260e3198291369249883dedaff028c4eb1fa","observation_id":"5aab33d7-fc39-46c2-b2d0-039504dbc521","resolution":{"observed_at":"2026-08-07T04:51:28.666040Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T04:51:35.474311Z","title":"Dico-nerf: Difference of cosine similarity for neural rendering of fisheye driving scenes","venue":null,"work_id":"d25ce722-eb23-42e4-b011-fa4a2f20074c","year":2024},"citing_paper":{"arxiv_id":"2506.09565","last_updated":"2025-06-13T08:30:38Z","snapshot_observed_at":"2026-08-09T06:12:36.730836Z","submitted_at":"2025-06-11T09:56:39Z","title":"SemanticSplat: Feed-Forward 3D Scene Understanding with Language-Aware Gaussian Fields","version":2},"reference_index":7,"source":"pdf_text","source_observed_at":"2026-08-07T04:51:28.751492Z"},"links":{"citing_paper":"/paper/2506.09565"},"observation_digest":"sha256:6e9e988765f9026ff85edd8c7b02fd2fa7a8e9b826cb4f13b265571d339c7056","observation_id":"7cfb13eb-f122-4884-bbb7-f5a031cc34b0","resolution":{"observed_at":"2026-08-07T04:51:35.479070Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T04:51:35.457668Z","title":"A space-sweep approach to true multi- image matching","venue":null,"work_id":"8dd14068-e7c0-4695-aea0-386dedeb1c75","year":1996},"citing_paper":{"arxiv_id":"2506.09565","last_updated":"2025-06-13T08:30:38Z","snapshot_observed_at":"2026-08-09T06:12:36.730836Z","submitted_at":"2025-06-11T09:56:39Z","title":"SemanticSplat: Feed-Forward 3D Scene Understanding with Language-Aware Gaussian Fields","version":2},"reference_index":8,"source":"pdf_text","source_observed_at":"2026-08-07T04:51:28.804286Z"},"links":{"citing_paper":"/paper/2506.09565"},"observation_digest":"sha256:649d692cda668301c1502deedbc59888b212b1b5329ea5a53ec02624693d5c4a","observation_id":"47844662-2386-40c3-a612-2556a9282dd3","resolution":{"observed_at":"2026-08-07T04:51:35.462882Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T04:51:28.905999Z","title":"Scannet: Richly-annotated 3d reconstructions of indoor scenes","venue":null,"work_id":null,"year":2017},"citing_paper":{"arxiv_id":"2506.09565","last_updated":"2025-06-13T08:30:38Z","snapshot_observed_at":"2026-08-09T06:12:36.730836Z","submitted_at":"2025-06-11T09:56:39Z","title":"SemanticSplat: Feed-Forward 3D Scene Understanding with Language-Aware Gaussian Fields","version":2},"reference_index":9,"source":"pdf_text","source_observed_at":"2026-08-07T04:51:28.905999Z"},"links":{"citing_paper":"/paper/2506.09565"},"observation_digest":"sha256:2e9925f955e340c1b6d7915299a2974e6cf2f803ce107026cfa3e567281a39ef","observation_id":"6af5b9da-af5b-4340-a72e-230f00990a2e","resolution":{"observed_at":"2026-08-07T04:51:28.905999Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2010.11929","last_updated":"2021-06-03T13:08:56Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2020-10-22T17:55:59Z","title":"An Image is Worth 16x16 Words: Transformers for Image Recognition at Scale","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2010.11929","snapshot_observed_at":"2026-08-07T04:51:29.015566Z","title":"An image is worth 16x16 words: Trans- formers for image recognition at scale.arXiv preprint arXiv:2010.11929, 2020","venue":null,"work_id":null,"year":2010},"citing_paper":{"arxiv_id":"2506.09565","last_updated":"2025-06-13T08:30:38Z","snapshot_observed_at":"2026-08-09T06:12:36.730836Z","submitted_at":"2025-06-11T09:56:39Z","title":"SemanticSplat: Feed-Forward 3D Scene Understanding with Language-Aware Gaussian Fields","version":2},"reference_index":10,"source":"pdf_text","source_observed_at":"2026-08-07T04:51:29.015566Z"},"links":{"cited_paper":"/paper/2010.11929","citing_paper":"/paper/2506.09565"},"observation_digest":"sha256:96cfbc4d34b3b8b240e01a6ed4bbfbbd2fef87eb35d43722cfad49a5000ca7bf","observation_id":"995ffc11-d0b0-4fb7-9c13-fd028e8bcb53","resolution":{"observed_at":"2026-08-07T04:51:29.015566Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2209.08776","last_updated":"2022-10-12T00:33:31Z","snapshot_observed_at":"2026-08-02T15:53:33.492056Z","submitted_at":"2022-09-19T06:03:17Z","title":"NeRF-SOS: Any-View Self-supervised Object Segmentation on Complex Scenes","version":6},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2209.08776","snapshot_observed_at":"2026-08-07T04:51:29.129239Z","title":"Nerf-sos: Any-view self- supervised object segmentation on complex scenes.arXiv preprint arXiv:2209.08776, 2022","venue":null,"work_id":null,"year":2022},"citing_paper":{"arxiv_id":"2506.09565","last_updated":"2025-06-13T08:30:38Z","snapshot_observed_at":"2026-08-09T06:12:36.730836Z","submitted_at":"2025-06-11T09:56:39Z","title":"SemanticSplat: Feed-Forward 3D Scene Understanding with Language-Aware Gaussian Fields","version":2},"reference_index":11,"source":"pdf_text","source_observed_at":"2026-08-07T04:51:29.129239Z"},"links":{"cited_paper":"/paper/2209.08776","citing_paper":"/paper/2506.09565"},"observation_digest":"sha256:1cc7b4cf075dc833c7c2bae0620b958b38aed73050c3009c88e959b6cdc331b0","observation_id":"30116496-1dab-4ad4-9717-18999f0f96d1","resolution":{"observed_at":"2026-08-07T04:51:29.129239Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T04:51:35.428467Z","title":"Large spatial model: End-to-end unposed images to semantic 3d.NeurIPS, 2024","venue":null,"work_id":"a6751ea3-c728-4b1a-ab02-a35ec5f4375a","year":2024},"citing_paper":{"arxiv_id":"2506.09565","last_updated":"2025-06-13T08:30:38Z","snapshot_observed_at":"2026-08-09T06:12:36.730836Z","submitted_at":"2025-06-11T09:56:39Z","title":"SemanticSplat: Feed-Forward 3D Scene Understanding with Language-Aware Gaussian Fields","version":2},"reference_index":12,"source":"pdf_text","source_observed_at":"2026-08-07T04:51:29.229348Z"},"links":{"citing_paper":"/paper/2506.09565"},"observation_digest":"sha256:a74ba1b037c524f6e77ea41ac8f23006f8b2563185e49299a92d547ae5dabe45","observation_id":"ad412090-908e-4cc3-a5af-6ea1cd1b488f","resolution":{"observed_at":"2026-08-07T04:51:35.433320Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2407.01220","last_updated":"2024-12-19T09:16:19Z","snapshot_observed_at":"2026-08-01T14:32:56.909671Z","submitted_at":"2024-07-01T12:07:26Z","title":"Fast and Efficient: Mask Neural Fields for 3D Scene Segmentation","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2407.01220","snapshot_observed_at":"2026-08-07T04:51:29.365175Z","title":"Fast and efficient: Mask neural fields for 3d scene segmentation.arXiv preprint arXiv:2407.01220, 2024","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2506.09565","last_updated":"2025-06-13T08:30:38Z","snapshot_observed_at":"2026-08-09T06:12:36.730836Z","submitted_at":"2025-06-11T09:56:39Z","title":"SemanticSplat: Feed-Forward 3D Scene Understanding with Language-Aware Gaussian Fields","version":2},"reference_index":13,"source":"pdf_text","source_observed_at":"2026-08-07T04:51:29.365175Z"},"links":{"cited_paper":"/paper/2407.01220","citing_paper":"/paper/2506.09565"},"observation_digest":"sha256:ab007f17a360a4d1071f9c72fbe25e1e039d28968a87f446f228930d8d2fac6b","observation_id":"92fb754c-c8f1-43c6-bd58-da9d7277fed6","resolution":{"observed_at":"2026-08-07T04:51:29.365175Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T04:51:35.410013Z","title":"Finding structure with randomness: Probabilistic algorithms for constructing approximate matrix decompositions.SIAM review, 53(2):217–288, 2011","venue":null,"work_id":"d054bf80-681c-4995-8626-983d3319ade5","year":2011},"citing_paper":{"arxiv_id":"2506.09565","last_updated":"2025-06-13T08:30:38Z","snapshot_observed_at":"2026-08-09T06:12:36.730836Z","submitted_at":"2025-06-11T09:56:39Z","title":"SemanticSplat: Feed-Forward 3D Scene Understanding with Language-Aware Gaussian Fields","version":2},"reference_index":14,"source":"pdf_text","source_observed_at":"2026-08-07T04:51:29.529794Z"},"links":{"citing_paper":"/paper/2506.09565"},"observation_digest":"sha256:b5ab5d021bf9d10e5883d967079a49588723e5fc215bbc8c2d619a455659f199","observation_id":"0d1b1dad-566e-42d8-918d-1af15e797bde","resolution":{"observed_at":"2026-08-07T04:51:35.415798Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T04:51:35.393828Z","title":"3d gaussian splatting for real-time radiance field rendering.ACM Transactions on Graphics, 42(4), 2023","venue":null,"work_id":"d840b073-9835-4ffa-ba07-731a8fbc4d5c","year":2023},"citing_paper":{"arxiv_id":"2506.09565","last_updated":"2025-06-13T08:30:38Z","snapshot_observed_at":"2026-08-09T06:12:36.730836Z","submitted_at":"2025-06-11T09:56:39Z","title":"SemanticSplat: Feed-Forward 3D Scene Understanding with Language-Aware Gaussian Fields","version":2},"reference_index":15,"source":"pdf_text","source_observed_at":"2026-08-07T04:51:29.682323Z"},"links":{"citing_paper":"/paper/2506.09565"},"observation_digest":"sha256:daf85eb1e36615cba2b31241ee4d4e8d4bb7fa43b503725b342190b29a24f94a","observation_id":"9550bdf7-9d0b-4147-ba7b-648dccda2c2e","resolution":{"observed_at":"2026-08-07T04:51:35.398647Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T04:51:35.376636Z","title":"Lerf: Language embedded radiance fields","venue":null,"work_id":"492fb63e-6cf1-4cad-bc52-2e7e088ed629","year":2023},"citing_paper":{"arxiv_id":"2506.09565","last_updated":"2025-06-13T08:30:38Z","snapshot_observed_at":"2026-08-09T06:12:36.730836Z","submitted_at":"2025-06-11T09:56:39Z","title":"SemanticSplat: Feed-Forward 3D Scene Understanding with Language-Aware Gaussian Fields","version":2},"reference_index":16,"source":"pdf_text","source_observed_at":"2026-08-07T04:51:29.828251Z"},"links":{"citing_paper":"/paper/2506.09565"},"observation_digest":"sha256:e2e5202d6701946441c633d9b1a4cdc8284d7a6725f048d58153ba33526df88d","observation_id":"b0a13ef5-fa30-401d-bbe1-cf8301470ebe","resolution":{"observed_at":"2026-08-07T04:51:35.381994Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T04:51:35.356747Z","title":"Lerf: Language embed- ded radiance fields","venue":null,"work_id":"bb4d9f9a-36ef-468c-af02-737dd1efa96b","year":2023},"citing_paper":{"arxiv_id":"2506.09565","last_updated":"2025-06-13T08:30:38Z","snapshot_observed_at":"2026-08-09T06:12:36.730836Z","submitted_at":"2025-06-11T09:56:39Z","title":"SemanticSplat: Feed-Forward 3D Scene Understanding with Language-Aware Gaussian Fields","version":2},"reference_index":17,"source":"pdf_text","source_observed_at":"2026-08-07T04:51:29.951827Z"},"links":{"citing_paper":"/paper/2506.09565"},"observation_digest":"sha256:6a46cbc8291a2f95adb62c407fc9d859965886459916abfe83b34a3503e25494","observation_id":"88cd8378-066b-4baa-b8da-c4ff31e6c1c6","resolution":{"observed_at":"2026-08-07T04:51:35.363611Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T04:51:35.339998Z","title":"Garfield: Group any- thing with radiance fields","venue":null,"work_id":"6ee8ba87-1dcc-41ac-b65c-bc9f285e7101","year":2024},"citing_paper":{"arxiv_id":"2506.09565","last_updated":"2025-06-13T08:30:38Z","snapshot_observed_at":"2026-08-09T06:12:36.730836Z","submitted_at":"2025-06-11T09:56:39Z","title":"SemanticSplat: Feed-Forward 3D Scene Understanding with Language-Aware Gaussian Fields","version":2},"reference_index":18,"source":"pdf_text","source_observed_at":"2026-08-07T04:51:30.101415Z"},"links":{"citing_paper":"/paper/2506.09565"},"observation_digest":"sha256:07d09c0eb65825614f26847529303e590c2900972e9a5d00ff3dec221525ad3f","observation_id":"b39d4120-f28f-4bd0-aa9e-fc91a689014d","resolution":{"observed_at":"2026-08-07T04:51:35.344583Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"1412.6980","last_updated":"2017-01-30T01:27:54Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2014-12-22T13:54:29Z","title":"Adam: A Method for Stochastic Optimization","version":9},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"1412.6980","snapshot_observed_at":"2026-08-07T04:51:30.243799Z","title":"Adam: A method for stochastic optimiza- tion.arXiv preprint arXiv:1412.6980, 2014","venue":null,"work_id":null,"year":2014},"citing_paper":{"arxiv_id":"2506.09565","last_updated":"2025-06-13T08:30:38Z","snapshot_observed_at":"2026-08-09T06:12:36.730836Z","submitted_at":"2025-06-11T09:56:39Z","title":"SemanticSplat: Feed-Forward 3D Scene Understanding with Language-Aware Gaussian Fields","version":2},"reference_index":19,"source":"pdf_text","source_observed_at":"2026-08-07T04:51:30.243799Z"},"links":{"cited_paper":"/paper/1412.6980","citing_paper":"/paper/2506.09565"},"observation_digest":"sha256:91fb350f5801992b60448863d4cf7421557c84dbef3fe371281e77427afeed04","observation_id":"b304be4e-4021-411d-bd3a-c38370374b98","resolution":{"observed_at":"2026-08-07T04:51:30.243799Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T04:51:35.324528Z","title":"Segment any- thing","venue":null,"work_id":"6bbbb8a7-85ea-4579-ae7c-734a7ac1fec2","year":2023},"citing_paper":{"arxiv_id":"2506.09565","last_updated":"2025-06-13T08:30:38Z","snapshot_observed_at":"2026-08-09T06:12:36.730836Z","submitted_at":"2025-06-11T09:56:39Z","title":"SemanticSplat: Feed-Forward 3D Scene Understanding with Language-Aware Gaussian Fields","version":2},"reference_index":20,"source":"pdf_text","source_observed_at":"2026-08-07T04:51:30.362021Z"},"links":{"citing_paper":"/paper/2506.09565"},"observation_digest":"sha256:8a2a9b67618c95e53985d912ec6907983b2479e64419e345579b62708584abbe","observation_id":"66d3601c-8fe7-47f2-ab2c-fd23b044a1e6","resolution":{"observed_at":"2026-08-07T04:51:35.328897Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T04:51:35.307564Z","title":"Rethinking open-vocabulary segmentation of radiance fields in 3d space","venue":null,"work_id":"f2bbe630-1e74-498d-a678-1e7decb08bf1","year":2025},"citing_paper":{"arxiv_id":"2506.09565","last_updated":"2025-06-13T08:30:38Z","snapshot_observed_at":"2026-08-09T06:12:36.730836Z","submitted_at":"2025-06-11T09:56:39Z","title":"SemanticSplat: Feed-Forward 3D Scene Understanding with Language-Aware Gaussian Fields","version":2},"reference_index":21,"source":"pdf_text","source_observed_at":"2026-08-07T04:51:30.473589Z"},"links":{"citing_paper":"/paper/2506.09565"},"observation_digest":"sha256:b7ad013fcba76fecac8e24a92305997783d87a625e472745dc8f2bb0dd8fe60a","observation_id":"3a1c8495-6ac3-4541-a7e5-8a481a39248a","resolution":{"observed_at":"2026-08-07T04:51:35.313500Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T04:51:35.287597Z","title":"Language-driven semantic seg- mentation","venue":null,"work_id":"28ac0389-90a8-417d-bef6-f27cbb131c81","year":2022},"citing_paper":{"arxiv_id":"2506.09565","last_updated":"2025-06-13T08:30:38Z","snapshot_observed_at":"2026-08-09T06:12:36.730836Z","submitted_at":"2025-06-11T09:56:39Z","title":"SemanticSplat: Feed-Forward 3D Scene Understanding with Language-Aware Gaussian Fields","version":2},"reference_index":22,"source":"pdf_text","source_observed_at":"2026-08-07T04:51:30.592153Z"},"links":{"citing_paper":"/paper/2506.09565"},"observation_digest":"sha256:39f5f847053f960780401ab9961b4298bad1ba55c5ad32b239a831aa4f75dd89","observation_id":"33cc8f46-9e44-472d-9060-da82626fae0d","resolution":{"observed_at":"2026-08-07T04:51:35.293453Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T04:51:30.684951Z","title":"Langsurf: Language- embedded surface gaussians for 3d scene understanding","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2506.09565","last_updated":"2025-06-13T08:30:38Z","snapshot_observed_at":"2026-08-09T06:12:36.730836Z","submitted_at":"2025-06-11T09:56:39Z","title":"SemanticSplat: Feed-Forward 3D Scene Understanding with Language-Aware Gaussian Fields","version":2},"reference_index":23,"source":"pdf_text","source_observed_at":"2026-08-07T04:51:30.684951Z"},"links":{"citing_paper":"/paper/2506.09565"},"observation_digest":"sha256:ebfa027f2881cf20fbedf4224fe6bb6bbd3240456fda3e974b24a5583e6afd0d","observation_id":"b243533b-1deb-42b0-a033-05128f4f0225","resolution":{"observed_at":"2026-08-07T04:51:30.684951Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2503.10437","last_updated":"2025-04-01T03:10:36Z","snapshot_observed_at":"2026-08-07T17:06:51.817878Z","submitted_at":"2025-03-13T14:58:22Z","title":"4D LangSplat: 4D Language Gaussian Splatting via Multimodal Large Language Models","version":2},"cited_work":{"arxiv_id":"2503.10437","doi":null,"metadata_source":"pith","pith_arxiv_id":"2503.10437","snapshot_observed_at":"2026-08-07T04:51:34.688869Z","title":"4D LangSplat: 4D Language Gaussian Splatting via Multimodal Large Language Models","venue":"cs.CV","work_id":"5087e628-cf63-4754-8697-229313848cb5","year":2025},"citing_paper":{"arxiv_id":"2506.09565","last_updated":"2025-06-13T08:30:38Z","snapshot_observed_at":"2026-08-09T06:12:36.730836Z","submitted_at":"2025-06-11T09:56:39Z","title":"SemanticSplat: Feed-Forward 3D Scene Understanding with Language-Aware Gaussian Fields","version":2},"reference_index":24,"source":"pdf_text","source_observed_at":"2026-08-07T04:51:30.789148Z"},"links":{"cited_paper":"/paper/2503.10437","citing_paper":"/paper/2506.09565"},"observation_digest":"sha256:65108970b049d9c04ff4b7281717d1b8df2e02fc5dd9e699ca34d1d3528945de","observation_id":"223f0773-3665-4495-8eb3-8e7876f8139b","resolution":{"observed_at":"2026-08-07T04:51:34.693956Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T04:51:35.266667Z","title":null,"venue":null,"work_id":"dc7083d9-841a-417b-8ac5-d3ac8a50236e","year":2024},"citing_paper":{"arxiv_id":"2506.09565","last_updated":"2025-06-13T08:30:38Z","snapshot_observed_at":"2026-08-09T06:12:36.730836Z","submitted_at":"2025-06-11T09:56:39Z","title":"SemanticSplat: Feed-Forward 3D Scene Understanding with Language-Aware Gaussian Fields","version":2},"reference_index":25,"source":"pdf_text","source_observed_at":"2026-08-07T04:51:30.940843Z"},"links":{"citing_paper":"/paper/2506.09565"},"observation_digest":"sha256:5e1a4610a5fda4f0eff9eade6c8e6229e723fe8a9524d6203484dc43a6d59e33","observation_id":"276e8e1a-0da4-4fe8-a919-f766fc0819fa","resolution":{"observed_at":"2026-08-07T04:51:35.272429Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T04:51:31.072127Z","title":"Focal loss for dense object detection","venue":null,"work_id":null,"year":2017},"citing_paper":{"arxiv_id":"2506.09565","last_updated":"2025-06-13T08:30:38Z","snapshot_observed_at":"2026-08-09T06:12:36.730836Z","submitted_at":"2025-06-11T09:56:39Z","title":"SemanticSplat: Feed-Forward 3D Scene Understanding with Language-Aware Gaussian Fields","version":2},"reference_index":26,"source":"pdf_text","source_observed_at":"2026-08-07T04:51:31.072127Z"},"links":{"citing_paper":"/paper/2506.09565"},"observation_digest":"sha256:c4cd91a59edab499daeb6794e3b62264efeb1b380bd4799edaa5c3c6f9a4e968","observation_id":"56f6696c-0f88-4b01-8017-bae2d60b0e2b","resolution":{"observed_at":"2026-08-07T04:51:31.072127Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2410.06014","last_updated":"2024-10-08T13:16:49Z","snapshot_observed_at":"2026-08-06T18:59:25.212327Z","submitted_at":"2024-10-08T13:16:49Z","title":"SplaTraj: Camera Trajectory Generation with Semantic Gaussian Splatting","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2410.06014","snapshot_observed_at":"2026-08-07T04:51:31.182990Z","title":"Splatraj: Camera trajectory generation with semantic gaussian splatting.arXiv preprint arXiv:2410.06014, 2024","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2506.09565","last_updated":"2025-06-13T08:30:38Z","snapshot_observed_at":"2026-08-09T06:12:36.730836Z","submitted_at":"2025-06-11T09:56:39Z","title":"SemanticSplat: Feed-Forward 3D Scene Understanding with Language-Aware Gaussian Fields","version":2},"reference_index":27,"source":"pdf_text","source_observed_at":"2026-08-07T04:51:31.182990Z"},"links":{"cited_paper":"/paper/2410.06014","citing_paper":"/paper/2506.09565"},"observation_digest":"sha256:969523c1d8bb7a74c129f43a4964db712d57ec54955ce1d55953a7f888c6995e","observation_id":"0db308ab-e439-4f6c-9570-682f9c0dd140","resolution":{"observed_at":"2026-08-07T04:51:31.182990Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T04:51:31.313886Z","title":"Swin transformer v2: Scaling up capacity and resolution","venue":null,"work_id":null,"year":2022},"citing_paper":{"arxiv_id":"2506.09565","last_updated":"2025-06-13T08:30:38Z","snapshot_observed_at":"2026-08-09T06:12:36.730836Z","submitted_at":"2025-06-11T09:56:39Z","title":"SemanticSplat: Feed-Forward 3D Scene Understanding with Language-Aware Gaussian Fields","version":2},"reference_index":28,"source":"pdf_text","source_observed_at":"2026-08-07T04:51:31.313886Z"},"links":{"citing_paper":"/paper/2506.09565"},"observation_digest":"sha256:e2e0b5f9ef1bf6150d2728d3159da885500bfa3c88fb6601ffdb8b2d56d340c6","observation_id":"fe78d7f4-0ebd-4496-8e0f-eca19822e7c5","resolution":{"observed_at":"2026-08-07T04:51:31.313886Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T04:51:35.223848Z","title":"Srinivasan, Matthew Tancik, Jonathan T","venue":null,"work_id":"95ea84da-6c7e-42af-b3ef-4854d55f83be","year":2020},"citing_paper":{"arxiv_id":"2506.09565","last_updated":"2025-06-13T08:30:38Z","snapshot_observed_at":"2026-08-09T06:12:36.730836Z","submitted_at":"2025-06-11T09:56:39Z","title":"SemanticSplat: Feed-Forward 3D Scene Understanding with Language-Aware Gaussian Fields","version":2},"reference_index":29,"source":"pdf_text","source_observed_at":"2026-08-07T04:51:31.444238Z"},"links":{"citing_paper":"/paper/2506.09565"},"observation_digest":"sha256:9b7baf42ddf3f6f888ec56a99b183d0de051d76e916d57587ca185c5f0d28d32","observation_id":"c6ad3fd5-a530-4d11-97db-5117b1ecdd84","resolution":{"observed_at":"2026-08-07T04:51:35.230033Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T04:51:31.557523Z","title":"V-net: Fully convolutional neural networks for volumetric medical image segmentation","venue":null,"work_id":null,"year":2016},"citing_paper":{"arxiv_id":"2506.09565","last_updated":"2025-06-13T08:30:38Z","snapshot_observed_at":"2026-08-09T06:12:36.730836Z","submitted_at":"2025-06-11T09:56:39Z","title":"SemanticSplat: Feed-Forward 3D Scene Understanding with Language-Aware Gaussian Fields","version":2},"reference_index":30,"source":"pdf_text","source_observed_at":"2026-08-07T04:51:31.557523Z"},"links":{"citing_paper":"/paper/2506.09565"},"observation_digest":"sha256:66495a3ed80488800489175c0bcb697792b4b9f12914708271247f5cfb3176b9","observation_id":"034f2b61-b07e-4e0f-83d9-217d64d1ad92","resolution":{"observed_at":"2026-08-07T04:51:31.557523Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T04:51:35.193767Z","title":"Scikit-learn: Machine learning in python.the Journal of machine Learning research, 12:2825–2830, 2011","venue":null,"work_id":"c5f53695-4bfc-4372-b83b-74f8802f0263","year":2011},"citing_paper":{"arxiv_id":"2506.09565","last_updated":"2025-06-13T08:30:38Z","snapshot_observed_at":"2026-08-09T06:12:36.730836Z","submitted_at":"2025-06-11T09:56:39Z","title":"SemanticSplat: Feed-Forward 3D Scene Understanding with Language-Aware Gaussian Fields","version":2},"reference_index":31,"source":"pdf_text","source_observed_at":"2026-08-07T04:51:31.677527Z"},"links":{"citing_paper":"/paper/2506.09565"},"observation_digest":"sha256:7e0bad808b5245a31537985d486e8df25fa4ddc7c81bf8742e63c18ab87808c0","observation_id":"f6a78f7d-1b44-4b95-a460-1d949883708d","resolution":{"observed_at":"2026-08-07T04:51:35.198635Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T04:51:35.174672Z","title":"Langsplat: 3d language gaussian splatting","venue":null,"work_id":"adece6a2-3f88-4e8e-acfe-6f191b3932c3","year":2024},"citing_paper":{"arxiv_id":"2506.09565","last_updated":"2025-06-13T08:30:38Z","snapshot_observed_at":"2026-08-09T06:12:36.730836Z","submitted_at":"2025-06-11T09:56:39Z","title":"SemanticSplat: Feed-Forward 3D Scene Understanding with Language-Aware Gaussian Fields","version":2},"reference_index":32,"source":"pdf_text","source_observed_at":"2026-08-07T04:51:31.766138Z"},"links":{"citing_paper":"/paper/2506.09565"},"observation_digest":"sha256:5e6da313c687d9de4b103b9676fb097630dfe59e2a771ba2e097c39752dc7c09","observation_id":"9b292a39-4bd6-457b-9223-57b6d526cc7f","resolution":{"observed_at":"2026-08-07T04:51:35.179420Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T04:51:31.845361Z","title":"Learning transferable visual models from natural language supervi- sion","venue":null,"work_id":null,"year":2021},"citing_paper":{"arxiv_id":"2506.09565","last_updated":"2025-06-13T08:30:38Z","snapshot_observed_at":"2026-08-09T06:12:36.730836Z","submitted_at":"2025-06-11T09:56:39Z","title":"SemanticSplat: Feed-Forward 3D Scene Understanding with Language-Aware Gaussian Fields","version":2},"reference_index":33,"source":"pdf_text","source_observed_at":"2026-08-07T04:51:31.845361Z"},"links":{"citing_paper":"/paper/2506.09565"},"observation_digest":"sha256:6b34f34213fec11a49602bf651b4472c360505368e9ee22abac2090e1f3c42ed","observation_id":"1186f828-3861-41a2-a691-2a4b3bfa1925","resolution":{"observed_at":"2026-08-07T04:51:31.845361Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T04:51:31.974727Z","title":"High-resolution image synthesis with latent diffusion models","venue":null,"work_id":null,"year":2022},"citing_paper":{"arxiv_id":"2506.09565","last_updated":"2025-06-13T08:30:38Z","snapshot_observed_at":"2026-08-09T06:12:36.730836Z","submitted_at":"2025-06-11T09:56:39Z","title":"SemanticSplat: Feed-Forward 3D Scene Understanding with Language-Aware Gaussian Fields","version":2},"reference_index":34,"source":"pdf_text","source_observed_at":"2026-08-07T04:51:31.974727Z"},"links":{"citing_paper":"/paper/2506.09565"},"observation_digest":"sha256:f6c5fd915a1ba4a2c1943bff1b316e1d0ded9c2efdfc03f8e265b008b2431dab","observation_id":"6de12d53-f76a-4f09-8e3e-6f54443914f9","resolution":{"observed_at":"2026-08-07T04:51:31.974727Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T04:51:35.135794Z","title":"U- net: Convolutional networks for biomedical image segmenta- tion","venue":null,"work_id":"be38b0c1-1287-46e4-9829-1fe83fc5d2b6","year":2015},"citing_paper":{"arxiv_id":"2506.09565","last_updated":"2025-06-13T08:30:38Z","snapshot_observed_at":"2026-08-09T06:12:36.730836Z","submitted_at":"2025-06-11T09:56:39Z","title":"SemanticSplat: Feed-Forward 3D Scene Understanding with Language-Aware Gaussian Fields","version":2},"reference_index":35,"source":"pdf_text","source_observed_at":"2026-08-07T04:51:32.047791Z"},"links":{"citing_paper":"/paper/2506.09565"},"observation_digest":"sha256:d097d4ecfb6fd5ee5abcdf21b0ad7ff8ff03cdc8d904145468c7993fbd13433d","observation_id":"d5e1ef4e-3753-4350-bd10-c3de2b0a220f","resolution":{"observed_at":"2026-08-07T04:51:35.140466Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2405.20310","last_updated":"2024-06-03T15:13:55Z","snapshot_observed_at":"2026-08-09T11:37:37.604099Z","submitted_at":"2024-05-30T17:52:52Z","title":"A Pixel Is Worth More Than One 3D Gaussians in Single-View 3D Reconstruction","version":3},"cited_work":{"arxiv_id":"2405.20310","doi":null,"metadata_source":"pith","pith_arxiv_id":"2405.20310","snapshot_observed_at":"2026-08-07T04:51:34.647939Z","title":"A Pixel Is Worth More Than One 3D Gaussians in Single-View 3D Reconstruction","venue":"cs.CV","work_id":"cb3a79c8-5e94-4653-92b2-ba19c51a55ac","year":2024},"citing_paper":{"arxiv_id":"2506.09565","last_updated":"2025-06-13T08:30:38Z","snapshot_observed_at":"2026-08-09T06:12:36.730836Z","submitted_at":"2025-06-11T09:56:39Z","title":"SemanticSplat: Feed-Forward 3D Scene Understanding with Language-Aware Gaussian Fields","version":2},"reference_index":36,"source":"pdf_text","source_observed_at":"2026-08-07T04:51:32.145643Z"},"links":{"cited_paper":"/paper/2405.20310","citing_paper":"/paper/2506.09565"},"observation_digest":"sha256:e64f974dc51ac54b72fbfb8dbfc2f2e7520b27c0f2442b37002a706ad280a46a","observation_id":"f1231d2b-da43-4ffe-a7fd-2dccbd621a54","resolution":{"observed_at":"2026-08-07T04:51:34.653540Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2403.18795","last_updated":"2024-05-24T18:43:28Z","snapshot_observed_at":"2026-07-06T17:52:05.486230Z","submitted_at":"2024-03-27T17:40:14Z","title":"Gamba: Marry Gaussian Splatting with Mamba for single view 3D reconstruction","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2403.18795","snapshot_observed_at":"2026-08-07T04:51:32.300655Z","title":"Gamba: Marry gaussian splatting with mamba for single view 3d reconstruc- tion.arXiv preprint arXiv:2403.18795, 2024","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2506.09565","last_updated":"2025-06-13T08:30:38Z","snapshot_observed_at":"2026-08-09T06:12:36.730836Z","submitted_at":"2025-06-11T09:56:39Z","title":"SemanticSplat: Feed-Forward 3D Scene Understanding with Language-Aware Gaussian Fields","version":2},"reference_index":37,"source":"pdf_text","source_observed_at":"2026-08-07T04:51:32.300655Z"},"links":{"cited_paper":"/paper/2403.18795","citing_paper":"/paper/2506.09565"},"observation_digest":"sha256:50bb81d790a325db9b412ad36679da54906e8c5af83192fbcb98c0ec5c2419e4","observation_id":"ea391e68-bf1a-4bb4-9591-70b4bcc23ad2","resolution":{"observed_at":"2026-08-07T04:51:32.300655Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T04:51:35.117971Z","title":"Distilled feature fields enable few-shot language-guided manipulation","venue":null,"work_id":"fd74a361-2c66-470c-bd1a-a67263fa05ac","year":2023},"citing_paper":{"arxiv_id":"2506.09565","last_updated":"2025-06-13T08:30:38Z","snapshot_observed_at":"2026-08-09T06:12:36.730836Z","submitted_at":"2025-06-11T09:56:39Z","title":"SemanticSplat: Feed-Forward 3D Scene Understanding with Language-Aware Gaussian Fields","version":2},"reference_index":38,"source":"pdf_text","source_observed_at":"2026-08-07T04:51:32.415364Z"},"links":{"citing_paper":"/paper/2506.09565"},"observation_digest":"sha256:1769be6113808cf1be35c252bba88047887a9e5a4372be60517bb45157ea7f8d","observation_id":"8a57950f-f9f0-453b-a8ac-571350b55e5e","resolution":{"observed_at":"2026-08-07T04:51:35.122876Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T04:51:35.102139Z","title":"Language embedded 3d gaussians for open-vocabulary scene understanding","venue":null,"work_id":"0cffbcae-e399-4e96-9511-b51d78bda4f2","year":2024},"citing_paper":{"arxiv_id":"2506.09565","last_updated":"2025-06-13T08:30:38Z","snapshot_observed_at":"2026-08-09T06:12:36.730836Z","submitted_at":"2025-06-11T09:56:39Z","title":"SemanticSplat: Feed-Forward 3D Scene Understanding with Language-Aware Gaussian Fields","version":2},"reference_index":39,"source":"pdf_text","source_observed_at":"2026-08-07T04:51:32.492336Z"},"links":{"citing_paper":"/paper/2506.09565"},"observation_digest":"sha256:81b8c4dae4fe35e7c66ce21073c5f28f8eddc1bc52372755c52aa26ab1b9e7b6","observation_id":"76de3596-b14c-4a5e-a7d7-839250c0d19d","resolution":{"observed_at":"2026-08-07T04:51:35.107357Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T04:51:35.087272Z","title":"Panoptic lifting for 3d scene understanding with neural fields","venue":null,"work_id":"98b0b8b9-295b-4282-b7c2-ab8ea18c2738","year":2023},"citing_paper":{"arxiv_id":"2506.09565","last_updated":"2025-06-13T08:30:38Z","snapshot_observed_at":"2026-08-09T06:12:36.730836Z","submitted_at":"2025-06-11T09:56:39Z","title":"SemanticSplat: Feed-Forward 3D Scene Understanding with Language-Aware Gaussian Fields","version":2},"reference_index":40,"source":"pdf_text","source_observed_at":"2026-08-07T04:51:32.567715Z"},"links":{"citing_paper":"/paper/2506.09565"},"observation_digest":"sha256:5e392475eec9afc46800d1939a5d7a5b3c859299ae57ef10707d8af5ed255156","observation_id":"476706b9-51a2-4807-b19c-eff667a85a2a","resolution":{"observed_at":"2026-08-07T04:51:35.091936Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T04:51:35.071909Z","title":"Splatter image: Ultra-fast single-view 3d recon- struction","venue":null,"work_id":"74d57523-ae47-4425-a3e2-934f8bb3017f","year":2024},"citing_paper":{"arxiv_id":"2506.09565","last_updated":"2025-06-13T08:30:38Z","snapshot_observed_at":"2026-08-09T06:12:36.730836Z","submitted_at":"2025-06-11T09:56:39Z","title":"SemanticSplat: Feed-Forward 3D Scene Understanding with Language-Aware Gaussian Fields","version":2},"reference_index":41,"source":"pdf_text","source_observed_at":"2026-08-07T04:51:32.712893Z"},"links":{"citing_paper":"/paper/2506.09565"},"observation_digest":"sha256:a8ba95ed6d8d38d851235da7820192c90d4daa862c66e95d8192852164fbb661","observation_id":"754a4510-c5b4-4903-a502-e99017bd9b4b","resolution":{"observed_at":"2026-08-07T04:51:35.076683Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T04:51:32.843416Z","title":"Lgm: Large multi-view gaussian model for high-resolution 3d content creation","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2506.09565","last_updated":"2025-06-13T08:30:38Z","snapshot_observed_at":"2026-08-09T06:12:36.730836Z","submitted_at":"2025-06-11T09:56:39Z","title":"SemanticSplat: Feed-Forward 3D Scene Understanding with Language-Aware Gaussian Fields","version":2},"reference_index":42,"source":"pdf_text","source_observed_at":"2026-08-07T04:51:32.843416Z"},"links":{"citing_paper":"/paper/2506.09565"},"observation_digest":"sha256:f082e48747398eb477b6b482382f694d30fc8a1a33482d3301d41b92fa6b6644","observation_id":"467edd44-c41b-44cb-8c03-14340b2e6ed7","resolution":{"observed_at":"2026-08-07T04:51:32.843416Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T04:51:35.046120Z","title":"Bovik, H.R","venue":null,"work_id":"099dbe66-92d2-46ff-931a-449a25231170","year":2004},"citing_paper":{"arxiv_id":"2506.09565","last_updated":"2025-06-13T08:30:38Z","snapshot_observed_at":"2026-08-09T06:12:36.730836Z","submitted_at":"2025-06-11T09:56:39Z","title":"SemanticSplat: Feed-Forward 3D Scene Understanding with Language-Aware Gaussian Fields","version":2},"reference_index":43,"source":"pdf_text","source_observed_at":"2026-08-07T04:51:32.956574Z"},"links":{"citing_paper":"/paper/2506.09565"},"observation_digest":"sha256:e12b19591d4836734a2227f6b14efd248ed32116e1a84bd196c5a330b909568d","observation_id":"a9572d4a-bac7-4513-92af-8c4be503812d","resolution":{"observed_at":"2026-08-07T04:51:35.050657Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T04:51:33.092394Z","title":"Gmflow: Learning optical flow via global matching","venue":null,"work_id":null,"year":2022},"citing_paper":{"arxiv_id":"2506.09565","last_updated":"2025-06-13T08:30:38Z","snapshot_observed_at":"2026-08-09T06:12:36.730836Z","submitted_at":"2025-06-11T09:56:39Z","title":"SemanticSplat: Feed-Forward 3D Scene Understanding with Language-Aware Gaussian Fields","version":2},"reference_index":44,"source":"pdf_text","source_observed_at":"2026-08-07T04:51:33.092394Z"},"links":{"citing_paper":"/paper/2506.09565"},"observation_digest":"sha256:4313bc81e00fa3d04b5d08ff8161df760f558652462ee72e7d5b0fd65fe97e4f","observation_id":"6278affd-9d73-43a7-a1a1-12953a65009c","resolution":{"observed_at":"2026-08-07T04:51:33.092394Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T04:51:33.215566Z","title":"Unifying flow, stereo and depth estimation.IEEE Transactions on Pattern Analysis and Machine Intelligence, 45(11):13941–13958, 2023","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2506.09565","last_updated":"2025-06-13T08:30:38Z","snapshot_observed_at":"2026-08-09T06:12:36.730836Z","submitted_at":"2025-06-11T09:56:39Z","title":"SemanticSplat: Feed-Forward 3D Scene Understanding with Language-Aware Gaussian Fields","version":2},"reference_index":45,"source":"pdf_text","source_observed_at":"2026-08-07T04:51:33.215566Z"},"links":{"citing_paper":"/paper/2506.09565"},"observation_digest":"sha256:b0eba4f46664f696f8d32639ffba40f044d8905d237fb5970ea97c3e2214e8ba","observation_id":"ace658ea-fb66-41bc-a193-e93c0d7e86ca","resolution":{"observed_at":"2026-08-07T04:51:33.215566Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T04:51:35.009935Z","title":"Grm: Large gaussian reconstruction model for efficient 3d recon- 10 struction and generation","venue":null,"work_id":"0e436372-feeb-40c8-8fab-f5e3069df9a4","year":2024},"citing_paper":{"arxiv_id":"2506.09565","last_updated":"2025-06-13T08:30:38Z","snapshot_observed_at":"2026-08-09T06:12:36.730836Z","submitted_at":"2025-06-11T09:56:39Z","title":"SemanticSplat: Feed-Forward 3D Scene Understanding with Language-Aware Gaussian Fields","version":2},"reference_index":46,"source":"pdf_text","source_observed_at":"2026-08-07T04:51:33.351815Z"},"links":{"citing_paper":"/paper/2506.09565"},"observation_digest":"sha256:a18e58b982056b049b061700420b61cd23d2e2bf7dca0b610ab69cc0f86ff8b4","observation_id":"a1baaffe-b180-45da-8f26-d4827b84fc77","resolution":{"observed_at":"2026-08-07T04:51:35.014700Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T04:51:34.994429Z","title":"Featurenerf: Learning generalizable nerfs by distilling foundation models","venue":null,"work_id":"2a5a971f-3764-4289-a423-c267ec7a8b0f","year":2023},"citing_paper":{"arxiv_id":"2506.09565","last_updated":"2025-06-13T08:30:38Z","snapshot_observed_at":"2026-08-09T06:12:36.730836Z","submitted_at":"2025-06-11T09:56:39Z","title":"SemanticSplat: Feed-Forward 3D Scene Understanding with Language-Aware Gaussian Fields","version":2},"reference_index":47,"source":"pdf_text","source_observed_at":"2026-08-07T04:51:33.463254Z"},"links":{"citing_paper":"/paper/2506.09565"},"observation_digest":"sha256:14232d03e808806b667185dfe61ddc09fd3d5c5955f09affce91d130afd761b8","observation_id":"53a31f79-7a5e-44c8-b5f3-34d639a33dcc","resolution":{"observed_at":"2026-08-07T04:51:34.999138Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T04:51:34.979597Z","title":"Gaus- sian grouping: Segment and edit anything in 3d scenes","venue":null,"work_id":"0a3b0757-72fb-4746-b90a-bafdbbaefb26","year":2025},"citing_paper":{"arxiv_id":"2506.09565","last_updated":"2025-06-13T08:30:38Z","snapshot_observed_at":"2026-08-09T06:12:36.730836Z","submitted_at":"2025-06-11T09:56:39Z","title":"SemanticSplat: Feed-Forward 3D Scene Understanding with Language-Aware Gaussian Fields","version":2},"reference_index":48,"source":"pdf_text","source_observed_at":"2026-08-07T04:51:33.563595Z"},"links":{"citing_paper":"/paper/2506.09565"},"observation_digest":"sha256:633fe05f93056962573b6e1b4fa223ad7e0c2b58f62f9f7732340ed94e69c34d","observation_id":"a2c0e3df-2e5f-499f-b020-104d7ae8ac88","resolution":{"observed_at":"2026-08-07T04:51:34.983970Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T04:51:34.963936Z","title":"Gs-lrm: Large recon- struction model for 3d gaussian splatting","venue":null,"work_id":"289aff08-da85-469e-b6c1-bb549e5d43a9","year":2024},"citing_paper":{"arxiv_id":"2506.09565","last_updated":"2025-06-13T08:30:38Z","snapshot_observed_at":"2026-08-09T06:12:36.730836Z","submitted_at":"2025-06-11T09:56:39Z","title":"SemanticSplat: Feed-Forward 3D Scene Understanding with Language-Aware Gaussian Fields","version":2},"reference_index":49,"source":"pdf_text","source_observed_at":"2026-08-07T04:51:33.649160Z"},"links":{"citing_paper":"/paper/2506.09565"},"observation_digest":"sha256:c844f7f238434f6561813fb264b69473099cfc052fcc1179dc0916fd779865b3","observation_id":"2f4bafa0-82d2-4d74-951e-f4e5cf824c04","resolution":{"observed_at":"2026-08-07T04:51:34.968383Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T04:51:33.804613Z","title":"The unreasonable effectiveness of deep features as a perceptual metric","venue":null,"work_id":null,"year":2018},"citing_paper":{"arxiv_id":"2506.09565","last_updated":"2025-06-13T08:30:38Z","snapshot_observed_at":"2026-08-09T06:12:36.730836Z","submitted_at":"2025-06-11T09:56:39Z","title":"SemanticSplat: Feed-Forward 3D Scene Understanding with Language-Aware Gaussian Fields","version":2},"reference_index":50,"source":"pdf_text","source_observed_at":"2026-08-07T04:51:33.804613Z"},"links":{"citing_paper":"/paper/2506.09565"},"observation_digest":"sha256:49c354dd554b11e4c5d25ededfb2a302c9bf3e1bcd0ed268bd10932c7f858f6b","observation_id":"0ff37b68-ab74-4158-9b6d-4ed84369f9f2","resolution":{"observed_at":"2026-08-07T04:51:33.804613Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T04:51:34.937430Z","title":null,"venue":null,"work_id":"e0816ee2-b312-45db-87f2-c54b240c838b","year":2021},"citing_paper":{"arxiv_id":"2506.09565","last_updated":"2025-06-13T08:30:38Z","snapshot_observed_at":"2026-08-09T06:12:36.730836Z","submitted_at":"2025-06-11T09:56:39Z","title":"SemanticSplat: Feed-Forward 3D Scene Understanding with Language-Aware Gaussian Fields","version":2},"reference_index":51,"source":"pdf_text","source_observed_at":"2026-08-07T04:51:33.912482Z"},"links":{"citing_paper":"/paper/2506.09565"},"observation_digest":"sha256:85ab1dbd242640b192a2c48a7b5796d3bd4f09c002e109d1500879e0339093f6","observation_id":"531b33bd-d8e6-4e4c-98fd-fd1f633ea66e","resolution":{"observed_at":"2026-08-07T04:51:34.941666Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T04:51:34.921725Z","title":"Feature 3dgs: Supercharging 3d gaussian splatting to enable distilled feature fields","venue":null,"work_id":"b6bf9121-fdcc-4f41-af78-ffdc52f32ea6","year":2024},"citing_paper":{"arxiv_id":"2506.09565","last_updated":"2025-06-13T08:30:38Z","snapshot_observed_at":"2026-08-09T06:12:36.730836Z","submitted_at":"2025-06-11T09:56:39Z","title":"SemanticSplat: Feed-Forward 3D Scene Understanding with Language-Aware Gaussian Fields","version":2},"reference_index":53,"source":"pdf_text","source_observed_at":"2026-08-07T04:51:34.203367Z"},"links":{"citing_paper":"/paper/2506.09565"},"observation_digest":"sha256:2451823cacf85d4b01542727ace4851f7a8495d718c1ac25d74376ce9074f391","observation_id":"ad2f8e55-aa73-4b67-816b-7f79432eb760","resolution":{"observed_at":"2026-08-07T04:51:34.926631Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2503.20776","last_updated":"2025-03-28T04:48:48Z","snapshot_observed_at":"2026-08-07T16:34:51.326881Z","submitted_at":"2025-03-26T17:56:16Z","title":"Feature4X: Bridging Any Monocular Video to 4D Agentic AI with Versatile Gaussian Feature Fields","version":2},"cited_work":{"arxiv_id":"2503.20776","doi":null,"metadata_source":"pith","pith_arxiv_id":"2503.20776","snapshot_observed_at":"2026-08-07T04:51:34.604761Z","title":"Feature4X: Bridging Any Monocular Video to 4D Agentic AI with Versatile Gaussian Feature Fields","venue":"cs.CV","work_id":"a19e8d76-48e6-417b-a7f9-3a2c50b41394","year":2025},"citing_paper":{"arxiv_id":"2506.09565","last_updated":"2025-06-13T08:30:38Z","snapshot_observed_at":"2026-08-09T06:12:36.730836Z","submitted_at":"2025-06-11T09:56:39Z","title":"SemanticSplat: Feed-Forward 3D Scene Understanding with Language-Aware Gaussian Fields","version":2},"reference_index":54,"source":"pdf_text","source_observed_at":"2026-08-07T04:51:34.330095Z"},"links":{"cited_paper":"/paper/2503.20776","citing_paper":"/paper/2506.09565"},"observation_digest":"sha256:68a87c1504b26c606276bfef591de3f1d46f8d67ac4c772b0d091a3d09d290c5","observation_id":"a6f5ed98-91aa-4cc1-89a4-220a4f8d4b15","resolution":{"observed_at":"2026-08-07T04:51:34.612448Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"1805.09817","last_updated":"2018-05-24T17:58:02Z","snapshot_observed_at":"2026-07-06T06:41:05.545426Z","submitted_at":"2018-05-24T17:58:02Z","title":"Stereo Magnification: Learning View Synthesis using Multiplane Images","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"1805.09817","snapshot_observed_at":"2026-08-07T04:51:34.452486Z","title":"Stereo magnification: Learning view synthesis using multiplane images.arXiv preprint arXiv:1805.09817, 2018","venue":null,"work_id":null,"year":2018},"citing_paper":{"arxiv_id":"2506.09565","last_updated":"2025-06-13T08:30:38Z","snapshot_observed_at":"2026-08-09T06:12:36.730836Z","submitted_at":"2025-06-11T09:56:39Z","title":"SemanticSplat: Feed-Forward 3D Scene Understanding with Language-Aware Gaussian Fields","version":2},"reference_index":55,"source":"pdf_text","source_observed_at":"2026-08-07T04:51:34.452486Z"},"links":{"cited_paper":"/paper/1805.09817","citing_paper":"/paper/2506.09565"},"observation_digest":"sha256:bbcee9ccaed7c6e3ce544dd49c1f71c2d650989fde3b352ea77977bff79884a3","observation_id":"1be5cd9c-038e-493f-ba7a-91fc85b2d4ec","resolution":{"observed_at":"2026-08-07T04:51:34.452486Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2401.01970","last_updated":"2024-05-03T23:33:07Z","snapshot_observed_at":"2026-08-01T21:55:32.669838Z","submitted_at":"2024-01-03T20:39:02Z","title":"FMGS: Foundation Model Embedded 3D Gaussian Splatting for Holistic 3D Scene Understanding","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2401.01970","snapshot_observed_at":"2026-08-07T04:51:34.522907Z","title":"Fmgs: Foundation model embedded 3d gaussian splatting for holistic 3d scene understanding.arXiv preprint arXiv:2401.01970, 2024","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2506.09565","last_updated":"2025-06-13T08:30:38Z","snapshot_observed_at":"2026-08-09T06:12:36.730836Z","submitted_at":"2025-06-11T09:56:39Z","title":"SemanticSplat: Feed-Forward 3D Scene Understanding with Language-Aware Gaussian Fields","version":2},"reference_index":56,"source":"pdf_text","source_observed_at":"2026-08-07T04:51:34.522907Z"},"links":{"cited_paper":"/paper/2401.01970","citing_paper":"/paper/2506.09565"},"observation_digest":"sha256:aea3f5b652050c140d494e97a9724be8101fc36f1a5976890dead5d5848445e5","observation_id":"2d30c059-2770-4276-bf29-c95422c8963b","resolution":{"observed_at":"2026-08-07T04:51:34.522907Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}}],"paper":{"arxiv_id":"2506.09565","last_updated":"2025-06-13T08:30:38Z","latest_version":2,"primary_category":"cs.CV","snapshot_observed_at":"2026-08-09T06:12:36.730836Z","submitted_at":"2025-06-11T09:56:39Z","title":"SemanticSplat: Feed-Forward 3D Scene Understanding with Language-Aware Gaussian Fields"},"reference_resolution":{"displayed":55,"state_counts":{"malformed_identifier":0,"metadata_mismatch":0,"parse_uncertain":0,"unresolved":24,"verified_exact":3,"verified_fuzzy":28},"total_outbound_references":55},"refusal":"A citation records a reference. It does not transfer a finding from one paper to another.","schema":"pith.paper-citation-record.v1","standing_sources":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"thesis":"As of 9 August 2026, this Paper Citation Record lists 55 of 55 outbound references and 9 inbound Pith citation observations for arXiv:2506.09565."}