{"as_of":"2026-08-08T07:57:00Z","caps":{"database_statements":6,"inbound":100,"outbound":100},"context_digest":"sha256:7fc6f021265c23bb8f051963fee335e9f508ea62df451e4c3df1a9fb6e962a45","coverage":[{"denominator":81,"lane":"reference_resolution","note":"Typed states for the displayed outbound observations.","records_observed":81,"source":"paper_references, paper_reference_links","source_observed_at":"2026-07-15T13:44:18.418529Z","state":"measured"},{"denominator":81,"lane":"standing_notices","note":"One-hop event checks from named stored sources.","records_observed":81,"source":"scholarly_work_events, retraction_status_cache","source_observed_at":"2026-08-08T06:32:00.761636+00:00","state":"measured"},{"denominator":0,"lane":"inbound_itemization","note":"Pith citing papers itemized under the disclosed page cap.","records_observed":0,"source":"paper_references, paper_reference_links","source_observed_at":null,"state":"measured"},{"denominator":1,"lane":"external_citation_measurements","note":"A source-named dated measurement, never combined with another source.","records_observed":0,"source":"cited_works","source_observed_at":null,"state":"measured"}],"external_citation_measurements":[],"inbound":[],"links":{"evidence":"/evidence","html":"/paper/2603.06576/citation-record","integrity":"/paper/2603.06576/integrity","json":"/paper/2603.06576/citation-record.json","paper":"/paper/2603.06576"},"outbound":[{"citation":{"cited_paper":{"arxiv_id":"2303.08774","last_updated":"2024-03-04T06:01:33Z","snapshot_observed_at":"2026-08-07T07:30:12.213965Z","submitted_at":"2023-03-15T17:15:04Z","title":"GPT-4 Technical Report","version":6},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2303.08774","snapshot_observed_at":"2026-07-15T13:44:18.418529Z","title":"arXiv preprint arXiv:2303.08774 (2023)","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2603.06576","last_updated":"2026-07-02T20:40:34Z","snapshot_observed_at":"2026-08-05T20:27:31.445015Z","submitted_at":"2026-03-06T18:59:55Z","title":"BEVLM: Distilling Semantic Knowledge from LLMs into Bird's-Eye View Representations","version":2},"reference_index":1,"source":"pdf_text","source_observed_at":"2026-07-15T13:44:18.418529Z"},"links":{"cited_paper":"/paper/2303.08774","citing_paper":"/paper/2603.06576"},"observation_digest":"sha256:6ccc927d3329a9cbfd3421a9c3d6567d251f67ec70dfb5b1178e3a4d1b41085a","observation_id":"f5b27618-3efe-47f2-9253-566617b4ecfc","resolution":{"observed_at":"2026-07-15T13:44:18.418529Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-07-15T13:44:18.418529Z","title":"In: Proceedings of the IEEE International Con- ference on Computer Vision (ICCV) (2015)","venue":null,"work_id":null,"year":2015},"citing_paper":{"arxiv_id":"2603.06576","last_updated":"2026-07-02T20:40:34Z","snapshot_observed_at":"2026-08-05T20:27:31.445015Z","submitted_at":"2026-03-06T18:59:55Z","title":"BEVLM: Distilling Semantic Knowledge from LLMs into Bird's-Eye View Representations","version":2},"reference_index":2,"source":"pdf_text","source_observed_at":"2026-07-15T13:44:18.418529Z"},"links":{"citing_paper":"/paper/2603.06576"},"observation_digest":"sha256:8c2074f6aa1fe5940470311a09282b0548333be7044bd0419b22f36ce7d4570e","observation_id":"56e67d05-19ba-4e35-8b75-16690db92dd6","resolution":{"observed_at":"2026-07-15T13:44:18.418529Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-07-15T13:44:18.418529Z","title":"In: Proceedings of the acl work- shop on intrinsic and extrinsic evaluation measures for machine translation and/or summarization","venue":null,"work_id":null,"year":2005},"citing_paper":{"arxiv_id":"2603.06576","last_updated":"2026-07-02T20:40:34Z","snapshot_observed_at":"2026-08-05T20:27:31.445015Z","submitted_at":"2026-03-06T18:59:55Z","title":"BEVLM: Distilling Semantic Knowledge from LLMs into Bird's-Eye View Representations","version":2},"reference_index":3,"source":"pdf_text","source_observed_at":"2026-07-15T13:44:18.418529Z"},"links":{"citing_paper":"/paper/2603.06576"},"observation_digest":"sha256:3d18d6e6a9285a41fc43bc1092ed1bebeb55f022f75a8cc620741533e98e47ca","observation_id":"fded5ec6-4664-4287-9c50-d5a8851afe09","resolution":{"observed_at":"2026-07-15T13:44:18.418529Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2507.19370","last_updated":"2025-07-25T15:22:56Z","snapshot_observed_at":"2026-08-06T14:20:26.136837Z","submitted_at":"2025-07-25T15:22:56Z","title":"BEV-LLM: Leveraging Multimodal BEV Maps for Scene Captioning in Autonomous Driving","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2507.19370","snapshot_observed_at":"2026-07-15T13:44:18.418529Z","title":"arXiv Preprint arXiv:2507.19370 (2025)","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2603.06576","last_updated":"2026-07-02T20:40:34Z","snapshot_observed_at":"2026-08-05T20:27:31.445015Z","submitted_at":"2026-03-06T18:59:55Z","title":"BEVLM: Distilling Semantic Knowledge from LLMs into Bird's-Eye View Representations","version":2},"reference_index":4,"source":"pdf_text","source_observed_at":"2026-07-15T13:44:18.418529Z"},"links":{"cited_paper":"/paper/2507.19370","citing_paper":"/paper/2603.06576"},"observation_digest":"sha256:88116fdb58b4cd06bd21b45cbd9f9260c854573b2d8848712d9a8e95c60ddf35","observation_id":"4f9c9b50-f0d3-482f-8feb-dc6ad8db4541","resolution":{"observed_at":"2026-07-15T13:44:18.418529Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-07-15T13:44:18.418529Z","title":"In: Proceedings of the IEEE/CVF Conference on Computer Vision and Pattern Recognition (CVPR) (2020)","venue":null,"work_id":null,"year":2020},"citing_paper":{"arxiv_id":"2603.06576","last_updated":"2026-07-02T20:40:34Z","snapshot_observed_at":"2026-08-05T20:27:31.445015Z","submitted_at":"2026-03-06T18:59:55Z","title":"BEVLM: Distilling Semantic Knowledge from LLMs into Bird's-Eye View Representations","version":2},"reference_index":5,"source":"pdf_text","source_observed_at":"2026-07-15T13:44:18.418529Z"},"links":{"citing_paper":"/paper/2603.06576"},"observation_digest":"sha256:daf6effcf7e520f796d65be31bdb0e05f4092125a4200c093fbbea4c8c0359f3","observation_id":"6c986c9a-6769-4e43-90c2-e3432b1df1ea","resolution":{"observed_at":"2026-07-15T13:44:18.418529Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-07-15T13:44:18.418529Z","title":"In: Conference on Robot Learning (CoRL) (2025)","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2603.06576","last_updated":"2026-07-02T20:40:34Z","snapshot_observed_at":"2026-08-05T20:27:31.445015Z","submitted_at":"2026-03-06T18:59:55Z","title":"BEVLM: Distilling Semantic Knowledge from LLMs into Bird's-Eye View Representations","version":2},"reference_index":6,"source":"pdf_text","source_observed_at":"2026-07-15T13:44:18.418529Z"},"links":{"citing_paper":"/paper/2603.06576"},"observation_digest":"sha256:cef2edc5c2886741c83a2665fd494ee3dcfc2f2cf830fabc58bc842c621201e6","observation_id":"427eac64-010f-4c5b-97f7-2aa0e23cc2d8","resolution":{"observed_at":"2026-07-15T13:44:18.418529Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-07-15T13:44:18.418529Z","title":"In: IEEE Symposium on Security and Privacy (SP)","venue":null,"work_id":null,"year":2021},"citing_paper":{"arxiv_id":"2603.06576","last_updated":"2026-07-02T20:40:34Z","snapshot_observed_at":"2026-08-05T20:27:31.445015Z","submitted_at":"2026-03-06T18:59:55Z","title":"BEVLM: Distilling Semantic Knowledge from LLMs into Bird's-Eye View Representations","version":2},"reference_index":7,"source":"pdf_text","source_observed_at":"2026-07-15T13:44:18.418529Z"},"links":{"citing_paper":"/paper/2603.06576"},"observation_digest":"sha256:545a435836242f2cf8e2701f752b5c9ca474c82272aa23ed39f74a2ba3ee534f","observation_id":"3c1f7167-46f1-45c5-836d-850b3c4fca37","resolution":{"observed_at":"2026-07-15T13:44:18.418529Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-07-15T13:44:18.418529Z","title":"In: IEEE/CVF Conference on Computer Vision and Pattern Recognition (CVPR)","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2603.06576","last_updated":"2026-07-02T20:40:34Z","snapshot_observed_at":"2026-08-05T20:27:31.445015Z","submitted_at":"2026-03-06T18:59:55Z","title":"BEVLM: Distilling Semantic Knowledge from LLMs into Bird's-Eye View Representations","version":2},"reference_index":8,"source":"pdf_text","source_observed_at":"2026-07-15T13:44:18.418529Z"},"links":{"citing_paper":"/paper/2603.06576"},"observation_digest":"sha256:a584f17ce9cf803dddccaf6a6a5b3ab665424388f9538b0722a626d932360255","observation_id":"deecc11f-441a-41df-88f0-43c4f2c210c4","resolution":{"observed_at":"2026-07-15T13:44:18.418529Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2412.18607","last_updated":"2024-12-24T18:59:37Z","snapshot_observed_at":"2026-07-06T20:12:54.710203Z","submitted_at":"2024-12-24T18:59:37Z","title":"DrivingGPT: Unifying Driving World Modeling and Planning with Multi-modal Autoregressive Transformers","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2412.18607","snapshot_observed_at":"2026-07-15T13:44:18.418529Z","title":"arXiv preprint arXiv:2412.18607 (2024)","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2603.06576","last_updated":"2026-07-02T20:40:34Z","snapshot_observed_at":"2026-08-05T20:27:31.445015Z","submitted_at":"2026-03-06T18:59:55Z","title":"BEVLM: Distilling Semantic Knowledge from LLMs into Bird's-Eye View Representations","version":2},"reference_index":9,"source":"pdf_text","source_observed_at":"2026-07-15T13:44:18.418529Z"},"links":{"cited_paper":"/paper/2412.18607","citing_paper":"/paper/2603.06576"},"observation_digest":"sha256:54318db0ae1efe5d25fc298c7f1cc2982e1296d6deeb4e135ec05b0166f1aca4","observation_id":"c78396ac-1a8e-4b07-a795-1c8e07197410","resolution":{"observed_at":"2026-07-15T13:44:18.418529Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-07-15T13:44:18.418529Z","title":"In: Proceedings of the IEEE Conference on Computer Vision and Pattern Recognition (CVPR) (2017)","venue":null,"work_id":null,"year":2017},"citing_paper":{"arxiv_id":"2603.06576","last_updated":"2026-07-02T20:40:34Z","snapshot_observed_at":"2026-08-05T20:27:31.445015Z","submitted_at":"2026-03-06T18:59:55Z","title":"BEVLM: Distilling Semantic Knowledge from LLMs into Bird's-Eye View Representations","version":2},"reference_index":10,"source":"pdf_text","source_observed_at":"2026-07-15T13:44:18.418529Z"},"links":{"citing_paper":"/paper/2603.06576"},"observation_digest":"sha256:c3dc3d2cf451660ae4d1c2b874285eb8e2469657ed23ac1a50ad28de0db6bbb2","observation_id":"c3458e9f-34dd-4cbf-a52f-9504f697d256","resolution":{"observed_at":"2026-07-15T13:44:18.418529Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2507.06261","last_updated":"2025-12-19T14:25:46Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2025-07-07T17:36:04Z","title":"Gemini 2.5: Pushing the Frontier with Advanced Reasoning, Multimodality, Long Context, and Next Generation Agentic Capabilities","version":6},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2507.06261","snapshot_observed_at":"2026-07-15T13:44:18.418529Z","title":"arXiv Preprint arXiv:2507.06261 (2025)","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2603.06576","last_updated":"2026-07-02T20:40:34Z","snapshot_observed_at":"2026-08-05T20:27:31.445015Z","submitted_at":"2026-03-06T18:59:55Z","title":"BEVLM: Distilling Semantic Knowledge from LLMs into Bird's-Eye View Representations","version":2},"reference_index":11,"source":"pdf_text","source_observed_at":"2026-07-15T13:44:18.418529Z"},"links":{"cited_paper":"/paper/2507.06261","citing_paper":"/paper/2603.06576"},"observation_digest":"sha256:485c78184d760dad7c46efab8400fe19a8fe7daad9bdead44ba736d845fd7e3a","observation_id":"81330417-4931-4746-8015-fb5755b96352","resolution":{"observed_at":"2026-07-15T13:44:18.418529Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-07-15T13:44:18.418529Z","title":"In: Conference on Robot Learn- ing","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2603.06576","last_updated":"2026-07-02T20:40:34Z","snapshot_observed_at":"2026-08-05T20:27:31.445015Z","submitted_at":"2026-03-06T18:59:55Z","title":"BEVLM: Distilling Semantic Knowledge from LLMs into Bird's-Eye View Representations","version":2},"reference_index":12,"source":"pdf_text","source_observed_at":"2026-07-15T13:44:18.418529Z"},"links":{"citing_paper":"/paper/2603.06576"},"observation_digest":"sha256:6a5ce6575bfb464c8d80d3000128180af3ae7f9b9298b24bebd62cfbf42c2ba1","observation_id":"f13017d9-709e-40b8-9158-0d927e840c51","resolution":{"observed_at":"2026-07-15T13:44:18.418529Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-07-15T13:44:18.418529Z","title":"Advances in Neural Information Processing Systems37, 28706–28719 (2024)","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2603.06576","last_updated":"2026-07-02T20:40:34Z","snapshot_observed_at":"2026-08-05T20:27:31.445015Z","submitted_at":"2026-03-06T18:59:55Z","title":"BEVLM: Distilling Semantic Knowledge from LLMs into Bird's-Eye View Representations","version":2},"reference_index":13,"source":"pdf_text","source_observed_at":"2026-07-15T13:44:18.418529Z"},"links":{"citing_paper":"/paper/2603.06576"},"observation_digest":"sha256:ca5fd53f87c39d06f735e07719628db6553fcf1d14523fdb7cea4266cc67f474","observation_id":"f346c399-58c8-49d9-a3ca-688051ede5d4","resolution":{"observed_at":"2026-07-15T13:44:18.418529Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-07-15T13:44:18.418529Z","title":"In: NeurIPS 2025 Workshop on Regulatable ML (2025)","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2603.06576","last_updated":"2026-07-02T20:40:34Z","snapshot_observed_at":"2026-08-05T20:27:31.445015Z","submitted_at":"2026-03-06T18:59:55Z","title":"BEVLM: Distilling Semantic Knowledge from LLMs into Bird's-Eye View Representations","version":2},"reference_index":14,"source":"pdf_text","source_observed_at":"2026-07-15T13:44:18.418529Z"},"links":{"citing_paper":"/paper/2603.06576"},"observation_digest":"sha256:60601226423583c10e5285f714632101da800dd9c12ea0988a352d4e260439cb","observation_id":"598be738-fc35-461c-b280-388034588063","resolution":{"observed_at":"2026-07-15T13:44:18.418529Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2010.11929","last_updated":"2021-06-03T13:08:56Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2020-10-22T17:55:59Z","title":"An Image is Worth 16x16 Words: Transformers for Image Recognition at Scale","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2010.11929","snapshot_observed_at":"2026-07-15T13:44:18.418529Z","title":"arXiv Preprint arXiv:2010.11929 (2020)","venue":null,"work_id":null,"year":2010},"citing_paper":{"arxiv_id":"2603.06576","last_updated":"2026-07-02T20:40:34Z","snapshot_observed_at":"2026-08-05T20:27:31.445015Z","submitted_at":"2026-03-06T18:59:55Z","title":"BEVLM: Distilling Semantic Knowledge from LLMs into Bird's-Eye View Representations","version":2},"reference_index":15,"source":"pdf_text","source_observed_at":"2026-07-15T13:44:18.418529Z"},"links":{"cited_paper":"/paper/2010.11929","citing_paper":"/paper/2603.06576"},"observation_digest":"sha256:cf9754ce44ec92d670e34fd5debac3ad4eaeeead069904e8d5f0ad1f78ba25a6","observation_id":"750480db-bebe-4b19-9968-a9dd6ba97053","resolution":{"observed_at":"2026-07-15T13:44:18.418529Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-07-15T13:44:18.418529Z","title":"In: Proceedings of the IEEE/CVF Winter Con- ference on Applications of Computer Vision (WACV) (2022)","venue":null,"work_id":null,"year":2022},"citing_paper":{"arxiv_id":"2603.06576","last_updated":"2026-07-02T20:40:34Z","snapshot_observed_at":"2026-08-05T20:27:31.445015Z","submitted_at":"2026-03-06T18:59:55Z","title":"BEVLM: Distilling Semantic Knowledge from LLMs into Bird's-Eye View Representations","version":2},"reference_index":16,"source":"pdf_text","source_observed_at":"2026-07-15T13:44:18.418529Z"},"links":{"citing_paper":"/paper/2603.06576"},"observation_digest":"sha256:f753201a86ff454077c4d904dc1c1da1799388d0ae3d446cbd4cb409e86948d2","observation_id":"0b9b8b27-ae19-43b5-88ab-3ec2aa2b9ad9","resolution":{"observed_at":"2026-07-15T13:44:18.418529Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2505.15925","last_updated":"2026-04-03T18:48:28Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2025-05-21T18:24:36Z","title":"VERDI: VLM-Embedded Reasoning for Autonomous Driving","version":4},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2505.15925","snapshot_observed_at":"2026-07-15T13:44:18.418529Z","title":"arXiv preprint arXiv:2505.15925 (2025)","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2603.06576","last_updated":"2026-07-02T20:40:34Z","snapshot_observed_at":"2026-08-05T20:27:31.445015Z","submitted_at":"2026-03-06T18:59:55Z","title":"BEVLM: Distilling Semantic Knowledge from LLMs into Bird's-Eye View Representations","version":2},"reference_index":17,"source":"pdf_text","source_observed_at":"2026-07-15T13:44:18.418529Z"},"links":{"cited_paper":"/paper/2505.15925","citing_paper":"/paper/2603.06576"},"observation_digest":"sha256:8acc09b7611f85d1d255bf01dac7d064368d8a344a09ba14d39adf933c6c0a10","observation_id":"59687166-f6d2-495c-8c93-9767e994ae80","resolution":{"observed_at":"2026-07-15T13:44:18.418529Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2503.19755","last_updated":"2025-03-25T15:18:43Z","snapshot_observed_at":"2026-08-05T21:57:39.132989Z","submitted_at":"2025-03-25T15:18:43Z","title":"ORION: A Holistic End-to-End Autonomous Driving Framework by Vision-Language Instructed Action Generation","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2503.19755","snapshot_observed_at":"2026-07-15T13:44:18.418529Z","title":"arXiv preprint arXiv:2503.19755 (2025)","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2603.06576","last_updated":"2026-07-02T20:40:34Z","snapshot_observed_at":"2026-08-05T20:27:31.445015Z","submitted_at":"2026-03-06T18:59:55Z","title":"BEVLM: Distilling Semantic Knowledge from LLMs into Bird's-Eye View Representations","version":2},"reference_index":18,"source":"pdf_text","source_observed_at":"2026-07-15T13:44:18.418529Z"},"links":{"cited_paper":"/paper/2503.19755","citing_paper":"/paper/2603.06576"},"observation_digest":"sha256:fad7d893ac2da7d690c1e7ed05a53eed13d137d48f0db562ff0d67783623cacc","observation_id":"d6867d28-8a58-4cd2-a720-765fd83f104a","resolution":{"observed_at":"2026-07-15T13:44:18.418529Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-07-15T13:44:18.418529Z","title":"arXiv preprint arXiv:2509.06266 (2025)","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2603.06576","last_updated":"2026-07-02T20:40:34Z","snapshot_observed_at":"2026-08-05T20:27:31.445015Z","submitted_at":"2026-03-06T18:59:55Z","title":"BEVLM: Distilling Semantic Knowledge from LLMs into Bird's-Eye View Representations","version":2},"reference_index":19,"source":"pdf_text","source_observed_at":"2026-07-15T13:44:18.418529Z"},"links":{"citing_paper":"/paper/2603.06576"},"observation_digest":"sha256:e1a9012a679a1309f69714a604fc95bcbef5f8fe5f38a802e004864249d0e9e5","observation_id":"d53b4d22-004b-47e4-9f1b-6a76c59771ce","resolution":{"observed_at":"2026-07-15T13:44:18.418529Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2502.20108","last_updated":"2025-03-01T23:17:26Z","snapshot_observed_at":"2026-08-07T17:42:52.004108Z","submitted_at":"2025-02-27T14:02:14Z","title":"VDT-Auto: End-to-end Autonomous Driving with VLM-Guided Diffusion Transformers","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2502.20108","snapshot_observed_at":"2026-07-15T13:44:18.418529Z","title":null,"venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2603.06576","last_updated":"2026-07-02T20:40:34Z","snapshot_observed_at":"2026-08-05T20:27:31.445015Z","submitted_at":"2026-03-06T18:59:55Z","title":"BEVLM: Distilling Semantic Knowledge from LLMs into Bird's-Eye View Representations","version":2},"reference_index":20,"source":"pdf_text","source_observed_at":"2026-07-15T13:44:18.418529Z"},"links":{"cited_paper":"/paper/2502.20108","citing_paper":"/paper/2603.06576"},"observation_digest":"sha256:835d8970324c7996d5e9019ef2772b2e8b1a87c41d319897c5d9837f229b571b","observation_id":"25c05b00-9658-49c7-a415-4069f1923d82","resolution":{"observed_at":"2026-07-15T13:44:18.418529Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-07-15T13:44:18.418529Z","title":"In: Proceedings of the IEEE Conference on Computer Vision and Pattern Recognition (CVPR) (2016)","venue":null,"work_id":null,"year":2016},"citing_paper":{"arxiv_id":"2603.06576","last_updated":"2026-07-02T20:40:34Z","snapshot_observed_at":"2026-08-05T20:27:31.445015Z","submitted_at":"2026-03-06T18:59:55Z","title":"BEVLM: Distilling Semantic Knowledge from LLMs into Bird's-Eye View Representations","version":2},"reference_index":21,"source":"pdf_text","source_observed_at":"2026-07-15T13:44:18.418529Z"},"links":{"citing_paper":"/paper/2603.06576"},"observation_digest":"sha256:10839c1a78b175e906b3a59a4de293dcb0264bba80f37a9e1b5f790cb149a984","observation_id":"5f360623-0c6e-4f1b-99af-c1ae14e401dc","resolution":{"observed_at":"2026-07-15T13:44:18.418529Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-07-15T13:44:18.418529Z","title":"In: Proceedings of the IEEE/CVF International Con- ference on Computer Vision (ICCV)","venue":null,"work_id":null,"year":1921},"citing_paper":{"arxiv_id":"2603.06576","last_updated":"2026-07-02T20:40:34Z","snapshot_observed_at":"2026-08-05T20:27:31.445015Z","submitted_at":"2026-03-06T18:59:55Z","title":"BEVLM: Distilling Semantic Knowledge from LLMs into Bird's-Eye View Representations","version":2},"reference_index":22,"source":"pdf_text","source_observed_at":"2026-07-15T13:44:18.418529Z"},"links":{"citing_paper":"/paper/2603.06576"},"observation_digest":"sha256:079eb2a56130febb50655ff51f0c7e3e1fab7890517c14e7072de64d204c9721","observation_id":"1868be13-a686-4497-8772-aa653e5d48c5","resolution":{"observed_at":"2026-07-15T13:44:18.418529Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"1503.02531","last_updated":"2015-03-09T15:44:49Z","snapshot_observed_at":"2026-07-06T04:11:24.157003Z","submitted_at":"2015-03-09T15:44:49Z","title":"Distilling the Knowledge in a Neural Network","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"1503.02531","snapshot_observed_at":"2026-07-15T13:44:18.418529Z","title":"arXiv preprint arXiv:1503.02531 (2015)","venue":null,"work_id":null,"year":2015},"citing_paper":{"arxiv_id":"2603.06576","last_updated":"2026-07-02T20:40:34Z","snapshot_observed_at":"2026-08-05T20:27:31.445015Z","submitted_at":"2026-03-06T18:59:55Z","title":"BEVLM: Distilling Semantic Knowledge from LLMs into Bird's-Eye View Representations","version":2},"reference_index":23,"source":"pdf_text","source_observed_at":"2026-07-15T13:44:18.418529Z"},"links":{"cited_paper":"/paper/1503.02531","citing_paper":"/paper/2603.06576"},"observation_digest":"sha256:1346db7114f1c8484d47a0cdaab66ba62cb5b2494dedce6a69cf015001bbc317","observation_id":"114d13f6-c92a-4a66-a954-3eba869f78a9","resolution":{"observed_at":"2026-07-15T13:44:18.418529Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-07-15T13:44:18.418529Z","title":"In: IEEE/CVF Conference on Computer Vision and Pattern Recognition (CVPR) (2023) BEVLM 17","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2603.06576","last_updated":"2026-07-02T20:40:34Z","snapshot_observed_at":"2026-08-05T20:27:31.445015Z","submitted_at":"2026-03-06T18:59:55Z","title":"BEVLM: Distilling Semantic Knowledge from LLMs into Bird's-Eye View Representations","version":2},"reference_index":24,"source":"pdf_text","source_observed_at":"2026-07-15T13:44:18.418529Z"},"links":{"citing_paper":"/paper/2603.06576"},"observation_digest":"sha256:27b13181f2e40afba8112e4ef71c9eae448874075f9b73c29ade69680440198e","observation_id":"f4f20e89-0c6b-46e5-8daa-8ecb39292e81","resolution":{"observed_at":"2026-07-15T13:44:18.418529Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2112.11790","last_updated":"2022-06-16T09:15:52Z","snapshot_observed_at":"2026-07-06T12:21:28.059661Z","submitted_at":"2021-12-22T10:48:06Z","title":"BEVDet: High-performance Multi-camera 3D Object Detection in Bird-Eye-View","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2112.11790","snapshot_observed_at":"2026-07-15T13:44:18.418529Z","title":"arXiv Preprint arXiv:2112.11790 (2021)","venue":null,"work_id":null,"year":2021},"citing_paper":{"arxiv_id":"2603.06576","last_updated":"2026-07-02T20:40:34Z","snapshot_observed_at":"2026-08-05T20:27:31.445015Z","submitted_at":"2026-03-06T18:59:55Z","title":"BEVLM: Distilling Semantic Knowledge from LLMs into Bird's-Eye View Representations","version":2},"reference_index":25,"source":"pdf_text","source_observed_at":"2026-07-15T13:44:18.418529Z"},"links":{"cited_paper":"/paper/2112.11790","citing_paper":"/paper/2603.06576"},"observation_digest":"sha256:cdc032cf2e8b087030ea9714d7c418b20e09d0b85d04b44a3addabfece9489a1","observation_id":"a49c0bf7-34e3-4843-9288-e1e4f55da90d","resolution":{"observed_at":"2026-07-15T13:44:18.418529Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-07-15T13:44:18.418529Z","title":"In: European Conference on Computer Vision (ECCV)","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2603.06576","last_updated":"2026-07-02T20:40:34Z","snapshot_observed_at":"2026-08-05T20:27:31.445015Z","submitted_at":"2026-03-06T18:59:55Z","title":"BEVLM: Distilling Semantic Knowledge from LLMs into Bird's-Eye View Representations","version":2},"reference_index":26,"source":"pdf_text","source_observed_at":"2026-07-15T13:44:18.418529Z"},"links":{"citing_paper":"/paper/2603.06576"},"observation_digest":"sha256:d86ec4068d8cf9846c157865e7b48d4789dff2098e318d84031cbab9a54ff7bb","observation_id":"9dd4eccc-6504-431f-b947-4c1721a8e6b5","resolution":{"observed_at":"2026-07-15T13:44:18.418529Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-07-15T13:44:18.418529Z","title":"Transactions on Machine Learning Research (2025),https://openreview.net/forum?id=kH3t5lmOU8","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2603.06576","last_updated":"2026-07-02T20:40:34Z","snapshot_observed_at":"2026-08-05T20:27:31.445015Z","submitted_at":"2026-03-06T18:59:55Z","title":"BEVLM: Distilling Semantic Knowledge from LLMs into Bird's-Eye View Representations","version":2},"reference_index":27,"source":"pdf_text","source_observed_at":"2026-07-15T13:44:18.418529Z"},"links":{"citing_paper":"/paper/2603.06576"},"observation_digest":"sha256:4890e935b70899c8711dd7ae81a98b2a99e83a0882284a661a2f0c16f217cbc5","observation_id":"583ce413-c923-44bf-91b1-00a7ec2b1142","resolution":{"observed_at":"2026-07-15T13:44:18.418529Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-07-15T13:44:18.418529Z","title":"Advances in Neu- ral Information Processing Systems37, 819–844 (2024)","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2603.06576","last_updated":"2026-07-02T20:40:34Z","snapshot_observed_at":"2026-08-05T20:27:31.445015Z","submitted_at":"2026-03-06T18:59:55Z","title":"BEVLM: Distilling Semantic Knowledge from LLMs into Bird's-Eye View Representations","version":2},"reference_index":28,"source":"pdf_text","source_observed_at":"2026-07-15T13:44:18.418529Z"},"links":{"citing_paper":"/paper/2603.06576"},"observation_digest":"sha256:e30b6ecfbf03c71e0f1197ba6637de248179356e2bda0b7747201b7a96e75a4c","observation_id":"9f50a375-8da1-4665-ab16-9cbc7441f6df","resolution":{"observed_at":"2026-07-15T13:44:18.418529Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2410.22313","last_updated":"2024-10-29T17:53:56Z","snapshot_observed_at":"2026-07-31T01:16:26.372370Z","submitted_at":"2024-10-29T17:53:56Z","title":"Senna: Bridging Large Vision-Language Models and End-to-End Autonomous Driving","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2410.22313","snapshot_observed_at":"2026-07-15T13:44:18.418529Z","title":"arXiv preprint arXiv:2410.22313 (2024)","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2603.06576","last_updated":"2026-07-02T20:40:34Z","snapshot_observed_at":"2026-08-05T20:27:31.445015Z","submitted_at":"2026-03-06T18:59:55Z","title":"BEVLM: Distilling Semantic Knowledge from LLMs into Bird's-Eye View Representations","version":2},"reference_index":29,"source":"pdf_text","source_observed_at":"2026-07-15T13:44:18.418529Z"},"links":{"cited_paper":"/paper/2410.22313","citing_paper":"/paper/2603.06576"},"observation_digest":"sha256:8dac88d4af5e071dca7c2c9337155c2c10cbc2bb43087bd02cdd16fe226295ef","observation_id":"9a414f1b-9019-4fd5-b4ca-00d59c29dbc3","resolution":{"observed_at":"2026-07-15T13:44:18.418529Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-07-15T13:44:18.418529Z","title":"In: IEEE/CVF International Conference on Computer Vision (ICCV) (2023)","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2603.06576","last_updated":"2026-07-02T20:40:34Z","snapshot_observed_at":"2026-08-05T20:27:31.445015Z","submitted_at":"2026-03-06T18:59:55Z","title":"BEVLM: Distilling Semantic Knowledge from LLMs into Bird's-Eye View Representations","version":2},"reference_index":30,"source":"pdf_text","source_observed_at":"2026-07-15T13:44:18.418529Z"},"links":{"citing_paper":"/paper/2603.06576"},"observation_digest":"sha256:e1580e71493bddd410aa414052cbd467c65370a9c376a477ffe30462fe419b33","observation_id":"f752117d-c68b-48cc-a5c3-8168e90d6e49","resolution":{"observed_at":"2026-07-15T13:44:18.418529Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-07-15T13:44:18.418529Z","title":"In: IEEE/CVF International Conference on Computer Vision (ICCV) (2023)","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2603.06576","last_updated":"2026-07-02T20:40:34Z","snapshot_observed_at":"2026-08-05T20:27:31.445015Z","submitted_at":"2026-03-06T18:59:55Z","title":"BEVLM: Distilling Semantic Knowledge from LLMs into Bird's-Eye View Representations","version":2},"reference_index":31,"source":"pdf_text","source_observed_at":"2026-07-15T13:44:18.418529Z"},"links":{"citing_paper":"/paper/2603.06576"},"observation_digest":"sha256:13a07a589545eb10ce6992e9b52314044e8916ba293b14a9fec8f342e89347bb","observation_id":"b9a5ba3b-94c0-4723-9216-131b05863a14","resolution":{"observed_at":"2026-07-15T13:44:18.418529Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-07-15T13:44:18.418529Z","title":"In: IEEE/CVF International Conference on Computer Vision (ICCV) (2023)","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2603.06576","last_updated":"2026-07-02T20:40:34Z","snapshot_observed_at":"2026-08-05T20:27:31.445015Z","submitted_at":"2026-03-06T18:59:55Z","title":"BEVLM: Distilling Semantic Knowledge from LLMs into Bird's-Eye View Representations","version":2},"reference_index":32,"source":"pdf_text","source_observed_at":"2026-07-15T13:44:18.418529Z"},"links":{"citing_paper":"/paper/2603.06576"},"observation_digest":"sha256:45d57c14010a55d2ef7d85c5609415cbab96ffd58c6197c87b80d920c3596c1d","observation_id":"85168dc6-1234-4d32-bede-e72f57d92d67","resolution":{"observed_at":"2026-07-15T13:44:18.418529Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-07-15T13:44:18.418529Z","title":"In: Inter- national Conference on Machine Learning (ICML)","venue":null,"work_id":null,"year":2022},"citing_paper":{"arxiv_id":"2603.06576","last_updated":"2026-07-02T20:40:34Z","snapshot_observed_at":"2026-08-05T20:27:31.445015Z","submitted_at":"2026-03-06T18:59:55Z","title":"BEVLM: Distilling Semantic Knowledge from LLMs into Bird's-Eye View Representations","version":2},"reference_index":33,"source":"pdf_text","source_observed_at":"2026-07-15T13:44:18.418529Z"},"links":{"citing_paper":"/paper/2603.06576"},"observation_digest":"sha256:a3c1f4969fd4190970a01980db023ec33cfbc1bd087472247bb9b37d52658eb6","observation_id":"52e13757-a135-46a4-b094-30a32cac5848","resolution":{"observed_at":"2026-07-15T13:44:18.418529Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-07-15T13:44:18.418529Z","title":"In: European Conference on Computer Vision (ECCV)","venue":null,"work_id":null,"year":2022},"citing_paper":{"arxiv_id":"2603.06576","last_updated":"2026-07-02T20:40:34Z","snapshot_observed_at":"2026-08-05T20:27:31.445015Z","submitted_at":"2026-03-06T18:59:55Z","title":"BEVLM: Distilling Semantic Knowledge from LLMs into Bird's-Eye View Representations","version":2},"reference_index":34,"source":"pdf_text","source_observed_at":"2026-07-15T13:44:18.418529Z"},"links":{"citing_paper":"/paper/2603.06576"},"observation_digest":"sha256:077a3eb62aec6f7db1d02730856845e9a4b48aa5e3c256b738518b576e65586a","observation_id":"777a778f-3a81-4029-b034-b32bfac75618","resolution":{"observed_at":"2026-07-15T13:44:18.418529Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-07-15T13:44:18.418529Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2603.06576","last_updated":"2026-07-02T20:40:34Z","snapshot_observed_at":"2026-08-05T20:27:31.445015Z","submitted_at":"2026-03-06T18:59:55Z","title":"BEVLM: Distilling Semantic Knowledge from LLMs into Bird's-Eye View Representations","version":2},"reference_index":35,"source":"pdf_text","source_observed_at":"2026-07-15T13:44:18.418529Z"},"links":{"citing_paper":"/paper/2603.06576"},"observation_digest":"sha256:6593ae619735c63817181d04fc51593172ac5e7584c7b78d3184247658d5eb10","observation_id":"5763abb4-07c1-4d33-a3d8-f07a2d65b2b0","resolution":{"observed_at":"2026-07-15T13:44:18.418529Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-07-15T13:44:18.418529Z","title":"In: The Eleventh International Conference on Learning Representations (ICLR) (2023)","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2603.06576","last_updated":"2026-07-02T20:40:34Z","snapshot_observed_at":"2026-08-05T20:27:31.445015Z","submitted_at":"2026-03-06T18:59:55Z","title":"BEVLM: Distilling Semantic Knowledge from LLMs into Bird's-Eye View Representations","version":2},"reference_index":36,"source":"pdf_text","source_observed_at":"2026-07-15T13:44:18.418529Z"},"links":{"citing_paper":"/paper/2603.06576"},"observation_digest":"sha256:ddae2dfc9b23f81bad317efb66624471848abcc85a1701ddbc2d42852b9254d5","observation_id":"710c8da8-18f5-4b18-8004-56be9ea0ad83","resolution":{"observed_at":"2026-07-15T13:44:18.418529Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-07-15T13:44:18.418529Z","title":"In: Text Summarization Branches Out","venue":null,"work_id":null,"year":2004},"citing_paper":{"arxiv_id":"2603.06576","last_updated":"2026-07-02T20:40:34Z","snapshot_observed_at":"2026-08-05T20:27:31.445015Z","submitted_at":"2026-03-06T18:59:55Z","title":"BEVLM: Distilling Semantic Knowledge from LLMs into Bird's-Eye View Representations","version":2},"reference_index":37,"source":"pdf_text","source_observed_at":"2026-07-15T13:44:18.418529Z"},"links":{"citing_paper":"/paper/2603.06576"},"observation_digest":"sha256:33a10a9aff933f88ca252958019e5a63cdb8a45759e5a9aad47cae71b4810d75","observation_id":"fdd19884-fae7-4888-b100-69a31801864d","resolution":{"observed_at":"2026-07-15T13:44:18.418529Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-07-15T13:44:18.418529Z","title":"In: Proceedings of the IEEE Conference on Computer Vision and Pattern Recognition (CVPR) (2017)","venue":null,"work_id":null,"year":2017},"citing_paper":{"arxiv_id":"2603.06576","last_updated":"2026-07-02T20:40:34Z","snapshot_observed_at":"2026-08-05T20:27:31.445015Z","submitted_at":"2026-03-06T18:59:55Z","title":"BEVLM: Distilling Semantic Knowledge from LLMs into Bird's-Eye View Representations","version":2},"reference_index":38,"source":"pdf_text","source_observed_at":"2026-07-15T13:44:18.418529Z"},"links":{"citing_paper":"/paper/2603.06576"},"observation_digest":"sha256:8beb8e4dc7602e79d4b7441647c266e14e5a1af5d86c211e4725d61a8ac27ced","observation_id":"40badd2f-d804-4475-b717-ee979f2f62d7","resolution":{"observed_at":"2026-07-15T13:44:18.418529Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-07-15T13:44:18.418529Z","title":"Advances in Neural Information Processing Systems36(2023)","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2603.06576","last_updated":"2026-07-02T20:40:34Z","snapshot_observed_at":"2026-08-05T20:27:31.445015Z","submitted_at":"2026-03-06T18:59:55Z","title":"BEVLM: Distilling Semantic Knowledge from LLMs into Bird's-Eye View Representations","version":2},"reference_index":39,"source":"pdf_text","source_observed_at":"2026-07-15T13:44:18.418529Z"},"links":{"citing_paper":"/paper/2603.06576"},"observation_digest":"sha256:aeb505fc3d96b72fa38f09afe018836018a36667593762ba6ff6956fcb5fbdd8","observation_id":"6d27de12-1c84-4db3-90e8-c4c495ef1cb4","resolution":{"observed_at":"2026-07-15T13:44:18.418529Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-07-15T13:44:18.418529Z","title":"In: 2023 IEEE International Conference on Robotics and Automation (ICRA)","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2603.06576","last_updated":"2026-07-02T20:40:34Z","snapshot_observed_at":"2026-08-05T20:27:31.445015Z","submitted_at":"2026-03-06T18:59:55Z","title":"BEVLM: Distilling Semantic Knowledge from LLMs into Bird's-Eye View Representations","version":2},"reference_index":40,"source":"pdf_text","source_observed_at":"2026-07-15T13:44:18.418529Z"},"links":{"citing_paper":"/paper/2603.06576"},"observation_digest":"sha256:07b7b0cb242e2bea733744ee1f045b37711f337aefcfbc8e7625265bb269764c","observation_id":"7471c53b-eeb2-4b10-8bfd-844b121bf9b9","resolution":{"observed_at":"2026-07-15T13:44:18.418529Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-07-15T13:44:18.418529Z","title":"In: European Conference on Computer Vision (ECCV)","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2603.06576","last_updated":"2026-07-02T20:40:34Z","snapshot_observed_at":"2026-08-05T20:27:31.445015Z","submitted_at":"2026-03-06T18:59:55Z","title":"BEVLM: Distilling Semantic Knowledge from LLMs into Bird's-Eye View Representations","version":2},"reference_index":41,"source":"pdf_text","source_observed_at":"2026-07-15T13:44:18.418529Z"},"links":{"citing_paper":"/paper/2603.06576"},"observation_digest":"sha256:2bad1b0bc58cb576b86b517eaeb800bb17fcebeff41d9a35bf71d38c07cb8d53","observation_id":"6a29d038-b268-4740-a6d7-8daf9207867b","resolution":{"observed_at":"2026-07-15T13:44:18.418529Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"1711.05101","last_updated":"2019-01-04T21:01:49Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2017-11-14T14:24:06Z","title":"Decoupled Weight Decay Regularization","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"1711.05101","snapshot_observed_at":"2026-07-15T13:44:18.418529Z","title":"arXiv preprint arXiv:1711.05101 (2017)","venue":null,"work_id":null,"year":2017},"citing_paper":{"arxiv_id":"2603.06576","last_updated":"2026-07-02T20:40:34Z","snapshot_observed_at":"2026-08-05T20:27:31.445015Z","submitted_at":"2026-03-06T18:59:55Z","title":"BEVLM: Distilling Semantic Knowledge from LLMs into Bird's-Eye View Representations","version":2},"reference_index":42,"source":"pdf_text","source_observed_at":"2026-07-15T13:44:18.418529Z"},"links":{"cited_paper":"/paper/1711.05101","citing_paper":"/paper/2603.06576"},"observation_digest":"sha256:abb374e89440827f5b5dcc55953edb4a9c6376b499b17b62e86b86676288a306","observation_id":"578ec04d-4f0b-46e1-ac93-5d233f79fd48","resolution":{"observed_at":"2026-07-15T13:44:18.418529Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-07-15T13:44:18.418529Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2603.06576","last_updated":"2026-07-02T20:40:34Z","snapshot_observed_at":"2026-08-05T20:27:31.445015Z","submitted_at":"2026-03-06T18:59:55Z","title":"BEVLM: Distilling Semantic Knowledge from LLMs into Bird's-Eye View Representations","version":2},"reference_index":43,"source":"pdf_text","source_observed_at":"2026-07-15T13:44:18.418529Z"},"links":{"citing_paper":"/paper/2603.06576"},"observation_digest":"sha256:1fd100bf819c24f850d078d0bd3f4a28f96a61326c85f8da5e18a4811ddd6e4a","observation_id":"c5889815-3906-4938-8f51-a80cc25ff597","resolution":{"observed_at":"2026-07-15T13:44:18.418529Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-07-15T13:44:18.418529Z","title":"IEEE Transac- tions on Pattern Analysis and Machine Intelligence (2024)","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2603.06576","last_updated":"2026-07-02T20:40:34Z","snapshot_observed_at":"2026-08-05T20:27:31.445015Z","submitted_at":"2026-03-06T18:59:55Z","title":"BEVLM: Distilling Semantic Knowledge from LLMs into Bird's-Eye View Representations","version":2},"reference_index":44,"source":"pdf_text","source_observed_at":"2026-07-15T13:44:18.418529Z"},"links":{"citing_paper":"/paper/2603.06576"},"observation_digest":"sha256:6a9e9237465d87b691ad5edf2a1bc5310cd84fa571aa8a067827ae31457b8742","observation_id":"12845b5b-9beb-4bcc-b351-70b9445326be","resolution":{"observed_at":"2026-07-15T13:44:18.418529Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-07-15T13:44:18.418529Z","title":"In: Proceedings of the IEEE/CVF Conference on Computer Vision and Pattern Recognition (CVPR) (2024)","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2603.06576","last_updated":"2026-07-02T20:40:34Z","snapshot_observed_at":"2026-08-05T20:27:31.445015Z","submitted_at":"2026-03-06T18:59:55Z","title":"BEVLM: Distilling Semantic Knowledge from LLMs into Bird's-Eye View Representations","version":2},"reference_index":45,"source":"pdf_text","source_observed_at":"2026-07-15T13:44:18.418529Z"},"links":{"citing_paper":"/paper/2603.06576"},"observation_digest":"sha256:df59c1acb82e02c88d7fa40a066cbd408aa83b4e2f356028299112751c772f09","observation_id":"40fb6eb5-3cc6-4258-912a-23f177512241","resolution":{"observed_at":"2026-07-15T13:44:18.418529Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-07-15T13:44:18.418529Z","title":"arXiv preprint arXiv:2503.13430 (2025)","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2603.06576","last_updated":"2026-07-02T20:40:34Z","snapshot_observed_at":"2026-08-05T20:27:31.445015Z","submitted_at":"2026-03-06T18:59:55Z","title":"BEVLM: Distilling Semantic Knowledge from LLMs into Bird's-Eye View Representations","version":2},"reference_index":46,"source":"pdf_text","source_observed_at":"2026-07-15T13:44:18.418529Z"},"links":{"citing_paper":"/paper/2603.06576"},"observation_digest":"sha256:2b9e797ce46b81759333c31370a6c97d1c9a220e9ccc9bf1a53f1a53e894cb94","observation_id":"4965b3ba-4de5-412e-9f49-ab41b54d09a6","resolution":{"observed_at":"2026-07-15T13:44:18.418529Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-07-15T13:44:18.418529Z","title":"In: 2025 IEEE/RSJ International Conference on Intelligent Robots and Systems (IROS)","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2603.06576","last_updated":"2026-07-02T20:40:34Z","snapshot_observed_at":"2026-08-05T20:27:31.445015Z","submitted_at":"2026-03-06T18:59:55Z","title":"BEVLM: Distilling Semantic Knowledge from LLMs into Bird's-Eye View Representations","version":2},"reference_index":47,"source":"pdf_text","source_observed_at":"2026-07-15T13:44:18.418529Z"},"links":{"citing_paper":"/paper/2603.06576"},"observation_digest":"sha256:b64184eb1894d8a14c3a17e4551ec2de5c8e18089117c788b90b1cad88c3518c","observation_id":"13d92c4e-f7b8-4321-85df-07074268ddf8","resolution":{"observed_at":"2026-07-15T13:44:18.418529Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-07-15T13:44:18.418529Z","title":"IEEE Transactions on Intelligent Transportation Systems22(7), 4316–4336 (2020)","venue":null,"work_id":null,"year":2020},"citing_paper":{"arxiv_id":"2603.06576","last_updated":"2026-07-02T20:40:34Z","snapshot_observed_at":"2026-08-05T20:27:31.445015Z","submitted_at":"2026-03-06T18:59:55Z","title":"BEVLM: Distilling Semantic Knowledge from LLMs into Bird's-Eye View Representations","version":2},"reference_index":48,"source":"pdf_text","source_observed_at":"2026-07-15T13:44:18.418529Z"},"links":{"citing_paper":"/paper/2603.06576"},"observation_digest":"sha256:78d51d0410dbc77c6af9cdd9d44e6885dac6283b73e543b729efb6038e041541","observation_id":"18214693-320b-4ee2-ac00-12e2472bb377","resolution":{"observed_at":"2026-07-15T13:44:18.418529Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-07-15T13:44:18.418529Z","title":"In: IEEE/CVF Confer- ence on Computer Vision and Pattern Recognition (CVPR) (2024)","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2603.06576","last_updated":"2026-07-02T20:40:34Z","snapshot_observed_at":"2026-08-05T20:27:31.445015Z","submitted_at":"2026-03-06T18:59:55Z","title":"BEVLM: Distilling Semantic Knowledge from LLMs into Bird's-Eye View Representations","version":2},"reference_index":49,"source":"pdf_text","source_observed_at":"2026-07-15T13:44:18.418529Z"},"links":{"citing_paper":"/paper/2603.06576"},"observation_digest":"sha256:62a01708ca0e24e5710da90dbe413bd557c3ddd9bc94caa51754d550fe715c34","observation_id":"f685bb5b-882d-4d0c-af76-392e21352153","resolution":{"observed_at":"2026-07-15T13:44:18.418529Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-07-15T13:44:18.418529Z","title":"In: Annual Meeting of the Association for Computational Linguistics","venue":null,"work_id":null,"year":2002},"citing_paper":{"arxiv_id":"2603.06576","last_updated":"2026-07-02T20:40:34Z","snapshot_observed_at":"2026-08-05T20:27:31.445015Z","submitted_at":"2026-03-06T18:59:55Z","title":"BEVLM: Distilling Semantic Knowledge from LLMs into Bird's-Eye View Representations","version":2},"reference_index":50,"source":"pdf_text","source_observed_at":"2026-07-15T13:44:18.418529Z"},"links":{"citing_paper":"/paper/2603.06576"},"observation_digest":"sha256:29aa889fc7a05be4c4e33029f999223f1e48a3bd3ea68f351a8bb2ea1f2af895","observation_id":"6184c05c-6d70-4371-8b00-8bd40ddc7fe4","resolution":{"observed_at":"2026-07-15T13:44:18.418529Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-07-15T13:44:18.418529Z","title":"In: AAAI Conference on Artificial Intelligence","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2603.06576","last_updated":"2026-07-02T20:40:34Z","snapshot_observed_at":"2026-08-05T20:27:31.445015Z","submitted_at":"2026-03-06T18:59:55Z","title":"BEVLM: Distilling Semantic Knowledge from LLMs into Bird's-Eye View Representations","version":2},"reference_index":51,"source":"pdf_text","source_observed_at":"2026-07-15T13:44:18.418529Z"},"links":{"citing_paper":"/paper/2603.06576"},"observation_digest":"sha256:f69380534b13256775e3ab3f5c501ae6588c11473c3bd992e67b13baa12fead5","observation_id":"209ec093-4b10-4a8b-8a45-6bde5cc9de11","resolution":{"observed_at":"2026-07-15T13:44:18.418529Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-07-15T13:44:18.418529Z","title":"In: International Conference on Machine Learning (ICML)","venue":null,"work_id":null,"year":2021},"citing_paper":{"arxiv_id":"2603.06576","last_updated":"2026-07-02T20:40:34Z","snapshot_observed_at":"2026-08-05T20:27:31.445015Z","submitted_at":"2026-03-06T18:59:55Z","title":"BEVLM: Distilling Semantic Knowledge from LLMs into Bird's-Eye View Representations","version":2},"reference_index":52,"source":"pdf_text","source_observed_at":"2026-07-15T13:44:18.418529Z"},"links":{"citing_paper":"/paper/2603.06576"},"observation_digest":"sha256:245d48bb3d1a5982f6ea2541aa1c33cb832c15b2d9070ff27fbbd59542b05f33","observation_id":"9fde8331-0d60-43de-8117-979ea610bc4f","resolution":{"observed_at":"2026-07-15T13:44:18.418529Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-07-15T13:44:18.418529Z","title":"In: IEEE/CVF Conference on Computer Vision and Pattern Recognition (CVPR)","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2603.06576","last_updated":"2026-07-02T20:40:34Z","snapshot_observed_at":"2026-08-05T20:27:31.445015Z","submitted_at":"2026-03-06T18:59:55Z","title":"BEVLM: Distilling Semantic Knowledge from LLMs into Bird's-Eye View Representations","version":2},"reference_index":53,"source":"pdf_text","source_observed_at":"2026-07-15T13:44:18.418529Z"},"links":{"citing_paper":"/paper/2603.06576"},"observation_digest":"sha256:12ec519cf08cdeadf562650937cd866501af93882220d87e1b4a19d133d474b6","observation_id":"e580cfb1-1f3a-48f0-9923-d4eb0f2c55a7","resolution":{"observed_at":"2026-07-15T13:44:18.418529Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-07-15T13:44:18.418529Z","title":"In: 30th USENIX Security Symposium (USENIX Security) (2021)","venue":null,"work_id":null,"year":2021},"citing_paper":{"arxiv_id":"2603.06576","last_updated":"2026-07-02T20:40:34Z","snapshot_observed_at":"2026-08-05T20:27:31.445015Z","submitted_at":"2026-03-06T18:59:55Z","title":"BEVLM: Distilling Semantic Knowledge from LLMs into Bird's-Eye View Representations","version":2},"reference_index":54,"source":"pdf_text","source_observed_at":"2026-07-15T13:44:18.418529Z"},"links":{"citing_paper":"/paper/2603.06576"},"observation_digest":"sha256:9697f758f09c5fa40040cc294ff28496841dfb5a4f797b9c2994e5e9cc3d23af","observation_id":"e2d46682-8bff-4479-a1db-8b007606da0d","resolution":{"observed_at":"2026-07-15T13:44:18.418529Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-07-15T13:44:18.418529Z","title":"In: IEEE/CVF Conference on Computer Vision and Pattern Recognition (CVPR) (2024)","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2603.06576","last_updated":"2026-07-02T20:40:34Z","snapshot_observed_at":"2026-08-05T20:27:31.445015Z","submitted_at":"2026-03-06T18:59:55Z","title":"BEVLM: Distilling Semantic Knowledge from LLMs into Bird's-Eye View Representations","version":2},"reference_index":55,"source":"pdf_text","source_observed_at":"2026-07-15T13:44:18.418529Z"},"links":{"citing_paper":"/paper/2603.06576"},"observation_digest":"sha256:93c1edbbfa2128ab742d00e16f5e214b697cd94a674037462e85c621eb68c8f1","observation_id":"60e33ce4-0004-4eea-8e76-a9417ea935d4","resolution":{"observed_at":"2026-07-15T13:44:18.418529Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-07-15T13:44:18.418529Z","title":"In: European Conference on Computer Vision (ECCV)","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2603.06576","last_updated":"2026-07-02T20:40:34Z","snapshot_observed_at":"2026-08-05T20:27:31.445015Z","submitted_at":"2026-03-06T18:59:55Z","title":"BEVLM: Distilling Semantic Knowledge from LLMs into Bird's-Eye View Representations","version":2},"reference_index":56,"source":"pdf_text","source_observed_at":"2026-07-15T13:44:18.418529Z"},"links":{"citing_paper":"/paper/2603.06576"},"observation_digest":"sha256:19cd9f78bd5060f06b991e831eb8fbdacc8e14da5d101617d436a40cc60de992","observation_id":"0a70ffa3-460f-4b99-ba73-1327affd6432","resolution":{"observed_at":"2026-07-15T13:44:18.418529Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-07-15T13:44:18.418529Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2603.06576","last_updated":"2026-07-02T20:40:34Z","snapshot_observed_at":"2026-08-05T20:27:31.445015Z","submitted_at":"2026-03-06T18:59:55Z","title":"BEVLM: Distilling Semantic Knowledge from LLMs into Bird's-Eye View Representations","version":2},"reference_index":57,"source":"pdf_text","source_observed_at":"2026-07-15T13:44:18.418529Z"},"links":{"citing_paper":"/paper/2603.06576"},"observation_digest":"sha256:1f0497f4fe95e5b8a2b0d603547bb4476ca1cbd1bed3dc209ad88bf75ff7ac2b","observation_id":"3efe3f48-7259-48b2-b016-ff1a5c2c95b6","resolution":{"observed_at":"2026-07-15T13:44:18.418529Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-07-15T13:44:18.418529Z","title":"In: IEEE/CVF Conference on Computer Vision and Pattern Recognition (CVPR)","venue":null,"work_id":null,"year":2015},"citing_paper":{"arxiv_id":"2603.06576","last_updated":"2026-07-02T20:40:34Z","snapshot_observed_at":"2026-08-05T20:27:31.445015Z","submitted_at":"2026-03-06T18:59:55Z","title":"BEVLM: Distilling Semantic Knowledge from LLMs into Bird's-Eye View Representations","version":2},"reference_index":58,"source":"pdf_text","source_observed_at":"2026-07-15T13:44:18.418529Z"},"links":{"citing_paper":"/paper/2603.06576"},"observation_digest":"sha256:d95bbbb0000c2fd907b10e94de879e5a94a09aa5cb3d18acb411bac92ba6e79d","observation_id":"8d8ef046-2e20-446e-bc65-2f52537d47d0","resolution":{"observed_at":"2026-07-15T13:44:18.418529Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-07-15T13:44:18.418529Z","title":"In: ISOC Network and Distributed Systems Security (NDSS) Symposium (2022)","venue":null,"work_id":null,"year":2022},"citing_paper":{"arxiv_id":"2603.06576","last_updated":"2026-07-02T20:40:34Z","snapshot_observed_at":"2026-08-05T20:27:31.445015Z","submitted_at":"2026-03-06T18:59:55Z","title":"BEVLM: Distilling Semantic Knowledge from LLMs into Bird's-Eye View Representations","version":2},"reference_index":59,"source":"pdf_text","source_observed_at":"2026-07-15T13:44:18.418529Z"},"links":{"citing_paper":"/paper/2603.06576"},"observation_digest":"sha256:0ebffdd1329b9c82dc6a0ae39c4fc26f65d4f8d84549d47a44afd69bb8410483","observation_id":"92644c19-48de-41f6-b8b4-a0930881ef90","resolution":{"observed_at":"2026-07-15T13:44:18.418529Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-07-15T13:44:18.418529Z","title":"In: IEEE/CVF Conference on Computer Vision and Pattern Recognition (CVPR) (2025)","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2603.06576","last_updated":"2026-07-02T20:40:34Z","snapshot_observed_at":"2026-08-05T20:27:31.445015Z","submitted_at":"2026-03-06T18:59:55Z","title":"BEVLM: Distilling Semantic Knowledge from LLMs into Bird's-Eye View Representations","version":2},"reference_index":60,"source":"pdf_text","source_observed_at":"2026-07-15T13:44:18.418529Z"},"links":{"citing_paper":"/paper/2603.06576"},"observation_digest":"sha256:c401e66b5cc09ab7ed631d4aaaed10a31a18fdb835238e97f778693054ec76b2","observation_id":"94d7996a-c65b-4f2f-ad8f-3623e2fce067","resolution":{"observed_at":"2026-07-15T13:44:18.418529Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2511.00088","last_updated":"2026-01-07T09:09:57Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2025-10-30T01:25:34Z","title":"Alpamayo-R1: Bridging Reasoning and Action Prediction for Generalizable Autonomous Driving in the Long Tail","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2511.00088","snapshot_observed_at":"2026-07-15T13:44:18.418529Z","title":"arXiv preprint arXiv:2511.00088 (2025)","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2603.06576","last_updated":"2026-07-02T20:40:34Z","snapshot_observed_at":"2026-08-05T20:27:31.445015Z","submitted_at":"2026-03-06T18:59:55Z","title":"BEVLM: Distilling Semantic Knowledge from LLMs into Bird's-Eye View Representations","version":2},"reference_index":61,"source":"pdf_text","source_observed_at":"2026-07-15T13:44:18.418529Z"},"links":{"cited_paper":"/paper/2511.00088","citing_paper":"/paper/2603.06576"},"observation_digest":"sha256:54c74eef17f737a4572a9bf745a9ad7a6e8c4da664e83471c07e42a1c23d0725","observation_id":"745b5932-12a0-4cd3-8307-06c09c4179e4","resolution":{"observed_at":"2026-07-15T13:44:18.418529Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2503.03074","last_updated":"2025-03-05T00:27:32Z","snapshot_observed_at":"2026-08-07T17:29:01.212505Z","submitted_at":"2025-03-05T00:27:32Z","title":"BEVDriver: Leveraging BEV Maps in LLMs for Robust Closed-Loop Driving","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2503.03074","snapshot_observed_at":"2026-07-15T13:44:18.418529Z","title":"arXiv Preprint arXiv:2503.03074 (2025)","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2603.06576","last_updated":"2026-07-02T20:40:34Z","snapshot_observed_at":"2026-08-05T20:27:31.445015Z","submitted_at":"2026-03-06T18:59:55Z","title":"BEVLM: Distilling Semantic Knowledge from LLMs into Bird's-Eye View Representations","version":2},"reference_index":62,"source":"pdf_text","source_observed_at":"2026-07-15T13:44:18.418529Z"},"links":{"cited_paper":"/paper/2503.03074","citing_paper":"/paper/2603.06576"},"observation_digest":"sha256:e3a0fff1373e57672e92a4a4a58f8d2d16282a0257cd23ff59fe5e51509ca311","observation_id":"f051b78a-4fa9-4985-9c9a-12b45f5dbaf0","resolution":{"observed_at":"2026-07-15T13:44:18.418529Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-07-15T13:44:18.418529Z","title":"In: Proceedings of the IEEE/CVF Conference on Computer Vision and Pattern Recognition (CVPR) (2020)","venue":null,"work_id":null,"year":2020},"citing_paper":{"arxiv_id":"2603.06576","last_updated":"2026-07-02T20:40:34Z","snapshot_observed_at":"2026-08-05T20:27:31.445015Z","submitted_at":"2026-03-06T18:59:55Z","title":"BEVLM: Distilling Semantic Knowledge from LLMs into Bird's-Eye View Representations","version":2},"reference_index":63,"source":"pdf_text","source_observed_at":"2026-07-15T13:44:18.418529Z"},"links":{"citing_paper":"/paper/2603.06576"},"observation_digest":"sha256:a7a8de5d22b7933bc27111bef2853dc01d460166e22bf1f110edeea36d07f159","observation_id":"d689c7bd-2309-4b9c-b1f7-ee8e9b9fda94","resolution":{"observed_at":"2026-07-15T13:44:18.418529Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-07-15T13:44:18.418529Z","title":"In: IEEE/CVF International Conference on Computer Vision (ICCV) (October 2025)","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2603.06576","last_updated":"2026-07-02T20:40:34Z","snapshot_observed_at":"2026-08-05T20:27:31.445015Z","submitted_at":"2026-03-06T18:59:55Z","title":"BEVLM: Distilling Semantic Knowledge from LLMs into Bird's-Eye View Representations","version":2},"reference_index":64,"source":"pdf_text","source_observed_at":"2026-07-15T13:44:18.418529Z"},"links":{"citing_paper":"/paper/2603.06576"},"observation_digest":"sha256:1041fb7f96ad632956857a8c48e9de0f50ae2b41696826f00db71c39d0cf8bb2","observation_id":"75ed3432-b3a3-4b0d-b0a3-76bd89073db3","resolution":{"observed_at":"2026-07-15T13:44:18.418529Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":"2025.353596","doi":"10.1109/tpami.2025.3535960","metadata_source":"arxiv_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"IEEE Transactions on Pattern Analysis and Machine Intelligence47(5) (2025)","venue":"IEEE Transactions on Pattern Analysis and Machine Intelligence","work_id":"cadf9d84-ea34-47e3-8b43-c1af698d6ddf","year":2025},"citing_paper":{"arxiv_id":"2603.06576","last_updated":"2026-07-02T20:40:34Z","snapshot_observed_at":"2026-08-05T20:27:31.445015Z","submitted_at":"2026-03-06T18:59:55Z","title":"BEVLM: Distilling Semantic Knowledge from LLMs into Bird's-Eye View Representations","version":2},"reference_index":65,"source":"pdf_text","source_observed_at":"2026-07-15T13:44:18.418529Z"},"links":{"citing_paper":"/paper/2603.06576"},"observation_digest":"sha256:f0fc3f0e22083a873422932256a0d445cb0daf1e8900af818c713031f47081d4","observation_id":"f9ca44e5-25ec-4a3e-b05b-1dfd76669da6","resolution":{"observed_at":"2026-07-15T13:51:31.175506Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-07-15T13:44:18.418529Z","title":"Transactions on Machine Learning Research (2024)","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2603.06576","last_updated":"2026-07-02T20:40:34Z","snapshot_observed_at":"2026-08-05T20:27:31.445015Z","submitted_at":"2026-03-06T18:59:55Z","title":"BEVLM: Distilling Semantic Knowledge from LLMs into Bird's-Eye View Representations","version":2},"reference_index":66,"source":"pdf_text","source_observed_at":"2026-07-15T13:44:18.418529Z"},"links":{"citing_paper":"/paper/2603.06576"},"observation_digest":"sha256:9765e422f2d496d35a6cd26a1ae4ac4dbb01f0a30f75fc3c50b4bdae68336bb2","observation_id":"2fe2fbe3-d470-409b-9fe5-365e604c3464","resolution":{"observed_at":"2026-07-15T13:44:18.418529Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-07-15T13:44:18.418529Z","title":"In: IEEE/CVF Conference on Computer Vision and Pattern Recognition (CVPR) (2025)","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2603.06576","last_updated":"2026-07-02T20:40:34Z","snapshot_observed_at":"2026-08-05T20:27:31.445015Z","submitted_at":"2026-03-06T18:59:55Z","title":"BEVLM: Distilling Semantic Knowledge from LLMs into Bird's-Eye View Representations","version":2},"reference_index":67,"source":"pdf_text","source_observed_at":"2026-07-15T13:44:18.418529Z"},"links":{"citing_paper":"/paper/2603.06576"},"observation_digest":"sha256:09cfa714443c6accd90c907d911ad167cfc311a4b83894c63af7db8f68fec458","observation_id":"bcc27adc-4caf-4c5d-9715-0c2fe3cc5cd6","resolution":{"observed_at":"2026-07-15T13:44:18.418529Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2412.14446","last_updated":"2025-08-29T20:50:08Z","snapshot_observed_at":"2026-08-02T03:23:17.398540Z","submitted_at":"2024-12-19T01:53:36Z","title":"VLM-AD: End-to-End Autonomous Driving through Vision-Language Model Supervision","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2412.14446","snapshot_observed_at":"2026-07-15T13:44:18.418529Z","title":"arXiv preprint arXiv:2412.14446 (2024)","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2603.06576","last_updated":"2026-07-02T20:40:34Z","snapshot_observed_at":"2026-08-05T20:27:31.445015Z","submitted_at":"2026-03-06T18:59:55Z","title":"BEVLM: Distilling Semantic Knowledge from LLMs into Bird's-Eye View Representations","version":2},"reference_index":68,"source":"pdf_text","source_observed_at":"2026-07-15T13:44:18.418529Z"},"links":{"cited_paper":"/paper/2412.14446","citing_paper":"/paper/2603.06576"},"observation_digest":"sha256:42714d43deb0323b6a00f3982a97c1497c353f1a2d949d8f61af4edbeb61a90f","observation_id":"c9a306e1-a05b-49bc-80f8-3723aa4fa9cd","resolution":{"observed_at":"2026-07-15T13:44:18.418529Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-07-15T13:44:18.418529Z","title":"IEEE Robotics and Automation Letters9(10) (2024)","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2603.06576","last_updated":"2026-07-02T20:40:34Z","snapshot_observed_at":"2026-08-05T20:27:31.445015Z","submitted_at":"2026-03-06T18:59:55Z","title":"BEVLM: Distilling Semantic Knowledge from LLMs into Bird's-Eye View Representations","version":2},"reference_index":69,"source":"pdf_text","source_observed_at":"2026-07-15T13:44:18.418529Z"},"links":{"citing_paper":"/paper/2603.06576"},"observation_digest":"sha256:2ece123ee6547c3f5919fbb32a835e4a02f6cd046c6b0218557e6cde4cec6020","observation_id":"a50e79ad-3300-42e1-9316-bc3092a5e15f","resolution":{"observed_at":"2026-07-15T13:44:18.418529Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2505.09388","last_updated":"2025-05-14T13:41:34Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2025-05-14T13:41:34Z","title":"Qwen3 Technical Report","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2505.09388","snapshot_observed_at":"2026-07-15T13:44:18.418529Z","title":"arXiv Preprint arXiv:2505.09388 (2025)","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2603.06576","last_updated":"2026-07-02T20:40:34Z","snapshot_observed_at":"2026-08-05T20:27:31.445015Z","submitted_at":"2026-03-06T18:59:55Z","title":"BEVLM: Distilling Semantic Knowledge from LLMs into Bird's-Eye View Representations","version":2},"reference_index":70,"source":"pdf_text","source_observed_at":"2026-07-15T13:44:18.418529Z"},"links":{"cited_paper":"/paper/2505.09388","citing_paper":"/paper/2603.06576"},"observation_digest":"sha256:76b9434b6311b69c66e47287d787f9df112b6464c8eb762926c4498f01d4082c","observation_id":"85fb5430-4248-4cf6-89f1-6478fa2276d2","resolution":{"observed_at":"2026-07-15T13:44:18.418529Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2412.15115","last_updated":"2025-01-03T02:18:21Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2024-12-19T17:56:09Z","title":"Qwen2.5 Technical Report","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2412.15115","snapshot_observed_at":"2026-07-15T13:44:18.418529Z","title":"Monninger, S","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2603.06576","last_updated":"2026-07-02T20:40:34Z","snapshot_observed_at":"2026-08-05T20:27:31.445015Z","submitted_at":"2026-03-06T18:59:55Z","title":"BEVLM: Distilling Semantic Knowledge from LLMs into Bird's-Eye View Representations","version":2},"reference_index":71,"source":"pdf_text","source_observed_at":"2026-07-15T13:44:18.418529Z"},"links":{"cited_paper":"/paper/2412.15115","citing_paper":"/paper/2603.06576"},"observation_digest":"sha256:7e4828b7a6a98e52fd8acea60aa456cff5f1040bdad540b75c800255dfe3f1b8","observation_id":"610447a8-9e33-42c3-a42f-8156ee85cf47","resolution":{"observed_at":"2026-07-15T13:44:18.418529Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-07-15T13:44:18.418529Z","title":"In: Proceedings of the IEEE/CVF Conference on Computer Vision and Pattern Recognition (CVPR) (2024)","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2603.06576","last_updated":"2026-07-02T20:40:34Z","snapshot_observed_at":"2026-08-05T20:27:31.445015Z","submitted_at":"2026-03-06T18:59:55Z","title":"BEVLM: Distilling Semantic Knowledge from LLMs into Bird's-Eye View Representations","version":2},"reference_index":72,"source":"pdf_text","source_observed_at":"2026-07-15T13:44:18.418529Z"},"links":{"citing_paper":"/paper/2603.06576"},"observation_digest":"sha256:13ee038949efb985235ba9652bd0036dd2f5c44d3c849acd34be65b6997ac176","observation_id":"5535a450-f1f6-4da5-9bea-abdf1e17cdc6","resolution":{"observed_at":"2026-07-15T13:44:18.418529Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2311.01043","last_updated":"2024-08-12T11:53:28Z","snapshot_observed_at":"2026-07-06T16:42:04.139453Z","submitted_at":"2023-11-02T07:23:33Z","title":"LLM4Drive: A Survey of Large Language Models for Autonomous Driving","version":4},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2311.01043","snapshot_observed_at":"2026-07-15T13:44:18.418529Z","title":"arXiv preprint arXiv:2311.01043 (2023)","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2603.06576","last_updated":"2026-07-02T20:40:34Z","snapshot_observed_at":"2026-08-05T20:27:31.445015Z","submitted_at":"2026-03-06T18:59:55Z","title":"BEVLM: Distilling Semantic Knowledge from LLMs into Bird's-Eye View Representations","version":2},"reference_index":73,"source":"pdf_text","source_observed_at":"2026-07-15T13:44:18.418529Z"},"links":{"cited_paper":"/paper/2311.01043","citing_paper":"/paper/2603.06576"},"observation_digest":"sha256:c58e93860ad231cb0bd1a9c339c6aab419eaa78a4e311924b432cd945418ea54","observation_id":"ea210f01-f590-44bb-96cd-bdfad9ac8ce5","resolution":{"observed_at":"2026-07-15T13:44:18.418529Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2505.17685","last_updated":"2025-11-11T01:31:25Z","snapshot_observed_at":"2026-08-02T20:06:32.255780Z","submitted_at":"2025-05-23T09:55:32Z","title":"FutureSightDrive: Thinking Visually with Spatio-Temporal CoT for Autonomous Driving","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2505.17685","snapshot_observed_at":"2026-07-15T13:44:18.418529Z","title":"arXiv preprint arXiv:2505.17685 (2025)","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2603.06576","last_updated":"2026-07-02T20:40:34Z","snapshot_observed_at":"2026-08-05T20:27:31.445015Z","submitted_at":"2026-03-06T18:59:55Z","title":"BEVLM: Distilling Semantic Knowledge from LLMs into Bird's-Eye View Representations","version":2},"reference_index":74,"source":"pdf_text","source_observed_at":"2026-07-15T13:44:18.418529Z"},"links":{"cited_paper":"/paper/2505.17685","citing_paper":"/paper/2603.06576"},"observation_digest":"sha256:6e4a9f3423072a42409c4204678d7cfee8d327c8a6ba32784c8beb3822a15a21","observation_id":"1e626683-a5b1-4cc4-ae44-67df716b5270","resolution":{"observed_at":"2026-07-15T13:44:18.418529Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-07-15T13:44:18.418529Z","title":"In: IEEE/CVF International Conference on Computer Vision (ICCV) (2023)","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2603.06576","last_updated":"2026-07-02T20:40:34Z","snapshot_observed_at":"2026-08-05T20:27:31.445015Z","submitted_at":"2026-03-06T18:59:55Z","title":"BEVLM: Distilling Semantic Knowledge from LLMs into Bird's-Eye View Representations","version":2},"reference_index":75,"source":"pdf_text","source_observed_at":"2026-07-15T13:44:18.418529Z"},"links":{"citing_paper":"/paper/2603.06576"},"observation_digest":"sha256:11ad4f1631c050ab12b36e8960514df6506c8223e82965f63ad1e1662e0004de","observation_id":"524ed3bc-e4e1-4656-a7c9-793de1e5ce7b","resolution":{"observed_at":"2026-07-15T13:44:18.418529Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2205.09743","last_updated":"2022-05-19T17:55:35Z","snapshot_observed_at":"2026-08-04T03:03:57.465654Z","submitted_at":"2022-05-19T17:55:35Z","title":"BEVerse: Unified Perception and Prediction in Birds-Eye-View for Vision-Centric Autonomous Driving","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2205.09743","snapshot_observed_at":"2026-07-15T13:44:18.418529Z","title":"arXiv Preprint arXiv:2205.09743 (2022)","venue":null,"work_id":null,"year":2022},"citing_paper":{"arxiv_id":"2603.06576","last_updated":"2026-07-02T20:40:34Z","snapshot_observed_at":"2026-08-05T20:27:31.445015Z","submitted_at":"2026-03-06T18:59:55Z","title":"BEVLM: Distilling Semantic Knowledge from LLMs into Bird's-Eye View Representations","version":2},"reference_index":76,"source":"pdf_text","source_observed_at":"2026-07-15T13:44:18.418529Z"},"links":{"cited_paper":"/paper/2205.09743","citing_paper":"/paper/2603.06576"},"observation_digest":"sha256:60a41a92e7d9ef9b1b7ad1db3fac86261bc16917257810e3f95231683a1645b9","observation_id":"7ab4407d-9937-4744-a85a-f4cbcd7479cc","resolution":{"observed_at":"2026-07-15T13:44:18.418529Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-07-15T13:44:18.418529Z","title":"In: IEEE/CVF International Conference on Com- puter Vision (ICCV) (October 2025)","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2603.06576","last_updated":"2026-07-02T20:40:34Z","snapshot_observed_at":"2026-08-05T20:27:31.445015Z","submitted_at":"2026-03-06T18:59:55Z","title":"BEVLM: Distilling Semantic Knowledge from LLMs into Bird's-Eye View Representations","version":2},"reference_index":77,"source":"pdf_text","source_observed_at":"2026-07-15T13:44:18.418529Z"},"links":{"citing_paper":"/paper/2603.06576"},"observation_digest":"sha256:7937207891f7491123ce64217bd748a030b8b585f8dd4c41072684650214c0af","observation_id":"ad2fc8f1-e569-488b-8dea-4cb047bca263","resolution":{"observed_at":"2026-07-15T13:44:18.418529Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-07-15T13:44:18.418529Z","title":"arXiv preprint arXiv:2503.23463 (2025)","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2603.06576","last_updated":"2026-07-02T20:40:34Z","snapshot_observed_at":"2026-08-05T20:27:31.445015Z","submitted_at":"2026-03-06T18:59:55Z","title":"BEVLM: Distilling Semantic Knowledge from LLMs into Bird's-Eye View Representations","version":2},"reference_index":78,"source":"pdf_text","source_observed_at":"2026-07-15T13:44:18.418529Z"},"links":{"citing_paper":"/paper/2603.06576"},"observation_digest":"sha256:ec419da234f93408d26dfae365e75288875c802767cf54581ceb3c2c0fbc20cd","observation_id":"38e10452-f958-477a-a2dc-8c44729c3575","resolution":{"observed_at":"2026-07-15T13:44:18.418529Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2506.13757","last_updated":"2025-11-05T23:46:20Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2025-06-16T17:58:50Z","title":"AutoVLA: A Vision-Language-Action Model for End-to-End Autonomous Driving with Adaptive Reasoning and Reinforcement Fine-Tuning","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2506.13757","snapshot_observed_at":"2026-07-15T13:44:18.418529Z","title":"arXiv preprint arXiv:2506.13757 (2025)","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2603.06576","last_updated":"2026-07-02T20:40:34Z","snapshot_observed_at":"2026-08-05T20:27:31.445015Z","submitted_at":"2026-03-06T18:59:55Z","title":"BEVLM: Distilling Semantic Knowledge from LLMs into Bird's-Eye View Representations","version":2},"reference_index":79,"source":"pdf_text","source_observed_at":"2026-07-15T13:44:18.418529Z"},"links":{"cited_paper":"/paper/2506.13757","citing_paper":"/paper/2603.06576"},"observation_digest":"sha256:cbb74d09ef45496c4f4c67f2d78af99b16524a4cf9809f46b4074143e6c4eca5","observation_id":"1aa880eb-14da-4067-bd94-9ea01dcc86e9","resolution":{"observed_at":"2026-07-15T13:44:18.418529Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-07-15T13:44:18.418529Z","title":"In: The Twelfth International Conference on Learning Representations (ICLR) (2024)","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2603.06576","last_updated":"2026-07-02T20:40:34Z","snapshot_observed_at":"2026-08-05T20:27:31.445015Z","submitted_at":"2026-03-06T18:59:55Z","title":"BEVLM: Distilling Semantic Knowledge from LLMs into Bird's-Eye View Representations","version":2},"reference_index":80,"source":"pdf_text","source_observed_at":"2026-07-15T13:44:18.418529Z"},"links":{"citing_paper":"/paper/2603.06576"},"observation_digest":"sha256:161ed9b4c117c29a02095a37fec10f65620010ff68f008e7d1132fb3e1867cf5","observation_id":"4f2affcf-acde-413a-a9c1-c87e07e39e14","resolution":{"observed_at":"2026-07-15T13:44:18.418529Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2504.10479","last_updated":"2025-04-19T03:47:21Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2025-04-14T17:59:25Z","title":"InternVL3: Exploring Advanced Training and Test-Time Recipes for Open-Source Multimodal Models","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2504.10479","snapshot_observed_at":"2026-07-15T13:44:18.418529Z","title":"object-centric","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2603.06576","last_updated":"2026-07-02T20:40:34Z","snapshot_observed_at":"2026-08-05T20:27:31.445015Z","submitted_at":"2026-03-06T18:59:55Z","title":"BEVLM: Distilling Semantic Knowledge from LLMs into Bird's-Eye View Representations","version":2},"reference_index":81,"source":"pdf_text","source_observed_at":"2026-07-15T13:44:18.418529Z"},"links":{"cited_paper":"/paper/2504.10479","citing_paper":"/paper/2603.06576"},"observation_digest":"sha256:f7e903d818965dd2f72ec04e8669cf6802f35141595f3b69d8f1092411f14b88","observation_id":"46a24873-beef-42d2-9c3a-42a8364ef63a","resolution":{"observed_at":"2026-07-15T13:44:18.418529Z","resolver_source":null,"status":"malformed_identifier"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}}],"paper":{"arxiv_id":"2603.06576","last_updated":"2026-07-02T20:40:34Z","latest_version":2,"primary_category":"cs.CV","snapshot_observed_at":"2026-08-05T20:27:31.445015Z","submitted_at":"2026-03-06T18:59:55Z","title":"BEVLM: Distilling Semantic Knowledge from LLMs into Bird's-Eye View Representations"},"reference_resolution":{"displayed":81,"state_counts":{"malformed_identifier":1,"metadata_mismatch":0,"parse_uncertain":0,"unresolved":79,"verified_exact":1,"verified_fuzzy":0},"total_outbound_references":81},"refusal":"A citation records a reference. It does not transfer a finding from one paper to another.","schema":"pith.paper-citation-record.v1","standing_sources":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"thesis":"As of 8 August 2026, this Paper Citation Record lists 81 of 81 outbound references and 0 inbound Pith citation observations for arXiv:2603.06576."}