{"bundle_type":"pith_open_graph_bundle","bundle_version":"1.0","pith_number":"pith:2025:T3TUDI2XSVEYGMQMT36Q7HUBCP","short_pith_number":"pith:T3TUDI2X","canonical_record":{"source":{"id":"2508.12404","kind":"arxiv","version":1},"metadata":{"license":"http://creativecommons.org/licenses/by/4.0/","primary_cat":"cs.CV","submitted_at":"2025-08-17T15:42:54Z","cross_cats_sorted":[],"title_canon_sha256":"fbdcc685ca29538d13ba1d6bb1dcae96f8a64e95d0c652de9f3f6f7138aa4063","abstract_canon_sha256":"c71c0f96e1b6afacdf5d6c538ec2b637dd390d37d17b6191b11041f71da6d5d2"},"schema_version":"1.0"},"canonical_sha256":"9ee741a357954983320c9efd0f9e8113f0f15f2d8fdf8124369ce9570d1fa06d","source":{"kind":"arxiv","id":"2508.12404","version":1},"source_aliases":[{"alias_kind":"arxiv","alias_value":"2508.12404","created_at":"2026-07-05T11:55:05Z"},{"alias_kind":"arxiv_version","alias_value":"2508.12404v1","created_at":"2026-07-05T11:55:05Z"},{"alias_kind":"doi","alias_value":"10.48550/arxiv.2508.12404","created_at":"2026-07-05T11:55:05Z"},{"alias_kind":"pith_short_12","alias_value":"T3TUDI2XSVEY","created_at":"2026-07-05T11:55:05Z"},{"alias_kind":"pith_short_16","alias_value":"T3TUDI2XSVEYGMQM","created_at":"2026-07-05T11:55:05Z"},{"alias_kind":"pith_short_8","alias_value":"T3TUDI2X","created_at":"2026-07-05T11:55:05Z"}],"events":[{"event_type":"record_created","subject_pith_number":"pith:2025:T3TUDI2XSVEYGMQMT36Q7HUBCP","target":"record","payload":{"canonical_record":{"source":{"id":"2508.12404","kind":"arxiv","version":1},"metadata":{"license":"http://creativecommons.org/licenses/by/4.0/","primary_cat":"cs.CV","submitted_at":"2025-08-17T15:42:54Z","cross_cats_sorted":[],"title_canon_sha256":"fbdcc685ca29538d13ba1d6bb1dcae96f8a64e95d0c652de9f3f6f7138aa4063","abstract_canon_sha256":"c71c0f96e1b6afacdf5d6c538ec2b637dd390d37d17b6191b11041f71da6d5d2"},"schema_version":"1.0"},"canonical_sha256":"9ee741a357954983320c9efd0f9e8113f0f15f2d8fdf8124369ce9570d1fa06d","receipt":{"kind":"pith_receipt","key_id":"pith-v1-2026-05","algorithm":"ed25519","signed_at":"2026-07-05T11:55:05.157525Z","signature_b64":"l1J2kBEyC7dpIUjZ41JSCEj1RnP8Ot8ZxSvef7r+ee7ecSRWT125wtpjoBTobkUNukNAvAzmepU36qcKlkVPCg==","signed_message":"canonical_sha256_bytes","builder_version":"pith-number-builder-2026-05-17-v1","receipt_version":"0.3","canonical_sha256":"9ee741a357954983320c9efd0f9e8113f0f15f2d8fdf8124369ce9570d1fa06d","last_reissued_at":"2026-07-05T11:55:05.157165Z","signature_status":"signed_v1","first_computed_at":"2026-07-05T11:55:05.157165Z","public_key_fingerprint":"8d4b5ee74e4693bcd1df2446408b0d54"},"source_kind":"arxiv","source_id":"2508.12404","source_version":1,"attestation_state":"computed"},"signer":{"signer_id":"pith.science","signer_type":"pith_registry","key_id":"pith-v1-2026-05","public_key_fingerprint":"8d4b5ee74e4693bcd1df2446408b0d54"},"created_at":"2026-07-05T11:55:05Z","supersedes":[],"prev_event":null,"signature":{"signature_status":"signed_v1","algorithm":"ed25519","key_id":"pith-v1-2026-05","public_key_fingerprint":"8d4b5ee74e4693bcd1df2446408b0d54","signature_b64":"5vLJsEGf8Lp0c//AurEpkSgV/JoGDAwJrroR0VIm+YVa3GoNdYzW/Cts3BsamrxuIowLv8YSqt20YjFH38QdDQ==","signed_message":"open_graph_event_sha256_bytes","signed_at":"2026-08-10T08:06:17.393940Z"},"content_sha256":"886223460754d794ef11215c095ec09a0494de7633f8dd0e5948715d8ef306cf","schema_version":"1.0","event_id":"sha256:886223460754d794ef11215c095ec09a0494de7633f8dd0e5948715d8ef306cf"},{"event_type":"graph_snapshot","subject_pith_number":"pith:2025:T3TUDI2XSVEYGMQMT36Q7HUBCP","target":"graph","payload":{"graph_snapshot":{"paper":{"title":"LMAD: Integrated End-to-End Vision-Language Model for Explainable Autonomous Driving","license":"http://creativecommons.org/licenses/by/4.0/","headline":"","cross_cats":[],"primary_cat":"cs.CV","authors_text":"Bozhou Zhang, Jiankang Deng, Li Zhang, Nan Song, Xiatian Zhu","submitted_at":"2025-08-17T15:42:54Z","abstract_excerpt":"Large vision-language models (VLMs) have shown promising capabilities in scene understanding, enhancing the explainability of driving behaviors and interactivity with users. Existing methods primarily fine-tune VLMs on on-board multi-view images and scene reasoning text, but this approach often lacks the holistic and nuanced scene recognition and powerful spatial awareness required for autonomous driving, especially in complex situations. To address this gap, we propose a novel vision-language framework tailored for autonomous driving, called LMAD. Our framework emulates modern end-to-end driv"},"claims":{"count":0,"items":[],"snapshot_sha256":"258153158e38e3291e3d48162225fcdb2d5a3ed65a07baac614ab91432fd4f57"},"source":{"id":"2508.12404","kind":"arxiv","version":1},"verdict":{"id":null,"model_set":{},"created_at":null,"strongest_claim":"","one_line_summary":"","pipeline_version":null,"weakest_assumption":"","pith_extraction_headline":""},"integrity":{"clean":true,"summary":{"advisory":0,"critical":0,"by_detector":{},"informational":0},"endpoint":"/pith/2508.12404/integrity.json","findings":[],"available":true,"detectors_run":[],"snapshot_sha256":"c28c3603d3b5d939e8dc4c7e95fa8dfce3d595e45f758748cecf8e644a296938"},"references":{"count":0,"sample":[],"resolved_work":0,"snapshot_sha256":"258153158e38e3291e3d48162225fcdb2d5a3ed65a07baac614ab91432fd4f57","internal_anchors":0},"formal_canon":{"evidence_count":0,"snapshot_sha256":"258153158e38e3291e3d48162225fcdb2d5a3ed65a07baac614ab91432fd4f57"},"author_claims":{"count":0,"strong_count":0,"snapshot_sha256":"258153158e38e3291e3d48162225fcdb2d5a3ed65a07baac614ab91432fd4f57"},"builder_version":"pith-number-builder-2026-05-17-v1"},"verdict_id":null},"signer":{"signer_id":"pith.science","signer_type":"pith_registry","key_id":"pith-v1-2026-05","public_key_fingerprint":"8d4b5ee74e4693bcd1df2446408b0d54"},"created_at":"2026-07-05T11:55:05Z","supersedes":[],"prev_event":null,"signature":{"signature_status":"signed_v1","algorithm":"ed25519","key_id":"pith-v1-2026-05","public_key_fingerprint":"8d4b5ee74e4693bcd1df2446408b0d54","signature_b64":"BBI1qC/Mt0GtTfyy9ASz84QSLMPeKZC847KJ0Fk6Aj0TR5CFqpEyu4fcmog55hNY9Ms+J9N5903xHREDETnXBg==","signed_message":"open_graph_event_sha256_bytes","signed_at":"2026-08-10T08:06:17.394918Z"},"content_sha256":"541ec6d272c9c6e25f3d2419be7215a907983040bbee727a94614114fe307f97","schema_version":"1.0","event_id":"sha256:541ec6d272c9c6e25f3d2419be7215a907983040bbee727a94614114fe307f97"}],"timestamp_proofs":[],"mirror_hints":[{"mirror_type":"https","name":"Pith Resolver","base_url":"https://pith.science","bundle_url":"https://pith.science/pith/T3TUDI2XSVEYGMQMT36Q7HUBCP/bundle.json","state_url":"https://pith.science/pith/T3TUDI2XSVEYGMQMT36Q7HUBCP/state.json","well_known_bundle_url":"https://pith.science/.well-known/pith/T3TUDI2XSVEYGMQMT36Q7HUBCP/bundle.json","status":"primary"}],"public_keys":[{"key_id":"pith-v1-2026-05","algorithm":"ed25519","format":"raw","public_key_b64":"stVStoiQhXFxp4s2pdzPNoqVNBMojDU/fJ2db5S3CbM=","public_key_hex":"b2d552b68890857171a78b36a5dccf368a953413288c353f7c9d9d6f94b709b3","fingerprint_sha256_b32_first128bits":"RVFV5Z2OI2J3ZUO7ERDEBCYNKS","fingerprint_sha256_hex":"8d4b5ee74e4693bcd1df2446408b0d54","rotates_at":null,"url":"https://pith.science/pith-signing-key.json","notes":"Pith uses this Ed25519 key to sign canonical record SHA-256 digests. Verify with: ed25519_verify(public_key, message=canonical_sha256_bytes, signature=base64decode(signature_b64))."}],"merge_version":"pith-open-graph-merge-v1","built_at":"2026-08-10T08:06:17Z","links":{"resolver":"https://pith.science/pith/T3TUDI2XSVEYGMQMT36Q7HUBCP","bundle":"https://pith.science/pith/T3TUDI2XSVEYGMQMT36Q7HUBCP/bundle.json","state":"https://pith.science/pith/T3TUDI2XSVEYGMQMT36Q7HUBCP/state.json","well_known_bundle":"https://pith.science/.well-known/pith/T3TUDI2XSVEYGMQMT36Q7HUBCP/bundle.json"},"state":{"state_type":"pith_open_graph_state","state_version":"1.0","pith_number":"pith:2025:T3TUDI2XSVEYGMQMT36Q7HUBCP","merge_version":"pith-open-graph-merge-v1","event_count":2,"valid_event_count":2,"invalid_event_count":0,"equivocation_count":0,"current":{"canonical_record":{"metadata":{"abstract_canon_sha256":"c71c0f96e1b6afacdf5d6c538ec2b637dd390d37d17b6191b11041f71da6d5d2","cross_cats_sorted":[],"license":"http://creativecommons.org/licenses/by/4.0/","primary_cat":"cs.CV","submitted_at":"2025-08-17T15:42:54Z","title_canon_sha256":"fbdcc685ca29538d13ba1d6bb1dcae96f8a64e95d0c652de9f3f6f7138aa4063"},"schema_version":"1.0","source":{"id":"2508.12404","kind":"arxiv","version":1}},"source_aliases":[{"alias_kind":"arxiv","alias_value":"2508.12404","created_at":"2026-07-05T11:55:05Z"},{"alias_kind":"arxiv_version","alias_value":"2508.12404v1","created_at":"2026-07-05T11:55:05Z"},{"alias_kind":"doi","alias_value":"10.48550/arxiv.2508.12404","created_at":"2026-07-05T11:55:05Z"},{"alias_kind":"pith_short_12","alias_value":"T3TUDI2XSVEY","created_at":"2026-07-05T11:55:05Z"},{"alias_kind":"pith_short_16","alias_value":"T3TUDI2XSVEYGMQM","created_at":"2026-07-05T11:55:05Z"},{"alias_kind":"pith_short_8","alias_value":"T3TUDI2X","created_at":"2026-07-05T11:55:05Z"}],"graph_snapshots":[{"event_id":"sha256:541ec6d272c9c6e25f3d2419be7215a907983040bbee727a94614114fe307f97","target":"graph","created_at":"2026-07-05T11:55:05Z","signer":{"key_id":"pith-v1-2026-05","public_key_fingerprint":"8d4b5ee74e4693bcd1df2446408b0d54","signer_id":"pith.science","signer_type":"pith_registry"},"payload":{"graph_snapshot":{"author_claims":{"count":0,"snapshot_sha256":"258153158e38e3291e3d48162225fcdb2d5a3ed65a07baac614ab91432fd4f57","strong_count":0},"builder_version":"pith-number-builder-2026-05-17-v1","claims":{"count":0,"items":[],"snapshot_sha256":"258153158e38e3291e3d48162225fcdb2d5a3ed65a07baac614ab91432fd4f57"},"formal_canon":{"evidence_count":0,"snapshot_sha256":"258153158e38e3291e3d48162225fcdb2d5a3ed65a07baac614ab91432fd4f57"},"integrity":{"available":true,"clean":true,"detectors_run":[],"endpoint":"/pith/2508.12404/integrity.json","findings":[],"snapshot_sha256":"c28c3603d3b5d939e8dc4c7e95fa8dfce3d595e45f758748cecf8e644a296938","summary":{"advisory":0,"by_detector":{},"critical":0,"informational":0}},"paper":{"abstract_excerpt":"Large vision-language models (VLMs) have shown promising capabilities in scene understanding, enhancing the explainability of driving behaviors and interactivity with users. Existing methods primarily fine-tune VLMs on on-board multi-view images and scene reasoning text, but this approach often lacks the holistic and nuanced scene recognition and powerful spatial awareness required for autonomous driving, especially in complex situations. To address this gap, we propose a novel vision-language framework tailored for autonomous driving, called LMAD. Our framework emulates modern end-to-end driv","authors_text":"Bozhou Zhang, Jiankang Deng, Li Zhang, Nan Song, Xiatian Zhu","cross_cats":[],"headline":"","license":"http://creativecommons.org/licenses/by/4.0/","primary_cat":"cs.CV","submitted_at":"2025-08-17T15:42:54Z","title":"LMAD: Integrated End-to-End Vision-Language Model for Explainable Autonomous Driving"},"references":{"count":0,"internal_anchors":0,"resolved_work":0,"sample":[],"snapshot_sha256":"258153158e38e3291e3d48162225fcdb2d5a3ed65a07baac614ab91432fd4f57"},"source":{"id":"2508.12404","kind":"arxiv","version":1},"verdict":{"created_at":null,"id":null,"model_set":{},"one_line_summary":"","pipeline_version":null,"pith_extraction_headline":"","strongest_claim":"","weakest_assumption":""}},"verdict_id":null}}],"author_attestations":[],"timestamp_anchors":[],"storage_attestations":[],"citation_signatures":[],"replication_records":[],"corrections":[],"mirror_hints":[],"record_created":{"event_id":"sha256:886223460754d794ef11215c095ec09a0494de7633f8dd0e5948715d8ef306cf","target":"record","created_at":"2026-07-05T11:55:05Z","signer":{"key_id":"pith-v1-2026-05","public_key_fingerprint":"8d4b5ee74e4693bcd1df2446408b0d54","signer_id":"pith.science","signer_type":"pith_registry"},"payload":{"attestation_state":"computed","canonical_record":{"metadata":{"abstract_canon_sha256":"c71c0f96e1b6afacdf5d6c538ec2b637dd390d37d17b6191b11041f71da6d5d2","cross_cats_sorted":[],"license":"http://creativecommons.org/licenses/by/4.0/","primary_cat":"cs.CV","submitted_at":"2025-08-17T15:42:54Z","title_canon_sha256":"fbdcc685ca29538d13ba1d6bb1dcae96f8a64e95d0c652de9f3f6f7138aa4063"},"schema_version":"1.0","source":{"id":"2508.12404","kind":"arxiv","version":1}},"canonical_sha256":"9ee741a357954983320c9efd0f9e8113f0f15f2d8fdf8124369ce9570d1fa06d","receipt":{"algorithm":"ed25519","builder_version":"pith-number-builder-2026-05-17-v1","canonical_sha256":"9ee741a357954983320c9efd0f9e8113f0f15f2d8fdf8124369ce9570d1fa06d","first_computed_at":"2026-07-05T11:55:05.157165Z","key_id":"pith-v1-2026-05","kind":"pith_receipt","last_reissued_at":"2026-07-05T11:55:05.157165Z","public_key_fingerprint":"8d4b5ee74e4693bcd1df2446408b0d54","receipt_version":"0.3","signature_b64":"l1J2kBEyC7dpIUjZ41JSCEj1RnP8Ot8ZxSvef7r+ee7ecSRWT125wtpjoBTobkUNukNAvAzmepU36qcKlkVPCg==","signature_status":"signed_v1","signed_at":"2026-07-05T11:55:05.157525Z","signed_message":"canonical_sha256_bytes"},"source_id":"2508.12404","source_kind":"arxiv","source_version":1}}},"equivocations":[],"invalid_events":[],"applied_event_ids":["sha256:886223460754d794ef11215c095ec09a0494de7633f8dd0e5948715d8ef306cf","sha256:541ec6d272c9c6e25f3d2419be7215a907983040bbee727a94614114fe307f97"],"state_sha256":"45d8b12aa2f7e19c41e421f45c7c88dc0573330eeb369e6306d7fc0616bfefbc"},"bundle_signature":{"signature_status":"signed_v1","algorithm":"ed25519","key_id":"pith-v1-2026-05","public_key_fingerprint":"8d4b5ee74e4693bcd1df2446408b0d54","signature_b64":"wmIv5AZVWfBRi5gqWSQj67jXv9cFQcQj6Dr71SnOHNt4dvywEBEHF3W+1B7Kh8ysk+0Je3TR6itpC+WqfsLbAQ==","signed_message":"bundle_sha256_bytes","signed_at":"2026-08-10T08:06:17.401517Z","bundle_sha256":"421fe2b6120b41f4af4d563d9368965d7e57976f4986f3446eb117c417c69aed"}}