{"record_type":"pith_number_record","schema_url":"https://pith.science/schemas/pith-number/v1.json","pith_number":"pith:2025:CMX4TGQTRB73HRU64CJASIWGNP","short_pith_number":"pith:CMX4TGQT","schema_version":"1.0","canonical_sha256":"132fc99a13887fb3c69ee0920922c66bde9ff09c0c6e5ad4705082e0a4a3d5ca","source":{"kind":"arxiv","id":"2502.07855","version":2},"attestation_state":"computed","paper":{"title":"Vision-Language Models for Edge Networks: A Comprehensive Survey","license":"http://creativecommons.org/licenses/by/4.0/","headline":"","cross_cats":["cs.AI","cs.CL"],"primary_cat":"cs.CV","authors_text":"Ahmed Sharshar, Latif U. Khan, Mohsen Guizani, Waseem Ullah","submitted_at":"2025-02-11T14:04:43Z","abstract_excerpt":"Vision Large Language Models (VLMs) combine visual understanding with natural language processing, enabling tasks like image captioning, visual question answering, and video analysis. While VLMs show impressive capabilities across domains such as autonomous vehicles, smart surveillance, and healthcare, their deployment on resource-constrained edge devices remains challenging due to processing power, memory, and energy limitations. This survey explores recent advancements in optimizing VLMs for edge environments, focusing on model compression techniques, including pruning, quantization, knowled"},"verification_status":{"content_addressed":true,"pith_receipt":true,"author_attested":false,"weak_author_claims":0,"strong_author_claims":0,"externally_anchored":false,"storage_verified":false,"citation_signatures":0,"replication_records":0,"graph_snapshot":true,"references_resolved":false,"formal_links_present":false},"canonical_record":{"source":{"id":"2502.07855","kind":"arxiv","version":2},"metadata":{"license":"http://creativecommons.org/licenses/by/4.0/","primary_cat":"cs.CV","submitted_at":"2025-02-11T14:04:43Z","cross_cats_sorted":["cs.AI","cs.CL"],"title_canon_sha256":"b46e86eb62f4f7bdd4768033b0be86141f3f9b5df10de632f8eb3d808826f427","abstract_canon_sha256":"be26c9bf00746be534e7779fbf56f183bf0a4929123e6e1d0c5ae49cdeccf960"},"schema_version":"1.0"},"receipt":{"kind":"pith_receipt","key_id":"pith-v1-2026-05","algorithm":"ed25519","signed_at":"2026-07-05T11:22:36.527568Z","signature_b64":"b/j1RiPd9PQoxE2tlBGMBz3bjB448tJBTPy3uA8ko36faktP8SWoXPELoqea/TRCuYIDwg8FxzpzvwRidn16AA==","signed_message":"canonical_sha256_bytes","builder_version":"pith-number-builder-2026-05-17-v1","receipt_version":"0.3","canonical_sha256":"132fc99a13887fb3c69ee0920922c66bde9ff09c0c6e5ad4705082e0a4a3d5ca","last_reissued_at":"2026-07-05T11:22:36.526869Z","signature_status":"signed_v1","first_computed_at":"2026-07-05T11:22:36.526869Z","public_key_fingerprint":"8d4b5ee74e4693bcd1df2446408b0d54"},"graph_snapshot":{"paper":{"title":"Vision-Language Models for Edge Networks: A Comprehensive Survey","license":"http://creativecommons.org/licenses/by/4.0/","headline":"","cross_cats":["cs.AI","cs.CL"],"primary_cat":"cs.CV","authors_text":"Ahmed Sharshar, Latif U. Khan, Mohsen Guizani, Waseem Ullah","submitted_at":"2025-02-11T14:04:43Z","abstract_excerpt":"Vision Large Language Models (VLMs) combine visual understanding with natural language processing, enabling tasks like image captioning, visual question answering, and video analysis. While VLMs show impressive capabilities across domains such as autonomous vehicles, smart surveillance, and healthcare, their deployment on resource-constrained edge devices remains challenging due to processing power, memory, and energy limitations. This survey explores recent advancements in optimizing VLMs for edge environments, focusing on model compression techniques, including pruning, quantization, knowled"},"claims":{"count":0,"items":[],"snapshot_sha256":"258153158e38e3291e3d48162225fcdb2d5a3ed65a07baac614ab91432fd4f57"},"source":{"id":"2502.07855","kind":"arxiv","version":2},"verdict":{"id":null,"model_set":{},"created_at":null,"strongest_claim":"","one_line_summary":"","pipeline_version":null,"weakest_assumption":"","pith_extraction_headline":""},"integrity":{"clean":true,"summary":{"advisory":0,"critical":0,"by_detector":{},"informational":0},"endpoint":"/pith/2502.07855/integrity.json","findings":[],"available":true,"detectors_run":[],"snapshot_sha256":"c28c3603d3b5d939e8dc4c7e95fa8dfce3d595e45f758748cecf8e644a296938"},"references":{"count":0,"sample":[],"resolved_work":0,"snapshot_sha256":"258153158e38e3291e3d48162225fcdb2d5a3ed65a07baac614ab91432fd4f57","internal_anchors":0},"formal_canon":{"evidence_count":0,"snapshot_sha256":"258153158e38e3291e3d48162225fcdb2d5a3ed65a07baac614ab91432fd4f57"},"author_claims":{"count":0,"strong_count":0,"snapshot_sha256":"258153158e38e3291e3d48162225fcdb2d5a3ed65a07baac614ab91432fd4f57"},"builder_version":"pith-number-builder-2026-05-17-v1"},"aliases":[{"alias_kind":"arxiv","alias_value":"2502.07855","created_at":"2026-07-05T11:22:36.526958+00:00"},{"alias_kind":"arxiv_version","alias_value":"2502.07855v2","created_at":"2026-07-05T11:22:36.526958+00:00"},{"alias_kind":"doi","alias_value":"10.48550/arxiv.2502.07855","created_at":"2026-07-05T11:22:36.526958+00:00"},{"alias_kind":"pith_short_12","alias_value":"CMX4TGQTRB73","created_at":"2026-07-05T11:22:36.526958+00:00"},{"alias_kind":"pith_short_16","alias_value":"CMX4TGQTRB73HRU6","created_at":"2026-07-05T11:22:36.526958+00:00"},{"alias_kind":"pith_short_8","alias_value":"CMX4TGQT","created_at":"2026-07-05T11:22:36.526958+00:00"}],"events":[],"event_summary":{},"paper_claims":[],"inbound_citations":{"count":1,"internal_anchor_count":0,"sample":[{"citing_arxiv_id":"2505.06907","citing_title":"A Survey on Foundation Models for Personalized Federated Intelligence","ref_index":182,"is_internal_anchor":false}]},"formal_canon":{"evidence_count":0,"sample":[],"anchors":[]},"links":{"html":"https://pith.science/pith/CMX4TGQTRB73HRU64CJASIWGNP","json":"https://pith.science/pith/CMX4TGQTRB73HRU64CJASIWGNP.json","graph_json":"https://pith.science/api/pith-number/CMX4TGQTRB73HRU64CJASIWGNP/graph.json","events_json":"https://pith.science/api/pith-number/CMX4TGQTRB73HRU64CJASIWGNP/events.json","paper":"https://pith.science/paper/CMX4TGQT"},"agent_actions":{"view_html":"https://pith.science/pith/CMX4TGQTRB73HRU64CJASIWGNP","download_json":"https://pith.science/pith/CMX4TGQTRB73HRU64CJASIWGNP.json","view_paper":"https://pith.science/paper/CMX4TGQT","resolve_alias":"https://pith.science/api/pith-number/resolve?arxiv=2502.07855&json=true","fetch_graph":"https://pith.science/api/pith-number/CMX4TGQTRB73HRU64CJASIWGNP/graph.json","fetch_events":"https://pith.science/api/pith-number/CMX4TGQTRB73HRU64CJASIWGNP/events.json","actions":{"anchor_timestamp":"https://pith.science/pith/CMX4TGQTRB73HRU64CJASIWGNP/action/timestamp_anchor","attest_storage":"https://pith.science/pith/CMX4TGQTRB73HRU64CJASIWGNP/action/storage_attestation","attest_author":"https://pith.science/pith/CMX4TGQTRB73HRU64CJASIWGNP/action/author_attestation","sign_citation":"https://pith.science/pith/CMX4TGQTRB73HRU64CJASIWGNP/action/citation_signature","submit_replication":"https://pith.science/pith/CMX4TGQTRB73HRU64CJASIWGNP/action/replication_record"}},"created_at":"2026-07-05T11:22:36.526958+00:00","updated_at":"2026-07-05T11:22:36.526958+00:00"}