{"record_type":"pith_number_record","schema_url":"https://pith.science/schemas/pith-number/v1.json","pith_number":"pith:2024:XZQ6GVCIDGIBL6XM2IJJ4ZBGDT","short_pith_number":"pith:XZQ6GVCI","schema_version":"1.0","canonical_sha256":"be61e35448199015faecd2129e64261cc111747846e1515603bd816f7adfb147","source":{"kind":"arxiv","id":"2408.01228","version":2},"attestation_state":"computed","paper":{"title":"The Phantom Menace: Unmasking Privacy Leakages in Vision-Language Models","license":"http://creativecommons.org/licenses/by-nc-sa/4.0/","headline":"","cross_cats":[],"primary_cat":"cs.CV","authors_text":"Elisa Ricci, Massimiliano Mancini, Rahaf Aljundi, Simone Caldarella","submitted_at":"2024-08-02T12:36:13Z","abstract_excerpt":"Vision-Language Models (VLMs) combine visual and textual understanding, rendering them well-suited for diverse tasks like generating image captions and answering visual questions across various domains. However, these capabilities are built upon training on large amount of uncurated data crawled from the web. The latter may include sensitive information that VLMs could memorize and leak, raising significant privacy concerns. In this paper, we assess whether these vulnerabilities exist, focusing on identity leakage. Our study leads to three key findings: (i) VLMs leak identity information, even"},"verification_status":{"content_addressed":true,"pith_receipt":true,"author_attested":false,"weak_author_claims":0,"strong_author_claims":0,"externally_anchored":false,"storage_verified":false,"citation_signatures":0,"replication_records":0,"graph_snapshot":true,"references_resolved":false,"formal_links_present":false},"canonical_record":{"source":{"id":"2408.01228","kind":"arxiv","version":2},"metadata":{"license":"http://creativecommons.org/licenses/by-nc-sa/4.0/","primary_cat":"cs.CV","submitted_at":"2024-08-02T12:36:13Z","cross_cats_sorted":[],"title_canon_sha256":"772db543dc1b00f7e4205f3304f6166cce60425227cca90c01866056a965f7a6","abstract_canon_sha256":"6de49261e10ef1f1e7a57ded0788dc4507a643e432b28bdc8bd1cb240cf37982"},"schema_version":"1.0"},"receipt":{"kind":"pith_receipt","key_id":"pith-v1-2026-05","algorithm":"ed25519","signed_at":"2026-07-05T08:56:48.328279Z","signature_b64":"5xvTc0bUqhCFG3DO07vqMX9oKmHzEvUtjNsZ+YAYbnWMBZuLB85IBKPAwWKqsMLUL1RJzTjkP7IOWHapkrIFCA==","signed_message":"canonical_sha256_bytes","builder_version":"pith-number-builder-2026-05-17-v1","receipt_version":"0.3","canonical_sha256":"be61e35448199015faecd2129e64261cc111747846e1515603bd816f7adfb147","last_reissued_at":"2026-07-05T08:56:48.327687Z","signature_status":"signed_v1","first_computed_at":"2026-07-05T08:56:48.327687Z","public_key_fingerprint":"8d4b5ee74e4693bcd1df2446408b0d54"},"graph_snapshot":{"paper":{"title":"The Phantom Menace: Unmasking Privacy Leakages in Vision-Language Models","license":"http://creativecommons.org/licenses/by-nc-sa/4.0/","headline":"","cross_cats":[],"primary_cat":"cs.CV","authors_text":"Elisa Ricci, Massimiliano Mancini, Rahaf Aljundi, Simone Caldarella","submitted_at":"2024-08-02T12:36:13Z","abstract_excerpt":"Vision-Language Models (VLMs) combine visual and textual understanding, rendering them well-suited for diverse tasks like generating image captions and answering visual questions across various domains. However, these capabilities are built upon training on large amount of uncurated data crawled from the web. The latter may include sensitive information that VLMs could memorize and leak, raising significant privacy concerns. In this paper, we assess whether these vulnerabilities exist, focusing on identity leakage. Our study leads to three key findings: (i) VLMs leak identity information, even"},"claims":{"count":0,"items":[],"snapshot_sha256":"258153158e38e3291e3d48162225fcdb2d5a3ed65a07baac614ab91432fd4f57"},"source":{"id":"2408.01228","kind":"arxiv","version":2},"verdict":{"id":null,"model_set":{},"created_at":null,"strongest_claim":"","one_line_summary":"","pipeline_version":null,"weakest_assumption":"","pith_extraction_headline":""},"integrity":{"clean":true,"summary":{"advisory":0,"critical":0,"by_detector":{},"informational":0},"endpoint":"/pith/2408.01228/integrity.json","findings":[],"available":true,"detectors_run":[],"snapshot_sha256":"c28c3603d3b5d939e8dc4c7e95fa8dfce3d595e45f758748cecf8e644a296938"},"references":{"count":0,"sample":[],"resolved_work":0,"snapshot_sha256":"258153158e38e3291e3d48162225fcdb2d5a3ed65a07baac614ab91432fd4f57","internal_anchors":0},"formal_canon":{"evidence_count":0,"snapshot_sha256":"258153158e38e3291e3d48162225fcdb2d5a3ed65a07baac614ab91432fd4f57"},"author_claims":{"count":0,"strong_count":0,"snapshot_sha256":"258153158e38e3291e3d48162225fcdb2d5a3ed65a07baac614ab91432fd4f57"},"builder_version":"pith-number-builder-2026-05-17-v1"},"aliases":[{"alias_kind":"arxiv","alias_value":"2408.01228","created_at":"2026-07-05T08:56:48.327757+00:00"},{"alias_kind":"arxiv_version","alias_value":"2408.01228v2","created_at":"2026-07-05T08:56:48.327757+00:00"},{"alias_kind":"doi","alias_value":"10.48550/arxiv.2408.01228","created_at":"2026-07-05T08:56:48.327757+00:00"},{"alias_kind":"pith_short_12","alias_value":"XZQ6GVCIDGIB","created_at":"2026-07-05T08:56:48.327757+00:00"},{"alias_kind":"pith_short_16","alias_value":"XZQ6GVCIDGIBL6XM","created_at":"2026-07-05T08:56:48.327757+00:00"},{"alias_kind":"pith_short_8","alias_value":"XZQ6GVCI","created_at":"2026-07-05T08:56:48.327757+00:00"}],"events":[],"event_summary":{},"paper_claims":[],"inbound_citations":{"count":2,"internal_anchor_count":0,"sample":[{"citing_arxiv_id":"2603.09002","citing_title":"Security Considerations for Multi-agent Systems","ref_index":216,"is_internal_anchor":false},{"citing_arxiv_id":"2604.05296","citing_title":"From Measurement to Mitigation: Quantifying and Reducing Identity Leakage in Image Representation Encoders with Linear Subspace Removal","ref_index":5,"is_internal_anchor":false}]},"formal_canon":{"evidence_count":0,"sample":[],"anchors":[]},"links":{"html":"https://pith.science/pith/XZQ6GVCIDGIBL6XM2IJJ4ZBGDT","json":"https://pith.science/pith/XZQ6GVCIDGIBL6XM2IJJ4ZBGDT.json","graph_json":"https://pith.science/api/pith-number/XZQ6GVCIDGIBL6XM2IJJ4ZBGDT/graph.json","events_json":"https://pith.science/api/pith-number/XZQ6GVCIDGIBL6XM2IJJ4ZBGDT/events.json","paper":"https://pith.science/paper/XZQ6GVCI"},"agent_actions":{"view_html":"https://pith.science/pith/XZQ6GVCIDGIBL6XM2IJJ4ZBGDT","download_json":"https://pith.science/pith/XZQ6GVCIDGIBL6XM2IJJ4ZBGDT.json","view_paper":"https://pith.science/paper/XZQ6GVCI","resolve_alias":"https://pith.science/api/pith-number/resolve?arxiv=2408.01228&json=true","fetch_graph":"https://pith.science/api/pith-number/XZQ6GVCIDGIBL6XM2IJJ4ZBGDT/graph.json","fetch_events":"https://pith.science/api/pith-number/XZQ6GVCIDGIBL6XM2IJJ4ZBGDT/events.json","actions":{"anchor_timestamp":"https://pith.science/pith/XZQ6GVCIDGIBL6XM2IJJ4ZBGDT/action/timestamp_anchor","attest_storage":"https://pith.science/pith/XZQ6GVCIDGIBL6XM2IJJ4ZBGDT/action/storage_attestation","attest_author":"https://pith.science/pith/XZQ6GVCIDGIBL6XM2IJJ4ZBGDT/action/author_attestation","sign_citation":"https://pith.science/pith/XZQ6GVCIDGIBL6XM2IJJ4ZBGDT/action/citation_signature","submit_replication":"https://pith.science/pith/XZQ6GVCIDGIBL6XM2IJJ4ZBGDT/action/replication_record"}},"created_at":"2026-07-05T08:56:48.327757+00:00","updated_at":"2026-07-05T08:56:48.327757+00:00"}