{"record_type":"pith_number_record","schema_url":"https://pith.science/schemas/pith-number/v1.json","pith_number":"pith:2024:GOEL2JTQHF2F5L46ELTGF2JVLV","short_pith_number":"pith:GOEL2JTQ","schema_version":"1.0","canonical_sha256":"3388bd267039745eaf9e22e662e9355d728eba5858d583f7b973edc023432066","source":{"kind":"arxiv","id":"2409.07353","version":1},"attestation_state":"computed","paper":{"title":"Securing Vision-Language Models with a Robust Encoder Against Jailbreak and Adversarial Attacks","license":"http://arxiv.org/licenses/nonexclusive-distrib/1.0/","headline":"","cross_cats":["cs.AI"],"primary_cat":"cs.CV","authors_text":"Ahmed Imteaj, Md Zarif Hossain","submitted_at":"2024-09-11T15:39:42Z","abstract_excerpt":"Large Vision-Language Models (LVLMs), trained on multimodal big datasets, have significantly advanced AI by excelling in vision-language tasks. However, these models remain vulnerable to adversarial attacks, particularly jailbreak attacks, which bypass safety protocols and cause the model to generate misleading or harmful responses. This vulnerability stems from both the inherent susceptibilities of LLMs and the expanded attack surface introduced by the visual modality. We propose Sim-CLIP+, a novel defense mechanism that adversarially fine-tunes the CLIP vision encoder by leveraging a Siamese"},"verification_status":{"content_addressed":true,"pith_receipt":true,"author_attested":false,"weak_author_claims":0,"strong_author_claims":0,"externally_anchored":false,"storage_verified":false,"citation_signatures":0,"replication_records":0,"graph_snapshot":true,"references_resolved":false,"formal_links_present":false},"canonical_record":{"source":{"id":"2409.07353","kind":"arxiv","version":1},"metadata":{"license":"http://arxiv.org/licenses/nonexclusive-distrib/1.0/","primary_cat":"cs.CV","submitted_at":"2024-09-11T15:39:42Z","cross_cats_sorted":["cs.AI"],"title_canon_sha256":"50c7a85c271ed2dda466b952c7bc28850a648b6c1bf43addd22b8e21a75d9baf","abstract_canon_sha256":"1a14696dbc720335e55a038e9105394fe85c97fb6d62bfec03803b133ced74db"},"schema_version":"1.0"},"receipt":{"kind":"pith_receipt","key_id":"pith-v1-2026-05","algorithm":"ed25519","signed_at":"2026-07-05T09:05:56.612758Z","signature_b64":"bWtdsyRJYb3QeZeVbcdW2pPMLQkdFrOmAB+3+Eo5gm3PBMJQiB+Cm+piMscbIxpHJN0jsHYQhnzVBa3VZSCLDw==","signed_message":"canonical_sha256_bytes","builder_version":"pith-number-builder-2026-05-17-v1","receipt_version":"0.3","canonical_sha256":"3388bd267039745eaf9e22e662e9355d728eba5858d583f7b973edc023432066","last_reissued_at":"2026-07-05T09:05:56.612295Z","signature_status":"signed_v1","first_computed_at":"2026-07-05T09:05:56.612295Z","public_key_fingerprint":"8d4b5ee74e4693bcd1df2446408b0d54"},"graph_snapshot":{"paper":{"title":"Securing Vision-Language Models with a Robust Encoder Against Jailbreak and Adversarial Attacks","license":"http://arxiv.org/licenses/nonexclusive-distrib/1.0/","headline":"","cross_cats":["cs.AI"],"primary_cat":"cs.CV","authors_text":"Ahmed Imteaj, Md Zarif Hossain","submitted_at":"2024-09-11T15:39:42Z","abstract_excerpt":"Large Vision-Language Models (LVLMs), trained on multimodal big datasets, have significantly advanced AI by excelling in vision-language tasks. However, these models remain vulnerable to adversarial attacks, particularly jailbreak attacks, which bypass safety protocols and cause the model to generate misleading or harmful responses. This vulnerability stems from both the inherent susceptibilities of LLMs and the expanded attack surface introduced by the visual modality. We propose Sim-CLIP+, a novel defense mechanism that adversarially fine-tunes the CLIP vision encoder by leveraging a Siamese"},"claims":{"count":0,"items":[],"snapshot_sha256":"258153158e38e3291e3d48162225fcdb2d5a3ed65a07baac614ab91432fd4f57"},"source":{"id":"2409.07353","kind":"arxiv","version":1},"verdict":{"id":null,"model_set":{},"created_at":null,"strongest_claim":"","one_line_summary":"","pipeline_version":null,"weakest_assumption":"","pith_extraction_headline":""},"integrity":{"clean":true,"summary":{"advisory":0,"critical":0,"by_detector":{},"informational":0},"endpoint":"/pith/2409.07353/integrity.json","findings":[],"available":true,"detectors_run":[],"snapshot_sha256":"c28c3603d3b5d939e8dc4c7e95fa8dfce3d595e45f758748cecf8e644a296938"},"references":{"count":0,"sample":[],"resolved_work":0,"snapshot_sha256":"258153158e38e3291e3d48162225fcdb2d5a3ed65a07baac614ab91432fd4f57","internal_anchors":0},"formal_canon":{"evidence_count":0,"snapshot_sha256":"258153158e38e3291e3d48162225fcdb2d5a3ed65a07baac614ab91432fd4f57"},"author_claims":{"count":0,"strong_count":0,"snapshot_sha256":"258153158e38e3291e3d48162225fcdb2d5a3ed65a07baac614ab91432fd4f57"},"builder_version":"pith-number-builder-2026-05-17-v1"},"aliases":[{"alias_kind":"arxiv","alias_value":"2409.07353","created_at":"2026-07-05T09:05:56.612361+00:00"},{"alias_kind":"arxiv_version","alias_value":"2409.07353v1","created_at":"2026-07-05T09:05:56.612361+00:00"},{"alias_kind":"doi","alias_value":"10.48550/arxiv.2409.07353","created_at":"2026-07-05T09:05:56.612361+00:00"},{"alias_kind":"pith_short_12","alias_value":"GOEL2JTQHF2F","created_at":"2026-07-05T09:05:56.612361+00:00"},{"alias_kind":"pith_short_16","alias_value":"GOEL2JTQHF2F5L46","created_at":"2026-07-05T09:05:56.612361+00:00"},{"alias_kind":"pith_short_8","alias_value":"GOEL2JTQ","created_at":"2026-07-05T09:05:56.612361+00:00"}],"events":[],"event_summary":{},"paper_claims":[],"inbound_citations":{"count":1,"internal_anchor_count":0,"sample":[{"citing_arxiv_id":"2606.03713","citing_title":"Investigating Adversarial Robustness of Multi-modal Large Language Models","ref_index":21,"is_internal_anchor":false}]},"formal_canon":{"evidence_count":0,"sample":[],"anchors":[]},"links":{"html":"https://pith.science/pith/GOEL2JTQHF2F5L46ELTGF2JVLV","json":"https://pith.science/pith/GOEL2JTQHF2F5L46ELTGF2JVLV.json","graph_json":"https://pith.science/api/pith-number/GOEL2JTQHF2F5L46ELTGF2JVLV/graph.json","events_json":"https://pith.science/api/pith-number/GOEL2JTQHF2F5L46ELTGF2JVLV/events.json","paper":"https://pith.science/paper/GOEL2JTQ"},"agent_actions":{"view_html":"https://pith.science/pith/GOEL2JTQHF2F5L46ELTGF2JVLV","download_json":"https://pith.science/pith/GOEL2JTQHF2F5L46ELTGF2JVLV.json","view_paper":"https://pith.science/paper/GOEL2JTQ","resolve_alias":"https://pith.science/api/pith-number/resolve?arxiv=2409.07353&json=true","fetch_graph":"https://pith.science/api/pith-number/GOEL2JTQHF2F5L46ELTGF2JVLV/graph.json","fetch_events":"https://pith.science/api/pith-number/GOEL2JTQHF2F5L46ELTGF2JVLV/events.json","actions":{"anchor_timestamp":"https://pith.science/pith/GOEL2JTQHF2F5L46ELTGF2JVLV/action/timestamp_anchor","attest_storage":"https://pith.science/pith/GOEL2JTQHF2F5L46ELTGF2JVLV/action/storage_attestation","attest_author":"https://pith.science/pith/GOEL2JTQHF2F5L46ELTGF2JVLV/action/author_attestation","sign_citation":"https://pith.science/pith/GOEL2JTQHF2F5L46ELTGF2JVLV/action/citation_signature","submit_replication":"https://pith.science/pith/GOEL2JTQHF2F5L46ELTGF2JVLV/action/replication_record"}},"created_at":"2026-07-05T09:05:56.612361+00:00","updated_at":"2026-07-05T09:05:56.612361+00:00"}