{"record_type":"pith_number_record","schema_url":"https://pith.science/schemas/pith-number/v1.json","pith_number":"pith:2024:TNVQXO2WZMLZGZ42PHN6NEDMGL","short_pith_number":"pith:TNVQXO2W","schema_version":"1.0","canonical_sha256":"9b6b0bbb56cb1793679a79dbe6906c32ccd75d7c31ef73fa3dd1dcb6092bed02","source":{"kind":"arxiv","id":"2408.12808","version":1},"attestation_state":"computed","paper":{"title":"VALE: A Multimodal Visual and Language Explanation Framework for Image Classifiers using eXplainable AI and Language Models","license":"http://creativecommons.org/licenses/by/4.0/","headline":"","cross_cats":["cs.AI","cs.CL","cs.LG"],"primary_cat":"cs.CV","authors_text":"Athira Nambiar, Purushothaman Natarajan","submitted_at":"2024-08-23T03:02:11Z","abstract_excerpt":"Deep Neural Networks (DNNs) have revolutionized various fields by enabling task automation and reducing human error. However, their internal workings and decision-making processes remain obscure due to their black box nature. Consequently, the lack of interpretability limits the application of these models in high-risk scenarios. To address this issue, the emerging field of eXplainable Artificial Intelligence (XAI) aims to explain and interpret the inner workings of DNNs. Despite advancements, XAI faces challenges such as the semantic gap between machine and human understanding, the trade-off "},"verification_status":{"content_addressed":true,"pith_receipt":true,"author_attested":false,"weak_author_claims":0,"strong_author_claims":0,"externally_anchored":false,"storage_verified":false,"citation_signatures":0,"replication_records":0,"graph_snapshot":true,"references_resolved":false,"formal_links_present":false},"canonical_record":{"source":{"id":"2408.12808","kind":"arxiv","version":1},"metadata":{"license":"http://creativecommons.org/licenses/by/4.0/","primary_cat":"cs.CV","submitted_at":"2024-08-23T03:02:11Z","cross_cats_sorted":["cs.AI","cs.CL","cs.LG"],"title_canon_sha256":"f80390afc366b89776de2320d35c0b647f1317ac5a1ea2809033630df42bc704","abstract_canon_sha256":"06bc8e42b76e5edb8a72e1b9ee4774b7a5dcc86f92ae408da1818835dc3af1c8"},"schema_version":"1.0"},"receipt":{"kind":"pith_receipt","key_id":"pith-v1-2026-05","algorithm":"ed25519","signed_at":"2026-07-05T08:58:30.757495Z","signature_b64":"oxfuWuSQ/AyFpG/3qFDgY2KBeTwz9896ka4xHkop/thCv2vSXUQSmzRK2I6Kynx7RQMZKXI0/TEDl+IDsLjqDA==","signed_message":"canonical_sha256_bytes","builder_version":"pith-number-builder-2026-05-17-v1","receipt_version":"0.3","canonical_sha256":"9b6b0bbb56cb1793679a79dbe6906c32ccd75d7c31ef73fa3dd1dcb6092bed02","last_reissued_at":"2026-07-05T08:58:30.756931Z","signature_status":"signed_v1","first_computed_at":"2026-07-05T08:58:30.756931Z","public_key_fingerprint":"8d4b5ee74e4693bcd1df2446408b0d54"},"graph_snapshot":{"paper":{"title":"VALE: A Multimodal Visual and Language Explanation Framework for Image Classifiers using eXplainable AI and Language Models","license":"http://creativecommons.org/licenses/by/4.0/","headline":"","cross_cats":["cs.AI","cs.CL","cs.LG"],"primary_cat":"cs.CV","authors_text":"Athira Nambiar, Purushothaman Natarajan","submitted_at":"2024-08-23T03:02:11Z","abstract_excerpt":"Deep Neural Networks (DNNs) have revolutionized various fields by enabling task automation and reducing human error. However, their internal workings and decision-making processes remain obscure due to their black box nature. Consequently, the lack of interpretability limits the application of these models in high-risk scenarios. To address this issue, the emerging field of eXplainable Artificial Intelligence (XAI) aims to explain and interpret the inner workings of DNNs. Despite advancements, XAI faces challenges such as the semantic gap between machine and human understanding, the trade-off "},"claims":{"count":0,"items":[],"snapshot_sha256":"258153158e38e3291e3d48162225fcdb2d5a3ed65a07baac614ab91432fd4f57"},"source":{"id":"2408.12808","kind":"arxiv","version":1},"verdict":{"id":null,"model_set":{},"created_at":null,"strongest_claim":"","one_line_summary":"","pipeline_version":null,"weakest_assumption":"","pith_extraction_headline":""},"integrity":{"clean":true,"summary":{"advisory":0,"critical":0,"by_detector":{},"informational":0},"endpoint":"/pith/2408.12808/integrity.json","findings":[],"available":true,"detectors_run":[],"snapshot_sha256":"c28c3603d3b5d939e8dc4c7e95fa8dfce3d595e45f758748cecf8e644a296938"},"references":{"count":0,"sample":[],"resolved_work":0,"snapshot_sha256":"258153158e38e3291e3d48162225fcdb2d5a3ed65a07baac614ab91432fd4f57","internal_anchors":0},"formal_canon":{"evidence_count":0,"snapshot_sha256":"258153158e38e3291e3d48162225fcdb2d5a3ed65a07baac614ab91432fd4f57"},"author_claims":{"count":0,"strong_count":0,"snapshot_sha256":"258153158e38e3291e3d48162225fcdb2d5a3ed65a07baac614ab91432fd4f57"},"builder_version":"pith-number-builder-2026-05-17-v1"},"aliases":[{"alias_kind":"arxiv","alias_value":"2408.12808","created_at":"2026-07-05T08:58:30.756990+00:00"},{"alias_kind":"arxiv_version","alias_value":"2408.12808v1","created_at":"2026-07-05T08:58:30.756990+00:00"},{"alias_kind":"doi","alias_value":"10.48550/arxiv.2408.12808","created_at":"2026-07-05T08:58:30.756990+00:00"},{"alias_kind":"pith_short_12","alias_value":"TNVQXO2WZMLZ","created_at":"2026-07-05T08:58:30.756990+00:00"},{"alias_kind":"pith_short_16","alias_value":"TNVQXO2WZMLZGZ42","created_at":"2026-07-05T08:58:30.756990+00:00"},{"alias_kind":"pith_short_8","alias_value":"TNVQXO2W","created_at":"2026-07-05T08:58:30.756990+00:00"}],"events":[],"event_summary":{},"paper_claims":[],"inbound_citations":{"count":1,"internal_anchor_count":1,"sample":[{"citing_arxiv_id":"2501.12203","citing_title":"Explainability for Vision Foundation Models: A Survey","ref_index":137,"is_internal_anchor":true}]},"formal_canon":{"evidence_count":0,"sample":[],"anchors":[]},"links":{"html":"https://pith.science/pith/TNVQXO2WZMLZGZ42PHN6NEDMGL","json":"https://pith.science/pith/TNVQXO2WZMLZGZ42PHN6NEDMGL.json","graph_json":"https://pith.science/api/pith-number/TNVQXO2WZMLZGZ42PHN6NEDMGL/graph.json","events_json":"https://pith.science/api/pith-number/TNVQXO2WZMLZGZ42PHN6NEDMGL/events.json","paper":"https://pith.science/paper/TNVQXO2W"},"agent_actions":{"view_html":"https://pith.science/pith/TNVQXO2WZMLZGZ42PHN6NEDMGL","download_json":"https://pith.science/pith/TNVQXO2WZMLZGZ42PHN6NEDMGL.json","view_paper":"https://pith.science/paper/TNVQXO2W","resolve_alias":"https://pith.science/api/pith-number/resolve?arxiv=2408.12808&json=true","fetch_graph":"https://pith.science/api/pith-number/TNVQXO2WZMLZGZ42PHN6NEDMGL/graph.json","fetch_events":"https://pith.science/api/pith-number/TNVQXO2WZMLZGZ42PHN6NEDMGL/events.json","actions":{"anchor_timestamp":"https://pith.science/pith/TNVQXO2WZMLZGZ42PHN6NEDMGL/action/timestamp_anchor","attest_storage":"https://pith.science/pith/TNVQXO2WZMLZGZ42PHN6NEDMGL/action/storage_attestation","attest_author":"https://pith.science/pith/TNVQXO2WZMLZGZ42PHN6NEDMGL/action/author_attestation","sign_citation":"https://pith.science/pith/TNVQXO2WZMLZGZ42PHN6NEDMGL/action/citation_signature","submit_replication":"https://pith.science/pith/TNVQXO2WZMLZGZ42PHN6NEDMGL/action/replication_record"}},"created_at":"2026-07-05T08:58:30.756990+00:00","updated_at":"2026-07-05T08:58:30.756990+00:00"}