{"record_type":"pith_number_record","schema_url":"https://pith.science/schemas/pith-number/v1.json","pith_number":"pith:2024:MMUVDXK63F6TOZQY2P6EPQRNLK","short_pith_number":"pith:MMUVDXK6","schema_version":"1.0","canonical_sha256":"632951dd5ed97d376618d3fc47c22d5a8af460b86805f8cd8dc21a6c945b5db8","source":{"kind":"arxiv","id":"2407.18129","version":2},"attestation_state":"computed","paper":{"title":"Dallah: A Dialect-Aware Multimodal Large Language Model for Arabic","license":"http://creativecommons.org/licenses/by-sa/4.0/","headline":"","cross_cats":["cs.AI"],"primary_cat":"cs.CL","authors_text":"Fakhraddin Alwajih, Gagan Bhatia, Muhammad Abdul-Mageed","submitted_at":"2024-07-25T15:36:48Z","abstract_excerpt":"Recent advancements have significantly enhanced the capabilities of Multimodal Large Language Models (MLLMs) in generating and understanding image-to-text content. Despite these successes, progress is predominantly limited to English due to the scarcity of high quality multimodal resources in other languages. This limitation impedes the development of competitive models in languages such as Arabic. To alleviate this situation, we introduce an efficient Arabic multimodal assistant, dubbed Dallah, that utilizes an advanced language model based on LLaMA-2 to facilitate multimodal interactions. Da"},"verification_status":{"content_addressed":true,"pith_receipt":true,"author_attested":false,"weak_author_claims":0,"strong_author_claims":0,"externally_anchored":false,"storage_verified":false,"citation_signatures":0,"replication_records":0,"graph_snapshot":true,"references_resolved":false,"formal_links_present":false},"canonical_record":{"source":{"id":"2407.18129","kind":"arxiv","version":2},"metadata":{"license":"http://creativecommons.org/licenses/by-sa/4.0/","primary_cat":"cs.CL","submitted_at":"2024-07-25T15:36:48Z","cross_cats_sorted":["cs.AI"],"title_canon_sha256":"7229b648c7ae4d0e76b23c0166f90b772590c3a191ea80921334e4ec73b02cde","abstract_canon_sha256":"749fc829f33b71e3bb8f0fd7fc02adff9d24b943bae7e136f711b526eef457f6"},"schema_version":"1.0"},"receipt":{"kind":"pith_receipt","key_id":"pith-v1-2026-05","algorithm":"ed25519","signed_at":"2026-07-05T08:48:46.723178Z","signature_b64":"ZCfTKJmvfCJMpyd2VG4MKyePwR7/f24C5Ul1iOgebuTqmsLiR6J0RmozmKp5NyhvgMNwIArNuQZbvHathzrcDA==","signed_message":"canonical_sha256_bytes","builder_version":"pith-number-builder-2026-05-17-v1","receipt_version":"0.3","canonical_sha256":"632951dd5ed97d376618d3fc47c22d5a8af460b86805f8cd8dc21a6c945b5db8","last_reissued_at":"2026-07-05T08:48:46.722746Z","signature_status":"signed_v1","first_computed_at":"2026-07-05T08:48:46.722746Z","public_key_fingerprint":"8d4b5ee74e4693bcd1df2446408b0d54"},"graph_snapshot":{"paper":{"title":"Dallah: A Dialect-Aware Multimodal Large Language Model for Arabic","license":"http://creativecommons.org/licenses/by-sa/4.0/","headline":"","cross_cats":["cs.AI"],"primary_cat":"cs.CL","authors_text":"Fakhraddin Alwajih, Gagan Bhatia, Muhammad Abdul-Mageed","submitted_at":"2024-07-25T15:36:48Z","abstract_excerpt":"Recent advancements have significantly enhanced the capabilities of Multimodal Large Language Models (MLLMs) in generating and understanding image-to-text content. Despite these successes, progress is predominantly limited to English due to the scarcity of high quality multimodal resources in other languages. This limitation impedes the development of competitive models in languages such as Arabic. To alleviate this situation, we introduce an efficient Arabic multimodal assistant, dubbed Dallah, that utilizes an advanced language model based on LLaMA-2 to facilitate multimodal interactions. Da"},"claims":{"count":0,"items":[],"snapshot_sha256":"258153158e38e3291e3d48162225fcdb2d5a3ed65a07baac614ab91432fd4f57"},"source":{"id":"2407.18129","kind":"arxiv","version":2},"verdict":{"id":null,"model_set":{},"created_at":null,"strongest_claim":"","one_line_summary":"","pipeline_version":null,"weakest_assumption":"","pith_extraction_headline":""},"integrity":{"clean":true,"summary":{"advisory":0,"critical":0,"by_detector":{},"informational":0},"endpoint":"/pith/2407.18129/integrity.json","findings":[],"available":true,"detectors_run":[],"snapshot_sha256":"c28c3603d3b5d939e8dc4c7e95fa8dfce3d595e45f758748cecf8e644a296938"},"references":{"count":0,"sample":[],"resolved_work":0,"snapshot_sha256":"258153158e38e3291e3d48162225fcdb2d5a3ed65a07baac614ab91432fd4f57","internal_anchors":0},"formal_canon":{"evidence_count":0,"snapshot_sha256":"258153158e38e3291e3d48162225fcdb2d5a3ed65a07baac614ab91432fd4f57"},"author_claims":{"count":0,"strong_count":0,"snapshot_sha256":"258153158e38e3291e3d48162225fcdb2d5a3ed65a07baac614ab91432fd4f57"},"builder_version":"pith-number-builder-2026-05-17-v1"},"aliases":[{"alias_kind":"arxiv","alias_value":"2407.18129","created_at":"2026-07-05T08:48:46.722804+00:00"},{"alias_kind":"arxiv_version","alias_value":"2407.18129v2","created_at":"2026-07-05T08:48:46.722804+00:00"},{"alias_kind":"doi","alias_value":"10.48550/arxiv.2407.18129","created_at":"2026-07-05T08:48:46.722804+00:00"},{"alias_kind":"pith_short_12","alias_value":"MMUVDXK63F6T","created_at":"2026-07-05T08:48:46.722804+00:00"},{"alias_kind":"pith_short_16","alias_value":"MMUVDXK63F6TOZQY","created_at":"2026-07-05T08:48:46.722804+00:00"},{"alias_kind":"pith_short_8","alias_value":"MMUVDXK6","created_at":"2026-07-05T08:48:46.722804+00:00"}],"events":[],"event_summary":{},"paper_claims":[],"inbound_citations":{"count":1,"internal_anchor_count":1,"sample":[{"citing_arxiv_id":"2505.00114","citing_title":"Fine-Tuning LLMs for Low-Resource Dialect Translation: The Case of Lebanese","ref_index":5,"is_internal_anchor":true}]},"formal_canon":{"evidence_count":0,"sample":[],"anchors":[]},"links":{"html":"https://pith.science/pith/MMUVDXK63F6TOZQY2P6EPQRNLK","json":"https://pith.science/pith/MMUVDXK63F6TOZQY2P6EPQRNLK.json","graph_json":"https://pith.science/api/pith-number/MMUVDXK63F6TOZQY2P6EPQRNLK/graph.json","events_json":"https://pith.science/api/pith-number/MMUVDXK63F6TOZQY2P6EPQRNLK/events.json","paper":"https://pith.science/paper/MMUVDXK6"},"agent_actions":{"view_html":"https://pith.science/pith/MMUVDXK63F6TOZQY2P6EPQRNLK","download_json":"https://pith.science/pith/MMUVDXK63F6TOZQY2P6EPQRNLK.json","view_paper":"https://pith.science/paper/MMUVDXK6","resolve_alias":"https://pith.science/api/pith-number/resolve?arxiv=2407.18129&json=true","fetch_graph":"https://pith.science/api/pith-number/MMUVDXK63F6TOZQY2P6EPQRNLK/graph.json","fetch_events":"https://pith.science/api/pith-number/MMUVDXK63F6TOZQY2P6EPQRNLK/events.json","actions":{"anchor_timestamp":"https://pith.science/pith/MMUVDXK63F6TOZQY2P6EPQRNLK/action/timestamp_anchor","attest_storage":"https://pith.science/pith/MMUVDXK63F6TOZQY2P6EPQRNLK/action/storage_attestation","attest_author":"https://pith.science/pith/MMUVDXK63F6TOZQY2P6EPQRNLK/action/author_attestation","sign_citation":"https://pith.science/pith/MMUVDXK63F6TOZQY2P6EPQRNLK/action/citation_signature","submit_replication":"https://pith.science/pith/MMUVDXK63F6TOZQY2P6EPQRNLK/action/replication_record"}},"created_at":"2026-07-05T08:48:46.722804+00:00","updated_at":"2026-07-05T08:48:46.722804+00:00"}