{"record_type":"pith_number_record","schema_url":"https://pith.science/schemas/pith-number/v1.json","pith_number":"pith:2025:HSAPFXM5WVBHJHK3D4RLQO7ROR","short_pith_number":"pith:HSAPFXM5","schema_version":"1.0","canonical_sha256":"3c80f2dd9db542749d5b1f22b83bf174730aee297781fd76fa73d7e9f4c6b6da","source":{"kind":"arxiv","id":"2508.12662","version":1},"attestation_state":"computed","paper":{"title":"Breaking Language Barriers: Equitable Performance in Multilingual Language Models","license":"http://creativecommons.org/licenses/by/4.0/","headline":"","cross_cats":["cs.AI"],"primary_cat":"cs.CL","authors_text":"Anna Sokol, Grigorii Khvatskii, Nitesh V. Chawla, Tanay Nagar","submitted_at":"2025-08-18T06:50:24Z","abstract_excerpt":"Cutting-edge LLMs have emerged as powerful tools for multilingual communication and understanding. However, LLMs perform worse in Common Sense Reasoning (CSR) tasks when prompted in low-resource languages (LRLs) like Hindi or Swahili compared to high-resource languages (HRLs) like English. Equalizing this inconsistent access to quality LLM outputs is crucial to ensure fairness for speakers of LRLs and across diverse linguistic communities. In this paper, we propose an approach to bridge this gap in LLM performance. Our approach involves fine-tuning an LLM on synthetic code-switched text genera"},"verification_status":{"content_addressed":true,"pith_receipt":true,"author_attested":false,"weak_author_claims":0,"strong_author_claims":0,"externally_anchored":false,"storage_verified":false,"citation_signatures":0,"replication_records":0,"graph_snapshot":true,"references_resolved":false,"formal_links_present":false},"canonical_record":{"source":{"id":"2508.12662","kind":"arxiv","version":1},"metadata":{"license":"http://creativecommons.org/licenses/by/4.0/","primary_cat":"cs.CL","submitted_at":"2025-08-18T06:50:24Z","cross_cats_sorted":["cs.AI"],"title_canon_sha256":"c805f0701f664e4b0a31a257a2856e730ba47471923063d2dbcb9f179a090603","abstract_canon_sha256":"c63f24e211b34121a8dd5b8c3bc6c599a27ea5ae32ebd32063e34867bbe431e4"},"schema_version":"1.0"},"receipt":{"kind":"pith_receipt","key_id":"pith-v1-2026-05","algorithm":"ed25519","signed_at":"2026-07-05T11:55:24.976383Z","signature_b64":"9cu1JGDOY1IOW/uU80pWPwexCpv88LrdzqwdxcLz6Xl7uzBMXqOpY1ccjkg0cYyuq+FV+Hj8BOD7f1BDLpY5Bg==","signed_message":"canonical_sha256_bytes","builder_version":"pith-number-builder-2026-05-17-v1","receipt_version":"0.3","canonical_sha256":"3c80f2dd9db542749d5b1f22b83bf174730aee297781fd76fa73d7e9f4c6b6da","last_reissued_at":"2026-07-05T11:55:24.975978Z","signature_status":"signed_v1","first_computed_at":"2026-07-05T11:55:24.975978Z","public_key_fingerprint":"8d4b5ee74e4693bcd1df2446408b0d54"},"graph_snapshot":{"paper":{"title":"Breaking Language Barriers: Equitable Performance in Multilingual Language Models","license":"http://creativecommons.org/licenses/by/4.0/","headline":"","cross_cats":["cs.AI"],"primary_cat":"cs.CL","authors_text":"Anna Sokol, Grigorii Khvatskii, Nitesh V. Chawla, Tanay Nagar","submitted_at":"2025-08-18T06:50:24Z","abstract_excerpt":"Cutting-edge LLMs have emerged as powerful tools for multilingual communication and understanding. However, LLMs perform worse in Common Sense Reasoning (CSR) tasks when prompted in low-resource languages (LRLs) like Hindi or Swahili compared to high-resource languages (HRLs) like English. Equalizing this inconsistent access to quality LLM outputs is crucial to ensure fairness for speakers of LRLs and across diverse linguistic communities. In this paper, we propose an approach to bridge this gap in LLM performance. Our approach involves fine-tuning an LLM on synthetic code-switched text genera"},"claims":{"count":0,"items":[],"snapshot_sha256":"258153158e38e3291e3d48162225fcdb2d5a3ed65a07baac614ab91432fd4f57"},"source":{"id":"2508.12662","kind":"arxiv","version":1},"verdict":{"id":null,"model_set":{},"created_at":null,"strongest_claim":"","one_line_summary":"","pipeline_version":null,"weakest_assumption":"","pith_extraction_headline":""},"integrity":{"clean":true,"summary":{"advisory":0,"critical":0,"by_detector":{},"informational":0},"endpoint":"/pith/2508.12662/integrity.json","findings":[],"available":true,"detectors_run":[],"snapshot_sha256":"c28c3603d3b5d939e8dc4c7e95fa8dfce3d595e45f758748cecf8e644a296938"},"references":{"count":0,"sample":[],"resolved_work":0,"snapshot_sha256":"258153158e38e3291e3d48162225fcdb2d5a3ed65a07baac614ab91432fd4f57","internal_anchors":0},"formal_canon":{"evidence_count":0,"snapshot_sha256":"258153158e38e3291e3d48162225fcdb2d5a3ed65a07baac614ab91432fd4f57"},"author_claims":{"count":0,"strong_count":0,"snapshot_sha256":"258153158e38e3291e3d48162225fcdb2d5a3ed65a07baac614ab91432fd4f57"},"builder_version":"pith-number-builder-2026-05-17-v1"},"aliases":[{"alias_kind":"arxiv","alias_value":"2508.12662","created_at":"2026-07-05T11:55:24.976040+00:00"},{"alias_kind":"arxiv_version","alias_value":"2508.12662v1","created_at":"2026-07-05T11:55:24.976040+00:00"},{"alias_kind":"doi","alias_value":"10.48550/arxiv.2508.12662","created_at":"2026-07-05T11:55:24.976040+00:00"},{"alias_kind":"pith_short_12","alias_value":"HSAPFXM5WVBH","created_at":"2026-07-05T11:55:24.976040+00:00"},{"alias_kind":"pith_short_16","alias_value":"HSAPFXM5WVBHJHK3","created_at":"2026-07-05T11:55:24.976040+00:00"},{"alias_kind":"pith_short_8","alias_value":"HSAPFXM5","created_at":"2026-07-05T11:55:24.976040+00:00"}],"events":[],"event_summary":{},"paper_claims":[],"inbound_citations":{"count":0,"internal_anchor_count":0,"sample":[]},"formal_canon":{"evidence_count":0,"sample":[],"anchors":[]},"links":{"html":"https://pith.science/pith/HSAPFXM5WVBHJHK3D4RLQO7ROR","json":"https://pith.science/pith/HSAPFXM5WVBHJHK3D4RLQO7ROR.json","graph_json":"https://pith.science/api/pith-number/HSAPFXM5WVBHJHK3D4RLQO7ROR/graph.json","events_json":"https://pith.science/api/pith-number/HSAPFXM5WVBHJHK3D4RLQO7ROR/events.json","paper":"https://pith.science/paper/HSAPFXM5"},"agent_actions":{"view_html":"https://pith.science/pith/HSAPFXM5WVBHJHK3D4RLQO7ROR","download_json":"https://pith.science/pith/HSAPFXM5WVBHJHK3D4RLQO7ROR.json","view_paper":"https://pith.science/paper/HSAPFXM5","resolve_alias":"https://pith.science/api/pith-number/resolve?arxiv=2508.12662&json=true","fetch_graph":"https://pith.science/api/pith-number/HSAPFXM5WVBHJHK3D4RLQO7ROR/graph.json","fetch_events":"https://pith.science/api/pith-number/HSAPFXM5WVBHJHK3D4RLQO7ROR/events.json","actions":{"anchor_timestamp":"https://pith.science/pith/HSAPFXM5WVBHJHK3D4RLQO7ROR/action/timestamp_anchor","attest_storage":"https://pith.science/pith/HSAPFXM5WVBHJHK3D4RLQO7ROR/action/storage_attestation","attest_author":"https://pith.science/pith/HSAPFXM5WVBHJHK3D4RLQO7ROR/action/author_attestation","sign_citation":"https://pith.science/pith/HSAPFXM5WVBHJHK3D4RLQO7ROR/action/citation_signature","submit_replication":"https://pith.science/pith/HSAPFXM5WVBHJHK3D4RLQO7ROR/action/replication_record"}},"created_at":"2026-07-05T11:55:24.976040+00:00","updated_at":"2026-07-05T11:55:24.976040+00:00"}