{"record_type":"pith_number_record","schema_url":"https://pith.science/schemas/pith-number/v1.json","pith_number":"pith:2023:AWZQ2MI4EGIKEZVGUCMM4ASKNS","short_pith_number":"pith:AWZQ2MI4","schema_version":"1.0","canonical_sha256":"05b30d311c2190a266a6a098ce024a6c8065b5cdd0ac77027c694d0cd4689629","source":{"kind":"arxiv","id":"2308.16884","version":2},"attestation_state":"computed","paper":{"title":"The Belebele Benchmark: a Parallel Reading Comprehension Dataset in 122 Language Variants","license":"http://creativecommons.org/licenses/by-sa/4.0/","headline":"","cross_cats":["cs.AI","cs.LG"],"primary_cat":"cs.CL","authors_text":"Abhinandan Krishnan, Benjamin Muller, Davis Liang, Donald Husa, Lucas Bandarkar, Luke Zettlemoyer, Madian Khabsa, Mikel Artetxe, Naman Goyal, Satya Narayan Shukla","submitted_at":"2023-08-31T17:43:08Z","abstract_excerpt":"We present Belebele, a multiple-choice machine reading comprehension (MRC) dataset spanning 122 language variants. Significantly expanding the language coverage of natural language understanding (NLU) benchmarks, this dataset enables the evaluation of text models in high-, medium-, and low-resource languages. Each question is based on a short passage from the Flores-200 dataset and has four multiple-choice answers. The questions were carefully curated to discriminate between models with different levels of general language comprehension. The English dataset on its own proves difficult enough t"},"verification_status":{"content_addressed":true,"pith_receipt":true,"author_attested":false,"weak_author_claims":0,"strong_author_claims":0,"externally_anchored":false,"storage_verified":false,"citation_signatures":0,"replication_records":0,"graph_snapshot":true,"references_resolved":false,"formal_links_present":false},"canonical_record":{"source":{"id":"2308.16884","kind":"arxiv","version":2},"metadata":{"license":"http://creativecommons.org/licenses/by-sa/4.0/","primary_cat":"cs.CL","submitted_at":"2023-08-31T17:43:08Z","cross_cats_sorted":["cs.AI","cs.LG"],"title_canon_sha256":"0be3615edcc0d8acf59ce6e0a598e8e88c9c02a979a1821cf2db189d8fddd863","abstract_canon_sha256":"13187c0f9aa88cf28e49a4dd4e832f86c62301a36e7ee6f3e9df6cbc8fe518c3"},"schema_version":"1.0"},"receipt":{"kind":"pith_receipt","key_id":"pith-v1-2026-05","algorithm":"ed25519","signed_at":"2026-07-05T09:14:27.910490Z","signature_b64":"wV/7hJ6gP5MVwbg6Lpio5lyJdjM9c0wmOpdmFV8kDThaPFsijzGYHt4EiXY86HceIRNoduGdnliPWr9YHbVoBg==","signed_message":"canonical_sha256_bytes","builder_version":"pith-number-builder-2026-05-17-v1","receipt_version":"0.3","canonical_sha256":"05b30d311c2190a266a6a098ce024a6c8065b5cdd0ac77027c694d0cd4689629","last_reissued_at":"2026-07-05T09:14:27.909972Z","signature_status":"signed_v1","first_computed_at":"2026-07-05T09:14:27.909972Z","public_key_fingerprint":"8d4b5ee74e4693bcd1df2446408b0d54"},"graph_snapshot":{"paper":{"title":"The Belebele Benchmark: a Parallel Reading Comprehension Dataset in 122 Language Variants","license":"http://creativecommons.org/licenses/by-sa/4.0/","headline":"","cross_cats":["cs.AI","cs.LG"],"primary_cat":"cs.CL","authors_text":"Abhinandan Krishnan, Benjamin Muller, Davis Liang, Donald Husa, Lucas Bandarkar, Luke Zettlemoyer, Madian Khabsa, Mikel Artetxe, Naman Goyal, Satya Narayan Shukla","submitted_at":"2023-08-31T17:43:08Z","abstract_excerpt":"We present Belebele, a multiple-choice machine reading comprehension (MRC) dataset spanning 122 language variants. Significantly expanding the language coverage of natural language understanding (NLU) benchmarks, this dataset enables the evaluation of text models in high-, medium-, and low-resource languages. Each question is based on a short passage from the Flores-200 dataset and has four multiple-choice answers. The questions were carefully curated to discriminate between models with different levels of general language comprehension. The English dataset on its own proves difficult enough t"},"claims":{"count":0,"items":[],"snapshot_sha256":"258153158e38e3291e3d48162225fcdb2d5a3ed65a07baac614ab91432fd4f57"},"source":{"id":"2308.16884","kind":"arxiv","version":2},"verdict":{"id":null,"model_set":{},"created_at":null,"strongest_claim":"","one_line_summary":"","pipeline_version":null,"weakest_assumption":"","pith_extraction_headline":""},"integrity":{"clean":true,"summary":{"advisory":0,"critical":0,"by_detector":{},"informational":0},"endpoint":"/pith/2308.16884/integrity.json","findings":[],"available":true,"detectors_run":[],"snapshot_sha256":"c28c3603d3b5d939e8dc4c7e95fa8dfce3d595e45f758748cecf8e644a296938"},"references":{"count":0,"sample":[],"resolved_work":0,"snapshot_sha256":"258153158e38e3291e3d48162225fcdb2d5a3ed65a07baac614ab91432fd4f57","internal_anchors":0},"formal_canon":{"evidence_count":0,"snapshot_sha256":"258153158e38e3291e3d48162225fcdb2d5a3ed65a07baac614ab91432fd4f57"},"author_claims":{"count":0,"strong_count":0,"snapshot_sha256":"258153158e38e3291e3d48162225fcdb2d5a3ed65a07baac614ab91432fd4f57"},"builder_version":"pith-number-builder-2026-05-17-v1"},"aliases":[{"alias_kind":"arxiv","alias_value":"2308.16884","created_at":"2026-07-05T09:14:27.910034+00:00"},{"alias_kind":"arxiv_version","alias_value":"2308.16884v2","created_at":"2026-07-05T09:14:27.910034+00:00"},{"alias_kind":"doi","alias_value":"10.48550/arxiv.2308.16884","created_at":"2026-07-05T09:14:27.910034+00:00"},{"alias_kind":"pith_short_12","alias_value":"AWZQ2MI4EGIK","created_at":"2026-07-05T09:14:27.910034+00:00"},{"alias_kind":"pith_short_16","alias_value":"AWZQ2MI4EGIKEZVG","created_at":"2026-07-05T09:14:27.910034+00:00"},{"alias_kind":"pith_short_8","alias_value":"AWZQ2MI4","created_at":"2026-07-05T09:14:27.910034+00:00"}],"events":[],"event_summary":{},"paper_claims":[],"inbound_citations":{"count":5,"internal_anchor_count":0,"sample":[{"citing_arxiv_id":"2412.15115","citing_title":"Qwen2.5 Technical Report","ref_index":7,"is_internal_anchor":false},{"citing_arxiv_id":"2507.08480","citing_title":"Improving Korean-English Cross-Lingual Retrieval: A Data-Centric Study of Language Composition and Model Merging","ref_index":1,"is_internal_anchor":false},{"citing_arxiv_id":"2605.18083","citing_title":"A Data-Efficient Path to Multilingual LLMs: Language Expansion via Post-training PARAM$\\Delta$ Integration into Upcycled MoE","ref_index":27,"is_internal_anchor":false},{"citing_arxiv_id":"2507.07847","citing_title":"From Ambiguity to Accuracy: The Transformative Effect of Coreference Resolution on Retrieval-Augmented Generation systems","ref_index":1,"is_internal_anchor":false},{"citing_arxiv_id":"2505.09388","citing_title":"Qwen3 Technical Report","ref_index":6,"is_internal_anchor":false}]},"formal_canon":{"evidence_count":0,"sample":[],"anchors":[]},"links":{"html":"https://pith.science/pith/AWZQ2MI4EGIKEZVGUCMM4ASKNS","json":"https://pith.science/pith/AWZQ2MI4EGIKEZVGUCMM4ASKNS.json","graph_json":"https://pith.science/api/pith-number/AWZQ2MI4EGIKEZVGUCMM4ASKNS/graph.json","events_json":"https://pith.science/api/pith-number/AWZQ2MI4EGIKEZVGUCMM4ASKNS/events.json","paper":"https://pith.science/paper/AWZQ2MI4"},"agent_actions":{"view_html":"https://pith.science/pith/AWZQ2MI4EGIKEZVGUCMM4ASKNS","download_json":"https://pith.science/pith/AWZQ2MI4EGIKEZVGUCMM4ASKNS.json","view_paper":"https://pith.science/paper/AWZQ2MI4","resolve_alias":"https://pith.science/api/pith-number/resolve?arxiv=2308.16884&json=true","fetch_graph":"https://pith.science/api/pith-number/AWZQ2MI4EGIKEZVGUCMM4ASKNS/graph.json","fetch_events":"https://pith.science/api/pith-number/AWZQ2MI4EGIKEZVGUCMM4ASKNS/events.json","actions":{"anchor_timestamp":"https://pith.science/pith/AWZQ2MI4EGIKEZVGUCMM4ASKNS/action/timestamp_anchor","attest_storage":"https://pith.science/pith/AWZQ2MI4EGIKEZVGUCMM4ASKNS/action/storage_attestation","attest_author":"https://pith.science/pith/AWZQ2MI4EGIKEZVGUCMM4ASKNS/action/author_attestation","sign_citation":"https://pith.science/pith/AWZQ2MI4EGIKEZVGUCMM4ASKNS/action/citation_signature","submit_replication":"https://pith.science/pith/AWZQ2MI4EGIKEZVGUCMM4ASKNS/action/replication_record"}},"created_at":"2026-07-05T09:14:27.910034+00:00","updated_at":"2026-07-05T09:14:27.910034+00:00"}