{"record_type":"pith_number_record","schema_url":"https://pith.science/schemas/pith-number/v1.json","pith_number":"pith:2023:T42VE62WITC6M6SGATY5RXVYVR","short_pith_number":"pith:T42VE62W","schema_version":"1.0","canonical_sha256":"9f35527b5644c5e67a4604f1d8deb8ac52315ad297362b00f218b341fd3d308b","source":{"kind":"arxiv","id":"2309.05472","version":2},"attestation_state":"computed","paper":{"title":"LeBenchmark 2.0: a Standardized, Replicable and Enhanced Framework for Self-supervised Representations of French Speech","license":"http://creativecommons.org/licenses/by/4.0/","headline":"","cross_cats":["cs.AI","cs.SD","eess.AS"],"primary_cat":"cs.CL","authors_text":"Adrien Pupier, Alexandre Allauzen, Benjamin Lecouteux, Didier Schwab, Fabien Ringeval, Francois Portet, Hang Le, Ha Nguyen, Jerome Goulian, Laurent Besacier, Marcely Zanon Boito, Marco Dinarelli, Maximin Coavoux, Mickael Rouvier, Natalia Tomashenko, Salima Mdhaffar, Shucong Zhang, Sina Alisamir, Solange Rossato, Solene Evain, Titouan Parcollet, Yannick Esteve","submitted_at":"2023-09-11T14:13:09Z","abstract_excerpt":"Self-supervised learning (SSL) is at the origin of unprecedented improvements in many different domains including computer vision and natural language processing. Speech processing drastically benefitted from SSL as most of the current domain-related tasks are now being approached with pre-trained models. This work introduces LeBenchmark 2.0 an open-source framework for assessing and building SSL-equipped French speech technologies. It includes documented, large-scale and heterogeneous corpora with up to 14,000 hours of heterogeneous speech, ten pre-trained SSL wav2vec 2.0 models containing fr"},"verification_status":{"content_addressed":true,"pith_receipt":true,"author_attested":false,"weak_author_claims":0,"strong_author_claims":0,"externally_anchored":false,"storage_verified":false,"citation_signatures":0,"replication_records":0,"graph_snapshot":true,"references_resolved":false,"formal_links_present":false},"canonical_record":{"source":{"id":"2309.05472","kind":"arxiv","version":2},"metadata":{"license":"http://creativecommons.org/licenses/by/4.0/","primary_cat":"cs.CL","submitted_at":"2023-09-11T14:13:09Z","cross_cats_sorted":["cs.AI","cs.SD","eess.AS"],"title_canon_sha256":"9d8a9f8c05d1414267f0897ae0228a3191137856ea68830d104043e7081e7a01","abstract_canon_sha256":"30887a5701adba17e13421be6b549b341ab55b75049dda88ac8051a2cf32e84f"},"schema_version":"1.0"},"receipt":{"kind":"pith_receipt","key_id":"pith-v1-2026-05","algorithm":"ed25519","signed_at":"2026-07-05T07:56:52.120433Z","signature_b64":"Ws+nlmC9R/9N07/rK48yyAWTKdpn1xz/rVCzkdLmDE8l8DT6Vq+dS5ixwrFqFWsbQyWPBK42mx41NLgwno8wCA==","signed_message":"canonical_sha256_bytes","builder_version":"pith-number-builder-2026-05-17-v1","receipt_version":"0.3","canonical_sha256":"9f35527b5644c5e67a4604f1d8deb8ac52315ad297362b00f218b341fd3d308b","last_reissued_at":"2026-07-05T07:56:52.119949Z","signature_status":"signed_v1","first_computed_at":"2026-07-05T07:56:52.119949Z","public_key_fingerprint":"8d4b5ee74e4693bcd1df2446408b0d54"},"graph_snapshot":{"paper":{"title":"LeBenchmark 2.0: a Standardized, Replicable and Enhanced Framework for Self-supervised Representations of French Speech","license":"http://creativecommons.org/licenses/by/4.0/","headline":"","cross_cats":["cs.AI","cs.SD","eess.AS"],"primary_cat":"cs.CL","authors_text":"Adrien Pupier, Alexandre Allauzen, Benjamin Lecouteux, Didier Schwab, Fabien Ringeval, Francois Portet, Hang Le, Ha Nguyen, Jerome Goulian, Laurent Besacier, Marcely Zanon Boito, Marco Dinarelli, Maximin Coavoux, Mickael Rouvier, Natalia Tomashenko, Salima Mdhaffar, Shucong Zhang, Sina Alisamir, Solange Rossato, Solene Evain, Titouan Parcollet, Yannick Esteve","submitted_at":"2023-09-11T14:13:09Z","abstract_excerpt":"Self-supervised learning (SSL) is at the origin of unprecedented improvements in many different domains including computer vision and natural language processing. Speech processing drastically benefitted from SSL as most of the current domain-related tasks are now being approached with pre-trained models. This work introduces LeBenchmark 2.0 an open-source framework for assessing and building SSL-equipped French speech technologies. It includes documented, large-scale and heterogeneous corpora with up to 14,000 hours of heterogeneous speech, ten pre-trained SSL wav2vec 2.0 models containing fr"},"claims":{"count":0,"items":[],"snapshot_sha256":"258153158e38e3291e3d48162225fcdb2d5a3ed65a07baac614ab91432fd4f57"},"source":{"id":"2309.05472","kind":"arxiv","version":2},"verdict":{"id":null,"model_set":{},"created_at":null,"strongest_claim":"","one_line_summary":"","pipeline_version":null,"weakest_assumption":"","pith_extraction_headline":""},"integrity":{"clean":true,"summary":{"advisory":0,"critical":0,"by_detector":{},"informational":0},"endpoint":"/pith/2309.05472/integrity.json","findings":[],"available":true,"detectors_run":[],"snapshot_sha256":"c28c3603d3b5d939e8dc4c7e95fa8dfce3d595e45f758748cecf8e644a296938"},"references":{"count":0,"sample":[],"resolved_work":0,"snapshot_sha256":"258153158e38e3291e3d48162225fcdb2d5a3ed65a07baac614ab91432fd4f57","internal_anchors":0},"formal_canon":{"evidence_count":0,"snapshot_sha256":"258153158e38e3291e3d48162225fcdb2d5a3ed65a07baac614ab91432fd4f57"},"author_claims":{"count":0,"strong_count":0,"snapshot_sha256":"258153158e38e3291e3d48162225fcdb2d5a3ed65a07baac614ab91432fd4f57"},"builder_version":"pith-number-builder-2026-05-17-v1"},"aliases":[{"alias_kind":"arxiv","alias_value":"2309.05472","created_at":"2026-07-05T07:56:52.120008+00:00"},{"alias_kind":"arxiv_version","alias_value":"2309.05472v2","created_at":"2026-07-05T07:56:52.120008+00:00"},{"alias_kind":"doi","alias_value":"10.48550/arxiv.2309.05472","created_at":"2026-07-05T07:56:52.120008+00:00"},{"alias_kind":"pith_short_12","alias_value":"T42VE62WITC6","created_at":"2026-07-05T07:56:52.120008+00:00"},{"alias_kind":"pith_short_16","alias_value":"T42VE62WITC6M6SG","created_at":"2026-07-05T07:56:52.120008+00:00"},{"alias_kind":"pith_short_8","alias_value":"T42VE62W","created_at":"2026-07-05T07:56:52.120008+00:00"}],"events":[],"event_summary":{},"paper_claims":[],"inbound_citations":{"count":2,"internal_anchor_count":0,"sample":[{"citing_arxiv_id":"2606.24910","citing_title":"End-to-End Voice Intent Recognition for Spontaneous Human-Drone Interaction with Naive Users","ref_index":45,"is_internal_anchor":false},{"citing_arxiv_id":"2606.08087","citing_title":"Assessing the Energy and Carbon Emissions of Neural Speaker Verification Model in Training and Inference","ref_index":19,"is_internal_anchor":false}]},"formal_canon":{"evidence_count":0,"sample":[],"anchors":[]},"links":{"html":"https://pith.science/pith/T42VE62WITC6M6SGATY5RXVYVR","json":"https://pith.science/pith/T42VE62WITC6M6SGATY5RXVYVR.json","graph_json":"https://pith.science/api/pith-number/T42VE62WITC6M6SGATY5RXVYVR/graph.json","events_json":"https://pith.science/api/pith-number/T42VE62WITC6M6SGATY5RXVYVR/events.json","paper":"https://pith.science/paper/T42VE62W"},"agent_actions":{"view_html":"https://pith.science/pith/T42VE62WITC6M6SGATY5RXVYVR","download_json":"https://pith.science/pith/T42VE62WITC6M6SGATY5RXVYVR.json","view_paper":"https://pith.science/paper/T42VE62W","resolve_alias":"https://pith.science/api/pith-number/resolve?arxiv=2309.05472&json=true","fetch_graph":"https://pith.science/api/pith-number/T42VE62WITC6M6SGATY5RXVYVR/graph.json","fetch_events":"https://pith.science/api/pith-number/T42VE62WITC6M6SGATY5RXVYVR/events.json","actions":{"anchor_timestamp":"https://pith.science/pith/T42VE62WITC6M6SGATY5RXVYVR/action/timestamp_anchor","attest_storage":"https://pith.science/pith/T42VE62WITC6M6SGATY5RXVYVR/action/storage_attestation","attest_author":"https://pith.science/pith/T42VE62WITC6M6SGATY5RXVYVR/action/author_attestation","sign_citation":"https://pith.science/pith/T42VE62WITC6M6SGATY5RXVYVR/action/citation_signature","submit_replication":"https://pith.science/pith/T42VE62WITC6M6SGATY5RXVYVR/action/replication_record"}},"created_at":"2026-07-05T07:56:52.120008+00:00","updated_at":"2026-07-05T07:56:52.120008+00:00"}