{"record_type":"pith_number_record","schema_url":"https://pith.science/schemas/pith-number/v1.json","pith_number":"pith:2022:GUY4XHJ4U7UBVFRGAHU7ATOINE","short_pith_number":"pith:GUY4XHJ4","schema_version":"1.0","canonical_sha256":"3531cb9d3ca7e81a962601e9f04dc8692faa9709f1bf376d0b08d4cefb5264c3","source":{"kind":"arxiv","id":"2203.03073","version":2},"attestation_state":"computed","paper":{"title":"ILDAE: Instance-Level Difficulty Analysis of Evaluation Data","license":"http://creativecommons.org/licenses/by-nc-sa/4.0/","headline":"","cross_cats":["cs.AI","cs.LG"],"primary_cat":"cs.CL","authors_text":"Chitta Baral, Neeraj Varshney, Swaroop Mishra","submitted_at":"2022-03-07T00:02:11Z","abstract_excerpt":"Knowledge of questions' difficulty level helps a teacher in several ways, such as estimating students' potential quickly by asking carefully selected questions and improving quality of examination by modifying trivial and hard questions. Can we extract such benefits of instance difficulty in NLP? To this end, we conduct Instance-Level Difficulty Analysis of Evaluation data (ILDAE) in a large-scale setup of 23 datasets and demonstrate its five novel applications: 1) conducting efficient-yet-accurate evaluations with fewer instances saving computational cost and time, 2) improving quality of exi"},"verification_status":{"content_addressed":true,"pith_receipt":true,"author_attested":false,"weak_author_claims":0,"strong_author_claims":0,"externally_anchored":false,"storage_verified":false,"citation_signatures":0,"replication_records":0,"graph_snapshot":true,"references_resolved":false,"formal_links_present":false},"canonical_record":{"source":{"id":"2203.03073","kind":"arxiv","version":2},"metadata":{"license":"http://creativecommons.org/licenses/by-nc-sa/4.0/","primary_cat":"cs.CL","submitted_at":"2022-03-07T00:02:11Z","cross_cats_sorted":["cs.AI","cs.LG"],"title_canon_sha256":"dd28e861b6cfad1a0dc9f4564d55bd9489ad1db904c7cb47ba816c3347554ee1","abstract_canon_sha256":"bbfe322446cada4bb5066cd28fba510fa1b72fdbc3ef9aff53ddc7a55dfa662d"},"schema_version":"1.0"},"receipt":{"kind":"pith_receipt","key_id":"pith-v1-2026-05","algorithm":"ed25519","signed_at":"2026-07-05T04:03:16.366869Z","signature_b64":"3MX0kyNevkxd5ZOEqcA+zLtSfbVwnTB2ztBdu8Jej2zv/qKBJ4HAiyPBlvRCRSErsgoSGMFp0D1Wwh+1SOFYCA==","signed_message":"canonical_sha256_bytes","builder_version":"pith-number-builder-2026-05-17-v1","receipt_version":"0.3","canonical_sha256":"3531cb9d3ca7e81a962601e9f04dc8692faa9709f1bf376d0b08d4cefb5264c3","last_reissued_at":"2026-07-05T04:03:16.366376Z","signature_status":"signed_v1","first_computed_at":"2026-07-05T04:03:16.366376Z","public_key_fingerprint":"8d4b5ee74e4693bcd1df2446408b0d54"},"graph_snapshot":{"paper":{"title":"ILDAE: Instance-Level Difficulty Analysis of Evaluation Data","license":"http://creativecommons.org/licenses/by-nc-sa/4.0/","headline":"","cross_cats":["cs.AI","cs.LG"],"primary_cat":"cs.CL","authors_text":"Chitta Baral, Neeraj Varshney, Swaroop Mishra","submitted_at":"2022-03-07T00:02:11Z","abstract_excerpt":"Knowledge of questions' difficulty level helps a teacher in several ways, such as estimating students' potential quickly by asking carefully selected questions and improving quality of examination by modifying trivial and hard questions. Can we extract such benefits of instance difficulty in NLP? To this end, we conduct Instance-Level Difficulty Analysis of Evaluation data (ILDAE) in a large-scale setup of 23 datasets and demonstrate its five novel applications: 1) conducting efficient-yet-accurate evaluations with fewer instances saving computational cost and time, 2) improving quality of exi"},"claims":{"count":0,"items":[],"snapshot_sha256":"258153158e38e3291e3d48162225fcdb2d5a3ed65a07baac614ab91432fd4f57"},"source":{"id":"2203.03073","kind":"arxiv","version":2},"verdict":{"id":null,"model_set":{},"created_at":null,"strongest_claim":"","one_line_summary":"","pipeline_version":null,"weakest_assumption":"","pith_extraction_headline":""},"integrity":{"clean":true,"summary":{"advisory":0,"critical":0,"by_detector":{},"informational":0},"endpoint":"/pith/2203.03073/integrity.json","findings":[],"available":true,"detectors_run":[],"snapshot_sha256":"c28c3603d3b5d939e8dc4c7e95fa8dfce3d595e45f758748cecf8e644a296938"},"references":{"count":0,"sample":[],"resolved_work":0,"snapshot_sha256":"258153158e38e3291e3d48162225fcdb2d5a3ed65a07baac614ab91432fd4f57","internal_anchors":0},"formal_canon":{"evidence_count":0,"snapshot_sha256":"258153158e38e3291e3d48162225fcdb2d5a3ed65a07baac614ab91432fd4f57"},"author_claims":{"count":0,"strong_count":0,"snapshot_sha256":"258153158e38e3291e3d48162225fcdb2d5a3ed65a07baac614ab91432fd4f57"},"builder_version":"pith-number-builder-2026-05-17-v1"},"aliases":[{"alias_kind":"arxiv","alias_value":"2203.03073","created_at":"2026-07-05T04:03:16.366437+00:00"},{"alias_kind":"arxiv_version","alias_value":"2203.03073v2","created_at":"2026-07-05T04:03:16.366437+00:00"},{"alias_kind":"doi","alias_value":"10.48550/arxiv.2203.03073","created_at":"2026-07-05T04:03:16.366437+00:00"},{"alias_kind":"pith_short_12","alias_value":"GUY4XHJ4U7UB","created_at":"2026-07-05T04:03:16.366437+00:00"},{"alias_kind":"pith_short_16","alias_value":"GUY4XHJ4U7UBVFRG","created_at":"2026-07-05T04:03:16.366437+00:00"},{"alias_kind":"pith_short_8","alias_value":"GUY4XHJ4","created_at":"2026-07-05T04:03:16.366437+00:00"}],"events":[],"event_summary":{},"paper_claims":[],"inbound_citations":{"count":1,"internal_anchor_count":1,"sample":[{"citing_arxiv_id":"2505.12844","citing_title":"AGI-Elo: How Far Are We From Mastering A Task?","ref_index":83,"is_internal_anchor":true}]},"formal_canon":{"evidence_count":0,"sample":[],"anchors":[]},"links":{"html":"https://pith.science/pith/GUY4XHJ4U7UBVFRGAHU7ATOINE","json":"https://pith.science/pith/GUY4XHJ4U7UBVFRGAHU7ATOINE.json","graph_json":"https://pith.science/api/pith-number/GUY4XHJ4U7UBVFRGAHU7ATOINE/graph.json","events_json":"https://pith.science/api/pith-number/GUY4XHJ4U7UBVFRGAHU7ATOINE/events.json","paper":"https://pith.science/paper/GUY4XHJ4"},"agent_actions":{"view_html":"https://pith.science/pith/GUY4XHJ4U7UBVFRGAHU7ATOINE","download_json":"https://pith.science/pith/GUY4XHJ4U7UBVFRGAHU7ATOINE.json","view_paper":"https://pith.science/paper/GUY4XHJ4","resolve_alias":"https://pith.science/api/pith-number/resolve?arxiv=2203.03073&json=true","fetch_graph":"https://pith.science/api/pith-number/GUY4XHJ4U7UBVFRGAHU7ATOINE/graph.json","fetch_events":"https://pith.science/api/pith-number/GUY4XHJ4U7UBVFRGAHU7ATOINE/events.json","actions":{"anchor_timestamp":"https://pith.science/pith/GUY4XHJ4U7UBVFRGAHU7ATOINE/action/timestamp_anchor","attest_storage":"https://pith.science/pith/GUY4XHJ4U7UBVFRGAHU7ATOINE/action/storage_attestation","attest_author":"https://pith.science/pith/GUY4XHJ4U7UBVFRGAHU7ATOINE/action/author_attestation","sign_citation":"https://pith.science/pith/GUY4XHJ4U7UBVFRGAHU7ATOINE/action/citation_signature","submit_replication":"https://pith.science/pith/GUY4XHJ4U7UBVFRGAHU7ATOINE/action/replication_record"}},"created_at":"2026-07-05T04:03:16.366437+00:00","updated_at":"2026-07-05T04:03:16.366437+00:00"}