{"record_type":"pith_number_record","schema_url":"https://pith.science/schemas/pith-number/v1.json","pith_number":"pith:2022:NISCMCUPQKZEHU3KJDT4FRXWOS","short_pith_number":"pith:NISCMCUP","schema_version":"1.0","canonical_sha256":"6a24260a8f82b243d36a48e7c2c6f67485020c4505bd9495b74efb2b12e15f00","source":{"kind":"arxiv","id":"2212.10769","version":1},"attestation_state":"computed","paper":{"title":"Uncontrolled Lexical Exposure Leads to Overestimation of Compositional Generalization in Pretrained Models","license":"http://creativecommons.org/licenses/by-nc-nd/4.0/","headline":"","cross_cats":[],"primary_cat":"cs.CL","authors_text":"Najoung Kim, Paul Smolensky, Tal Linzen","submitted_at":"2022-12-21T05:02:08Z","abstract_excerpt":"Human linguistic capacity is often characterized by compositionality and the generalization it enables -- human learners can produce and comprehend novel complex expressions by composing known parts. Several benchmarks exploit distributional control across training and test to gauge compositional generalization, where certain lexical items only occur in limited contexts during training. While recent work using these benchmarks suggests that pretrained models achieve impressive generalization performance, we argue that exposure to pretraining data may break the aforementioned distributional con"},"verification_status":{"content_addressed":true,"pith_receipt":true,"author_attested":false,"weak_author_claims":0,"strong_author_claims":0,"externally_anchored":false,"storage_verified":false,"citation_signatures":0,"replication_records":0,"graph_snapshot":true,"references_resolved":false,"formal_links_present":false},"canonical_record":{"source":{"id":"2212.10769","kind":"arxiv","version":1},"metadata":{"license":"http://creativecommons.org/licenses/by-nc-nd/4.0/","primary_cat":"cs.CL","submitted_at":"2022-12-21T05:02:08Z","cross_cats_sorted":[],"title_canon_sha256":"13e4a23a1d46c16c5279177ba5fa2d5d4bdee54e25362413069871633b62bed0","abstract_canon_sha256":"0c5336f9d6b0c28d64064ddfdd1c62b9ac4729d7ca5340580275dd5850b9c43d"},"schema_version":"1.0"},"receipt":{"kind":"pith_receipt","key_id":"pith-v1-2026-05","algorithm":"ed25519","signed_at":"2026-07-05T05:27:22.681031Z","signature_b64":"HztyMCPRnnVrQtaTmiiScuF8wT6bNvBdczsomoGkjA8ZEnsAP03N8/QYJE0hmpTVKqBLqdy67J+3EY8UodO+CA==","signed_message":"canonical_sha256_bytes","builder_version":"pith-number-builder-2026-05-17-v1","receipt_version":"0.3","canonical_sha256":"6a24260a8f82b243d36a48e7c2c6f67485020c4505bd9495b74efb2b12e15f00","last_reissued_at":"2026-07-05T05:27:22.680600Z","signature_status":"signed_v1","first_computed_at":"2026-07-05T05:27:22.680600Z","public_key_fingerprint":"8d4b5ee74e4693bcd1df2446408b0d54"},"graph_snapshot":{"paper":{"title":"Uncontrolled Lexical Exposure Leads to Overestimation of Compositional Generalization in Pretrained Models","license":"http://creativecommons.org/licenses/by-nc-nd/4.0/","headline":"","cross_cats":[],"primary_cat":"cs.CL","authors_text":"Najoung Kim, Paul Smolensky, Tal Linzen","submitted_at":"2022-12-21T05:02:08Z","abstract_excerpt":"Human linguistic capacity is often characterized by compositionality and the generalization it enables -- human learners can produce and comprehend novel complex expressions by composing known parts. Several benchmarks exploit distributional control across training and test to gauge compositional generalization, where certain lexical items only occur in limited contexts during training. While recent work using these benchmarks suggests that pretrained models achieve impressive generalization performance, we argue that exposure to pretraining data may break the aforementioned distributional con"},"claims":{"count":0,"items":[],"snapshot_sha256":"258153158e38e3291e3d48162225fcdb2d5a3ed65a07baac614ab91432fd4f57"},"source":{"id":"2212.10769","kind":"arxiv","version":1},"verdict":{"id":null,"model_set":{},"created_at":null,"strongest_claim":"","one_line_summary":"","pipeline_version":null,"weakest_assumption":"","pith_extraction_headline":""},"integrity":{"clean":true,"summary":{"advisory":0,"critical":0,"by_detector":{},"informational":0},"endpoint":"/pith/2212.10769/integrity.json","findings":[],"available":true,"detectors_run":[],"snapshot_sha256":"c28c3603d3b5d939e8dc4c7e95fa8dfce3d595e45f758748cecf8e644a296938"},"references":{"count":0,"sample":[],"resolved_work":0,"snapshot_sha256":"258153158e38e3291e3d48162225fcdb2d5a3ed65a07baac614ab91432fd4f57","internal_anchors":0},"formal_canon":{"evidence_count":0,"snapshot_sha256":"258153158e38e3291e3d48162225fcdb2d5a3ed65a07baac614ab91432fd4f57"},"author_claims":{"count":0,"strong_count":0,"snapshot_sha256":"258153158e38e3291e3d48162225fcdb2d5a3ed65a07baac614ab91432fd4f57"},"builder_version":"pith-number-builder-2026-05-17-v1"},"aliases":[{"alias_kind":"arxiv","alias_value":"2212.10769","created_at":"2026-07-05T05:27:22.680656+00:00"},{"alias_kind":"arxiv_version","alias_value":"2212.10769v1","created_at":"2026-07-05T05:27:22.680656+00:00"},{"alias_kind":"doi","alias_value":"10.48550/arxiv.2212.10769","created_at":"2026-07-05T05:27:22.680656+00:00"},{"alias_kind":"pith_short_12","alias_value":"NISCMCUPQKZE","created_at":"2026-07-05T05:27:22.680656+00:00"},{"alias_kind":"pith_short_16","alias_value":"NISCMCUPQKZEHU3K","created_at":"2026-07-05T05:27:22.680656+00:00"},{"alias_kind":"pith_short_8","alias_value":"NISCMCUP","created_at":"2026-07-05T05:27:22.680656+00:00"}],"events":[],"event_summary":{},"paper_claims":[],"inbound_citations":{"count":1,"internal_anchor_count":1,"sample":[{"citing_arxiv_id":"2506.03916","citing_title":"Compositional Generalisation for Explainable Hate Speech Detection","ref_index":10,"is_internal_anchor":true}]},"formal_canon":{"evidence_count":0,"sample":[],"anchors":[]},"links":{"html":"https://pith.science/pith/NISCMCUPQKZEHU3KJDT4FRXWOS","json":"https://pith.science/pith/NISCMCUPQKZEHU3KJDT4FRXWOS.json","graph_json":"https://pith.science/api/pith-number/NISCMCUPQKZEHU3KJDT4FRXWOS/graph.json","events_json":"https://pith.science/api/pith-number/NISCMCUPQKZEHU3KJDT4FRXWOS/events.json","paper":"https://pith.science/paper/NISCMCUP"},"agent_actions":{"view_html":"https://pith.science/pith/NISCMCUPQKZEHU3KJDT4FRXWOS","download_json":"https://pith.science/pith/NISCMCUPQKZEHU3KJDT4FRXWOS.json","view_paper":"https://pith.science/paper/NISCMCUP","resolve_alias":"https://pith.science/api/pith-number/resolve?arxiv=2212.10769&json=true","fetch_graph":"https://pith.science/api/pith-number/NISCMCUPQKZEHU3KJDT4FRXWOS/graph.json","fetch_events":"https://pith.science/api/pith-number/NISCMCUPQKZEHU3KJDT4FRXWOS/events.json","actions":{"anchor_timestamp":"https://pith.science/pith/NISCMCUPQKZEHU3KJDT4FRXWOS/action/timestamp_anchor","attest_storage":"https://pith.science/pith/NISCMCUPQKZEHU3KJDT4FRXWOS/action/storage_attestation","attest_author":"https://pith.science/pith/NISCMCUPQKZEHU3KJDT4FRXWOS/action/author_attestation","sign_citation":"https://pith.science/pith/NISCMCUPQKZEHU3KJDT4FRXWOS/action/citation_signature","submit_replication":"https://pith.science/pith/NISCMCUPQKZEHU3KJDT4FRXWOS/action/replication_record"}},"created_at":"2026-07-05T05:27:22.680656+00:00","updated_at":"2026-07-05T05:27:22.680656+00:00"}