{"bundle_type":"pith_open_graph_bundle","bundle_version":"1.0","pith_number":"pith:2019:YAD46YSPDR324DNFSO5T7WSINO","short_pith_number":"pith:YAD46YSP","canonical_record":{"source":{"id":"1903.07435","kind":"arxiv","version":2},"metadata":{"license":"http://arxiv.org/licenses/nonexclusive-distrib/1.0/","primary_cat":"cs.CL","submitted_at":"2019-03-18T13:38:54Z","cross_cats_sorted":[],"title_canon_sha256":"7d7f614b3457d885852482e14131f1bb114f989a8ebe205bee21a681c9363b59","abstract_canon_sha256":"24e6a046b43ff4c093c241b8959f77acb32126d6d8f26a6c855a64344e64c9bf"},"schema_version":"1.0"},"canonical_sha256":"c007cf624f1c77ae0da593bb3fda486ba110b80437a63ccff1e3bcfbe0d91c4f","source":{"kind":"arxiv","id":"1903.07435","version":2},"source_aliases":[{"alias_kind":"arxiv","alias_value":"1903.07435","created_at":"2026-05-17T23:49:38Z"},{"alias_kind":"arxiv_version","alias_value":"1903.07435v2","created_at":"2026-05-17T23:49:38Z"},{"alias_kind":"doi","alias_value":"10.48550/arxiv.1903.07435","created_at":"2026-05-17T23:49:38Z"},{"alias_kind":"pith_short_12","alias_value":"YAD46YSPDR32","created_at":"2026-05-18T12:33:33Z"},{"alias_kind":"pith_short_16","alias_value":"YAD46YSPDR324DNF","created_at":"2026-05-18T12:33:33Z"},{"alias_kind":"pith_short_8","alias_value":"YAD46YSP","created_at":"2026-05-18T12:33:33Z"}],"events":[{"event_type":"record_created","subject_pith_number":"pith:2019:YAD46YSPDR324DNFSO5T7WSINO","target":"record","payload":{"canonical_record":{"source":{"id":"1903.07435","kind":"arxiv","version":2},"metadata":{"license":"http://arxiv.org/licenses/nonexclusive-distrib/1.0/","primary_cat":"cs.CL","submitted_at":"2019-03-18T13:38:54Z","cross_cats_sorted":[],"title_canon_sha256":"7d7f614b3457d885852482e14131f1bb114f989a8ebe205bee21a681c9363b59","abstract_canon_sha256":"24e6a046b43ff4c093c241b8959f77acb32126d6d8f26a6c855a64344e64c9bf"},"schema_version":"1.0"},"canonical_sha256":"c007cf624f1c77ae0da593bb3fda486ba110b80437a63ccff1e3bcfbe0d91c4f","receipt":{"kind":"pith_receipt","key_id":"pith-v1-2026-05","algorithm":"ed25519","signed_at":"2026-05-17T23:49:38.504102Z","signature_b64":"/tiuk3TrH7CztLi4LqZvlBc05QgIEy+r074E8DjhcrTSrAnPkRPGJD0WIQQL79w980Yi8fAM+prXYQMEvxicBA==","signed_message":"canonical_sha256_bytes","builder_version":"pith-number-builder-2026-05-17-v1","receipt_version":"0.3","canonical_sha256":"c007cf624f1c77ae0da593bb3fda486ba110b80437a63ccff1e3bcfbe0d91c4f","last_reissued_at":"2026-05-17T23:49:38.503556Z","signature_status":"signed_v1","first_computed_at":"2026-05-17T23:49:38.503556Z","public_key_fingerprint":"8d4b5ee74e4693bcd1df2446408b0d54"},"source_kind":"arxiv","source_id":"1903.07435","source_version":2,"attestation_state":"computed"},"signer":{"signer_id":"pith.science","signer_type":"pith_registry","key_id":"pith-v1-2026-05","public_key_fingerprint":"8d4b5ee74e4693bcd1df2446408b0d54"},"created_at":"2026-05-17T23:49:38Z","supersedes":[],"prev_event":null,"signature":{"signature_status":"signed_v1","algorithm":"ed25519","key_id":"pith-v1-2026-05","public_key_fingerprint":"8d4b5ee74e4693bcd1df2446408b0d54","signature_b64":"685B0epnHPttybwI5aZLFdP/DwT5M6de3BzK+6KvAlSS9LC5C0VcPoT1sEHJf6TiALMD7qlvFQLwRvdsnxfZAw==","signed_message":"open_graph_event_sha256_bytes","signed_at":"2026-06-03T20:38:27.277565Z"},"content_sha256":"e2576d8a5260ac9b997ab62b891f523b13ec5d9a5f2553eaaf54987baa894512","schema_version":"1.0","event_id":"sha256:e2576d8a5260ac9b997ab62b891f523b13ec5d9a5f2553eaaf54987baa894512"},{"event_type":"graph_snapshot","subject_pith_number":"pith:2019:YAD46YSPDR324DNFSO5T7WSINO","target":"graph","payload":{"graph_snapshot":{"paper":{"title":"The emergence of number and syntax units in LSTM language models","license":"http://arxiv.org/licenses/nonexclusive-distrib/1.0/","headline":"","cross_cats":[],"primary_cat":"cs.CL","authors_text":"Dieuwke Hupkes, German Kruszewski, Marco Baroni, Stanislas Dehaene, Theo Desbordes, Yair Lakretz","submitted_at":"2019-03-18T13:38:54Z","abstract_excerpt":"Recent work has shown that LSTMs trained on a generic language modeling objective capture syntax-sensitive generalizations such as long-distance number agreement. We have however no mechanistic understanding of how they accomplish this remarkable feat. Some have conjectured it depends on heuristics that do not truly take hierarchical structure into account. We present here a detailed study of the inner mechanics of number tracking in LSTMs at the single neuron level. We discover that long-distance number information is largely managed by two `number units'. Importantly, the behaviour of these "},"claims":{"count":0,"items":[],"snapshot_sha256":"258153158e38e3291e3d48162225fcdb2d5a3ed65a07baac614ab91432fd4f57"},"source":{"id":"1903.07435","kind":"arxiv","version":2},"verdict":{"id":null,"model_set":{},"created_at":null,"strongest_claim":"","one_line_summary":"","pipeline_version":null,"weakest_assumption":"","pith_extraction_headline":""},"references":{"count":0,"sample":[],"resolved_work":0,"snapshot_sha256":"258153158e38e3291e3d48162225fcdb2d5a3ed65a07baac614ab91432fd4f57","internal_anchors":0},"formal_canon":{"evidence_count":0,"snapshot_sha256":"258153158e38e3291e3d48162225fcdb2d5a3ed65a07baac614ab91432fd4f57"},"author_claims":{"count":0,"strong_count":0,"snapshot_sha256":"258153158e38e3291e3d48162225fcdb2d5a3ed65a07baac614ab91432fd4f57"},"builder_version":"pith-number-builder-2026-05-17-v1"},"verdict_id":null},"signer":{"signer_id":"pith.science","signer_type":"pith_registry","key_id":"pith-v1-2026-05","public_key_fingerprint":"8d4b5ee74e4693bcd1df2446408b0d54"},"created_at":"2026-05-17T23:49:38Z","supersedes":[],"prev_event":null,"signature":{"signature_status":"signed_v1","algorithm":"ed25519","key_id":"pith-v1-2026-05","public_key_fingerprint":"8d4b5ee74e4693bcd1df2446408b0d54","signature_b64":"uWgsChJEV1Q8L7zlpLTqhsR/p3GB0nHM32BwcbcMQlV1qkCJq+EMtkA6rlCCIy5vvlHYPd3sGQV4KMzg8OiYAQ==","signed_message":"open_graph_event_sha256_bytes","signed_at":"2026-06-03T20:38:27.277915Z"},"content_sha256":"c725b4a1b4b32ebec1b07b9768b41ad1fb66358619d3a53e2c69c3ebe480f9c7","schema_version":"1.0","event_id":"sha256:c725b4a1b4b32ebec1b07b9768b41ad1fb66358619d3a53e2c69c3ebe480f9c7"}],"timestamp_proofs":[],"mirror_hints":[{"mirror_type":"https","name":"Pith Resolver","base_url":"https://pith.science","bundle_url":"https://pith.science/pith/YAD46YSPDR324DNFSO5T7WSINO/bundle.json","state_url":"https://pith.science/pith/YAD46YSPDR324DNFSO5T7WSINO/state.json","well_known_bundle_url":"https://pith.science/.well-known/pith/YAD46YSPDR324DNFSO5T7WSINO/bundle.json","status":"primary"}],"public_keys":[{"key_id":"pith-v1-2026-05","algorithm":"ed25519","format":"raw","public_key_b64":"stVStoiQhXFxp4s2pdzPNoqVNBMojDU/fJ2db5S3CbM=","public_key_hex":"b2d552b68890857171a78b36a5dccf368a953413288c353f7c9d9d6f94b709b3","fingerprint_sha256_b32_first128bits":"RVFV5Z2OI2J3ZUO7ERDEBCYNKS","fingerprint_sha256_hex":"8d4b5ee74e4693bcd1df2446408b0d54","rotates_at":null,"url":"https://pith.science/pith-signing-key.json","notes":"Pith uses this Ed25519 key to sign canonical record SHA-256 digests. Verify with: ed25519_verify(public_key, message=canonical_sha256_bytes, signature=base64decode(signature_b64))."}],"merge_version":"pith-open-graph-merge-v1","built_at":"2026-06-03T20:38:27Z","links":{"resolver":"https://pith.science/pith/YAD46YSPDR324DNFSO5T7WSINO","bundle":"https://pith.science/pith/YAD46YSPDR324DNFSO5T7WSINO/bundle.json","state":"https://pith.science/pith/YAD46YSPDR324DNFSO5T7WSINO/state.json","well_known_bundle":"https://pith.science/.well-known/pith/YAD46YSPDR324DNFSO5T7WSINO/bundle.json"},"state":{"state_type":"pith_open_graph_state","state_version":"1.0","pith_number":"pith:2019:YAD46YSPDR324DNFSO5T7WSINO","merge_version":"pith-open-graph-merge-v1","event_count":2,"valid_event_count":2,"invalid_event_count":0,"equivocation_count":0,"current":{"canonical_record":{"metadata":{"abstract_canon_sha256":"24e6a046b43ff4c093c241b8959f77acb32126d6d8f26a6c855a64344e64c9bf","cross_cats_sorted":[],"license":"http://arxiv.org/licenses/nonexclusive-distrib/1.0/","primary_cat":"cs.CL","submitted_at":"2019-03-18T13:38:54Z","title_canon_sha256":"7d7f614b3457d885852482e14131f1bb114f989a8ebe205bee21a681c9363b59"},"schema_version":"1.0","source":{"id":"1903.07435","kind":"arxiv","version":2}},"source_aliases":[{"alias_kind":"arxiv","alias_value":"1903.07435","created_at":"2026-05-17T23:49:38Z"},{"alias_kind":"arxiv_version","alias_value":"1903.07435v2","created_at":"2026-05-17T23:49:38Z"},{"alias_kind":"doi","alias_value":"10.48550/arxiv.1903.07435","created_at":"2026-05-17T23:49:38Z"},{"alias_kind":"pith_short_12","alias_value":"YAD46YSPDR32","created_at":"2026-05-18T12:33:33Z"},{"alias_kind":"pith_short_16","alias_value":"YAD46YSPDR324DNF","created_at":"2026-05-18T12:33:33Z"},{"alias_kind":"pith_short_8","alias_value":"YAD46YSP","created_at":"2026-05-18T12:33:33Z"}],"graph_snapshots":[{"event_id":"sha256:c725b4a1b4b32ebec1b07b9768b41ad1fb66358619d3a53e2c69c3ebe480f9c7","target":"graph","created_at":"2026-05-17T23:49:38Z","signer":{"key_id":"pith-v1-2026-05","public_key_fingerprint":"8d4b5ee74e4693bcd1df2446408b0d54","signer_id":"pith.science","signer_type":"pith_registry"},"payload":{"graph_snapshot":{"author_claims":{"count":0,"snapshot_sha256":"258153158e38e3291e3d48162225fcdb2d5a3ed65a07baac614ab91432fd4f57","strong_count":0},"builder_version":"pith-number-builder-2026-05-17-v1","claims":{"count":0,"items":[],"snapshot_sha256":"258153158e38e3291e3d48162225fcdb2d5a3ed65a07baac614ab91432fd4f57"},"formal_canon":{"evidence_count":0,"snapshot_sha256":"258153158e38e3291e3d48162225fcdb2d5a3ed65a07baac614ab91432fd4f57"},"paper":{"abstract_excerpt":"Recent work has shown that LSTMs trained on a generic language modeling objective capture syntax-sensitive generalizations such as long-distance number agreement. We have however no mechanistic understanding of how they accomplish this remarkable feat. Some have conjectured it depends on heuristics that do not truly take hierarchical structure into account. We present here a detailed study of the inner mechanics of number tracking in LSTMs at the single neuron level. We discover that long-distance number information is largely managed by two `number units'. Importantly, the behaviour of these ","authors_text":"Dieuwke Hupkes, German Kruszewski, Marco Baroni, Stanislas Dehaene, Theo Desbordes, Yair Lakretz","cross_cats":[],"headline":"","license":"http://arxiv.org/licenses/nonexclusive-distrib/1.0/","primary_cat":"cs.CL","submitted_at":"2019-03-18T13:38:54Z","title":"The emergence of number and syntax units in LSTM language models"},"references":{"count":0,"internal_anchors":0,"resolved_work":0,"sample":[],"snapshot_sha256":"258153158e38e3291e3d48162225fcdb2d5a3ed65a07baac614ab91432fd4f57"},"source":{"id":"1903.07435","kind":"arxiv","version":2},"verdict":{"created_at":null,"id":null,"model_set":{},"one_line_summary":"","pipeline_version":null,"pith_extraction_headline":"","strongest_claim":"","weakest_assumption":""}},"verdict_id":null}}],"author_attestations":[],"timestamp_anchors":[],"storage_attestations":[],"citation_signatures":[],"replication_records":[],"corrections":[],"mirror_hints":[],"record_created":{"event_id":"sha256:e2576d8a5260ac9b997ab62b891f523b13ec5d9a5f2553eaaf54987baa894512","target":"record","created_at":"2026-05-17T23:49:38Z","signer":{"key_id":"pith-v1-2026-05","public_key_fingerprint":"8d4b5ee74e4693bcd1df2446408b0d54","signer_id":"pith.science","signer_type":"pith_registry"},"payload":{"attestation_state":"computed","canonical_record":{"metadata":{"abstract_canon_sha256":"24e6a046b43ff4c093c241b8959f77acb32126d6d8f26a6c855a64344e64c9bf","cross_cats_sorted":[],"license":"http://arxiv.org/licenses/nonexclusive-distrib/1.0/","primary_cat":"cs.CL","submitted_at":"2019-03-18T13:38:54Z","title_canon_sha256":"7d7f614b3457d885852482e14131f1bb114f989a8ebe205bee21a681c9363b59"},"schema_version":"1.0","source":{"id":"1903.07435","kind":"arxiv","version":2}},"canonical_sha256":"c007cf624f1c77ae0da593bb3fda486ba110b80437a63ccff1e3bcfbe0d91c4f","receipt":{"algorithm":"ed25519","builder_version":"pith-number-builder-2026-05-17-v1","canonical_sha256":"c007cf624f1c77ae0da593bb3fda486ba110b80437a63ccff1e3bcfbe0d91c4f","first_computed_at":"2026-05-17T23:49:38.503556Z","key_id":"pith-v1-2026-05","kind":"pith_receipt","last_reissued_at":"2026-05-17T23:49:38.503556Z","public_key_fingerprint":"8d4b5ee74e4693bcd1df2446408b0d54","receipt_version":"0.3","signature_b64":"/tiuk3TrH7CztLi4LqZvlBc05QgIEy+r074E8DjhcrTSrAnPkRPGJD0WIQQL79w980Yi8fAM+prXYQMEvxicBA==","signature_status":"signed_v1","signed_at":"2026-05-17T23:49:38.504102Z","signed_message":"canonical_sha256_bytes"},"source_id":"1903.07435","source_kind":"arxiv","source_version":2}}},"equivocations":[],"invalid_events":[],"applied_event_ids":["sha256:e2576d8a5260ac9b997ab62b891f523b13ec5d9a5f2553eaaf54987baa894512","sha256:c725b4a1b4b32ebec1b07b9768b41ad1fb66358619d3a53e2c69c3ebe480f9c7"],"state_sha256":"f64e528f8ec35e37931d53cb7b74c8aa6f2ff9636b557a3ceca41403335f75e4"},"bundle_signature":{"signature_status":"signed_v1","algorithm":"ed25519","key_id":"pith-v1-2026-05","public_key_fingerprint":"8d4b5ee74e4693bcd1df2446408b0d54","signature_b64":"Icey+Lr/03dVCcTnusffLeZT+xndB8eDy/v6uFkfGEKaGwr2L99YVp4zPqbncitXcGpjGXlAnBe7PLgGFk47BQ==","signed_message":"bundle_sha256_bytes","signed_at":"2026-06-03T20:38:27.280243Z","bundle_sha256":"2e9ab28180407e23df4c8d1039c133805ba7fbf8d99e09fcff6b30af959aa949"}}