{"record_type":"pith_number_record","schema_url":"https://pith.science/schemas/pith-number/v1.json","pith_number":"pith:2019:YKGAFO2O5M7PYG25BUY6LCPUT3","short_pith_number":"pith:YKGAFO2O","schema_version":"1.0","canonical_sha256":"c28c02bb4eeb3efc1b5d0d31e589f49ed95a90725d709abf5bbca4a3ff6412a8","source":{"kind":"arxiv","id":"1903.00724","version":1},"attestation_state":"computed","paper":{"title":"Predicting and interpreting embeddings for out of vocabulary words in downstream tasks","license":"http://arxiv.org/licenses/nonexclusive-distrib/1.0/","headline":"","cross_cats":["cs.LG"],"primary_cat":"cs.CL","authors_text":"Jean-Samuel Leboeuf, Luc Lamontagne, Nicolas Garneau","submitted_at":"2019-03-02T15:32:39Z","abstract_excerpt":"We propose a novel way to handle out of vocabulary (OOV) words in downstream natural language processing (NLP) tasks. We implement a network that predicts useful embeddings for OOV words based on their morphology and on the context in which they appear. Our model also incorporates an attention mechanism indicating the focus allocated to the left context words, the right context words or the word's characters, hence making the prediction more interpretable. The model is a ``drop-in'' module that is jointly trained with the downstream task's neural network, thus producing embeddings specialized "},"verification_status":{"content_addressed":true,"pith_receipt":true,"author_attested":false,"weak_author_claims":0,"strong_author_claims":0,"externally_anchored":false,"storage_verified":false,"citation_signatures":0,"replication_records":0,"graph_snapshot":true,"references_resolved":false,"formal_links_present":false},"canonical_record":{"source":{"id":"1903.00724","kind":"arxiv","version":1},"metadata":{"license":"http://arxiv.org/licenses/nonexclusive-distrib/1.0/","primary_cat":"cs.CL","submitted_at":"2019-03-02T15:32:39Z","cross_cats_sorted":["cs.LG"],"title_canon_sha256":"3d515e920c0f9be699669e1b968c92551bc65e61a02fb180dde111bc62dcf666","abstract_canon_sha256":"d07994b49a67a66f5cde6313a01666be102a9724ee8b2e59576a00f2870196ed"},"schema_version":"1.0"},"receipt":{"kind":"pith_receipt","key_id":"pith-v1-2026-05","algorithm":"ed25519","signed_at":"2026-05-17T23:52:13.494557Z","signature_b64":"GFG9FIyoRufOnle5jtvYKPvfjQC+1Hed8lsbudbVamBxekcuNQ75uUXKAQrB+DKnXlLUhrEbX4vgfvg5YCTnDA==","signed_message":"canonical_sha256_bytes","builder_version":"pith-number-builder-2026-05-17-v1","receipt_version":"0.3","canonical_sha256":"c28c02bb4eeb3efc1b5d0d31e589f49ed95a90725d709abf5bbca4a3ff6412a8","last_reissued_at":"2026-05-17T23:52:13.493822Z","signature_status":"signed_v1","first_computed_at":"2026-05-17T23:52:13.493822Z","public_key_fingerprint":"8d4b5ee74e4693bcd1df2446408b0d54"},"graph_snapshot":{"paper":{"title":"Predicting and interpreting embeddings for out of vocabulary words in downstream tasks","license":"http://arxiv.org/licenses/nonexclusive-distrib/1.0/","headline":"","cross_cats":["cs.LG"],"primary_cat":"cs.CL","authors_text":"Jean-Samuel Leboeuf, Luc Lamontagne, Nicolas Garneau","submitted_at":"2019-03-02T15:32:39Z","abstract_excerpt":"We propose a novel way to handle out of vocabulary (OOV) words in downstream natural language processing (NLP) tasks. We implement a network that predicts useful embeddings for OOV words based on their morphology and on the context in which they appear. Our model also incorporates an attention mechanism indicating the focus allocated to the left context words, the right context words or the word's characters, hence making the prediction more interpretable. The model is a ``drop-in'' module that is jointly trained with the downstream task's neural network, thus producing embeddings specialized "},"claims":{"count":0,"items":[],"snapshot_sha256":"258153158e38e3291e3d48162225fcdb2d5a3ed65a07baac614ab91432fd4f57"},"source":{"id":"1903.00724","kind":"arxiv","version":1},"verdict":{"id":null,"model_set":{},"created_at":null,"strongest_claim":"","one_line_summary":"","pipeline_version":null,"weakest_assumption":"","pith_extraction_headline":""},"references":{"count":0,"sample":[],"resolved_work":0,"snapshot_sha256":"258153158e38e3291e3d48162225fcdb2d5a3ed65a07baac614ab91432fd4f57","internal_anchors":0},"formal_canon":{"evidence_count":0,"snapshot_sha256":"258153158e38e3291e3d48162225fcdb2d5a3ed65a07baac614ab91432fd4f57"},"author_claims":{"count":0,"strong_count":0,"snapshot_sha256":"258153158e38e3291e3d48162225fcdb2d5a3ed65a07baac614ab91432fd4f57"},"builder_version":"pith-number-builder-2026-05-17-v1"},"aliases":[{"alias_kind":"arxiv","alias_value":"1903.00724","created_at":"2026-05-17T23:52:13.493945+00:00"},{"alias_kind":"arxiv_version","alias_value":"1903.00724v1","created_at":"2026-05-17T23:52:13.493945+00:00"},{"alias_kind":"doi","alias_value":"10.48550/arxiv.1903.00724","created_at":"2026-05-17T23:52:13.493945+00:00"},{"alias_kind":"pith_short_12","alias_value":"YKGAFO2O5M7P","created_at":"2026-05-18T12:33:33.725879+00:00"},{"alias_kind":"pith_short_16","alias_value":"YKGAFO2O5M7PYG25","created_at":"2026-05-18T12:33:33.725879+00:00"},{"alias_kind":"pith_short_8","alias_value":"YKGAFO2O","created_at":"2026-05-18T12:33:33.725879+00:00"}],"events":[],"event_summary":{},"paper_claims":[],"inbound_citations":{"count":0,"internal_anchor_count":0,"sample":[]},"formal_canon":{"evidence_count":0,"sample":[],"anchors":[]},"links":{"html":"https://pith.science/pith/YKGAFO2O5M7PYG25BUY6LCPUT3","json":"https://pith.science/pith/YKGAFO2O5M7PYG25BUY6LCPUT3.json","graph_json":"https://pith.science/api/pith-number/YKGAFO2O5M7PYG25BUY6LCPUT3/graph.json","events_json":"https://pith.science/api/pith-number/YKGAFO2O5M7PYG25BUY6LCPUT3/events.json","paper":"https://pith.science/paper/YKGAFO2O"},"agent_actions":{"view_html":"https://pith.science/pith/YKGAFO2O5M7PYG25BUY6LCPUT3","download_json":"https://pith.science/pith/YKGAFO2O5M7PYG25BUY6LCPUT3.json","view_paper":"https://pith.science/paper/YKGAFO2O","resolve_alias":"https://pith.science/api/pith-number/resolve?arxiv=1903.00724&json=true","fetch_graph":"https://pith.science/api/pith-number/YKGAFO2O5M7PYG25BUY6LCPUT3/graph.json","fetch_events":"https://pith.science/api/pith-number/YKGAFO2O5M7PYG25BUY6LCPUT3/events.json","actions":{"anchor_timestamp":"https://pith.science/pith/YKGAFO2O5M7PYG25BUY6LCPUT3/action/timestamp_anchor","attest_storage":"https://pith.science/pith/YKGAFO2O5M7PYG25BUY6LCPUT3/action/storage_attestation","attest_author":"https://pith.science/pith/YKGAFO2O5M7PYG25BUY6LCPUT3/action/author_attestation","sign_citation":"https://pith.science/pith/YKGAFO2O5M7PYG25BUY6LCPUT3/action/citation_signature","submit_replication":"https://pith.science/pith/YKGAFO2O5M7PYG25BUY6LCPUT3/action/replication_record"}},"created_at":"2026-05-17T23:52:13.493945+00:00","updated_at":"2026-05-17T23:52:13.493945+00:00"}