{"bundle_type":"pith_open_graph_bundle","bundle_version":"1.0","pith_number":"pith:2019:5HPMQPEJI5FFGVQVBPS4GPHBA4","short_pith_number":"pith:5HPMQPEJ","canonical_record":{"source":{"id":"1902.09476","kind":"arxiv","version":1},"metadata":{"license":"http://arxiv.org/licenses/nonexclusive-distrib/1.0/","primary_cat":"cs.CL","submitted_at":"2019-02-25T17:53:20Z","cross_cats_sorted":["cs.LG"],"title_canon_sha256":"be2ca30d9f4b8dc10954bad9c861e52b20935e8bb12fa8148bfc6378505c9d90","abstract_canon_sha256":"231dd17789409a9b5fafcdc2afec6aa8c3c81167594ae23cce913d05ffaafa03"},"schema_version":"1.0"},"canonical_sha256":"e9dec83c89474a5356150be5c33ce1072dbe63cebb177964dca451a773125f92","source":{"kind":"arxiv","id":"1902.09476","version":1},"source_aliases":[{"alias_kind":"arxiv","alias_value":"1902.09476","created_at":"2026-05-17T23:52:45Z"},{"alias_kind":"arxiv_version","alias_value":"1902.09476v1","created_at":"2026-05-17T23:52:45Z"},{"alias_kind":"doi","alias_value":"10.48550/arxiv.1902.09476","created_at":"2026-05-17T23:52:45Z"},{"alias_kind":"pith_short_12","alias_value":"5HPMQPEJI5FF","created_at":"2026-05-18T12:33:10Z"},{"alias_kind":"pith_short_16","alias_value":"5HPMQPEJI5FFGVQV","created_at":"2026-05-18T12:33:10Z"},{"alias_kind":"pith_short_8","alias_value":"5HPMQPEJ","created_at":"2026-05-18T12:33:10Z"}],"events":[{"event_type":"record_created","subject_pith_number":"pith:2019:5HPMQPEJI5FFGVQVBPS4GPHBA4","target":"record","payload":{"canonical_record":{"source":{"id":"1902.09476","kind":"arxiv","version":1},"metadata":{"license":"http://arxiv.org/licenses/nonexclusive-distrib/1.0/","primary_cat":"cs.CL","submitted_at":"2019-02-25T17:53:20Z","cross_cats_sorted":["cs.LG"],"title_canon_sha256":"be2ca30d9f4b8dc10954bad9c861e52b20935e8bb12fa8148bfc6378505c9d90","abstract_canon_sha256":"231dd17789409a9b5fafcdc2afec6aa8c3c81167594ae23cce913d05ffaafa03"},"schema_version":"1.0"},"canonical_sha256":"e9dec83c89474a5356150be5c33ce1072dbe63cebb177964dca451a773125f92","receipt":{"kind":"pith_receipt","key_id":"pith-v1-2026-05","algorithm":"ed25519","signed_at":"2026-05-17T23:52:45.654311Z","signature_b64":"1yM018jnqNupmhbUghn6nXRQReXGvDSSiL0OO8LDrga4fVvLvDlUFglkWXdjppbQthvu7HjA7XWOWQXfTKJOBg==","signed_message":"canonical_sha256_bytes","builder_version":"pith-number-builder-2026-05-17-v1","receipt_version":"0.3","canonical_sha256":"e9dec83c89474a5356150be5c33ce1072dbe63cebb177964dca451a773125f92","last_reissued_at":"2026-05-17T23:52:45.653703Z","signature_status":"signed_v1","first_computed_at":"2026-05-17T23:52:45.653703Z","public_key_fingerprint":"8d4b5ee74e4693bcd1df2446408b0d54"},"source_kind":"arxiv","source_id":"1902.09476","source_version":1,"attestation_state":"computed"},"signer":{"signer_id":"pith.science","signer_type":"pith_registry","key_id":"pith-v1-2026-05","public_key_fingerprint":"8d4b5ee74e4693bcd1df2446408b0d54"},"created_at":"2026-05-17T23:52:45Z","supersedes":[],"prev_event":null,"signature":{"signature_status":"signed_v1","algorithm":"ed25519","key_id":"pith-v1-2026-05","public_key_fingerprint":"8d4b5ee74e4693bcd1df2446408b0d54","signature_b64":"GXV/NdlZDpZ5Y0P9njWc5KsfTgsfCdDXdz+uRvrD3fin4vD5iRwdff8dRCxQtuIqcl0OSsuoogcgE2If5dbvAg==","signed_message":"open_graph_event_sha256_bytes","signed_at":"2026-05-25T12:22:45.777855Z"},"content_sha256":"240ae223f7d43ce76abd962e4ee91a9d58fa59af0ee7eb0c6b609d235fd1e0fa","schema_version":"1.0","event_id":"sha256:240ae223f7d43ce76abd962e4ee91a9d58fa59af0ee7eb0c6b609d235fd1e0fa"},{"event_type":"graph_snapshot","subject_pith_number":"pith:2019:5HPMQPEJI5FFGVQVBPS4GPHBA4","target":"graph","payload":{"graph_snapshot":{"paper":{"title":"MedMentions: A Large Biomedical Corpus Annotated with UMLS Concepts","license":"http://arxiv.org/licenses/nonexclusive-distrib/1.0/","headline":"","cross_cats":["cs.LG"],"primary_cat":"cs.CL","authors_text":"Donghui Li, Sunil Mohan","submitted_at":"2019-02-25T17:53:20Z","abstract_excerpt":"This paper presents the formal release of MedMentions, a new manually annotated resource for the recognition of biomedical concepts. What distinguishes MedMentions from other annotated biomedical corpora is its size (over 4,000 abstracts and over 350,000 linked mentions), as well as the size of the concept ontology (over 3 million concepts from UMLS 2017) and its broad coverage of biomedical disciplines. In addition to the full corpus, a sub-corpus of MedMentions is also presented, comprising annotations for a subset of UMLS 2017 targeted towards document retrieval. To encourage research in Bi"},"claims":{"count":0,"items":[],"snapshot_sha256":"258153158e38e3291e3d48162225fcdb2d5a3ed65a07baac614ab91432fd4f57"},"source":{"id":"1902.09476","kind":"arxiv","version":1},"verdict":{"id":null,"model_set":{},"created_at":null,"strongest_claim":"","one_line_summary":"","pipeline_version":null,"weakest_assumption":"","pith_extraction_headline":""},"references":{"count":0,"sample":[],"resolved_work":0,"snapshot_sha256":"258153158e38e3291e3d48162225fcdb2d5a3ed65a07baac614ab91432fd4f57","internal_anchors":0},"formal_canon":{"evidence_count":0,"snapshot_sha256":"258153158e38e3291e3d48162225fcdb2d5a3ed65a07baac614ab91432fd4f57"},"author_claims":{"count":0,"strong_count":0,"snapshot_sha256":"258153158e38e3291e3d48162225fcdb2d5a3ed65a07baac614ab91432fd4f57"},"builder_version":"pith-number-builder-2026-05-17-v1"},"verdict_id":null},"signer":{"signer_id":"pith.science","signer_type":"pith_registry","key_id":"pith-v1-2026-05","public_key_fingerprint":"8d4b5ee74e4693bcd1df2446408b0d54"},"created_at":"2026-05-17T23:52:45Z","supersedes":[],"prev_event":null,"signature":{"signature_status":"signed_v1","algorithm":"ed25519","key_id":"pith-v1-2026-05","public_key_fingerprint":"8d4b5ee74e4693bcd1df2446408b0d54","signature_b64":"kfi+WBTzNXzYcV/QfcfHkGA7aYthLjaOs9Lgv6PUSjVsat6bY+UCCF0+1YpbgXN8p92Uv3x+aW4rs5OrD+nsDw==","signed_message":"open_graph_event_sha256_bytes","signed_at":"2026-05-25T12:22:45.778432Z"},"content_sha256":"6c0e58629ee224c1d79f101e26a7f630542e72cb3adc6fe017935f9d2d5d7216","schema_version":"1.0","event_id":"sha256:6c0e58629ee224c1d79f101e26a7f630542e72cb3adc6fe017935f9d2d5d7216"}],"timestamp_proofs":[],"mirror_hints":[{"mirror_type":"https","name":"Pith Resolver","base_url":"https://pith.science","bundle_url":"https://pith.science/pith/5HPMQPEJI5FFGVQVBPS4GPHBA4/bundle.json","state_url":"https://pith.science/pith/5HPMQPEJI5FFGVQVBPS4GPHBA4/state.json","well_known_bundle_url":"https://pith.science/.well-known/pith/5HPMQPEJI5FFGVQVBPS4GPHBA4/bundle.json","status":"primary"}],"public_keys":[{"key_id":"pith-v1-2026-05","algorithm":"ed25519","format":"raw","public_key_b64":"stVStoiQhXFxp4s2pdzPNoqVNBMojDU/fJ2db5S3CbM=","public_key_hex":"b2d552b68890857171a78b36a5dccf368a953413288c353f7c9d9d6f94b709b3","fingerprint_sha256_b32_first128bits":"RVFV5Z2OI2J3ZUO7ERDEBCYNKS","fingerprint_sha256_hex":"8d4b5ee74e4693bcd1df2446408b0d54","rotates_at":null,"url":"https://pith.science/pith-signing-key.json","notes":"Pith uses this Ed25519 key to sign canonical record SHA-256 digests. Verify with: ed25519_verify(public_key, message=canonical_sha256_bytes, signature=base64decode(signature_b64))."}],"merge_version":"pith-open-graph-merge-v1","built_at":"2026-05-25T12:22:45Z","links":{"resolver":"https://pith.science/pith/5HPMQPEJI5FFGVQVBPS4GPHBA4","bundle":"https://pith.science/pith/5HPMQPEJI5FFGVQVBPS4GPHBA4/bundle.json","state":"https://pith.science/pith/5HPMQPEJI5FFGVQVBPS4GPHBA4/state.json","well_known_bundle":"https://pith.science/.well-known/pith/5HPMQPEJI5FFGVQVBPS4GPHBA4/bundle.json"},"state":{"state_type":"pith_open_graph_state","state_version":"1.0","pith_number":"pith:2019:5HPMQPEJI5FFGVQVBPS4GPHBA4","merge_version":"pith-open-graph-merge-v1","event_count":2,"valid_event_count":2,"invalid_event_count":0,"equivocation_count":0,"current":{"canonical_record":{"metadata":{"abstract_canon_sha256":"231dd17789409a9b5fafcdc2afec6aa8c3c81167594ae23cce913d05ffaafa03","cross_cats_sorted":["cs.LG"],"license":"http://arxiv.org/licenses/nonexclusive-distrib/1.0/","primary_cat":"cs.CL","submitted_at":"2019-02-25T17:53:20Z","title_canon_sha256":"be2ca30d9f4b8dc10954bad9c861e52b20935e8bb12fa8148bfc6378505c9d90"},"schema_version":"1.0","source":{"id":"1902.09476","kind":"arxiv","version":1}},"source_aliases":[{"alias_kind":"arxiv","alias_value":"1902.09476","created_at":"2026-05-17T23:52:45Z"},{"alias_kind":"arxiv_version","alias_value":"1902.09476v1","created_at":"2026-05-17T23:52:45Z"},{"alias_kind":"doi","alias_value":"10.48550/arxiv.1902.09476","created_at":"2026-05-17T23:52:45Z"},{"alias_kind":"pith_short_12","alias_value":"5HPMQPEJI5FF","created_at":"2026-05-18T12:33:10Z"},{"alias_kind":"pith_short_16","alias_value":"5HPMQPEJI5FFGVQV","created_at":"2026-05-18T12:33:10Z"},{"alias_kind":"pith_short_8","alias_value":"5HPMQPEJ","created_at":"2026-05-18T12:33:10Z"}],"graph_snapshots":[{"event_id":"sha256:6c0e58629ee224c1d79f101e26a7f630542e72cb3adc6fe017935f9d2d5d7216","target":"graph","created_at":"2026-05-17T23:52:45Z","signer":{"key_id":"pith-v1-2026-05","public_key_fingerprint":"8d4b5ee74e4693bcd1df2446408b0d54","signer_id":"pith.science","signer_type":"pith_registry"},"payload":{"graph_snapshot":{"author_claims":{"count":0,"snapshot_sha256":"258153158e38e3291e3d48162225fcdb2d5a3ed65a07baac614ab91432fd4f57","strong_count":0},"builder_version":"pith-number-builder-2026-05-17-v1","claims":{"count":0,"items":[],"snapshot_sha256":"258153158e38e3291e3d48162225fcdb2d5a3ed65a07baac614ab91432fd4f57"},"formal_canon":{"evidence_count":0,"snapshot_sha256":"258153158e38e3291e3d48162225fcdb2d5a3ed65a07baac614ab91432fd4f57"},"paper":{"abstract_excerpt":"This paper presents the formal release of MedMentions, a new manually annotated resource for the recognition of biomedical concepts. What distinguishes MedMentions from other annotated biomedical corpora is its size (over 4,000 abstracts and over 350,000 linked mentions), as well as the size of the concept ontology (over 3 million concepts from UMLS 2017) and its broad coverage of biomedical disciplines. In addition to the full corpus, a sub-corpus of MedMentions is also presented, comprising annotations for a subset of UMLS 2017 targeted towards document retrieval. To encourage research in Bi","authors_text":"Donghui Li, Sunil Mohan","cross_cats":["cs.LG"],"headline":"","license":"http://arxiv.org/licenses/nonexclusive-distrib/1.0/","primary_cat":"cs.CL","submitted_at":"2019-02-25T17:53:20Z","title":"MedMentions: A Large Biomedical Corpus Annotated with UMLS Concepts"},"references":{"count":0,"internal_anchors":0,"resolved_work":0,"sample":[],"snapshot_sha256":"258153158e38e3291e3d48162225fcdb2d5a3ed65a07baac614ab91432fd4f57"},"source":{"id":"1902.09476","kind":"arxiv","version":1},"verdict":{"created_at":null,"id":null,"model_set":{},"one_line_summary":"","pipeline_version":null,"pith_extraction_headline":"","strongest_claim":"","weakest_assumption":""}},"verdict_id":null}}],"author_attestations":[],"timestamp_anchors":[],"storage_attestations":[],"citation_signatures":[],"replication_records":[],"corrections":[],"mirror_hints":[],"record_created":{"event_id":"sha256:240ae223f7d43ce76abd962e4ee91a9d58fa59af0ee7eb0c6b609d235fd1e0fa","target":"record","created_at":"2026-05-17T23:52:45Z","signer":{"key_id":"pith-v1-2026-05","public_key_fingerprint":"8d4b5ee74e4693bcd1df2446408b0d54","signer_id":"pith.science","signer_type":"pith_registry"},"payload":{"attestation_state":"computed","canonical_record":{"metadata":{"abstract_canon_sha256":"231dd17789409a9b5fafcdc2afec6aa8c3c81167594ae23cce913d05ffaafa03","cross_cats_sorted":["cs.LG"],"license":"http://arxiv.org/licenses/nonexclusive-distrib/1.0/","primary_cat":"cs.CL","submitted_at":"2019-02-25T17:53:20Z","title_canon_sha256":"be2ca30d9f4b8dc10954bad9c861e52b20935e8bb12fa8148bfc6378505c9d90"},"schema_version":"1.0","source":{"id":"1902.09476","kind":"arxiv","version":1}},"canonical_sha256":"e9dec83c89474a5356150be5c33ce1072dbe63cebb177964dca451a773125f92","receipt":{"algorithm":"ed25519","builder_version":"pith-number-builder-2026-05-17-v1","canonical_sha256":"e9dec83c89474a5356150be5c33ce1072dbe63cebb177964dca451a773125f92","first_computed_at":"2026-05-17T23:52:45.653703Z","key_id":"pith-v1-2026-05","kind":"pith_receipt","last_reissued_at":"2026-05-17T23:52:45.653703Z","public_key_fingerprint":"8d4b5ee74e4693bcd1df2446408b0d54","receipt_version":"0.3","signature_b64":"1yM018jnqNupmhbUghn6nXRQReXGvDSSiL0OO8LDrga4fVvLvDlUFglkWXdjppbQthvu7HjA7XWOWQXfTKJOBg==","signature_status":"signed_v1","signed_at":"2026-05-17T23:52:45.654311Z","signed_message":"canonical_sha256_bytes"},"source_id":"1902.09476","source_kind":"arxiv","source_version":1}}},"equivocations":[],"invalid_events":[],"applied_event_ids":["sha256:240ae223f7d43ce76abd962e4ee91a9d58fa59af0ee7eb0c6b609d235fd1e0fa","sha256:6c0e58629ee224c1d79f101e26a7f630542e72cb3adc6fe017935f9d2d5d7216"],"state_sha256":"4a1ff697e2004b6eb3414492e9c6b1a53dd80d60ae477d6fc7ac7f02dd0675fe"},"bundle_signature":{"signature_status":"signed_v1","algorithm":"ed25519","key_id":"pith-v1-2026-05","public_key_fingerprint":"8d4b5ee74e4693bcd1df2446408b0d54","signature_b64":"8BuNmizAo7uXztxjTslCzRZEYeQWH0bDa1/dxFyR64FwVn5qty63mdwduSjd6HcWDS/RnA8YMralNW2JUs/fDw==","signed_message":"bundle_sha256_bytes","signed_at":"2026-05-25T12:22:45.781580Z","bundle_sha256":"172d0937d27cd64536c4af0ab97b0e199d2774c55d534b50492849aec3887ccc"}}