{"bundle_type":"pith_open_graph_bundle","bundle_version":"1.0","pith_number":"pith:2018:L3HXCROWLPZPHQ5XZ2Y4CGZAK7","short_pith_number":"pith:L3HXCROW","canonical_record":{"source":{"id":"1807.06638","kind":"arxiv","version":1},"metadata":{"license":"http://arxiv.org/licenses/nonexclusive-distrib/1.0/","primary_cat":"cs.CL","submitted_at":"2018-07-17T19:40:28Z","cross_cats_sorted":["cs.IR"],"title_canon_sha256":"3908e13d74aecd32fcf2ffc441f910249a53dde52f6cfad61e0f22addaccfdcb","abstract_canon_sha256":"cd997ccc1667e984a3a78d7ef6779456458fa1b8acb17bc17546717667a06415"},"schema_version":"1.0"},"canonical_sha256":"5ecf7145d65bf2f3c3b7ceb1c11b2057c3414146845d6a7a328c0658e3e3277b","source":{"kind":"arxiv","id":"1807.06638","version":1},"source_aliases":[{"alias_kind":"arxiv","alias_value":"1807.06638","created_at":"2026-05-18T00:10:27Z"},{"alias_kind":"arxiv_version","alias_value":"1807.06638v1","created_at":"2026-05-18T00:10:27Z"},{"alias_kind":"doi","alias_value":"10.48550/arxiv.1807.06638","created_at":"2026-05-18T00:10:27Z"},{"alias_kind":"pith_short_12","alias_value":"L3HXCROWLPZP","created_at":"2026-05-18T12:32:33Z"},{"alias_kind":"pith_short_16","alias_value":"L3HXCROWLPZPHQ5X","created_at":"2026-05-18T12:32:33Z"},{"alias_kind":"pith_short_8","alias_value":"L3HXCROW","created_at":"2026-05-18T12:32:33Z"}],"events":[{"event_type":"record_created","subject_pith_number":"pith:2018:L3HXCROWLPZPHQ5XZ2Y4CGZAK7","target":"record","payload":{"canonical_record":{"source":{"id":"1807.06638","kind":"arxiv","version":1},"metadata":{"license":"http://arxiv.org/licenses/nonexclusive-distrib/1.0/","primary_cat":"cs.CL","submitted_at":"2018-07-17T19:40:28Z","cross_cats_sorted":["cs.IR"],"title_canon_sha256":"3908e13d74aecd32fcf2ffc441f910249a53dde52f6cfad61e0f22addaccfdcb","abstract_canon_sha256":"cd997ccc1667e984a3a78d7ef6779456458fa1b8acb17bc17546717667a06415"},"schema_version":"1.0"},"canonical_sha256":"5ecf7145d65bf2f3c3b7ceb1c11b2057c3414146845d6a7a328c0658e3e3277b","receipt":{"kind":"pith_receipt","key_id":"pith-v1-2026-05","algorithm":"ed25519","signed_at":"2026-05-18T00:10:27.314698Z","signature_b64":"F1jwh5T7q6zO9bRgv5wWVaf5A3ChCYqrUXi42hTAu5cbRAwzsbrTfi+NlVmr6CwK0O/x3KIrfm1mJXgRgxS9Cg==","signed_message":"canonical_sha256_bytes","builder_version":"pith-number-builder-2026-05-17-v1","receipt_version":"0.3","canonical_sha256":"5ecf7145d65bf2f3c3b7ceb1c11b2057c3414146845d6a7a328c0658e3e3277b","last_reissued_at":"2026-05-18T00:10:27.314135Z","signature_status":"signed_v1","first_computed_at":"2026-05-18T00:10:27.314135Z","public_key_fingerprint":"8d4b5ee74e4693bcd1df2446408b0d54"},"source_kind":"arxiv","source_id":"1807.06638","source_version":1,"attestation_state":"computed"},"signer":{"signer_id":"pith.science","signer_type":"pith_registry","key_id":"pith-v1-2026-05","public_key_fingerprint":"8d4b5ee74e4693bcd1df2446408b0d54"},"created_at":"2026-05-18T00:10:27Z","supersedes":[],"prev_event":null,"signature":{"signature_status":"signed_v1","algorithm":"ed25519","key_id":"pith-v1-2026-05","public_key_fingerprint":"8d4b5ee74e4693bcd1df2446408b0d54","signature_b64":"B7y7IcYvc+rIeT6w/IXQ7nbHDUuSmUtCzQTn9I7Nf49XT43m+O+OP5ugAmpoH7wOBb1GMSl18qTymCZHtx+VCw==","signed_message":"open_graph_event_sha256_bytes","signed_at":"2026-06-01T21:28:38.298301Z"},"content_sha256":"7a1cd94e542fc6f988f30e6c7de33ca217caedcedd23812222a2ff7e1c110b6a","schema_version":"1.0","event_id":"sha256:7a1cd94e542fc6f988f30e6c7de33ca217caedcedd23812222a2ff7e1c110b6a"},{"event_type":"graph_snapshot","subject_pith_number":"pith:2018:L3HXCROWLPZPHQ5XZ2Y4CGZAK7","target":"graph","payload":{"graph_snapshot":{"paper":{"title":"Developing a Portable Natural Language Processing Based Phenotyping System","license":"http://arxiv.org/licenses/nonexclusive-distrib/1.0/","headline":"","cross_cats":["cs.IR"],"primary_cat":"cs.CL","authors_text":"Chengsheng Mao, Guoqian Jiang, Haleh Vatani, Himanshu Sharma, Jyotishman Pathak, Liang Yao, Luke Rasmussen, Yizhen Zhang, Yizhen Zhong, Yuan Luo","submitted_at":"2018-07-17T19:40:28Z","abstract_excerpt":"This paper presents a portable phenotyping system that is capable of integrating both rule-based and statistical machine learning based approaches. Our system utilizes UMLS to extract clinically relevant features from the unstructured text and then facilitates portability across different institutions and data systems by incorporating OHDSI's OMOP Common Data Model (CDM) to standardize necessary data elements. Our system can also store the key components of rule-based systems (e.g., regular expression matches) in the format of OMOP CDM, thus enabling the reuse, adaptation and extension of many"},"claims":{"count":0,"items":[],"snapshot_sha256":"258153158e38e3291e3d48162225fcdb2d5a3ed65a07baac614ab91432fd4f57"},"source":{"id":"1807.06638","kind":"arxiv","version":1},"verdict":{"id":null,"model_set":{},"created_at":null,"strongest_claim":"","one_line_summary":"","pipeline_version":null,"weakest_assumption":"","pith_extraction_headline":""},"references":{"count":0,"sample":[],"resolved_work":0,"snapshot_sha256":"258153158e38e3291e3d48162225fcdb2d5a3ed65a07baac614ab91432fd4f57","internal_anchors":0},"formal_canon":{"evidence_count":0,"snapshot_sha256":"258153158e38e3291e3d48162225fcdb2d5a3ed65a07baac614ab91432fd4f57"},"author_claims":{"count":0,"strong_count":0,"snapshot_sha256":"258153158e38e3291e3d48162225fcdb2d5a3ed65a07baac614ab91432fd4f57"},"builder_version":"pith-number-builder-2026-05-17-v1"},"verdict_id":null},"signer":{"signer_id":"pith.science","signer_type":"pith_registry","key_id":"pith-v1-2026-05","public_key_fingerprint":"8d4b5ee74e4693bcd1df2446408b0d54"},"created_at":"2026-05-18T00:10:27Z","supersedes":[],"prev_event":null,"signature":{"signature_status":"signed_v1","algorithm":"ed25519","key_id":"pith-v1-2026-05","public_key_fingerprint":"8d4b5ee74e4693bcd1df2446408b0d54","signature_b64":"jDd08o2Rl133UReJYQBaDTWaP9xD+xbzDcMcwHVgMjhC1bxfRAi60zUMPcJdW1UynZFLLoPhP4C1WIpxcKxQBw==","signed_message":"open_graph_event_sha256_bytes","signed_at":"2026-06-01T21:28:38.298676Z"},"content_sha256":"ecebb978ea22b2039c01446d0244981182d85a40e8a7749d6ececa6ef0a39b65","schema_version":"1.0","event_id":"sha256:ecebb978ea22b2039c01446d0244981182d85a40e8a7749d6ececa6ef0a39b65"}],"timestamp_proofs":[],"mirror_hints":[{"mirror_type":"https","name":"Pith Resolver","base_url":"https://pith.science","bundle_url":"https://pith.science/pith/L3HXCROWLPZPHQ5XZ2Y4CGZAK7/bundle.json","state_url":"https://pith.science/pith/L3HXCROWLPZPHQ5XZ2Y4CGZAK7/state.json","well_known_bundle_url":"https://pith.science/.well-known/pith/L3HXCROWLPZPHQ5XZ2Y4CGZAK7/bundle.json","status":"primary"}],"public_keys":[{"key_id":"pith-v1-2026-05","algorithm":"ed25519","format":"raw","public_key_b64":"stVStoiQhXFxp4s2pdzPNoqVNBMojDU/fJ2db5S3CbM=","public_key_hex":"b2d552b68890857171a78b36a5dccf368a953413288c353f7c9d9d6f94b709b3","fingerprint_sha256_b32_first128bits":"RVFV5Z2OI2J3ZUO7ERDEBCYNKS","fingerprint_sha256_hex":"8d4b5ee74e4693bcd1df2446408b0d54","rotates_at":null,"url":"https://pith.science/pith-signing-key.json","notes":"Pith uses this Ed25519 key to sign canonical record SHA-256 digests. Verify with: ed25519_verify(public_key, message=canonical_sha256_bytes, signature=base64decode(signature_b64))."}],"merge_version":"pith-open-graph-merge-v1","built_at":"2026-06-01T21:28:38Z","links":{"resolver":"https://pith.science/pith/L3HXCROWLPZPHQ5XZ2Y4CGZAK7","bundle":"https://pith.science/pith/L3HXCROWLPZPHQ5XZ2Y4CGZAK7/bundle.json","state":"https://pith.science/pith/L3HXCROWLPZPHQ5XZ2Y4CGZAK7/state.json","well_known_bundle":"https://pith.science/.well-known/pith/L3HXCROWLPZPHQ5XZ2Y4CGZAK7/bundle.json"},"state":{"state_type":"pith_open_graph_state","state_version":"1.0","pith_number":"pith:2018:L3HXCROWLPZPHQ5XZ2Y4CGZAK7","merge_version":"pith-open-graph-merge-v1","event_count":2,"valid_event_count":2,"invalid_event_count":0,"equivocation_count":0,"current":{"canonical_record":{"metadata":{"abstract_canon_sha256":"cd997ccc1667e984a3a78d7ef6779456458fa1b8acb17bc17546717667a06415","cross_cats_sorted":["cs.IR"],"license":"http://arxiv.org/licenses/nonexclusive-distrib/1.0/","primary_cat":"cs.CL","submitted_at":"2018-07-17T19:40:28Z","title_canon_sha256":"3908e13d74aecd32fcf2ffc441f910249a53dde52f6cfad61e0f22addaccfdcb"},"schema_version":"1.0","source":{"id":"1807.06638","kind":"arxiv","version":1}},"source_aliases":[{"alias_kind":"arxiv","alias_value":"1807.06638","created_at":"2026-05-18T00:10:27Z"},{"alias_kind":"arxiv_version","alias_value":"1807.06638v1","created_at":"2026-05-18T00:10:27Z"},{"alias_kind":"doi","alias_value":"10.48550/arxiv.1807.06638","created_at":"2026-05-18T00:10:27Z"},{"alias_kind":"pith_short_12","alias_value":"L3HXCROWLPZP","created_at":"2026-05-18T12:32:33Z"},{"alias_kind":"pith_short_16","alias_value":"L3HXCROWLPZPHQ5X","created_at":"2026-05-18T12:32:33Z"},{"alias_kind":"pith_short_8","alias_value":"L3HXCROW","created_at":"2026-05-18T12:32:33Z"}],"graph_snapshots":[{"event_id":"sha256:ecebb978ea22b2039c01446d0244981182d85a40e8a7749d6ececa6ef0a39b65","target":"graph","created_at":"2026-05-18T00:10:27Z","signer":{"key_id":"pith-v1-2026-05","public_key_fingerprint":"8d4b5ee74e4693bcd1df2446408b0d54","signer_id":"pith.science","signer_type":"pith_registry"},"payload":{"graph_snapshot":{"author_claims":{"count":0,"snapshot_sha256":"258153158e38e3291e3d48162225fcdb2d5a3ed65a07baac614ab91432fd4f57","strong_count":0},"builder_version":"pith-number-builder-2026-05-17-v1","claims":{"count":0,"items":[],"snapshot_sha256":"258153158e38e3291e3d48162225fcdb2d5a3ed65a07baac614ab91432fd4f57"},"formal_canon":{"evidence_count":0,"snapshot_sha256":"258153158e38e3291e3d48162225fcdb2d5a3ed65a07baac614ab91432fd4f57"},"paper":{"abstract_excerpt":"This paper presents a portable phenotyping system that is capable of integrating both rule-based and statistical machine learning based approaches. Our system utilizes UMLS to extract clinically relevant features from the unstructured text and then facilitates portability across different institutions and data systems by incorporating OHDSI's OMOP Common Data Model (CDM) to standardize necessary data elements. Our system can also store the key components of rule-based systems (e.g., regular expression matches) in the format of OMOP CDM, thus enabling the reuse, adaptation and extension of many","authors_text":"Chengsheng Mao, Guoqian Jiang, Haleh Vatani, Himanshu Sharma, Jyotishman Pathak, Liang Yao, Luke Rasmussen, Yizhen Zhang, Yizhen Zhong, Yuan Luo","cross_cats":["cs.IR"],"headline":"","license":"http://arxiv.org/licenses/nonexclusive-distrib/1.0/","primary_cat":"cs.CL","submitted_at":"2018-07-17T19:40:28Z","title":"Developing a Portable Natural Language Processing Based Phenotyping System"},"references":{"count":0,"internal_anchors":0,"resolved_work":0,"sample":[],"snapshot_sha256":"258153158e38e3291e3d48162225fcdb2d5a3ed65a07baac614ab91432fd4f57"},"source":{"id":"1807.06638","kind":"arxiv","version":1},"verdict":{"created_at":null,"id":null,"model_set":{},"one_line_summary":"","pipeline_version":null,"pith_extraction_headline":"","strongest_claim":"","weakest_assumption":""}},"verdict_id":null}}],"author_attestations":[],"timestamp_anchors":[],"storage_attestations":[],"citation_signatures":[],"replication_records":[],"corrections":[],"mirror_hints":[],"record_created":{"event_id":"sha256:7a1cd94e542fc6f988f30e6c7de33ca217caedcedd23812222a2ff7e1c110b6a","target":"record","created_at":"2026-05-18T00:10:27Z","signer":{"key_id":"pith-v1-2026-05","public_key_fingerprint":"8d4b5ee74e4693bcd1df2446408b0d54","signer_id":"pith.science","signer_type":"pith_registry"},"payload":{"attestation_state":"computed","canonical_record":{"metadata":{"abstract_canon_sha256":"cd997ccc1667e984a3a78d7ef6779456458fa1b8acb17bc17546717667a06415","cross_cats_sorted":["cs.IR"],"license":"http://arxiv.org/licenses/nonexclusive-distrib/1.0/","primary_cat":"cs.CL","submitted_at":"2018-07-17T19:40:28Z","title_canon_sha256":"3908e13d74aecd32fcf2ffc441f910249a53dde52f6cfad61e0f22addaccfdcb"},"schema_version":"1.0","source":{"id":"1807.06638","kind":"arxiv","version":1}},"canonical_sha256":"5ecf7145d65bf2f3c3b7ceb1c11b2057c3414146845d6a7a328c0658e3e3277b","receipt":{"algorithm":"ed25519","builder_version":"pith-number-builder-2026-05-17-v1","canonical_sha256":"5ecf7145d65bf2f3c3b7ceb1c11b2057c3414146845d6a7a328c0658e3e3277b","first_computed_at":"2026-05-18T00:10:27.314135Z","key_id":"pith-v1-2026-05","kind":"pith_receipt","last_reissued_at":"2026-05-18T00:10:27.314135Z","public_key_fingerprint":"8d4b5ee74e4693bcd1df2446408b0d54","receipt_version":"0.3","signature_b64":"F1jwh5T7q6zO9bRgv5wWVaf5A3ChCYqrUXi42hTAu5cbRAwzsbrTfi+NlVmr6CwK0O/x3KIrfm1mJXgRgxS9Cg==","signature_status":"signed_v1","signed_at":"2026-05-18T00:10:27.314698Z","signed_message":"canonical_sha256_bytes"},"source_id":"1807.06638","source_kind":"arxiv","source_version":1}}},"equivocations":[],"invalid_events":[],"applied_event_ids":["sha256:7a1cd94e542fc6f988f30e6c7de33ca217caedcedd23812222a2ff7e1c110b6a","sha256:ecebb978ea22b2039c01446d0244981182d85a40e8a7749d6ececa6ef0a39b65"],"state_sha256":"cd60d7c86787f9c76e1a1cab3949600406dded1e09eeb2b72fcee5e20966722c"},"bundle_signature":{"signature_status":"signed_v1","algorithm":"ed25519","key_id":"pith-v1-2026-05","public_key_fingerprint":"8d4b5ee74e4693bcd1df2446408b0d54","signature_b64":"hFKyWnRPOgeAWIbhaSDYCRjN2vZXNKOPWtg2jABHZt+K8djR4cm+vP1JHBL0YBFhy5l1l+B+LsCfwJdM51b+CA==","signed_message":"bundle_sha256_bytes","signed_at":"2026-06-01T21:28:38.300670Z","bundle_sha256":"b9c6c6dfc55f51875be6d754ebc8372fcb1bb7e11df20602c4424484b25677fc"}}