{"bundle_type":"pith_open_graph_bundle","bundle_version":"1.0","pith_number":"pith:2018:XFO7QEPUW7ZXADOIF4KRFOTRTK","short_pith_number":"pith:XFO7QEPU","canonical_record":{"source":{"id":"1808.06907","kind":"arxiv","version":2},"metadata":{"license":"http://creativecommons.org/licenses/by/4.0/","primary_cat":"cs.DB","submitted_at":"2018-08-17T18:03:26Z","cross_cats_sorted":["cs.AI"],"title_canon_sha256":"1c73efbf81a9c8fa9b45043d90d8ef997de50a69942d124c2314de6b042a8fc5","abstract_canon_sha256":"2b422c461944fd250d2aa2ec209739f2574fcc646844950f3e3db2576aea1b27"},"schema_version":"1.0"},"canonical_sha256":"b95df811f4b7f3700dc82f1512ba719ab41fe242837b0d810d31c4cfe7b5378e","source":{"kind":"arxiv","id":"1808.06907","version":2},"source_aliases":[{"alias_kind":"arxiv","alias_value":"1808.06907","created_at":"2026-05-17T23:52:41Z"},{"alias_kind":"arxiv_version","alias_value":"1808.06907v2","created_at":"2026-05-17T23:52:41Z"},{"alias_kind":"doi","alias_value":"10.48550/arxiv.1808.06907","created_at":"2026-05-17T23:52:41Z"},{"alias_kind":"pith_short_12","alias_value":"XFO7QEPUW7ZX","created_at":"2026-05-18T12:33:01Z"},{"alias_kind":"pith_short_16","alias_value":"XFO7QEPUW7ZXADOI","created_at":"2026-05-18T12:33:01Z"},{"alias_kind":"pith_short_8","alias_value":"XFO7QEPU","created_at":"2026-05-18T12:33:01Z"}],"events":[{"event_type":"record_created","subject_pith_number":"pith:2018:XFO7QEPUW7ZXADOIF4KRFOTRTK","target":"record","payload":{"canonical_record":{"source":{"id":"1808.06907","kind":"arxiv","version":2},"metadata":{"license":"http://creativecommons.org/licenses/by/4.0/","primary_cat":"cs.DB","submitted_at":"2018-08-17T18:03:26Z","cross_cats_sorted":["cs.AI"],"title_canon_sha256":"1c73efbf81a9c8fa9b45043d90d8ef997de50a69942d124c2314de6b042a8fc5","abstract_canon_sha256":"2b422c461944fd250d2aa2ec209739f2574fcc646844950f3e3db2576aea1b27"},"schema_version":"1.0"},"canonical_sha256":"b95df811f4b7f3700dc82f1512ba719ab41fe242837b0d810d31c4cfe7b5378e","receipt":{"kind":"pith_receipt","key_id":"pith-v1-2026-05","algorithm":"ed25519","signed_at":"2026-05-17T23:52:41.208051Z","signature_b64":"q70Nij5QQu+qV/AkqwZlORZVm65DQFlsesaV9lkccGS273Is6p/mIpv6trGcSHbuwimPKCuCZvhenneifpWdAw==","signed_message":"canonical_sha256_bytes","builder_version":"pith-number-builder-2026-05-17-v1","receipt_version":"0.3","canonical_sha256":"b95df811f4b7f3700dc82f1512ba719ab41fe242837b0d810d31c4cfe7b5378e","last_reissued_at":"2026-05-17T23:52:41.207380Z","signature_status":"signed_v1","first_computed_at":"2026-05-17T23:52:41.207380Z","public_key_fingerprint":"8d4b5ee74e4693bcd1df2446408b0d54"},"source_kind":"arxiv","source_id":"1808.06907","source_version":2,"attestation_state":"computed"},"signer":{"signer_id":"pith.science","signer_type":"pith_registry","key_id":"pith-v1-2026-05","public_key_fingerprint":"8d4b5ee74e4693bcd1df2446408b0d54"},"created_at":"2026-05-17T23:52:41Z","supersedes":[],"prev_event":null,"signature":{"signature_status":"signed_v1","algorithm":"ed25519","key_id":"pith-v1-2026-05","public_key_fingerprint":"8d4b5ee74e4693bcd1df2446408b0d54","signature_b64":"4/PMFMB+12zoYqIaYzu5Y1gCofxzPiEY10899j0KXsYXUOGN0zjBjHLDymYKb6xkWBEiLY4baOOYfz5d/p29BA==","signed_message":"open_graph_event_sha256_bytes","signed_at":"2026-06-04T15:24:33.614229Z"},"content_sha256":"b33c8acac8ce61495f66b4bb853a1fe99e4b04182798994795fb8512bac6e3ab","schema_version":"1.0","event_id":"sha256:b33c8acac8ce61495f66b4bb853a1fe99e4b04182798994795fb8512bac6e3ab"},{"event_type":"graph_snapshot","subject_pith_number":"pith:2018:XFO7QEPUW7ZXADOIF4KRFOTRTK","target":"graph","payload":{"graph_snapshot":{"paper":{"title":"The variable quality of metadata about biological samples used in biomedical experiments","license":"http://creativecommons.org/licenses/by/4.0/","headline":"","cross_cats":["cs.AI"],"primary_cat":"cs.DB","authors_text":"Mark A. Musen, Rafael S. Gon\\c{c}alves","submitted_at":"2018-08-17T18:03:26Z","abstract_excerpt":"We present an analytical study of the quality of metadata about samples used in biomedical experiments. The metadata under analysis are stored in two well-known databases: BioSample---a repository managed by the National Center for Biotechnology Information (NCBI), and BioSamples---a repository managed by the European Bioinformatics Institute (EBI). We tested whether 11.4M sample metadata records in the two repositories are populated with values that fulfill the stated requirements for such values. Our study revealed multiple anomalies in the metadata. Most metadata field names and their value"},"claims":{"count":0,"items":[],"snapshot_sha256":"258153158e38e3291e3d48162225fcdb2d5a3ed65a07baac614ab91432fd4f57"},"source":{"id":"1808.06907","kind":"arxiv","version":2},"verdict":{"id":null,"model_set":{},"created_at":null,"strongest_claim":"","one_line_summary":"","pipeline_version":null,"weakest_assumption":"","pith_extraction_headline":""},"references":{"count":0,"sample":[],"resolved_work":0,"snapshot_sha256":"258153158e38e3291e3d48162225fcdb2d5a3ed65a07baac614ab91432fd4f57","internal_anchors":0},"formal_canon":{"evidence_count":0,"snapshot_sha256":"258153158e38e3291e3d48162225fcdb2d5a3ed65a07baac614ab91432fd4f57"},"author_claims":{"count":0,"strong_count":0,"snapshot_sha256":"258153158e38e3291e3d48162225fcdb2d5a3ed65a07baac614ab91432fd4f57"},"builder_version":"pith-number-builder-2026-05-17-v1"},"verdict_id":null},"signer":{"signer_id":"pith.science","signer_type":"pith_registry","key_id":"pith-v1-2026-05","public_key_fingerprint":"8d4b5ee74e4693bcd1df2446408b0d54"},"created_at":"2026-05-17T23:52:41Z","supersedes":[],"prev_event":null,"signature":{"signature_status":"signed_v1","algorithm":"ed25519","key_id":"pith-v1-2026-05","public_key_fingerprint":"8d4b5ee74e4693bcd1df2446408b0d54","signature_b64":"z01bqEXvqV+VfTqeI8UOIx8PigbdvbGO+UWyTmKpY293NIxjhvRbmVbO0Dku1Qv2L7UEZkjByr+CcSDBujzEBQ==","signed_message":"open_graph_event_sha256_bytes","signed_at":"2026-06-04T15:24:33.614600Z"},"content_sha256":"04f60be41a716213858863cd9a275fc7caba3417923dbf344e812c853749841e","schema_version":"1.0","event_id":"sha256:04f60be41a716213858863cd9a275fc7caba3417923dbf344e812c853749841e"}],"timestamp_proofs":[],"mirror_hints":[{"mirror_type":"https","name":"Pith Resolver","base_url":"https://pith.science","bundle_url":"https://pith.science/pith/XFO7QEPUW7ZXADOIF4KRFOTRTK/bundle.json","state_url":"https://pith.science/pith/XFO7QEPUW7ZXADOIF4KRFOTRTK/state.json","well_known_bundle_url":"https://pith.science/.well-known/pith/XFO7QEPUW7ZXADOIF4KRFOTRTK/bundle.json","status":"primary"}],"public_keys":[{"key_id":"pith-v1-2026-05","algorithm":"ed25519","format":"raw","public_key_b64":"stVStoiQhXFxp4s2pdzPNoqVNBMojDU/fJ2db5S3CbM=","public_key_hex":"b2d552b68890857171a78b36a5dccf368a953413288c353f7c9d9d6f94b709b3","fingerprint_sha256_b32_first128bits":"RVFV5Z2OI2J3ZUO7ERDEBCYNKS","fingerprint_sha256_hex":"8d4b5ee74e4693bcd1df2446408b0d54","rotates_at":null,"url":"https://pith.science/pith-signing-key.json","notes":"Pith uses this Ed25519 key to sign canonical record SHA-256 digests. Verify with: ed25519_verify(public_key, message=canonical_sha256_bytes, signature=base64decode(signature_b64))."}],"merge_version":"pith-open-graph-merge-v1","built_at":"2026-06-04T15:24:33Z","links":{"resolver":"https://pith.science/pith/XFO7QEPUW7ZXADOIF4KRFOTRTK","bundle":"https://pith.science/pith/XFO7QEPUW7ZXADOIF4KRFOTRTK/bundle.json","state":"https://pith.science/pith/XFO7QEPUW7ZXADOIF4KRFOTRTK/state.json","well_known_bundle":"https://pith.science/.well-known/pith/XFO7QEPUW7ZXADOIF4KRFOTRTK/bundle.json"},"state":{"state_type":"pith_open_graph_state","state_version":"1.0","pith_number":"pith:2018:XFO7QEPUW7ZXADOIF4KRFOTRTK","merge_version":"pith-open-graph-merge-v1","event_count":2,"valid_event_count":2,"invalid_event_count":0,"equivocation_count":0,"current":{"canonical_record":{"metadata":{"abstract_canon_sha256":"2b422c461944fd250d2aa2ec209739f2574fcc646844950f3e3db2576aea1b27","cross_cats_sorted":["cs.AI"],"license":"http://creativecommons.org/licenses/by/4.0/","primary_cat":"cs.DB","submitted_at":"2018-08-17T18:03:26Z","title_canon_sha256":"1c73efbf81a9c8fa9b45043d90d8ef997de50a69942d124c2314de6b042a8fc5"},"schema_version":"1.0","source":{"id":"1808.06907","kind":"arxiv","version":2}},"source_aliases":[{"alias_kind":"arxiv","alias_value":"1808.06907","created_at":"2026-05-17T23:52:41Z"},{"alias_kind":"arxiv_version","alias_value":"1808.06907v2","created_at":"2026-05-17T23:52:41Z"},{"alias_kind":"doi","alias_value":"10.48550/arxiv.1808.06907","created_at":"2026-05-17T23:52:41Z"},{"alias_kind":"pith_short_12","alias_value":"XFO7QEPUW7ZX","created_at":"2026-05-18T12:33:01Z"},{"alias_kind":"pith_short_16","alias_value":"XFO7QEPUW7ZXADOI","created_at":"2026-05-18T12:33:01Z"},{"alias_kind":"pith_short_8","alias_value":"XFO7QEPU","created_at":"2026-05-18T12:33:01Z"}],"graph_snapshots":[{"event_id":"sha256:04f60be41a716213858863cd9a275fc7caba3417923dbf344e812c853749841e","target":"graph","created_at":"2026-05-17T23:52:41Z","signer":{"key_id":"pith-v1-2026-05","public_key_fingerprint":"8d4b5ee74e4693bcd1df2446408b0d54","signer_id":"pith.science","signer_type":"pith_registry"},"payload":{"graph_snapshot":{"author_claims":{"count":0,"snapshot_sha256":"258153158e38e3291e3d48162225fcdb2d5a3ed65a07baac614ab91432fd4f57","strong_count":0},"builder_version":"pith-number-builder-2026-05-17-v1","claims":{"count":0,"items":[],"snapshot_sha256":"258153158e38e3291e3d48162225fcdb2d5a3ed65a07baac614ab91432fd4f57"},"formal_canon":{"evidence_count":0,"snapshot_sha256":"258153158e38e3291e3d48162225fcdb2d5a3ed65a07baac614ab91432fd4f57"},"paper":{"abstract_excerpt":"We present an analytical study of the quality of metadata about samples used in biomedical experiments. The metadata under analysis are stored in two well-known databases: BioSample---a repository managed by the National Center for Biotechnology Information (NCBI), and BioSamples---a repository managed by the European Bioinformatics Institute (EBI). We tested whether 11.4M sample metadata records in the two repositories are populated with values that fulfill the stated requirements for such values. Our study revealed multiple anomalies in the metadata. Most metadata field names and their value","authors_text":"Mark A. Musen, Rafael S. Gon\\c{c}alves","cross_cats":["cs.AI"],"headline":"","license":"http://creativecommons.org/licenses/by/4.0/","primary_cat":"cs.DB","submitted_at":"2018-08-17T18:03:26Z","title":"The variable quality of metadata about biological samples used in biomedical experiments"},"references":{"count":0,"internal_anchors":0,"resolved_work":0,"sample":[],"snapshot_sha256":"258153158e38e3291e3d48162225fcdb2d5a3ed65a07baac614ab91432fd4f57"},"source":{"id":"1808.06907","kind":"arxiv","version":2},"verdict":{"created_at":null,"id":null,"model_set":{},"one_line_summary":"","pipeline_version":null,"pith_extraction_headline":"","strongest_claim":"","weakest_assumption":""}},"verdict_id":null}}],"author_attestations":[],"timestamp_anchors":[],"storage_attestations":[],"citation_signatures":[],"replication_records":[],"corrections":[],"mirror_hints":[],"record_created":{"event_id":"sha256:b33c8acac8ce61495f66b4bb853a1fe99e4b04182798994795fb8512bac6e3ab","target":"record","created_at":"2026-05-17T23:52:41Z","signer":{"key_id":"pith-v1-2026-05","public_key_fingerprint":"8d4b5ee74e4693bcd1df2446408b0d54","signer_id":"pith.science","signer_type":"pith_registry"},"payload":{"attestation_state":"computed","canonical_record":{"metadata":{"abstract_canon_sha256":"2b422c461944fd250d2aa2ec209739f2574fcc646844950f3e3db2576aea1b27","cross_cats_sorted":["cs.AI"],"license":"http://creativecommons.org/licenses/by/4.0/","primary_cat":"cs.DB","submitted_at":"2018-08-17T18:03:26Z","title_canon_sha256":"1c73efbf81a9c8fa9b45043d90d8ef997de50a69942d124c2314de6b042a8fc5"},"schema_version":"1.0","source":{"id":"1808.06907","kind":"arxiv","version":2}},"canonical_sha256":"b95df811f4b7f3700dc82f1512ba719ab41fe242837b0d810d31c4cfe7b5378e","receipt":{"algorithm":"ed25519","builder_version":"pith-number-builder-2026-05-17-v1","canonical_sha256":"b95df811f4b7f3700dc82f1512ba719ab41fe242837b0d810d31c4cfe7b5378e","first_computed_at":"2026-05-17T23:52:41.207380Z","key_id":"pith-v1-2026-05","kind":"pith_receipt","last_reissued_at":"2026-05-17T23:52:41.207380Z","public_key_fingerprint":"8d4b5ee74e4693bcd1df2446408b0d54","receipt_version":"0.3","signature_b64":"q70Nij5QQu+qV/AkqwZlORZVm65DQFlsesaV9lkccGS273Is6p/mIpv6trGcSHbuwimPKCuCZvhenneifpWdAw==","signature_status":"signed_v1","signed_at":"2026-05-17T23:52:41.208051Z","signed_message":"canonical_sha256_bytes"},"source_id":"1808.06907","source_kind":"arxiv","source_version":2}}},"equivocations":[],"invalid_events":[],"applied_event_ids":["sha256:b33c8acac8ce61495f66b4bb853a1fe99e4b04182798994795fb8512bac6e3ab","sha256:04f60be41a716213858863cd9a275fc7caba3417923dbf344e812c853749841e"],"state_sha256":"d5ca316968bd081b943b02a71724d2e643cf45f3163a443bf62719564ef3aca5"},"bundle_signature":{"signature_status":"signed_v1","algorithm":"ed25519","key_id":"pith-v1-2026-05","public_key_fingerprint":"8d4b5ee74e4693bcd1df2446408b0d54","signature_b64":"ldWUv/I+YarpBP/ONOIq5BkI0xXfpFWYU9uEEGxsf4/y8GiVH6UyJMRsKHWr9FRcsPpS4rXUUP/fYoHMUgrCAg==","signed_message":"bundle_sha256_bytes","signed_at":"2026-06-04T15:24:33.616642Z","bundle_sha256":"b7c704148e76baa3e24b8ac9886584fb71f093b5997b76c8dfcd29e4daf53f52"}}