{"bundle_type":"pith_open_graph_bundle","bundle_version":"1.0","pith_number":"pith:2026:DPQ22QU65CCC35XGTUC3WCVQJP","short_pith_number":"pith:DPQ22QU6","canonical_record":{"source":{"id":"2604.15145","kind":"arxiv","version":2},"metadata":{"license":"http://creativecommons.org/licenses/by-sa/4.0/","primary_cat":"cs.AI","submitted_at":"2026-04-16T15:19:58Z","cross_cats_sorted":["cs.DL"],"title_canon_sha256":"216e4c00bf4aa30bae8db2a97ac360162f870f288d5668cc136ced87ce5f7431","abstract_canon_sha256":"c6a6a475293d0e53d672b0cb62644f844b6d05901d4f06aef1d4d2aaa8155c6c"},"schema_version":"1.0"},"canonical_sha256":"1be1ad429ee8842df6e69d05bb0ab04bf8a1b317054782e7c1918eaaed45c728","source":{"kind":"arxiv","id":"2604.15145","version":2},"source_aliases":[{"alias_kind":"arxiv","alias_value":"2604.15145","created_at":"2026-08-07T00:45:50Z"},{"alias_kind":"arxiv_version","alias_value":"2604.15145v2","created_at":"2026-08-07T00:45:50Z"},{"alias_kind":"doi","alias_value":"10.48550/arxiv.2604.15145","created_at":"2026-08-07T00:45:50Z"},{"alias_kind":"pith_short_12","alias_value":"DPQ22QU65CCC","created_at":"2026-08-07T00:45:50Z"},{"alias_kind":"pith_short_16","alias_value":"DPQ22QU65CCC35XG","created_at":"2026-08-07T00:45:50Z"},{"alias_kind":"pith_short_8","alias_value":"DPQ22QU6","created_at":"2026-08-07T00:45:50Z"}],"events":[{"event_type":"record_created","subject_pith_number":"pith:2026:DPQ22QU65CCC35XGTUC3WCVQJP","target":"record","payload":{"canonical_record":{"source":{"id":"2604.15145","kind":"arxiv","version":2},"metadata":{"license":"http://creativecommons.org/licenses/by-sa/4.0/","primary_cat":"cs.AI","submitted_at":"2026-04-16T15:19:58Z","cross_cats_sorted":["cs.DL"],"title_canon_sha256":"216e4c00bf4aa30bae8db2a97ac360162f870f288d5668cc136ced87ce5f7431","abstract_canon_sha256":"c6a6a475293d0e53d672b0cb62644f844b6d05901d4f06aef1d4d2aaa8155c6c"},"schema_version":"1.0"},"canonical_sha256":"1be1ad429ee8842df6e69d05bb0ab04bf8a1b317054782e7c1918eaaed45c728","receipt":{"kind":"pith_receipt","key_id":"pith-v1-2026-05","algorithm":"ed25519","signed_at":"2026-08-07T00:45:50.715232Z","signature_b64":"3RYw6rhMgDajKZFS5ClU1weKQvMoajD8TpSuOrOpez5SKcUsn631HZoicnzW0SysTi6+Sz8C9dkTavLR9YsvBA==","signed_message":"canonical_sha256_bytes","builder_version":"pith-number-builder-2026-05-17-v1","receipt_version":"0.3","canonical_sha256":"1be1ad429ee8842df6e69d05bb0ab04bf8a1b317054782e7c1918eaaed45c728","last_reissued_at":"2026-08-07T00:45:50.713682Z","signature_status":"signed_v1","first_computed_at":"2026-08-07T00:45:50.713682Z","public_key_fingerprint":"8d4b5ee74e4693bcd1df2446408b0d54"},"source_kind":"arxiv","source_id":"2604.15145","source_version":2,"attestation_state":"computed"},"signer":{"signer_id":"pith.science","signer_type":"pith_registry","key_id":"pith-v1-2026-05","public_key_fingerprint":"8d4b5ee74e4693bcd1df2446408b0d54"},"created_at":"2026-08-07T00:45:50Z","supersedes":[],"prev_event":null,"signature":{"signature_status":"signed_v1","algorithm":"ed25519","key_id":"pith-v1-2026-05","public_key_fingerprint":"8d4b5ee74e4693bcd1df2446408b0d54","signature_b64":"A93NvtWlIAY3Y1yAiBwAMUl776+8R+ExEAMot9Aryd8fIC+LmiRDQnusfpCZ86+oMF3BmdC3C7ekX53YtZa/BQ==","signed_message":"open_graph_event_sha256_bytes","signed_at":"2026-08-20T10:06:06.249999Z"},"content_sha256":"f73736e87e3544987b28d857397c6aa98b9fc2c60142b8cebc0784d2b58aab4e","schema_version":"1.0","event_id":"sha256:f73736e87e3544987b28d857397c6aa98b9fc2c60142b8cebc0784d2b58aab4e"},{"event_type":"graph_snapshot","subject_pith_number":"pith:2026:DPQ22QU65CCC35XGTUC3WCVQJP","target":"graph","payload":{"graph_snapshot":{"paper":{"title":"An Axiomatic Benchmark for Evaluation of Scientific Novelty Metrics","license":"http://creativecommons.org/licenses/by-sa/4.0/","headline":"No single metric satisfies all axioms for scientific novelty, but combining complementary architectures reaches 90.1 percent compliance.","cross_cats":["cs.DL"],"primary_cat":"cs.AI","authors_text":"ChengXiang Zhai, Miri Liu","submitted_at":"2026-04-16T15:19:58Z","abstract_excerpt":"The rigorous evaluation of the novelty of a scientific paper is, even for human scientists, a challenging task. With the increasing interest in AI scientists, it is becoming more and more important that this task be automatable and reliable, lest attention and compute be wasted on ideas that have already been explored. Due to the challenge of quantifying ground-truth novelty, however, existing novelty metrics generally validate against noisy, confounded signals such as citation counts or peer review scores. We introduce a benchmark that compares novelty metrics without requiring explicit novel"},"claims":{"count":4,"items":[{"kind":"strongest_claim","text":"Combining metrics of complementary architectures leads to consistent improvements on the benchmark, with per-axiom weighting achieving 90.1% versus 71.5% for the best individual metric.","source":"verdict.strongest_claim","status":"machine_extracted","claim_id":"C1","attestation":"unclaimed"},{"kind":"weakest_assumption","text":"The axioms defined capture the essential aspects of human scientific norms for novelty, and the ten tasks across three AI domains sufficiently represent the general problem of evaluating scientific novelty.","source":"verdict.weakest_assumption","status":"machine_extracted","claim_id":"C2","attestation":"unclaimed"},{"kind":"one_line_summary","text":"An axiomatic benchmark shows no single novelty metric satisfies all desired properties consistently, but combining complementary metrics reaches 90.1% performance versus 71.5% for the best individual metric.","source":"verdict.one_line_summary","status":"machine_extracted","claim_id":"C3","attestation":"unclaimed"},{"kind":"headline","text":"No single metric satisfies all axioms for scientific novelty, but combining complementary architectures reaches 90.1 percent compliance.","source":"verdict.pith_extraction.headline","status":"machine_extracted","claim_id":"C4","attestation":"unclaimed"}],"snapshot_sha256":"0bd77ca5a717a571e0dbc5ed10056c40c33d29a6869f64c4ed61b2fa1394a84c"},"source":{"id":"2604.15145","kind":"arxiv","version":2},"verdict":{"id":"7c0dd367-2a57-4add-8b43-173d079d7342","model_set":{"reader":"grok-4.3"},"created_at":"2026-05-10T11:18:46.444284Z","strongest_claim":"Combining metrics of complementary architectures leads to consistent improvements on the benchmark, with per-axiom weighting achieving 90.1% versus 71.5% for the best individual metric.","one_line_summary":"An axiomatic benchmark shows no single novelty metric satisfies all desired properties consistently, but combining complementary metrics reaches 90.1% performance versus 71.5% for the best individual metric.","pipeline_version":"pith-pipeline@v0.9.0","weakest_assumption":"The axioms defined capture the essential aspects of human scientific norms for novelty, and the ten tasks across three AI domains sufficiently represent the general problem of evaluating scientific novelty.","pith_extraction_headline":"No single metric satisfies all axioms for scientific novelty, but combining complementary architectures reaches 90.1 percent compliance."},"integrity":{"clean":true,"summary":{"advisory":0,"critical":0,"by_detector":{},"informational":0},"endpoint":"/pith/2604.15145/integrity.json","findings":[],"available":true,"detectors_run":[],"snapshot_sha256":"c28c3603d3b5d939e8dc4c7e95fa8dfce3d595e45f758748cecf8e644a296938"},"references":{"count":0,"sample":[],"resolved_work":0,"snapshot_sha256":"258153158e38e3291e3d48162225fcdb2d5a3ed65a07baac614ab91432fd4f57","internal_anchors":0},"formal_canon":{"evidence_count":0,"snapshot_sha256":"258153158e38e3291e3d48162225fcdb2d5a3ed65a07baac614ab91432fd4f57"},"author_claims":{"count":0,"strong_count":0,"snapshot_sha256":"258153158e38e3291e3d48162225fcdb2d5a3ed65a07baac614ab91432fd4f57"},"builder_version":"pith-number-builder-2026-05-17-v1"},"verdict_id":"7c0dd367-2a57-4add-8b43-173d079d7342"},"signer":{"signer_id":"pith.science","signer_type":"pith_registry","key_id":"pith-v1-2026-05","public_key_fingerprint":"8d4b5ee74e4693bcd1df2446408b0d54"},"created_at":"2026-08-07T00:45:50Z","supersedes":[],"prev_event":null,"signature":{"signature_status":"signed_v1","algorithm":"ed25519","key_id":"pith-v1-2026-05","public_key_fingerprint":"8d4b5ee74e4693bcd1df2446408b0d54","signature_b64":"Xf6F5lhJcQ3Xnmog1zR3LsDDiYP+SeteZLz1EDWyy0AZXbw4KzOa8fqQYYpbAkZwePWYG/a02xxpy22WRkD3Dg==","signed_message":"open_graph_event_sha256_bytes","signed_at":"2026-08-20T10:06:06.250633Z"},"content_sha256":"65086449a75c1fb4f36b138a907380c86a1dc5ae9a61b6f8e56a7ea86dbcebdf","schema_version":"1.0","event_id":"sha256:65086449a75c1fb4f36b138a907380c86a1dc5ae9a61b6f8e56a7ea86dbcebdf"}],"timestamp_proofs":[],"mirror_hints":[{"mirror_type":"https","name":"Pith Resolver","base_url":"https://pith.science","bundle_url":"https://pith.science/pith/DPQ22QU65CCC35XGTUC3WCVQJP/bundle.json","state_url":"https://pith.science/pith/DPQ22QU65CCC35XGTUC3WCVQJP/state.json","well_known_bundle_url":"https://pith.science/.well-known/pith/DPQ22QU65CCC35XGTUC3WCVQJP/bundle.json","status":"primary"}],"public_keys":[{"key_id":"pith-v1-2026-05","algorithm":"ed25519","format":"raw","public_key_b64":"stVStoiQhXFxp4s2pdzPNoqVNBMojDU/fJ2db5S3CbM=","public_key_hex":"b2d552b68890857171a78b36a5dccf368a953413288c353f7c9d9d6f94b709b3","fingerprint_sha256_b32_first128bits":"RVFV5Z2OI2J3ZUO7ERDEBCYNKS","fingerprint_sha256_hex":"8d4b5ee74e4693bcd1df2446408b0d54","rotates_at":null,"url":"https://pith.science/pith-signing-key.json","notes":"Pith uses this Ed25519 key to sign canonical record SHA-256 digests. Verify with: ed25519_verify(public_key, message=canonical_sha256_bytes, signature=base64decode(signature_b64))."}],"merge_version":"pith-open-graph-merge-v1","built_at":"2026-08-20T10:06:06Z","links":{"resolver":"https://pith.science/pith/DPQ22QU65CCC35XGTUC3WCVQJP","bundle":"https://pith.science/pith/DPQ22QU65CCC35XGTUC3WCVQJP/bundle.json","state":"https://pith.science/pith/DPQ22QU65CCC35XGTUC3WCVQJP/state.json","well_known_bundle":"https://pith.science/.well-known/pith/DPQ22QU65CCC35XGTUC3WCVQJP/bundle.json"},"state":{"state_type":"pith_open_graph_state","state_version":"1.0","pith_number":"pith:2026:DPQ22QU65CCC35XGTUC3WCVQJP","merge_version":"pith-open-graph-merge-v1","event_count":2,"valid_event_count":2,"invalid_event_count":0,"equivocation_count":0,"current":{"canonical_record":{"metadata":{"abstract_canon_sha256":"c6a6a475293d0e53d672b0cb62644f844b6d05901d4f06aef1d4d2aaa8155c6c","cross_cats_sorted":["cs.DL"],"license":"http://creativecommons.org/licenses/by-sa/4.0/","primary_cat":"cs.AI","submitted_at":"2026-04-16T15:19:58Z","title_canon_sha256":"216e4c00bf4aa30bae8db2a97ac360162f870f288d5668cc136ced87ce5f7431"},"schema_version":"1.0","source":{"id":"2604.15145","kind":"arxiv","version":2}},"source_aliases":[{"alias_kind":"arxiv","alias_value":"2604.15145","created_at":"2026-08-07T00:45:50Z"},{"alias_kind":"arxiv_version","alias_value":"2604.15145v2","created_at":"2026-08-07T00:45:50Z"},{"alias_kind":"doi","alias_value":"10.48550/arxiv.2604.15145","created_at":"2026-08-07T00:45:50Z"},{"alias_kind":"pith_short_12","alias_value":"DPQ22QU65CCC","created_at":"2026-08-07T00:45:50Z"},{"alias_kind":"pith_short_16","alias_value":"DPQ22QU65CCC35XG","created_at":"2026-08-07T00:45:50Z"},{"alias_kind":"pith_short_8","alias_value":"DPQ22QU6","created_at":"2026-08-07T00:45:50Z"}],"graph_snapshots":[{"event_id":"sha256:65086449a75c1fb4f36b138a907380c86a1dc5ae9a61b6f8e56a7ea86dbcebdf","target":"graph","created_at":"2026-08-07T00:45:50Z","signer":{"key_id":"pith-v1-2026-05","public_key_fingerprint":"8d4b5ee74e4693bcd1df2446408b0d54","signer_id":"pith.science","signer_type":"pith_registry"},"payload":{"graph_snapshot":{"author_claims":{"count":0,"snapshot_sha256":"258153158e38e3291e3d48162225fcdb2d5a3ed65a07baac614ab91432fd4f57","strong_count":0},"builder_version":"pith-number-builder-2026-05-17-v1","claims":{"count":4,"items":[{"attestation":"unclaimed","claim_id":"C1","kind":"strongest_claim","source":"verdict.strongest_claim","status":"machine_extracted","text":"Combining metrics of complementary architectures leads to consistent improvements on the benchmark, with per-axiom weighting achieving 90.1% versus 71.5% for the best individual metric."},{"attestation":"unclaimed","claim_id":"C2","kind":"weakest_assumption","source":"verdict.weakest_assumption","status":"machine_extracted","text":"The axioms defined capture the essential aspects of human scientific norms for novelty, and the ten tasks across three AI domains sufficiently represent the general problem of evaluating scientific novelty."},{"attestation":"unclaimed","claim_id":"C3","kind":"one_line_summary","source":"verdict.one_line_summary","status":"machine_extracted","text":"An axiomatic benchmark shows no single novelty metric satisfies all desired properties consistently, but combining complementary metrics reaches 90.1% performance versus 71.5% for the best individual metric."},{"attestation":"unclaimed","claim_id":"C4","kind":"headline","source":"verdict.pith_extraction.headline","status":"machine_extracted","text":"No single metric satisfies all axioms for scientific novelty, but combining complementary architectures reaches 90.1 percent compliance."}],"snapshot_sha256":"0bd77ca5a717a571e0dbc5ed10056c40c33d29a6869f64c4ed61b2fa1394a84c"},"formal_canon":{"evidence_count":0,"snapshot_sha256":"258153158e38e3291e3d48162225fcdb2d5a3ed65a07baac614ab91432fd4f57"},"integrity":{"available":true,"clean":true,"detectors_run":[],"endpoint":"/pith/2604.15145/integrity.json","findings":[],"snapshot_sha256":"c28c3603d3b5d939e8dc4c7e95fa8dfce3d595e45f758748cecf8e644a296938","summary":{"advisory":0,"by_detector":{},"critical":0,"informational":0}},"paper":{"abstract_excerpt":"The rigorous evaluation of the novelty of a scientific paper is, even for human scientists, a challenging task. With the increasing interest in AI scientists, it is becoming more and more important that this task be automatable and reliable, lest attention and compute be wasted on ideas that have already been explored. Due to the challenge of quantifying ground-truth novelty, however, existing novelty metrics generally validate against noisy, confounded signals such as citation counts or peer review scores. We introduce a benchmark that compares novelty metrics without requiring explicit novel","authors_text":"ChengXiang Zhai, Miri Liu","cross_cats":["cs.DL"],"headline":"No single metric satisfies all axioms for scientific novelty, but combining complementary architectures reaches 90.1 percent compliance.","license":"http://creativecommons.org/licenses/by-sa/4.0/","primary_cat":"cs.AI","submitted_at":"2026-04-16T15:19:58Z","title":"An Axiomatic Benchmark for Evaluation of Scientific Novelty Metrics"},"references":{"count":0,"internal_anchors":0,"resolved_work":0,"sample":[],"snapshot_sha256":"258153158e38e3291e3d48162225fcdb2d5a3ed65a07baac614ab91432fd4f57"},"source":{"id":"2604.15145","kind":"arxiv","version":2},"verdict":{"created_at":"2026-05-10T11:18:46.444284Z","id":"7c0dd367-2a57-4add-8b43-173d079d7342","model_set":{"reader":"grok-4.3"},"one_line_summary":"An axiomatic benchmark shows no single novelty metric satisfies all desired properties consistently, but combining complementary metrics reaches 90.1% performance versus 71.5% for the best individual metric.","pipeline_version":"pith-pipeline@v0.9.0","pith_extraction_headline":"No single metric satisfies all axioms for scientific novelty, but combining complementary architectures reaches 90.1 percent compliance.","strongest_claim":"Combining metrics of complementary architectures leads to consistent improvements on the benchmark, with per-axiom weighting achieving 90.1% versus 71.5% for the best individual metric.","weakest_assumption":"The axioms defined capture the essential aspects of human scientific norms for novelty, and the ten tasks across three AI domains sufficiently represent the general problem of evaluating scientific novelty."}},"verdict_id":"7c0dd367-2a57-4add-8b43-173d079d7342"}}],"author_attestations":[],"timestamp_anchors":[],"storage_attestations":[],"citation_signatures":[],"replication_records":[],"corrections":[],"mirror_hints":[],"record_created":{"event_id":"sha256:f73736e87e3544987b28d857397c6aa98b9fc2c60142b8cebc0784d2b58aab4e","target":"record","created_at":"2026-08-07T00:45:50Z","signer":{"key_id":"pith-v1-2026-05","public_key_fingerprint":"8d4b5ee74e4693bcd1df2446408b0d54","signer_id":"pith.science","signer_type":"pith_registry"},"payload":{"attestation_state":"computed","canonical_record":{"metadata":{"abstract_canon_sha256":"c6a6a475293d0e53d672b0cb62644f844b6d05901d4f06aef1d4d2aaa8155c6c","cross_cats_sorted":["cs.DL"],"license":"http://creativecommons.org/licenses/by-sa/4.0/","primary_cat":"cs.AI","submitted_at":"2026-04-16T15:19:58Z","title_canon_sha256":"216e4c00bf4aa30bae8db2a97ac360162f870f288d5668cc136ced87ce5f7431"},"schema_version":"1.0","source":{"id":"2604.15145","kind":"arxiv","version":2}},"canonical_sha256":"1be1ad429ee8842df6e69d05bb0ab04bf8a1b317054782e7c1918eaaed45c728","receipt":{"algorithm":"ed25519","builder_version":"pith-number-builder-2026-05-17-v1","canonical_sha256":"1be1ad429ee8842df6e69d05bb0ab04bf8a1b317054782e7c1918eaaed45c728","first_computed_at":"2026-08-07T00:45:50.713682Z","key_id":"pith-v1-2026-05","kind":"pith_receipt","last_reissued_at":"2026-08-07T00:45:50.713682Z","public_key_fingerprint":"8d4b5ee74e4693bcd1df2446408b0d54","receipt_version":"0.3","signature_b64":"3RYw6rhMgDajKZFS5ClU1weKQvMoajD8TpSuOrOpez5SKcUsn631HZoicnzW0SysTi6+Sz8C9dkTavLR9YsvBA==","signature_status":"signed_v1","signed_at":"2026-08-07T00:45:50.715232Z","signed_message":"canonical_sha256_bytes"},"source_id":"2604.15145","source_kind":"arxiv","source_version":2}}},"equivocations":[],"invalid_events":[],"applied_event_ids":["sha256:f73736e87e3544987b28d857397c6aa98b9fc2c60142b8cebc0784d2b58aab4e","sha256:65086449a75c1fb4f36b138a907380c86a1dc5ae9a61b6f8e56a7ea86dbcebdf"],"state_sha256":"57983bfd23bc2cf5d607d8858e1e62e2049ab1621e1a87a75179271625b6db8d"},"bundle_signature":{"signature_status":"signed_v1","algorithm":"ed25519","key_id":"pith-v1-2026-05","public_key_fingerprint":"8d4b5ee74e4693bcd1df2446408b0d54","signature_b64":"H1wclYDWusPLtlN5Jrg5stmF4RuXJ/wXc0w6xHGptbSYWgFXpCXdHAUwlq1pIf/LkpSlRdMXY2uVeSn49NtfCA==","signed_message":"bundle_sha256_bytes","signed_at":"2026-08-20T10:06:06.254630Z","bundle_sha256":"4e89ea89721bf530deaa2116bed5f9c04f47c2ec66d10e748bf8582d910da499"}}