{"record_type":"pith_number_record","schema_url":"https://pith.science/schemas/pith-number/v1.json","pith_number":"pith:2025:GQHJFRVRKAQQDJA3UXKSWFSMGP","short_pith_number":"pith:GQHJFRVR","schema_version":"1.0","canonical_sha256":"340e92c6b1502101a41ba5d52b164c33d3c57f1c60425e36b29541ca6f9d4644","source":{"kind":"arxiv","id":"2505.05541","version":1},"attestation_state":"computed","paper":{"title":"Safety by Measurement: A Systematic Literature Review of AI Safety Evaluation Methods","license":"http://creativecommons.org/licenses/by-nc-sa/4.0/","headline":"","cross_cats":[],"primary_cat":"cs.AI","authors_text":"Charbel-Rapha\\\"el Segerie, Markov Grey","submitted_at":"2025-05-08T16:55:07Z","abstract_excerpt":"As frontier AI systems advance toward transformative capabilities, we need a parallel transformation in how we measure and evaluate these systems to ensure safety and inform governance. While benchmarks have been the primary method for estimating model capabilities, they often fail to establish true upper bounds or predict deployment behavior. This literature review consolidates the rapidly evolving field of AI safety evaluations, proposing a systematic taxonomy around three dimensions: what properties we measure, how we measure them, and how these measurements integrate into frameworks. We sh"},"verification_status":{"content_addressed":true,"pith_receipt":true,"author_attested":false,"weak_author_claims":0,"strong_author_claims":0,"externally_anchored":false,"storage_verified":false,"citation_signatures":0,"replication_records":0,"graph_snapshot":true,"references_resolved":false,"formal_links_present":false},"canonical_record":{"source":{"id":"2505.05541","kind":"arxiv","version":1},"metadata":{"license":"http://creativecommons.org/licenses/by-nc-sa/4.0/","primary_cat":"cs.AI","submitted_at":"2025-05-08T16:55:07Z","cross_cats_sorted":[],"title_canon_sha256":"db6648453b5542dd2db655baacf7597bb3d52e0f67a84d8eb020daadd135db0b","abstract_canon_sha256":"2a5424e9587c647de2aef9f4b0ee7edae8b245d57e88fe4d4a5a16cb08dc6009"},"schema_version":"1.0"},"receipt":{"kind":"pith_receipt","key_id":"pith-v1-2026-05","algorithm":"ed25519","signed_at":"2026-07-05T11:00:27.672062Z","signature_b64":"z9FeevnmoUqOtL2UN86YIDQLFXz7PCw8/ShL66Ry+fbEi3o5WiSdXVecifhJJQ4vkyljAKT5+wA5eRY752odBA==","signed_message":"canonical_sha256_bytes","builder_version":"pith-number-builder-2026-05-17-v1","receipt_version":"0.3","canonical_sha256":"340e92c6b1502101a41ba5d52b164c33d3c57f1c60425e36b29541ca6f9d4644","last_reissued_at":"2026-07-05T11:00:27.671528Z","signature_status":"signed_v1","first_computed_at":"2026-07-05T11:00:27.671528Z","public_key_fingerprint":"8d4b5ee74e4693bcd1df2446408b0d54"},"graph_snapshot":{"paper":{"title":"Safety by Measurement: A Systematic Literature Review of AI Safety Evaluation Methods","license":"http://creativecommons.org/licenses/by-nc-sa/4.0/","headline":"","cross_cats":[],"primary_cat":"cs.AI","authors_text":"Charbel-Rapha\\\"el Segerie, Markov Grey","submitted_at":"2025-05-08T16:55:07Z","abstract_excerpt":"As frontier AI systems advance toward transformative capabilities, we need a parallel transformation in how we measure and evaluate these systems to ensure safety and inform governance. While benchmarks have been the primary method for estimating model capabilities, they often fail to establish true upper bounds or predict deployment behavior. This literature review consolidates the rapidly evolving field of AI safety evaluations, proposing a systematic taxonomy around three dimensions: what properties we measure, how we measure them, and how these measurements integrate into frameworks. We sh"},"claims":{"count":0,"items":[],"snapshot_sha256":"258153158e38e3291e3d48162225fcdb2d5a3ed65a07baac614ab91432fd4f57"},"source":{"id":"2505.05541","kind":"arxiv","version":1},"verdict":{"id":null,"model_set":{},"created_at":null,"strongest_claim":"","one_line_summary":"","pipeline_version":null,"weakest_assumption":"","pith_extraction_headline":""},"integrity":{"clean":true,"summary":{"advisory":0,"critical":0,"by_detector":{},"informational":0},"endpoint":"/pith/2505.05541/integrity.json","findings":[],"available":true,"detectors_run":[],"snapshot_sha256":"c28c3603d3b5d939e8dc4c7e95fa8dfce3d595e45f758748cecf8e644a296938"},"references":{"count":0,"sample":[],"resolved_work":0,"snapshot_sha256":"258153158e38e3291e3d48162225fcdb2d5a3ed65a07baac614ab91432fd4f57","internal_anchors":0},"formal_canon":{"evidence_count":0,"snapshot_sha256":"258153158e38e3291e3d48162225fcdb2d5a3ed65a07baac614ab91432fd4f57"},"author_claims":{"count":0,"strong_count":0,"snapshot_sha256":"258153158e38e3291e3d48162225fcdb2d5a3ed65a07baac614ab91432fd4f57"},"builder_version":"pith-number-builder-2026-05-17-v1"},"aliases":[{"alias_kind":"arxiv","alias_value":"2505.05541","created_at":"2026-07-05T11:00:27.671588+00:00"},{"alias_kind":"arxiv_version","alias_value":"2505.05541v1","created_at":"2026-07-05T11:00:27.671588+00:00"},{"alias_kind":"doi","alias_value":"10.48550/arxiv.2505.05541","created_at":"2026-07-05T11:00:27.671588+00:00"},{"alias_kind":"pith_short_12","alias_value":"GQHJFRVRKAQQ","created_at":"2026-07-05T11:00:27.671588+00:00"},{"alias_kind":"pith_short_16","alias_value":"GQHJFRVRKAQQDJA3","created_at":"2026-07-05T11:00:27.671588+00:00"},{"alias_kind":"pith_short_8","alias_value":"GQHJFRVR","created_at":"2026-07-05T11:00:27.671588+00:00"}],"events":[],"event_summary":{},"paper_claims":[],"inbound_citations":{"count":3,"internal_anchor_count":0,"sample":[{"citing_arxiv_id":"2508.19932","citing_title":"CASE: An Agentic AI Framework for Enhancing Scam Intelligence in Digital Payments","ref_index":11,"is_internal_anchor":false},{"citing_arxiv_id":"2512.10687","citing_title":"Safe for Whom? Rethinking How We Evaluate the Safety of LLMs for Real Users","ref_index":1,"is_internal_anchor":false},{"citing_arxiv_id":"2604.27911","citing_title":"Physical Foundation Models: Fixed hardware implementations of large-scale neural networks","ref_index":126,"is_internal_anchor":false}]},"formal_canon":{"evidence_count":0,"sample":[],"anchors":[]},"links":{"html":"https://pith.science/pith/GQHJFRVRKAQQDJA3UXKSWFSMGP","json":"https://pith.science/pith/GQHJFRVRKAQQDJA3UXKSWFSMGP.json","graph_json":"https://pith.science/api/pith-number/GQHJFRVRKAQQDJA3UXKSWFSMGP/graph.json","events_json":"https://pith.science/api/pith-number/GQHJFRVRKAQQDJA3UXKSWFSMGP/events.json","paper":"https://pith.science/paper/GQHJFRVR"},"agent_actions":{"view_html":"https://pith.science/pith/GQHJFRVRKAQQDJA3UXKSWFSMGP","download_json":"https://pith.science/pith/GQHJFRVRKAQQDJA3UXKSWFSMGP.json","view_paper":"https://pith.science/paper/GQHJFRVR","resolve_alias":"https://pith.science/api/pith-number/resolve?arxiv=2505.05541&json=true","fetch_graph":"https://pith.science/api/pith-number/GQHJFRVRKAQQDJA3UXKSWFSMGP/graph.json","fetch_events":"https://pith.science/api/pith-number/GQHJFRVRKAQQDJA3UXKSWFSMGP/events.json","actions":{"anchor_timestamp":"https://pith.science/pith/GQHJFRVRKAQQDJA3UXKSWFSMGP/action/timestamp_anchor","attest_storage":"https://pith.science/pith/GQHJFRVRKAQQDJA3UXKSWFSMGP/action/storage_attestation","attest_author":"https://pith.science/pith/GQHJFRVRKAQQDJA3UXKSWFSMGP/action/author_attestation","sign_citation":"https://pith.science/pith/GQHJFRVRKAQQDJA3UXKSWFSMGP/action/citation_signature","submit_replication":"https://pith.science/pith/GQHJFRVRKAQQDJA3UXKSWFSMGP/action/replication_record"}},"created_at":"2026-07-05T11:00:27.671588+00:00","updated_at":"2026-07-05T11:00:27.671588+00:00"}