{"bundle_type":"pith_open_graph_bundle","bundle_version":"1.0","pith_number":"pith:2026:KX2IWKYPRIHDEITF2Z2YNG6BRU","short_pith_number":"pith:KX2IWKYP","canonical_record":{"source":{"id":"2605.21834","kind":"arxiv","version":1},"metadata":{"license":"http://creativecommons.org/licenses/by/4.0/","primary_cat":"cs.LG","submitted_at":"2026-05-20T23:56:40Z","cross_cats_sorted":[],"title_canon_sha256":"719fa8b5629ebb57d79e35501f0f3a16dbf94baece04c9ddc0ccee68bec2dc14","abstract_canon_sha256":"15c5a41bc7370321fad940044980887a162238ba7dea48c1207dd06104e415d3"},"schema_version":"1.0"},"canonical_sha256":"55f48b2b0f8a0e322265d675869bc18d3765e84f8bfd5c641fd81b6199e7c39e","source":{"kind":"arxiv","id":"2605.21834","version":1},"source_aliases":[{"alias_kind":"arxiv","alias_value":"2605.21834","created_at":"2026-05-22T01:04:10Z"},{"alias_kind":"arxiv_version","alias_value":"2605.21834v1","created_at":"2026-05-22T01:04:10Z"},{"alias_kind":"doi","alias_value":"10.48550/arxiv.2605.21834","created_at":"2026-05-22T01:04:10Z"},{"alias_kind":"pith_short_12","alias_value":"KX2IWKYPRIHD","created_at":"2026-05-22T01:04:10Z"},{"alias_kind":"pith_short_16","alias_value":"KX2IWKYPRIHDEITF","created_at":"2026-05-22T01:04:10Z"},{"alias_kind":"pith_short_8","alias_value":"KX2IWKYP","created_at":"2026-05-22T01:04:10Z"}],"events":[{"event_type":"record_created","subject_pith_number":"pith:2026:KX2IWKYPRIHDEITF2Z2YNG6BRU","target":"record","payload":{"canonical_record":{"source":{"id":"2605.21834","kind":"arxiv","version":1},"metadata":{"license":"http://creativecommons.org/licenses/by/4.0/","primary_cat":"cs.LG","submitted_at":"2026-05-20T23:56:40Z","cross_cats_sorted":[],"title_canon_sha256":"719fa8b5629ebb57d79e35501f0f3a16dbf94baece04c9ddc0ccee68bec2dc14","abstract_canon_sha256":"15c5a41bc7370321fad940044980887a162238ba7dea48c1207dd06104e415d3"},"schema_version":"1.0"},"canonical_sha256":"55f48b2b0f8a0e322265d675869bc18d3765e84f8bfd5c641fd81b6199e7c39e","receipt":{"kind":"pith_receipt","key_id":"pith-v1-2026-05","algorithm":"ed25519","signed_at":"2026-05-22T01:04:10.168627Z","signature_b64":"vVMRF0IQlLu8uqlvki0QAog45riqY0LwJ+WeduPCOY9LsN//Xjq15ME8b6wVLD18WiHl4DBLjVoq67s8ebUGCw==","signed_message":"canonical_sha256_bytes","builder_version":"pith-number-builder-2026-05-17-v1","receipt_version":"0.3","canonical_sha256":"55f48b2b0f8a0e322265d675869bc18d3765e84f8bfd5c641fd81b6199e7c39e","last_reissued_at":"2026-05-22T01:04:10.167818Z","signature_status":"signed_v1","first_computed_at":"2026-05-22T01:04:10.167818Z","public_key_fingerprint":"8d4b5ee74e4693bcd1df2446408b0d54"},"source_kind":"arxiv","source_id":"2605.21834","source_version":1,"attestation_state":"computed"},"signer":{"signer_id":"pith.science","signer_type":"pith_registry","key_id":"pith-v1-2026-05","public_key_fingerprint":"8d4b5ee74e4693bcd1df2446408b0d54"},"created_at":"2026-05-22T01:04:10Z","supersedes":[],"prev_event":null,"signature":{"signature_status":"signed_v1","algorithm":"ed25519","key_id":"pith-v1-2026-05","public_key_fingerprint":"8d4b5ee74e4693bcd1df2446408b0d54","signature_b64":"V5JEXjE4FXuoQcCxCS/dBZNGjGLMMrctRzeJUOxFyJwhRtb8Mba8mSTMV/3U0EtwQ1oNT7mxp9niHUCCHD6fAA==","signed_message":"open_graph_event_sha256_bytes","signed_at":"2026-05-25T22:52:53.028885Z"},"content_sha256":"4f1830821d2d1d043eb6b2e51691f1f1e759a262c5a1861f356ce07ba24cc4b5","schema_version":"1.0","event_id":"sha256:4f1830821d2d1d043eb6b2e51691f1f1e759a262c5a1861f356ce07ba24cc4b5"},{"event_type":"graph_snapshot","subject_pith_number":"pith:2026:KX2IWKYPRIHDEITF2Z2YNG6BRU","target":"graph","payload":{"graph_snapshot":{"paper":{"title":"On-Policy Consistency Training Improves LLM Safety with Minimal Capability Degradation","license":"http://creativecommons.org/licenses/by/4.0/","headline":"","cross_cats":[],"primary_cat":"cs.LG","authors_text":"Andy Han, Avidan Shah, Chen Yueh-Han, Ilia Sucholutsky, Kai Xu, Kiet Nguyen, Kristina Fujimoto, Rico Angell","submitted_at":"2026-05-20T23:56:40Z","abstract_excerpt":"Aligned models can misbehave in several ways: they are often sycophantic, fall victim to jailbreaks, or fail to include appropriate safety warnings. Consistency training is a promising new alignment paradigm to mitigate such failures by training invariants into the model using contrastive input pairs. Existing consistency training procedures generate the supervision signal once, offline, and use supervised fine-tuning (SFT) to update the model. Unfortunately, the resulting models tend to merely memorize the surface forms of the training distribution and thus generalize poorly and regress in th"},"claims":{"count":0,"items":[],"snapshot_sha256":"258153158e38e3291e3d48162225fcdb2d5a3ed65a07baac614ab91432fd4f57"},"source":{"id":"2605.21834","kind":"arxiv","version":1},"verdict":{"id":null,"model_set":{},"created_at":null,"strongest_claim":"","one_line_summary":"","pipeline_version":null,"weakest_assumption":"","pith_extraction_headline":""},"integrity":{"clean":true,"summary":{"advisory":0,"critical":0,"by_detector":{},"informational":0},"endpoint":"/pith/2605.21834/integrity.json","findings":[],"available":true,"detectors_run":[],"snapshot_sha256":"c28c3603d3b5d939e8dc4c7e95fa8dfce3d595e45f758748cecf8e644a296938"},"references":{"count":0,"sample":[],"resolved_work":0,"snapshot_sha256":"258153158e38e3291e3d48162225fcdb2d5a3ed65a07baac614ab91432fd4f57","internal_anchors":0},"formal_canon":{"evidence_count":0,"snapshot_sha256":"258153158e38e3291e3d48162225fcdb2d5a3ed65a07baac614ab91432fd4f57"},"author_claims":{"count":0,"strong_count":0,"snapshot_sha256":"258153158e38e3291e3d48162225fcdb2d5a3ed65a07baac614ab91432fd4f57"},"builder_version":"pith-number-builder-2026-05-17-v1"},"verdict_id":null},"signer":{"signer_id":"pith.science","signer_type":"pith_registry","key_id":"pith-v1-2026-05","public_key_fingerprint":"8d4b5ee74e4693bcd1df2446408b0d54"},"created_at":"2026-05-22T01:04:10Z","supersedes":[],"prev_event":null,"signature":{"signature_status":"signed_v1","algorithm":"ed25519","key_id":"pith-v1-2026-05","public_key_fingerprint":"8d4b5ee74e4693bcd1df2446408b0d54","signature_b64":"jRkj1NRsig/stJs0wnMv8rfXz63qalgqvlGWaymdP3s2x7Ar17ByvX7e4A7poSofLaemR8itx6VWutIj36/qAA==","signed_message":"open_graph_event_sha256_bytes","signed_at":"2026-05-25T22:52:53.029561Z"},"content_sha256":"43df767ab778431fea32b0759338b645d0170e68ea16e6fa292ec751cefdc609","schema_version":"1.0","event_id":"sha256:43df767ab778431fea32b0759338b645d0170e68ea16e6fa292ec751cefdc609"}],"timestamp_proofs":[],"mirror_hints":[{"mirror_type":"https","name":"Pith Resolver","base_url":"https://pith.science","bundle_url":"https://pith.science/pith/KX2IWKYPRIHDEITF2Z2YNG6BRU/bundle.json","state_url":"https://pith.science/pith/KX2IWKYPRIHDEITF2Z2YNG6BRU/state.json","well_known_bundle_url":"https://pith.science/.well-known/pith/KX2IWKYPRIHDEITF2Z2YNG6BRU/bundle.json","status":"primary"}],"public_keys":[{"key_id":"pith-v1-2026-05","algorithm":"ed25519","format":"raw","public_key_b64":"stVStoiQhXFxp4s2pdzPNoqVNBMojDU/fJ2db5S3CbM=","public_key_hex":"b2d552b68890857171a78b36a5dccf368a953413288c353f7c9d9d6f94b709b3","fingerprint_sha256_b32_first128bits":"RVFV5Z2OI2J3ZUO7ERDEBCYNKS","fingerprint_sha256_hex":"8d4b5ee74e4693bcd1df2446408b0d54","rotates_at":null,"url":"https://pith.science/pith-signing-key.json","notes":"Pith uses this Ed25519 key to sign canonical record SHA-256 digests. Verify with: ed25519_verify(public_key, message=canonical_sha256_bytes, signature=base64decode(signature_b64))."}],"merge_version":"pith-open-graph-merge-v1","built_at":"2026-05-25T22:52:53Z","links":{"resolver":"https://pith.science/pith/KX2IWKYPRIHDEITF2Z2YNG6BRU","bundle":"https://pith.science/pith/KX2IWKYPRIHDEITF2Z2YNG6BRU/bundle.json","state":"https://pith.science/pith/KX2IWKYPRIHDEITF2Z2YNG6BRU/state.json","well_known_bundle":"https://pith.science/.well-known/pith/KX2IWKYPRIHDEITF2Z2YNG6BRU/bundle.json"},"state":{"state_type":"pith_open_graph_state","state_version":"1.0","pith_number":"pith:2026:KX2IWKYPRIHDEITF2Z2YNG6BRU","merge_version":"pith-open-graph-merge-v1","event_count":2,"valid_event_count":2,"invalid_event_count":0,"equivocation_count":0,"current":{"canonical_record":{"metadata":{"abstract_canon_sha256":"15c5a41bc7370321fad940044980887a162238ba7dea48c1207dd06104e415d3","cross_cats_sorted":[],"license":"http://creativecommons.org/licenses/by/4.0/","primary_cat":"cs.LG","submitted_at":"2026-05-20T23:56:40Z","title_canon_sha256":"719fa8b5629ebb57d79e35501f0f3a16dbf94baece04c9ddc0ccee68bec2dc14"},"schema_version":"1.0","source":{"id":"2605.21834","kind":"arxiv","version":1}},"source_aliases":[{"alias_kind":"arxiv","alias_value":"2605.21834","created_at":"2026-05-22T01:04:10Z"},{"alias_kind":"arxiv_version","alias_value":"2605.21834v1","created_at":"2026-05-22T01:04:10Z"},{"alias_kind":"doi","alias_value":"10.48550/arxiv.2605.21834","created_at":"2026-05-22T01:04:10Z"},{"alias_kind":"pith_short_12","alias_value":"KX2IWKYPRIHD","created_at":"2026-05-22T01:04:10Z"},{"alias_kind":"pith_short_16","alias_value":"KX2IWKYPRIHDEITF","created_at":"2026-05-22T01:04:10Z"},{"alias_kind":"pith_short_8","alias_value":"KX2IWKYP","created_at":"2026-05-22T01:04:10Z"}],"graph_snapshots":[{"event_id":"sha256:43df767ab778431fea32b0759338b645d0170e68ea16e6fa292ec751cefdc609","target":"graph","created_at":"2026-05-22T01:04:10Z","signer":{"key_id":"pith-v1-2026-05","public_key_fingerprint":"8d4b5ee74e4693bcd1df2446408b0d54","signer_id":"pith.science","signer_type":"pith_registry"},"payload":{"graph_snapshot":{"author_claims":{"count":0,"snapshot_sha256":"258153158e38e3291e3d48162225fcdb2d5a3ed65a07baac614ab91432fd4f57","strong_count":0},"builder_version":"pith-number-builder-2026-05-17-v1","claims":{"count":0,"items":[],"snapshot_sha256":"258153158e38e3291e3d48162225fcdb2d5a3ed65a07baac614ab91432fd4f57"},"formal_canon":{"evidence_count":0,"snapshot_sha256":"258153158e38e3291e3d48162225fcdb2d5a3ed65a07baac614ab91432fd4f57"},"integrity":{"available":true,"clean":true,"detectors_run":[],"endpoint":"/pith/2605.21834/integrity.json","findings":[],"snapshot_sha256":"c28c3603d3b5d939e8dc4c7e95fa8dfce3d595e45f758748cecf8e644a296938","summary":{"advisory":0,"by_detector":{},"critical":0,"informational":0}},"paper":{"abstract_excerpt":"Aligned models can misbehave in several ways: they are often sycophantic, fall victim to jailbreaks, or fail to include appropriate safety warnings. Consistency training is a promising new alignment paradigm to mitigate such failures by training invariants into the model using contrastive input pairs. Existing consistency training procedures generate the supervision signal once, offline, and use supervised fine-tuning (SFT) to update the model. Unfortunately, the resulting models tend to merely memorize the surface forms of the training distribution and thus generalize poorly and regress in th","authors_text":"Andy Han, Avidan Shah, Chen Yueh-Han, Ilia Sucholutsky, Kai Xu, Kiet Nguyen, Kristina Fujimoto, Rico Angell","cross_cats":[],"headline":"","license":"http://creativecommons.org/licenses/by/4.0/","primary_cat":"cs.LG","submitted_at":"2026-05-20T23:56:40Z","title":"On-Policy Consistency Training Improves LLM Safety with Minimal Capability Degradation"},"references":{"count":0,"internal_anchors":0,"resolved_work":0,"sample":[],"snapshot_sha256":"258153158e38e3291e3d48162225fcdb2d5a3ed65a07baac614ab91432fd4f57"},"source":{"id":"2605.21834","kind":"arxiv","version":1},"verdict":{"created_at":null,"id":null,"model_set":{},"one_line_summary":"","pipeline_version":null,"pith_extraction_headline":"","strongest_claim":"","weakest_assumption":""}},"verdict_id":null}}],"author_attestations":[],"timestamp_anchors":[],"storage_attestations":[],"citation_signatures":[],"replication_records":[],"corrections":[],"mirror_hints":[],"record_created":{"event_id":"sha256:4f1830821d2d1d043eb6b2e51691f1f1e759a262c5a1861f356ce07ba24cc4b5","target":"record","created_at":"2026-05-22T01:04:10Z","signer":{"key_id":"pith-v1-2026-05","public_key_fingerprint":"8d4b5ee74e4693bcd1df2446408b0d54","signer_id":"pith.science","signer_type":"pith_registry"},"payload":{"attestation_state":"computed","canonical_record":{"metadata":{"abstract_canon_sha256":"15c5a41bc7370321fad940044980887a162238ba7dea48c1207dd06104e415d3","cross_cats_sorted":[],"license":"http://creativecommons.org/licenses/by/4.0/","primary_cat":"cs.LG","submitted_at":"2026-05-20T23:56:40Z","title_canon_sha256":"719fa8b5629ebb57d79e35501f0f3a16dbf94baece04c9ddc0ccee68bec2dc14"},"schema_version":"1.0","source":{"id":"2605.21834","kind":"arxiv","version":1}},"canonical_sha256":"55f48b2b0f8a0e322265d675869bc18d3765e84f8bfd5c641fd81b6199e7c39e","receipt":{"algorithm":"ed25519","builder_version":"pith-number-builder-2026-05-17-v1","canonical_sha256":"55f48b2b0f8a0e322265d675869bc18d3765e84f8bfd5c641fd81b6199e7c39e","first_computed_at":"2026-05-22T01:04:10.167818Z","key_id":"pith-v1-2026-05","kind":"pith_receipt","last_reissued_at":"2026-05-22T01:04:10.167818Z","public_key_fingerprint":"8d4b5ee74e4693bcd1df2446408b0d54","receipt_version":"0.3","signature_b64":"vVMRF0IQlLu8uqlvki0QAog45riqY0LwJ+WeduPCOY9LsN//Xjq15ME8b6wVLD18WiHl4DBLjVoq67s8ebUGCw==","signature_status":"signed_v1","signed_at":"2026-05-22T01:04:10.168627Z","signed_message":"canonical_sha256_bytes"},"source_id":"2605.21834","source_kind":"arxiv","source_version":1}}},"equivocations":[],"invalid_events":[],"applied_event_ids":["sha256:4f1830821d2d1d043eb6b2e51691f1f1e759a262c5a1861f356ce07ba24cc4b5","sha256:43df767ab778431fea32b0759338b645d0170e68ea16e6fa292ec751cefdc609"],"state_sha256":"bdd61bc82700467148d80059b89128caa47a04ee04d67a3c8592d69abb00ca29"},"bundle_signature":{"signature_status":"signed_v1","algorithm":"ed25519","key_id":"pith-v1-2026-05","public_key_fingerprint":"8d4b5ee74e4693bcd1df2446408b0d54","signature_b64":"aR6lNQXEzG3c42QGo7Ja2mmiLOc6WnOMTCI52djQTecyln/afEP5zh0WJStqMhguUXAyGa0jUJFLIen9avz0CQ==","signed_message":"bundle_sha256_bytes","signed_at":"2026-05-25T22:52:53.032740Z","bundle_sha256":"1efd94403cf8b335667b1b9e4b029473c0956a02c50b53b1996b631696b003f9"}}