{"as_of":"2026-08-04T09:57:00Z","caps":{"database_statements":6,"inbound":100,"outbound":100},"context_digest":"sha256:512a0e6aae709a9c7bbccf76259dee87f811b36ed4da951750c4a7657d7083d7","coverage":[{"denominator":42,"lane":"reference_resolution","note":"Typed states for the displayed outbound observations.","records_observed":42,"source":"paper_references, paper_reference_links","source_observed_at":"2026-05-19T16:34:47.856606Z","state":"measured"},{"denominator":42,"lane":"standing_notices","note":"One-hop event checks from named stored sources.","records_observed":42,"source":"scholarly_work_events, retraction_status_cache","source_observed_at":"2026-08-04T06:34:03.388597+00:00","state":"measured"},{"denominator":0,"lane":"inbound_itemization","note":"Pith citing papers itemized under the disclosed page cap.","records_observed":0,"source":"paper_references, paper_reference_links","source_observed_at":null,"state":"measured"},{"denominator":1,"lane":"external_citation_measurements","note":"A source-named dated measurement, never combined with another source.","records_observed":0,"source":"cited_works","source_observed_at":null,"state":"measured"}],"external_citation_measurements":[],"inbound":[],"links":{"evidence":"/evidence","html":"/paper/2605.15239/citation-record","integrity":"/paper/2605.15239/integrity","json":"/paper/2605.15239/citation-record.json","paper":"/paper/2605.15239"},"outbound":[{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-07-09T17:56:26.096579Z","title":"Advances in neural information processing systems , volume=","venue":null,"work_id":"c25e8154-fab2-455c-8a26-56e40aed5d2b","year":null},"citing_paper":{"arxiv_id":"2605.15239","last_updated":"2026-05-14T03:40:07Z","snapshot_observed_at":"2026-07-06T23:26:32.979566Z","submitted_at":"2026-05-14T03:40:07Z","title":"Reducing the Safety Tax in LLM Safety Alignment with On-Policy Self-Distillation","version":1},"reference_index":1,"source":"arxiv_source","source_observed_at":"2026-05-19T16:34:47.856606Z"},"links":{"citing_paper":"/paper/2605.15239"},"observation_digest":"sha256:9da2ea57dcdc7d23fbb0acadb66af612b8b439e41d4e1aca3f30aa27eb734d55","observation_id":"1647b194-fba2-4184-bf6b-0ec83b0b12aa","resolution":{"observed_at":"2026-05-19T16:37:40.520366Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-04T06:34:03.388597+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-04T06:34:03.388597+00:00","source":"crossref"},{"observed_at":"2026-08-04T06:33:57.428241+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2204.05862","last_updated":"2022-04-12T15:02:38Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2022-04-12T15:02:38Z","title":"Training a Helpful and Harmless Assistant with Reinforcement Learning from Human Feedback","version":1},"cited_work":{"arxiv_id":"2204.05862","doi":"10.1016/j.respol.2005.01.014","metadata_source":"pith","pith_arxiv_id":"2204.05862","snapshot_observed_at":"2026-07-11T11:50:26.030339Z","title":"Training a Helpful and Harmless Assistant with Reinforcement Learning from Human Feedback","venue":"cs.CL","work_id":"a1f2574b-a899-4713-be60-c87ba332656c","year":2022},"citing_paper":{"arxiv_id":"2605.15239","last_updated":"2026-05-14T03:40:07Z","snapshot_observed_at":"2026-07-06T23:26:32.979566Z","submitted_at":"2026-05-14T03:40:07Z","title":"Reducing the Safety Tax in LLM Safety Alignment with On-Policy Self-Distillation","version":1},"reference_index":2,"source":"arxiv_source","source_observed_at":"2026-05-19T16:34:47.856606Z"},"links":{"cited_paper":"/paper/2204.05862","citing_paper":"/paper/2605.15239"},"observation_digest":"sha256:f437ee7c6f27f12280b39fd029edf044f2f7cc1cdfa650000ede2cea8285eac7","observation_id":"b3801cd0-768d-4307-8c7b-0f136a89ba46","resolution":{"observed_at":"2026-05-19T16:37:39.902681Z","resolver_source":"local_arxiv","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-04T06:34:03.388597+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-04T06:34:03.388597+00:00","source":"crossref"},{"observed_at":"2026-08-04T06:33:57.428241+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2407.04295","last_updated":"2024-08-30T11:57:47Z","snapshot_observed_at":"2026-08-03T06:32:58.433190Z","submitted_at":"2024-07-05T06:57:30Z","title":"Jailbreak Attacks and Defenses Against Large Language Models: A Survey","version":2},"cited_work":{"arxiv_id":"2407.04295","doi":"10.48550/arxiv.2407.04295","metadata_source":"pith","pith_arxiv_id":"2407.04295","snapshot_observed_at":"2026-07-10T06:15:00.866473Z","title":"Jailbreak Attacks and Defenses Against Large Language Models: A Survey","venue":"cs.CR","work_id":"0ee7fc45-ae61-432b-83ac-f1d93ccd88fb","year":2024},"citing_paper":{"arxiv_id":"2605.15239","last_updated":"2026-05-14T03:40:07Z","snapshot_observed_at":"2026-07-06T23:26:32.979566Z","submitted_at":"2026-05-14T03:40:07Z","title":"Reducing the Safety Tax in LLM Safety Alignment with On-Policy Self-Distillation","version":1},"reference_index":3,"source":"arxiv_source","source_observed_at":"2026-05-19T16:34:47.856606Z"},"links":{"cited_paper":"/paper/2407.04295","citing_paper":"/paper/2605.15239"},"observation_digest":"sha256:2762236f5ebad1c24b475289b0730241a6410b05678c23bace662955d89d14c7","observation_id":"35f5cc9c-5bdc-484e-8b07-ee502238aa0f","resolution":{"observed_at":"2026-05-19T16:37:39.882900Z","resolver_source":"local_arxiv","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-04T06:34:03.388597+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-04T06:34:03.388597+00:00","source":"crossref"},{"observed_at":"2026-05-24T00:53:04.67981+00:00","source":"crossref_status_cache"},{"observed_at":"2026-05-24T00:53:04.67981+00:00","source":"openalex_status_cache"},{"observed_at":"2026-08-04T06:33:57.428241+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Findings of the Association for Computational Linguistics: ACL 2025 , pages=","venue":null,"work_id":"923339bc-bc78-421a-9e00-8feeec5339ad","year":2025},"citing_paper":{"arxiv_id":"2605.15239","last_updated":"2026-05-14T03:40:07Z","snapshot_observed_at":"2026-07-06T23:26:32.979566Z","submitted_at":"2026-05-14T03:40:07Z","title":"Reducing the Safety Tax in LLM Safety Alignment with On-Policy Self-Distillation","version":1},"reference_index":4,"source":"arxiv_source","source_observed_at":"2026-05-19T16:34:47.856606Z"},"links":{"citing_paper":"/paper/2605.15239"},"observation_digest":"sha256:001b2111254a68911d879a296956fb630c718df01410f08610e1469d26033d52","observation_id":"188ceb5e-da97-47ec-8113-59e95365f06f","resolution":{"observed_at":"2026-05-19T16:37:40.518240Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-04T06:34:03.388597+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-04T06:34:03.388597+00:00","source":"crossref"},{"observed_at":"2026-08-04T06:33:57.428241+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Proceedings of the AAAI Conference on Artificial Intelligence , volume=","venue":null,"work_id":"e125f746-28c3-4ec3-a3a8-dc26bdb4d96b","year":null},"citing_paper":{"arxiv_id":"2605.15239","last_updated":"2026-05-14T03:40:07Z","snapshot_observed_at":"2026-07-06T23:26:32.979566Z","submitted_at":"2026-05-14T03:40:07Z","title":"Reducing the Safety Tax in LLM Safety Alignment with On-Policy Self-Distillation","version":1},"reference_index":5,"source":"arxiv_source","source_observed_at":"2026-05-19T16:34:47.856606Z"},"links":{"citing_paper":"/paper/2605.15239"},"observation_digest":"sha256:87a7138a5b33b445e8ed8c31eb8d8026d9ec937bf5fb6e58baddb0eef9b0c15f","observation_id":"4dcd14a2-c564-49a4-bcfc-c5eb5f3526c6","resolution":{"observed_at":"2026-05-19T16:37:40.512706Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-04T06:34:03.388597+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-04T06:34:03.388597+00:00","source":"crossref"},{"observed_at":"2026-08-04T06:33:57.428241+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-07-06T16:42:39.969253Z","title":"Proceedings of the 2025 Conference on Empirical Methods in Natural Language Processing , pages=","venue":null,"work_id":"89d2a70a-9313-4909-acde-2d404fd99285","year":2025},"citing_paper":{"arxiv_id":"2605.15239","last_updated":"2026-05-14T03:40:07Z","snapshot_observed_at":"2026-07-06T23:26:32.979566Z","submitted_at":"2026-05-14T03:40:07Z","title":"Reducing the Safety Tax in LLM Safety Alignment with On-Policy Self-Distillation","version":1},"reference_index":6,"source":"arxiv_source","source_observed_at":"2026-05-19T16:34:47.856606Z"},"links":{"citing_paper":"/paper/2605.15239"},"observation_digest":"sha256:2b6bcf7e49dce9df88311f7e1b60b8726e2d81eac9be13eee6bfafb75492c0bc","observation_id":"624b080a-b7e1-435e-98a5-69f03b29e923","resolution":{"observed_at":"2026-05-19T16:37:40.514665Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-04T06:34:03.388597+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-04T06:34:03.388597+00:00","source":"crossref"},{"observed_at":"2026-08-04T06:33:57.428241+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2505.09427","last_updated":"2025-05-15T07:22:20Z","snapshot_observed_at":"2026-07-06T21:23:51.965886Z","submitted_at":"2025-05-14T14:28:24Z","title":"SafePath: Conformal Prediction for Safe LLM-Based Autonomous Navigation","version":2},"cited_work":{"arxiv_id":"2505.09427","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2505.09427","snapshot_observed_at":"2026-07-03T23:29:02.512774Z","title":"arXiv preprint arXiv:2505.09427 , year=","venue":null,"work_id":"2b79d9d0-49f7-4a74-9acc-38ad36ced805","year":2025},"citing_paper":{"arxiv_id":"2605.15239","last_updated":"2026-05-14T03:40:07Z","snapshot_observed_at":"2026-07-06T23:26:32.979566Z","submitted_at":"2026-05-14T03:40:07Z","title":"Reducing the Safety Tax in LLM Safety Alignment with On-Policy Self-Distillation","version":1},"reference_index":7,"source":"arxiv_source","source_observed_at":"2026-05-19T16:34:47.856606Z"},"links":{"cited_paper":"/paper/2505.09427","citing_paper":"/paper/2605.15239"},"observation_digest":"sha256:a55015c5d64be4c9d36794310442d382fef348f0f06549df0f29cd876265bd7d","observation_id":"30fbb816-b491-4503-9819-297fb8ac689a","resolution":{"observed_at":"2026-05-19T16:37:39.906065Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-04T06:34:03.388597+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-04T06:34:03.388597+00:00","source":"crossref"},{"observed_at":"2026-08-04T06:33:57.428241+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2309.07875","last_updated":"2024-03-19T16:50:50Z","snapshot_observed_at":"2026-07-06T16:18:37.399390Z","submitted_at":"2023-09-14T17:23:37Z","title":"Safety-Tuned LLaMAs: Lessons From Improving the Safety of Large Language Models that Follow Instructions","version":3},"cited_work":{"arxiv_id":"2309.07875","doi":"10.48550/arxiv.2309.07875","metadata_source":"pith","pith_arxiv_id":"2309.07875","snapshot_observed_at":"2026-07-10T06:15:00.866473Z","title":"Safety-tuned llamas: Lessons from improving the safety of large language models that follow instructions","venue":"cs.CL","work_id":"4f3956ef-5c1a-4536-8f08-ab1c9a6ce6d5","year":2023},"citing_paper":{"arxiv_id":"2605.15239","last_updated":"2026-05-14T03:40:07Z","snapshot_observed_at":"2026-07-06T23:26:32.979566Z","submitted_at":"2026-05-14T03:40:07Z","title":"Reducing the Safety Tax in LLM Safety Alignment with On-Policy Self-Distillation","version":1},"reference_index":8,"source":"arxiv_source","source_observed_at":"2026-05-19T16:34:47.856606Z"},"links":{"cited_paper":"/paper/2309.07875","citing_paper":"/paper/2605.15239"},"observation_digest":"sha256:92c0232443437a6cb88b3736610a21df4812fd5556937145e52d8c85dec7bfb5","observation_id":"7fb23de7-ab3c-457a-863f-1c9f3eab6c0a","resolution":{"observed_at":"2026-05-19T16:37:39.913018Z","resolver_source":"arxiv_id","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-04T06:34:03.388597+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-04T06:34:03.388597+00:00","source":"crossref"},{"observed_at":"2026-08-04T06:33:57.428241+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2310.03693","last_updated":"2023-10-05T17:12:17Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2023-10-05T17:12:17Z","title":"Fine-tuning Aligned Language Models Compromises Safety, Even When Users Do Not Intend To!","version":1},"cited_work":{"arxiv_id":"2310.03693","doi":"10.48550/arxiv.2310.03693","metadata_source":"pith","pith_arxiv_id":"2310.03693","snapshot_observed_at":"2026-07-11T00:27:50.654851Z","title":"Fine-tuning Aligned Language Models Compromises Safety, Even When Users Do Not Intend To!","venue":"cs.CL","work_id":"8b07137a-7175-4dd1-a8d9-570493d3f404","year":2023},"citing_paper":{"arxiv_id":"2605.15239","last_updated":"2026-05-14T03:40:07Z","snapshot_observed_at":"2026-07-06T23:26:32.979566Z","submitted_at":"2026-05-14T03:40:07Z","title":"Reducing the Safety Tax in LLM Safety Alignment with On-Policy Self-Distillation","version":1},"reference_index":9,"source":"arxiv_source","source_observed_at":"2026-05-19T16:34:47.856606Z"},"links":{"cited_paper":"/paper/2310.03693","citing_paper":"/paper/2605.15239"},"observation_digest":"sha256:02a50a37c9f5ead00b42ce0173593e8518e59e113c6180bf55a38b71bc6c599c","observation_id":"66d46433-fcf6-4cf6-86b4-75736e1d4b03","resolution":{"observed_at":"2026-05-19T16:37:39.891487Z","resolver_source":"local_arxiv","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-04T06:34:03.388597+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-04T06:34:03.388597+00:00","source":"crossref"},{"observed_at":"2026-08-04T06:33:57.428241+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2502.09673","last_updated":"2025-02-21T03:12:17Z","snapshot_observed_at":"2026-07-06T20:36:17.029825Z","submitted_at":"2025-02-13T06:37:28Z","title":"Are Smarter LLMs Safer? Exploring Safety-Reasoning Trade-offs in Prompting and Fine-Tuning","version":2},"cited_work":{"arxiv_id":"2502.09673","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2502.09673","snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Are Smarter LLMs Safer? Exploring Safety- Reasoning Trade-offs in Prompting and Fine-Tuning.CoRR abs/2502.09673","venue":null,"work_id":"87390412-811f-482c-998b-04197e99da8f","year":2025},"citing_paper":{"arxiv_id":"2605.15239","last_updated":"2026-05-14T03:40:07Z","snapshot_observed_at":"2026-07-06T23:26:32.979566Z","submitted_at":"2026-05-14T03:40:07Z","title":"Reducing the Safety Tax in LLM Safety Alignment with On-Policy Self-Distillation","version":1},"reference_index":10,"source":"arxiv_source","source_observed_at":"2026-05-19T16:34:47.856606Z"},"links":{"cited_paper":"/paper/2502.09673","citing_paper":"/paper/2605.15239"},"observation_digest":"sha256:c1f0e2fc2f599134c963694848359ec4c5b74fc6a3b15f81d1040ac9adcab0b8","observation_id":"e3658083-daf6-45fe-875e-33ddb4234433","resolution":{"observed_at":"2026-05-19T16:37:39.908935Z","resolver_source":"arxiv_id","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-04T06:34:03.388597+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-04T06:34:03.388597+00:00","source":"crossref"},{"observed_at":"2026-08-04T06:33:57.428241+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2503.00555","last_updated":"2025-06-05T03:20:54Z","snapshot_observed_at":"2026-08-01T20:19:04.242454Z","submitted_at":"2025-03-01T16:42:01Z","title":"Safety Tax: Safety Alignment Makes Your Large Reasoning Models Less Reasonable","version":2},"cited_work":{"arxiv_id":"2503.00555","doi":"10.48550/arxiv.2503.00555","metadata_source":"arxiv_reference","pith_arxiv_id":"2503.00555","snapshot_observed_at":"2026-07-10T12:15:01.137692Z","title":"Safety tax: Safety alignment makes your large reasoning models less reasonable","venue":null,"work_id":"b0e63685-280f-43d8-adf7-2603cb1c2798","year":2025},"citing_paper":{"arxiv_id":"2605.15239","last_updated":"2026-05-14T03:40:07Z","snapshot_observed_at":"2026-07-06T23:26:32.979566Z","submitted_at":"2026-05-14T03:40:07Z","title":"Reducing the Safety Tax in LLM Safety Alignment with On-Policy Self-Distillation","version":1},"reference_index":11,"source":"arxiv_source","source_observed_at":"2026-05-19T16:34:47.856606Z"},"links":{"cited_paper":"/paper/2503.00555","citing_paper":"/paper/2605.15239"},"observation_digest":"sha256:b9dce100f51fda9913ec880f18b8d63f3f727e2c97e66c3c642cad9ef5ef5db4","observation_id":"c6a46327-1a16-44ef-bf8b-ca4471668f44","resolution":{"observed_at":"2026-05-19T16:37:39.880192Z","resolver_source":"arxiv_id","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-04T06:34:03.388597+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-04T06:34:03.388597+00:00","source":"crossref"},{"observed_at":"2026-08-04T06:33:57.428241+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2601.23143","last_updated":"2026-05-13T02:16:49Z","snapshot_observed_at":"2026-07-06T22:43:45.998263Z","submitted_at":"2026-01-30T16:31:02Z","title":"THINKSAFE: Self-Generated Safety Alignment for Reasoning Models","version":4},"cited_work":{"arxiv_id":"2601.23143","doi":null,"metadata_source":"pith","pith_arxiv_id":"2601.23143","snapshot_observed_at":"2026-07-04T17:30:00.640646Z","title":"THINKSAFE: Self-Generated Safety Alignment for Reasoning Models","venue":"cs.AI","work_id":"3f512943-c116-4d07-922b-8bde1756c0ff","year":2026},"citing_paper":{"arxiv_id":"2605.15239","last_updated":"2026-05-14T03:40:07Z","snapshot_observed_at":"2026-07-06T23:26:32.979566Z","submitted_at":"2026-05-14T03:40:07Z","title":"Reducing the Safety Tax in LLM Safety Alignment with On-Policy Self-Distillation","version":1},"reference_index":12,"source":"arxiv_source","source_observed_at":"2026-05-19T16:34:47.856606Z"},"links":{"cited_paper":"/paper/2601.23143","citing_paper":"/paper/2605.15239"},"observation_digest":"sha256:26b4fba10a7c9bfd55bc93bf1106d0e905e7d66a047d11c1bb3c426119ec7f24","observation_id":"4bf50f87-b1f6-4e6d-ba63-84ef7768e83e","resolution":{"observed_at":"2026-05-19T16:37:39.899800Z","resolver_source":"local_arxiv","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-04T06:34:03.388597+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-04T06:34:03.388597+00:00","source":"crossref"},{"observed_at":"2026-08-04T06:33:57.428241+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2406.05946","last_updated":"2024-06-10T00:35:23Z","snapshot_observed_at":"2026-07-06T18:27:54.379381Z","submitted_at":"2024-06-10T00:35:23Z","title":"Safety Alignment Should Be Made More Than Just a Few Tokens Deep","version":1},"cited_work":{"arxiv_id":"2406.05946","doi":"10.48550/arxiv.2406.05946","metadata_source":"arxiv_reference","pith_arxiv_id":"2406.05946","snapshot_observed_at":"2026-07-10T06:15:00.866473Z","title":"Safety alignment should be made more than just a few tokens deep","venue":null,"work_id":"40539ea6-b6ba-4195-971b-53c52efa9597","year":2024},"citing_paper":{"arxiv_id":"2605.15239","last_updated":"2026-05-14T03:40:07Z","snapshot_observed_at":"2026-07-06T23:26:32.979566Z","submitted_at":"2026-05-14T03:40:07Z","title":"Reducing the Safety Tax in LLM Safety Alignment with On-Policy Self-Distillation","version":1},"reference_index":13,"source":"arxiv_source","source_observed_at":"2026-05-19T16:34:47.856606Z"},"links":{"cited_paper":"/paper/2406.05946","citing_paper":"/paper/2605.15239"},"observation_digest":"sha256:83f843de2057e19b0656e8f33aba46845538c9007d69e1d5e88583023a13833e","observation_id":"b8910e1c-87c9-4d1a-ab7a-1ba54351b815","resolution":{"observed_at":"2026-05-19T16:37:39.886052Z","resolver_source":"arxiv_id","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-04T06:34:03.388597+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-04T06:34:03.388597+00:00","source":"crossref"},{"observed_at":"2026-08-04T06:33:57.428241+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2601.18734","last_updated":"2026-03-20T15:40:19Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2026-01-26T17:56:50Z","title":"Self-Distilled Reasoner: On-Policy Self-Distillation for Large Language Models","version":3},"cited_work":{"arxiv_id":"2601.18734","doi":"10.18653/v1/2025.emnlp-main.125.https://aclanthology.org/2025.emnlp-main.125/","metadata_source":"pith","pith_arxiv_id":"2601.18734","snapshot_observed_at":"2026-07-11T11:50:26.030339Z","title":"Self-Distilled Reasoner: On-Policy Self-Distillation for Large Language Models","venue":"cs.LG","work_id":"bae00e84-9b0d-433d-a066-20b951f0b4d0","year":2026},"citing_paper":{"arxiv_id":"2605.15239","last_updated":"2026-05-14T03:40:07Z","snapshot_observed_at":"2026-07-06T23:26:32.979566Z","submitted_at":"2026-05-14T03:40:07Z","title":"Reducing the Safety Tax in LLM Safety Alignment with On-Policy Self-Distillation","version":1},"reference_index":14,"source":"arxiv_source","source_observed_at":"2026-05-19T16:34:47.856606Z"},"links":{"cited_paper":"/paper/2601.18734","citing_paper":"/paper/2605.15239"},"observation_digest":"sha256:8d650b7abc926e9e104baa846fe6374b4d2dfefd15bc723cc915c64d6034c234","observation_id":"645a3552-6d92-4e53-ada0-4ce09a3283af","resolution":{"observed_at":"2026-05-19T16:37:39.896866Z","resolver_source":"local_arxiv","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-04T06:34:03.388597+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-04T06:34:03.388597+00:00","source":"crossref"},{"observed_at":"2026-08-04T06:33:57.428241+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Proceedings of the 62nd Annual Meeting of the Association for Computational Linguistics (Volume 1: Long Papers) , pages=","venue":null,"work_id":"8a4431e8-1f54-453e-8f38-d92f67695016","year":null},"citing_paper":{"arxiv_id":"2605.15239","last_updated":"2026-05-14T03:40:07Z","snapshot_observed_at":"2026-07-06T23:26:32.979566Z","submitted_at":"2026-05-14T03:40:07Z","title":"Reducing the Safety Tax in LLM Safety Alignment with On-Policy Self-Distillation","version":1},"reference_index":15,"source":"arxiv_source","source_observed_at":"2026-05-19T16:34:47.856606Z"},"links":{"citing_paper":"/paper/2605.15239"},"observation_digest":"sha256:3cc665c3507241fc27ab46d982ec6756b440010d7aba077d09ebcf5945272425","observation_id":"17c18425-8db0-4b32-bb2b-066e5455572e","resolution":{"observed_at":"2026-05-19T16:37:40.522211Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-04T06:34:03.388597+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-04T06:34:03.388597+00:00","source":"crossref"},{"observed_at":"2026-08-04T06:33:57.428241+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2402.03300","last_updated":"2024-04-27T15:25:53Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2024-02-05T18:55:32Z","title":"DeepSeekMath: Pushing the Limits of Mathematical Reasoning in Open Language Models","version":3},"cited_work":{"arxiv_id":"2402.03300","doi":"10.1016/0004-3702(73)90011-8","metadata_source":"pith","pith_arxiv_id":"2402.03300","snapshot_observed_at":"2026-07-11T11:50:26.030339Z","title":"DeepSeekMath: Pushing the Limits of Mathematical Reasoning in Open Language Models","venue":"cs.CL","work_id":"c5006563-f3ec-438a-9e35-b7b484f34828","year":2024},"citing_paper":{"arxiv_id":"2605.15239","last_updated":"2026-05-14T03:40:07Z","snapshot_observed_at":"2026-07-06T23:26:32.979566Z","submitted_at":"2026-05-14T03:40:07Z","title":"Reducing the Safety Tax in LLM Safety Alignment with On-Policy Self-Distillation","version":1},"reference_index":16,"source":"arxiv_source","source_observed_at":"2026-05-19T16:34:47.856606Z"},"links":{"cited_paper":"/paper/2402.03300","citing_paper":"/paper/2605.15239"},"observation_digest":"sha256:4d7ec295df67f207a9ad9d04435db6c069620be58c5f41cb30ff8d398489bbee","observation_id":"0f0c2af1-26f7-4db9-a58d-1711ff60a4a9","resolution":{"observed_at":"2026-05-19T16:37:39.888812Z","resolver_source":"local_arxiv","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-04T06:34:03.388597+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-04T06:34:03.388597+00:00","source":"crossref"},{"observed_at":"2026-08-04T06:33:57.428241+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2505.09388","last_updated":"2025-05-14T13:41:34Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2025-05-14T13:41:34Z","title":"Qwen3 Technical Report","version":1},"cited_work":{"arxiv_id":"2505.09388","doi":"10.1016/j.aiopen.2022.12","metadata_source":"pith","pith_arxiv_id":"2505.09388","snapshot_observed_at":"2026-07-11T11:50:26.030339Z","title":"Qwen3 Technical Report","venue":"cs.CL","work_id":"25a4e30c-1232-48e7-9925-02fa12ba7c9e","year":2025},"citing_paper":{"arxiv_id":"2605.15239","last_updated":"2026-05-14T03:40:07Z","snapshot_observed_at":"2026-07-06T23:26:32.979566Z","submitted_at":"2026-05-14T03:40:07Z","title":"Reducing the Safety Tax in LLM Safety Alignment with On-Policy Self-Distillation","version":1},"reference_index":17,"source":"arxiv_source","source_observed_at":"2026-05-19T16:34:47.856606Z"},"links":{"cited_paper":"/paper/2505.09388","citing_paper":"/paper/2605.15239"},"observation_digest":"sha256:737880ce5f6f9f9035726ba379c4ac31efcb8a8b69a503633744747fde49ca7d","observation_id":"1219f55d-52ca-4f48-8102-6a6bf30fa058","resolution":{"observed_at":"2026-05-19T16:37:39.894056Z","resolver_source":"local_arxiv","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-04T06:34:03.388597+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-04T06:34:03.388597+00:00","source":"crossref"},{"observed_at":"2026-08-04T06:33:57.428241+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2501.12948","last_updated":"2026-01-04T03:57:36Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2025-01-22T15:19:35Z","title":"DeepSeek-R1: Incentivizing Reasoning Capability in LLMs via Reinforcement Learning","version":2},"cited_work":{"arxiv_id":"2501.12948","doi":"10.1016/j.artmed.2024.103001","metadata_source":"pith","pith_arxiv_id":"2501.12948","snapshot_observed_at":"2026-07-11T11:50:26.030339Z","title":"DeepSeek-R1: Incentivizing Reasoning Capability in LLMs via Reinforcement Learning","venue":"cs.CL","work_id":"e6b75ad5-2877-4168-97c8-710407094d20","year":2025},"citing_paper":{"arxiv_id":"2605.15239","last_updated":"2026-05-14T03:40:07Z","snapshot_observed_at":"2026-07-06T23:26:32.979566Z","submitted_at":"2026-05-14T03:40:07Z","title":"Reducing the Safety Tax in LLM Safety Alignment with On-Policy Self-Distillation","version":1},"reference_index":18,"source":"arxiv_source","source_observed_at":"2026-05-19T16:34:47.856606Z"},"links":{"cited_paper":"/paper/2501.12948","citing_paper":"/paper/2605.15239"},"observation_digest":"sha256:b846adf916e10f23ed0a5869e9379ec0b7ba989f7d2028e91ab2106cb4a241f2","observation_id":"0e47f25c-cc7b-4631-a008-3c3da78768ff","resolution":{"observed_at":"2026-05-19T16:37:39.848530Z","resolver_source":"local_arxiv","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-04T06:34:03.388597+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-04T06:34:03.388597+00:00","source":"crossref"},{"observed_at":"2026-08-04T06:33:57.428241+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"1711.05101","last_updated":"2019-01-04T21:01:49Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2017-11-14T14:24:06Z","title":"Decoupled Weight Decay Regularization","version":3},"cited_work":{"arxiv_id":"1711.05101","doi":"10.1137/1.9781611972825.47","metadata_source":"pith","pith_arxiv_id":"1711.05101","snapshot_observed_at":"2026-07-11T11:50:26.030339Z","title":"Decoupled Weight Decay Regularization","venue":"cs.LG","work_id":"07ef7360-d385-4033-83f7-8384a6325204","year":2017},"citing_paper":{"arxiv_id":"2605.15239","last_updated":"2026-05-14T03:40:07Z","snapshot_observed_at":"2026-07-06T23:26:32.979566Z","submitted_at":"2026-05-14T03:40:07Z","title":"Reducing the Safety Tax in LLM Safety Alignment with On-Policy Self-Distillation","version":1},"reference_index":19,"source":"arxiv_source","source_observed_at":"2026-05-19T16:34:47.856606Z"},"links":{"cited_paper":"/paper/1711.05101","citing_paper":"/paper/2605.15239"},"observation_digest":"sha256:634f31354d74e50308d31d0987533e0a6416f66713f10d0311f1c8da9f481f75","observation_id":"bada1515-6f07-4203-8d60-a993c15b7806","resolution":{"observed_at":"2026-05-19T16:37:39.877288Z","resolver_source":"local_arxiv","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-04T06:34:03.388597+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-04T06:34:03.388597+00:00","source":"crossref"},{"observed_at":"2026-08-04T06:33:57.428241+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"2025 , note =","venue":null,"work_id":"575768cc-aad5-4674-8c43-db455dcffef5","year":2025},"citing_paper":{"arxiv_id":"2605.15239","last_updated":"2026-05-14T03:40:07Z","snapshot_observed_at":"2026-07-06T23:26:32.979566Z","submitted_at":"2026-05-14T03:40:07Z","title":"Reducing the Safety Tax in LLM Safety Alignment with On-Policy Self-Distillation","version":1},"reference_index":20,"source":"arxiv_source","source_observed_at":"2026-05-19T16:34:47.856606Z"},"links":{"citing_paper":"/paper/2605.15239"},"observation_digest":"sha256:4adb90c31c78bbc9caa4b347a88cdb8088b68b35323f6a565a56f7320fb06578","observation_id":"8c980e12-c1f3-4935-a25d-d519ec05b14e","resolution":{"observed_at":"2026-05-19T16:37:40.516440Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-04T06:34:03.388597+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-04T06:34:03.388597+00:00","source":"crossref"},{"observed_at":"2026-08-04T06:33:57.428241+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2312.06674","last_updated":"2023-12-07T19:40:50Z","snapshot_observed_at":"2026-07-06T17:00:00.321552Z","submitted_at":"2023-12-07T19:40:50Z","title":"Llama Guard: LLM-based Input-Output Safeguard for Human-AI Conversations","version":1},"cited_work":{"arxiv_id":"2312.06674","doi":"10.3390/info16050365","metadata_source":"pith","pith_arxiv_id":"2312.06674","snapshot_observed_at":"2026-07-11T11:50:26.030339Z","title":"Llama Guard: LLM-based Input-Output Safeguard for Human-AI Conversations","venue":"cs.CL","work_id":"93844332-869b-448c-a1be-35466150b1b2","year":2023},"citing_paper":{"arxiv_id":"2605.15239","last_updated":"2026-05-14T03:40:07Z","snapshot_observed_at":"2026-07-06T23:26:32.979566Z","submitted_at":"2026-05-14T03:40:07Z","title":"Reducing the Safety Tax in LLM Safety Alignment with On-Policy Self-Distillation","version":1},"reference_index":21,"source":"arxiv_source","source_observed_at":"2026-05-19T16:34:47.856606Z"},"links":{"cited_paper":"/paper/2312.06674","citing_paper":"/paper/2605.15239"},"observation_digest":"sha256:8ed1b5b1fe706beb02c22fa40987bb519638139f6199823cf281d005bc231b09","observation_id":"c6ac4c7b-8d16-4daa-9e07-761ad7d08e2d","resolution":{"observed_at":"2026-05-19T16:37:39.851629Z","resolver_source":"local_arxiv","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-04T06:34:03.388597+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-04T06:34:03.388597+00:00","source":"crossref"},{"observed_at":"2026-08-04T06:33:57.428241+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2402.04249","last_updated":"2024-02-27T04:43:08Z","snapshot_observed_at":"2026-07-06T17:26:23.067923Z","submitted_at":"2024-02-06T18:59:08Z","title":"HarmBench: A Standardized Evaluation Framework for Automated Red Teaming and Robust Refusal","version":2},"cited_work":{"arxiv_id":"2402.04249","doi":"10.48550/arxiv.2402.04249","metadata_source":"pith","pith_arxiv_id":"2402.04249","snapshot_observed_at":"2026-07-10T15:27:20.028353Z","title":"HarmBench: A Standardized Evaluation Framework for Automated Red Teaming and Robust Refusal","venue":"cs.LG","work_id":"b0b0303f-2444-4789-a979-8153624312ff","year":2024},"citing_paper":{"arxiv_id":"2605.15239","last_updated":"2026-05-14T03:40:07Z","snapshot_observed_at":"2026-07-06T23:26:32.979566Z","submitted_at":"2026-05-14T03:40:07Z","title":"Reducing the Safety Tax in LLM Safety Alignment with On-Policy Self-Distillation","version":1},"reference_index":22,"source":"arxiv_source","source_observed_at":"2026-05-19T16:34:47.856606Z"},"links":{"cited_paper":"/paper/2402.04249","citing_paper":"/paper/2605.15239"},"observation_digest":"sha256:dc295ee397fb698531d49e31276a3dfe93fa78bfb62ad972bccf2f1adbb9408b","observation_id":"54563ba4-806e-4092-bebe-f25238cc18fe","resolution":{"observed_at":"2026-05-19T16:37:39.866341Z","resolver_source":"local_arxiv","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-04T06:34:03.388597+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-04T06:34:03.388597+00:00","source":"crossref"},{"observed_at":"2026-08-04T06:33:57.428241+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Advances in Neural Information Processing Systems , volume=","venue":null,"work_id":"87a2d2aa-c798-4c39-a51c-54a3805e5884","year":null},"citing_paper":{"arxiv_id":"2605.15239","last_updated":"2026-05-14T03:40:07Z","snapshot_observed_at":"2026-07-06T23:26:32.979566Z","submitted_at":"2026-05-14T03:40:07Z","title":"Reducing the Safety Tax in LLM Safety Alignment with On-Policy Self-Distillation","version":1},"reference_index":23,"source":"arxiv_source","source_observed_at":"2026-05-19T16:34:47.856606Z"},"links":{"citing_paper":"/paper/2605.15239"},"observation_digest":"sha256:337360e997e9a0495d4349bf5c39339d691cf56459f14ce3f3a26552f04c900f","observation_id":"450f1e63-e795-4757-aea0-14fe2f96541c","resolution":{"observed_at":"2026-05-19T16:37:40.498831Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-04T06:34:03.388597+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-04T06:34:03.388597+00:00","source":"crossref"},{"observed_at":"2026-08-04T06:33:57.428241+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Advances in Neural Information Processing Systems , volume=","venue":null,"work_id":"c16b2972-c1c5-47e7-85dc-7dc8a54f32d2","year":null},"citing_paper":{"arxiv_id":"2605.15239","last_updated":"2026-05-14T03:40:07Z","snapshot_observed_at":"2026-07-06T23:26:32.979566Z","submitted_at":"2026-05-14T03:40:07Z","title":"Reducing the Safety Tax in LLM Safety Alignment with On-Policy Self-Distillation","version":1},"reference_index":24,"source":"arxiv_source","source_observed_at":"2026-05-19T16:34:47.856606Z"},"links":{"citing_paper":"/paper/2605.15239"},"observation_digest":"sha256:d1afada17015bd9d09b3034540af4e7bfa8ea22d0973e6fa2ba389efecd23421","observation_id":"b489457d-4b6b-4fd7-a74a-fd7424122851","resolution":{"observed_at":"2026-05-19T16:37:40.500770Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-04T06:34:03.388597+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-04T06:34:03.388597+00:00","source":"crossref"},{"observed_at":"2026-08-04T06:33:57.428241+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-07-05T06:50:48.645747Z","title":"Proceedings of the 2024 Conference of the North American Chapter of the Association for Computational Linguistics: Human Language Technologies (Volume 1: Long Papers) , pages=","venue":null,"work_id":"3c05557c-2471-4d17-b026-259d07bff219","year":2024},"citing_paper":{"arxiv_id":"2605.15239","last_updated":"2026-05-14T03:40:07Z","snapshot_observed_at":"2026-07-06T23:26:32.979566Z","submitted_at":"2026-05-14T03:40:07Z","title":"Reducing the Safety Tax in LLM Safety Alignment with On-Policy Self-Distillation","version":1},"reference_index":25,"source":"arxiv_source","source_observed_at":"2026-05-19T16:34:47.856606Z"},"links":{"citing_paper":"/paper/2605.15239"},"observation_digest":"sha256:7471c6bdc6619e1e33282b58a8d37a2d7fea10a1c0caee657718f3ee5ac5f15e","observation_id":"489b2817-5b56-4c72-a72c-117c41ce3542","resolution":{"observed_at":"2026-05-19T16:37:40.503054Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-04T06:34:03.388597+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-04T06:34:03.388597+00:00","source":"crossref"},{"observed_at":"2026-08-04T06:33:57.428241+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-07-05T23:21:35.999640Z","title":"Advances in neural information processing systems , volume=","venue":null,"work_id":"15496c9a-1638-467a-86bd-c04d39eafbfd","year":null},"citing_paper":{"arxiv_id":"2605.15239","last_updated":"2026-05-14T03:40:07Z","snapshot_observed_at":"2026-07-06T23:26:32.979566Z","submitted_at":"2026-05-14T03:40:07Z","title":"Reducing the Safety Tax in LLM Safety Alignment with On-Policy Self-Distillation","version":1},"reference_index":26,"source":"arxiv_source","source_observed_at":"2026-05-19T16:34:47.856606Z"},"links":{"citing_paper":"/paper/2605.15239"},"observation_digest":"sha256:2a58c16a32e6ccd612549da240569232c5073ccf98201b9deaede5ab74bc2e22","observation_id":"828a5639-82de-4708-bd71-951b7b1521c0","resolution":{"observed_at":"2026-05-19T16:37:40.504884Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-04T06:34:03.388597+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-04T06:34:03.388597+00:00","source":"crossref"},{"observed_at":"2026-08-04T06:33:57.428241+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2110.14168","last_updated":"2021-11-18T00:23:45Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2021-10-27T04:49:45Z","title":"Training Verifiers to Solve Math Word Problems","version":2},"cited_work":{"arxiv_id":"2110.14168","doi":"10.1002/j.1545-","metadata_source":"pith","pith_arxiv_id":"2110.14168","snapshot_observed_at":"2026-07-11T11:50:26.030339Z","title":"Training Verifiers to Solve Math Word Problems","venue":"cs.LG","work_id":"acab1aa8-b4d6-40e0-a3ee-25341701dca2","year":2021},"citing_paper":{"arxiv_id":"2605.15239","last_updated":"2026-05-14T03:40:07Z","snapshot_observed_at":"2026-07-06T23:26:32.979566Z","submitted_at":"2026-05-14T03:40:07Z","title":"Reducing the Safety Tax in LLM Safety Alignment with On-Policy Self-Distillation","version":1},"reference_index":27,"source":"arxiv_source","source_observed_at":"2026-05-19T16:34:47.856606Z"},"links":{"cited_paper":"/paper/2110.14168","citing_paper":"/paper/2605.15239"},"observation_digest":"sha256:43d115d78eb49b4546c7d3170ddd370f867e9e8b6f0b2370e4120c1f3cc68b8c","observation_id":"cf113726-6fe7-40f6-a9b8-c829c7694ea2","resolution":{"observed_at":"2026-05-19T16:37:39.863129Z","resolver_source":"local_arxiv","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-04T06:34:03.388597+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-04T06:34:03.388597+00:00","source":"crossref"},{"observed_at":"2026-08-04T06:33:57.428241+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"The twelfth international conference on learning representations , year=","venue":null,"work_id":"9706546a-6b95-4fb4-9188-69bf7c2228f9","year":null},"citing_paper":{"arxiv_id":"2605.15239","last_updated":"2026-05-14T03:40:07Z","snapshot_observed_at":"2026-07-06T23:26:32.979566Z","submitted_at":"2026-05-14T03:40:07Z","title":"Reducing the Safety Tax in LLM Safety Alignment with On-Policy Self-Distillation","version":1},"reference_index":28,"source":"arxiv_source","source_observed_at":"2026-05-19T16:34:47.856606Z"},"links":{"citing_paper":"/paper/2605.15239"},"observation_digest":"sha256:9ad4321aa61e752513469683ea59c5ee03ac53f5cb77a38bd7f8a92acacdba10","observation_id":"17eaba70-8051-40eb-b008-0df4b2303c26","resolution":{"observed_at":"2026-05-19T16:37:40.506689Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-04T06:34:03.388597+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-04T06:34:03.388597+00:00","source":"crossref"},{"observed_at":"2026-08-04T06:33:57.428241+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2311.12022","last_updated":"2023-11-20T18:57:34Z","snapshot_observed_at":"2026-08-02T23:46:38.854124Z","submitted_at":"2023-11-20T18:57:34Z","title":"GPQA: A Graduate-Level Google-Proof Q&A Benchmark","version":1},"cited_work":{"arxiv_id":"2311.12022","doi":"10.48550/arxiv.2311.12022","metadata_source":"pith","pith_arxiv_id":"2311.12022","snapshot_observed_at":"2026-07-10T14:47:14.590391Z","title":"GPQA: A Graduate-Level Google-Proof Q&A Benchmark","venue":"cs.AI","work_id":"9e2a976b-f5ad-4aee-af5c-243fe0fe75d2","year":2023},"citing_paper":{"arxiv_id":"2605.15239","last_updated":"2026-05-14T03:40:07Z","snapshot_observed_at":"2026-07-06T23:26:32.979566Z","submitted_at":"2026-05-14T03:40:07Z","title":"Reducing the Safety Tax in LLM Safety Alignment with On-Policy Self-Distillation","version":1},"reference_index":29,"source":"arxiv_source","source_observed_at":"2026-05-19T16:34:47.856606Z"},"links":{"cited_paper":"/paper/2311.12022","citing_paper":"/paper/2605.15239"},"observation_digest":"sha256:18683aa8f3d18853d36b744e810dd44a266e9c1eb7b633e8ab06447f14d176e0","observation_id":"d13beaa8-fe31-4f86-b159-8da330e3692d","resolution":{"observed_at":"2026-05-19T16:37:39.872110Z","resolver_source":"local_arxiv","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-04T06:34:03.388597+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-04T06:34:03.388597+00:00","source":"crossref"},{"observed_at":"2026-08-01T04:38:46.941438+00:00","source":"crossref_status_cache"},{"observed_at":"2026-08-01T04:38:46.941438+00:00","source":"openalex_status_cache"},{"observed_at":"2026-08-04T06:33:57.428241+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2107.03374","last_updated":"2021-07-14T17:16:02Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2021-07-07T17:41:24Z","title":"Evaluating Large Language Models Trained on Code","version":2},"cited_work":{"arxiv_id":"2107.03374","doi":"10.48550/arxiv.2107.03374","metadata_source":"pith","pith_arxiv_id":"2107.03374","snapshot_observed_at":"2026-07-11T03:27:46.202228Z","title":"Evaluating Large Language Models Trained on Code","venue":"cs.LG","work_id":"042493e9-b26f-4b4e-bbde-382072ca9b08","year":2021},"citing_paper":{"arxiv_id":"2605.15239","last_updated":"2026-05-14T03:40:07Z","snapshot_observed_at":"2026-07-06T23:26:32.979566Z","submitted_at":"2026-05-14T03:40:07Z","title":"Reducing the Safety Tax in LLM Safety Alignment with On-Policy Self-Distillation","version":1},"reference_index":30,"source":"arxiv_source","source_observed_at":"2026-05-19T16:34:47.856606Z"},"links":{"cited_paper":"/paper/2107.03374","citing_paper":"/paper/2605.15239"},"observation_digest":"sha256:87ef2e23832c80f009f3496e870a027598766eb277713bf79176e1fd4f153992","observation_id":"12e4ec2f-4a44-4581-b8cf-6da983dfe552","resolution":{"observed_at":"2026-05-19T16:37:39.845186Z","resolver_source":"local_arxiv","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-04T06:34:03.388597+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-04T06:34:03.388597+00:00","source":"crossref"},{"observed_at":"2026-08-01T08:08:23.404839+00:00","source":"crossref_status_cache"},{"observed_at":"2026-08-01T08:08:23.404839+00:00","source":"openalex_status_cache"},{"observed_at":"2026-08-04T06:33:57.428241+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2108.07732","last_updated":"2021-08-16T03:57:30Z","snapshot_observed_at":"2026-08-02T19:23:53.535075Z","submitted_at":"2021-08-16T03:57:30Z","title":"Program Synthesis with Large Language Models","version":1},"cited_work":{"arxiv_id":"2108.07732","doi":"10.1007/s11390-025-5518-5","metadata_source":"pith","pith_arxiv_id":"2108.07732","snapshot_observed_at":"2026-07-11T11:50:26.030339Z","title":"Program Synthesis with Large Language Models","venue":"cs.PL","work_id":"fd241a05-03b9-4de2-9588-9d77ce176125","year":2021},"citing_paper":{"arxiv_id":"2605.15239","last_updated":"2026-05-14T03:40:07Z","snapshot_observed_at":"2026-07-06T23:26:32.979566Z","submitted_at":"2026-05-14T03:40:07Z","title":"Reducing the Safety Tax in LLM Safety Alignment with On-Policy Self-Distillation","version":1},"reference_index":31,"source":"arxiv_source","source_observed_at":"2026-05-19T16:34:47.856606Z"},"links":{"cited_paper":"/paper/2108.07732","citing_paper":"/paper/2605.15239"},"observation_digest":"sha256:c0ab41b5de40cfbd4f026384454b96d806d70cd9a8ffcfa3cdcf3f5d1dbc0d9a","observation_id":"3654fbbd-a491-4118-bb0b-58bdd1d616fd","resolution":{"observed_at":"2026-05-19T16:37:39.854356Z","resolver_source":"local_arxiv","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-04T06:34:03.388597+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-04T06:34:03.388597+00:00","source":"crossref"},{"observed_at":"2026-08-04T06:33:57.428241+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2312.12321","last_updated":"2024-05-17T05:27:25Z","snapshot_observed_at":"2026-07-06T17:05:22.992895Z","submitted_at":"2023-12-19T16:47:12Z","title":"Bypassing the Safety Training of Open-Source LLMs with Priming Attacks","version":2},"cited_work":{"arxiv_id":"2312.12321","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2312.12321","snapshot_observed_at":"2026-07-01T01:15:13.251025Z","title":"arXiv preprint arXiv:2312.12321 , year=","venue":null,"work_id":"14f5dd62-0d53-4a37-ad39-b6ae346f04c9","year":null},"citing_paper":{"arxiv_id":"2605.15239","last_updated":"2026-05-14T03:40:07Z","snapshot_observed_at":"2026-07-06T23:26:32.979566Z","submitted_at":"2026-05-14T03:40:07Z","title":"Reducing the Safety Tax in LLM Safety Alignment with On-Policy Self-Distillation","version":1},"reference_index":32,"source":"arxiv_source","source_observed_at":"2026-05-19T16:34:47.856606Z"},"links":{"cited_paper":"/paper/2312.12321","citing_paper":"/paper/2605.15239"},"observation_digest":"sha256:93295c81a3427e33978120e7cb482890d4e258d7596eda8badd9752bae8188d7","observation_id":"935d416b-ec31-4a3f-8205-bcbc4af244cc","resolution":{"observed_at":"2026-05-19T16:37:39.857254Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-04T06:34:03.388597+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-04T06:34:03.388597+00:00","source":"crossref"},{"observed_at":"2026-08-04T06:33:57.428241+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"do anything now","venue":null,"work_id":"7203142b-c1e9-4c5f-8165-2d5cbe646439","year":2024},"citing_paper":{"arxiv_id":"2605.15239","last_updated":"2026-05-14T03:40:07Z","snapshot_observed_at":"2026-07-06T23:26:32.979566Z","submitted_at":"2026-05-14T03:40:07Z","title":"Reducing the Safety Tax in LLM Safety Alignment with On-Policy Self-Distillation","version":1},"reference_index":33,"source":"arxiv_source","source_observed_at":"2026-05-19T16:34:47.856606Z"},"links":{"citing_paper":"/paper/2605.15239"},"observation_digest":"sha256:fde562d530e29b66c704e916c0ef11f1a0652cd8b52ad18f64dc8d7d98ccaf48","observation_id":"a3a0dd48-b639-496f-ab77-2e8c2aee7266","resolution":{"observed_at":"2026-05-19T16:37:40.492981Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-04T06:34:03.388597+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-04T06:34:03.388597+00:00","source":"crossref"},{"observed_at":"2026-08-04T06:33:57.428241+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-07-06T16:42:40.001878Z","title":"Proceedings of the 62nd Annual Meeting of the Association for Computational Linguistics (Volume 1: Long Papers) , pages=","venue":null,"work_id":"5a7e6354-0582-413c-a0b6-ed5573b425de","year":null},"citing_paper":{"arxiv_id":"2605.15239","last_updated":"2026-05-14T03:40:07Z","snapshot_observed_at":"2026-07-06T23:26:32.979566Z","submitted_at":"2026-05-14T03:40:07Z","title":"Reducing the Safety Tax in LLM Safety Alignment with On-Policy Self-Distillation","version":1},"reference_index":34,"source":"arxiv_source","source_observed_at":"2026-05-19T16:34:47.856606Z"},"links":{"citing_paper":"/paper/2605.15239"},"observation_digest":"sha256:fcd5067ef3c0591539edfe78bd3fd7d5e07a7760a7c3055f902e41334e7fc832","observation_id":"1ee8f1ee-c151-40b2-a2e9-2eea915cbacc","resolution":{"observed_at":"2026-05-19T16:37:40.496642Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-04T06:34:03.388597+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-04T06:34:03.388597+00:00","source":"crossref"},{"observed_at":"2026-08-04T06:33:57.428241+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"2025 IEEE Conference on Secure and Trustworthy Machine Learning (SaTML) , pages=","venue":null,"work_id":"ba7524ec-ad30-45f9-90cd-2a9dd6b78a80","year":2025},"citing_paper":{"arxiv_id":"2605.15239","last_updated":"2026-05-14T03:40:07Z","snapshot_observed_at":"2026-07-06T23:26:32.979566Z","submitted_at":"2026-05-14T03:40:07Z","title":"Reducing the Safety Tax in LLM Safety Alignment with On-Policy Self-Distillation","version":1},"reference_index":35,"source":"arxiv_source","source_observed_at":"2026-05-19T16:34:47.856606Z"},"links":{"citing_paper":"/paper/2605.15239"},"observation_digest":"sha256:a03a0dd7f4c8f5e243d4c29a0094e45d1c7b35e721f38dad85d2ff7afc211433","observation_id":"0f7ab785-80f3-43d2-9c24-5d2b8724f155","resolution":{"observed_at":"2026-05-19T16:37:40.508591Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-04T06:34:03.388597+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-04T06:34:03.388597+00:00","source":"crossref"},{"observed_at":"2026-08-04T06:33:57.428241+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"2026 , eprint=","venue":null,"work_id":"e06f395e-6cc0-4a4b-926e-46210815a4d4","year":2026},"citing_paper":{"arxiv_id":"2605.15239","last_updated":"2026-05-14T03:40:07Z","snapshot_observed_at":"2026-07-06T23:26:32.979566Z","submitted_at":"2026-05-14T03:40:07Z","title":"Reducing the Safety Tax in LLM Safety Alignment with On-Policy Self-Distillation","version":1},"reference_index":36,"source":"arxiv_source","source_observed_at":"2026-05-19T16:34:47.856606Z"},"links":{"citing_paper":"/paper/2605.15239"},"observation_digest":"sha256:21611f97b23e06bd8202a3e2a25f12838c1fb363e5fbc1a82c12e53c9e695405","observation_id":"14ad3a0a-248b-40de-82a0-5051fcd8c429","resolution":{"observed_at":"2026-05-19T16:37:40.489565Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-04T06:34:03.388597+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-04T06:34:03.388597+00:00","source":"crossref"},{"observed_at":"2026-08-04T06:33:57.428241+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"2026 , eprint=","venue":null,"work_id":"0517ff45-041f-428e-9348-43465b609ccc","year":2026},"citing_paper":{"arxiv_id":"2605.15239","last_updated":"2026-05-14T03:40:07Z","snapshot_observed_at":"2026-07-06T23:26:32.979566Z","submitted_at":"2026-05-14T03:40:07Z","title":"Reducing the Safety Tax in LLM Safety Alignment with On-Policy Self-Distillation","version":1},"reference_index":37,"source":"arxiv_source","source_observed_at":"2026-05-19T16:34:47.856606Z"},"links":{"citing_paper":"/paper/2605.15239"},"observation_digest":"sha256:9a2bc2f493a0edd1e8999e9132f9196bca47bf045327d381bb7b9c5f8c6a58f0","observation_id":"f61e7f5e-796d-4f68-92f2-4a33b0e074ba","resolution":{"observed_at":"2026-05-19T16:37:40.487888Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-04T06:34:03.388597+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-04T06:34:03.388597+00:00","source":"crossref"},{"observed_at":"2026-08-04T06:33:57.428241+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"The twelfth international conference on learning representations , year=","venue":null,"work_id":"f24cb3a7-1269-49d3-ba3c-d42e0d9a2f99","year":null},"citing_paper":{"arxiv_id":"2605.15239","last_updated":"2026-05-14T03:40:07Z","snapshot_observed_at":"2026-07-06T23:26:32.979566Z","submitted_at":"2026-05-14T03:40:07Z","title":"Reducing the Safety Tax in LLM Safety Alignment with On-Policy Self-Distillation","version":1},"reference_index":38,"source":"arxiv_source","source_observed_at":"2026-05-19T16:34:47.856606Z"},"links":{"citing_paper":"/paper/2605.15239"},"observation_digest":"sha256:417aa1c8eff147271b10a20a74d8c376e109ca0d64c9a10f3df4260ad6886fed","observation_id":"505ac6d2-801b-4454-909e-4d28d5cbe0b1","resolution":{"observed_at":"2026-05-19T16:37:40.491320Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-04T06:34:03.388597+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-04T06:34:03.388597+00:00","source":"crossref"},{"observed_at":"2026-08-04T06:33:57.428241+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-07-09T13:26:17.298990Z","title":"The twelfth international conference on learning representations , year=","venue":null,"work_id":"41a49c2e-ddd8-4600-832d-1c9739f8163f","year":null},"citing_paper":{"arxiv_id":"2605.15239","last_updated":"2026-05-14T03:40:07Z","snapshot_observed_at":"2026-07-06T23:26:32.979566Z","submitted_at":"2026-05-14T03:40:07Z","title":"Reducing the Safety Tax in LLM Safety Alignment with On-Policy Self-Distillation","version":1},"reference_index":39,"source":"arxiv_source","source_observed_at":"2026-05-19T16:34:47.856606Z"},"links":{"citing_paper":"/paper/2605.15239"},"observation_digest":"sha256:1188b417b4c9f93c622670ab63820cdd9a41a350c857d4290aa7e97486078135","observation_id":"8b9c96eb-4e73-447f-8edb-bb1d9c7602b1","resolution":{"observed_at":"2026-05-19T16:37:40.510899Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-04T06:34:03.388597+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-04T06:34:03.388597+00:00","source":"crossref"},{"observed_at":"2026-08-04T06:33:57.428241+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2602.12275","last_updated":"2026-03-23T13:11:10Z","snapshot_observed_at":"2026-07-06T22:45:39.557598Z","submitted_at":"2026-02-12T18:58:28Z","title":"On-Policy Context Distillation for Language Models","version":2},"cited_work":{"arxiv_id":"2602.12275","doi":null,"metadata_source":"pith","pith_arxiv_id":"2602.12275","snapshot_observed_at":"2026-07-10T00:26:39.252472Z","title":"On-Policy Context Distillation for Language Models","venue":"cs.CL","work_id":"b56a7e15-d864-43f4-9212-59bc7ec70d21","year":2026},"citing_paper":{"arxiv_id":"2605.15239","last_updated":"2026-05-14T03:40:07Z","snapshot_observed_at":"2026-07-06T23:26:32.979566Z","submitted_at":"2026-05-14T03:40:07Z","title":"Reducing the Safety Tax in LLM Safety Alignment with On-Policy Self-Distillation","version":1},"reference_index":40,"source":"arxiv_source","source_observed_at":"2026-05-19T16:34:47.856606Z"},"links":{"cited_paper":"/paper/2602.12275","citing_paper":"/paper/2605.15239"},"observation_digest":"sha256:a0654eb3f088dd1215883521a0c983d9a649b67842a140e432d25641aa9f7058","observation_id":"856fc587-1272-452a-a13d-cc869e55fd9d","resolution":{"observed_at":"2026-05-19T16:37:39.874777Z","resolver_source":"local_arxiv","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-04T06:34:03.388597+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-04T06:34:03.388597+00:00","source":"crossref"},{"observed_at":"2026-08-04T06:33:57.428241+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2603.05433","last_updated":"2026-07-03T04:38:21Z","snapshot_observed_at":"2026-07-15T14:32:24.776402Z","submitted_at":"2026-03-05T17:54:40Z","title":"CRISP: Compressed Reasoning via Iterative Self-Policy Distillation","version":7},"cited_work":{"arxiv_id":"2603.05433","doi":null,"metadata_source":"pith","pith_arxiv_id":"2603.05433","snapshot_observed_at":"2026-07-03T20:18:56.074332Z","title":"CRISP: Compressed Reasoning via Iterative Self-Policy Distillation","venue":"cs.LG","work_id":"c5a99022-15b6-4d77-9850-23036df7a073","year":2026},"citing_paper":{"arxiv_id":"2605.15239","last_updated":"2026-05-14T03:40:07Z","snapshot_observed_at":"2026-07-06T23:26:32.979566Z","submitted_at":"2026-05-14T03:40:07Z","title":"Reducing the Safety Tax in LLM Safety Alignment with On-Policy Self-Distillation","version":1},"reference_index":41,"source":"arxiv_source","source_observed_at":"2026-05-19T16:34:47.856606Z"},"links":{"cited_paper":"/paper/2603.05433","citing_paper":"/paper/2605.15239"},"observation_digest":"sha256:73badd560c19b4fcee5f46b0a6a9b7adaa2983ad25a74d1be67f895908e4b42d","observation_id":"2a1119cb-e280-4965-b015-66c67fbee098","resolution":{"observed_at":"2026-05-19T16:37:39.869050Z","resolver_source":"local_arxiv","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-04T06:34:03.388597+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-04T06:34:03.388597+00:00","source":"crossref"},{"observed_at":"2026-08-04T06:33:57.428241+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2307.15043","last_updated":"2023-12-20T20:48:57Z","snapshot_observed_at":"2026-07-06T15:59:23.019044Z","submitted_at":"2023-07-27T17:49:12Z","title":"Universal and Transferable Adversarial Attacks on Aligned Language Models","version":2},"cited_work":{"arxiv_id":"2307.15043","doi":"10.48550/arxiv.2307.15043","metadata_source":"pith","pith_arxiv_id":"2307.15043","snapshot_observed_at":"2026-07-11T02:47:49.837877Z","title":"Universal and Transferable Adversarial Attacks on Aligned Language Models","venue":"cs.CL","work_id":"3322fa86-1768-4677-8425-dd326b45e078","year":2023},"citing_paper":{"arxiv_id":"2605.15239","last_updated":"2026-05-14T03:40:07Z","snapshot_observed_at":"2026-07-06T23:26:32.979566Z","submitted_at":"2026-05-14T03:40:07Z","title":"Reducing the Safety Tax in LLM Safety Alignment with On-Policy Self-Distillation","version":1},"reference_index":42,"source":"arxiv_source","source_observed_at":"2026-05-19T16:34:47.856606Z"},"links":{"cited_paper":"/paper/2307.15043","citing_paper":"/paper/2605.15239"},"observation_digest":"sha256:15cee7aa6262b2e4b36dba2a2747341bdc93ba4ca973e0f0382eb57c2ba824a6","observation_id":"eeafd913-7895-4fe5-9c18-a1ca37f899e4","resolution":{"observed_at":"2026-05-19T16:37:39.860061Z","resolver_source":"local_arxiv","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-04T06:34:03.388597+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-04T06:34:03.388597+00:00","source":"crossref"},{"observed_at":"2026-07-15T23:50:40.271168+00:00","source":"crossref_status_cache"},{"observed_at":"2026-07-15T23:50:40.271168+00:00","source":"openalex_status_cache"},{"observed_at":"2026-08-04T06:33:57.428241+00:00","source":"retraction_watch"}],"state":"measured"}}],"paper":{"arxiv_id":"2605.15239","last_updated":"2026-05-14T03:40:07Z","latest_version":1,"primary_category":"cs.LG","snapshot_observed_at":"2026-07-06T23:26:32.979566Z","submitted_at":"2026-05-14T03:40:07Z","title":"Reducing the Safety Tax in LLM Safety Alignment with On-Policy Self-Distillation"},"reference_resolution":{"displayed":42,"state_counts":{"malformed_identifier":0,"metadata_mismatch":22,"parse_uncertain":0,"unresolved":0,"verified_exact":2,"verified_fuzzy":18},"total_outbound_references":42},"refusal":"A citation records a reference. It does not transfer a finding from one paper to another.","schema":"pith.paper-citation-record.v1","standing_sources":[{"observed_at":"2026-08-04T06:34:03.388597+00:00","source":"crossref"},{"observed_at":"2026-08-04T06:33:57.428241+00:00","source":"retraction_watch"}],"thesis":"As of 4 August 2026, this Paper Citation Record lists 42 of 42 outbound references and 0 inbound Pith citation observations for arXiv:2605.15239."}