{"as_of":"2026-08-08T08:03:00Z","caps":{"database_statements":6,"inbound":100,"outbound":100},"context_digest":"sha256:df8b12fdf17b6defcd83199a837129312a2a0c176b188f8f145a8641795fedb6","coverage":[{"denominator":29,"lane":"reference_resolution","note":"Typed states for the displayed outbound observations.","records_observed":29,"source":"paper_references, paper_reference_links","source_observed_at":"2026-08-06T10:18:26.686397Z","state":"measured"},{"denominator":30,"lane":"standing_notices","note":"One-hop event checks from named stored sources.","records_observed":30,"source":"scholarly_work_events, retraction_status_cache","source_observed_at":"2026-08-08T06:32:00.761636+00:00","state":"measured"},{"denominator":1,"lane":"inbound_itemization","note":"Pith citing papers itemized under the disclosed page cap.","records_observed":1,"source":"paper_references, paper_reference_links","source_observed_at":"2026-06-25T23:37:49.412578Z","state":"measured"},{"denominator":1,"lane":"external_citation_measurements","note":"A source-named dated measurement, never combined with another source.","records_observed":0,"source":"arxiv_reference","source_observed_at":"2026-07-04T17:30:00.622652Z","state":"measured"}],"external_citation_measurements":[],"inbound":[{"citation":{"cited_paper":{"arxiv_id":"2508.00324","last_updated":"2025-08-01T05:14:13Z","snapshot_observed_at":"2026-08-07T21:57:02.982200Z","submitted_at":"2025-08-01T05:14:13Z","title":"R1-ACT: Efficient Reasoning Model Safety Alignment by Activating Safety Knowledge","version":1},"cited_work":{"arxiv_id":"2508.00324","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2508.00324","snapshot_observed_at":"2026-07-04T17:30:00.622652Z","title":"2508.00324 , archivePrefix =","venue":null,"work_id":"adef8db7-d37e-4d27-ad26-bba253437d59","year":null},"citing_paper":{"arxiv_id":"2606.25013","last_updated":"2026-06-23T17:59:01Z","snapshot_observed_at":"2026-08-03T16:39:59.297765Z","submitted_at":"2026-06-23T17:59:01Z","title":"Do Thinking Tokens Help with Safety?","version":1},"reference_index":47,"source":"arxiv_source","source_observed_at":"2026-06-25T23:37:49.412578Z"},"links":{"cited_paper":"/paper/2508.00324","citing_paper":"/paper/2606.25013"},"observation_digest":"sha256:1870ab5db6343cd843c0f8fa7d82141967acdefa744d757660eaf2d9980afae6","observation_id":"5487ee9c-be80-44d9-85ea-e29142fd2aab","resolution":{"observed_at":"2026-07-04T17:30:00.624437Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}}],"links":{"evidence":"/evidence","html":"/paper/2508.00324/citation-record","integrity":"/paper/2508.00324/integrity","json":"/paper/2508.00324/citation-record.json","paper":"/paper/2508.00324"},"outbound":[{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T10:18:27.883886Z","title":null,"venue":null,"work_id":"c0c978dd-7b39-4282-a54f-3a6b66885a0e","year":2024},"citing_paper":{"arxiv_id":"2508.00324","last_updated":"2025-08-01T05:14:13Z","snapshot_observed_at":"2026-08-07T21:57:02.982200Z","submitted_at":"2025-08-01T05:14:13Z","title":"R1-ACT: Efficient Reasoning Model Safety Alignment by Activating Safety Knowledge","version":1},"reference_index":1,"source":"arxiv_source","source_observed_at":"2026-08-06T10:18:24.201071Z"},"links":{"citing_paper":"/paper/2508.00324"},"observation_digest":"sha256:eefac44b5685fa7873028b693fa124f7ad730f4dbf85f58b820def233796a837","observation_id":"c2690ba0-7111-4ebd-b3f0-b6c72531ff4d","resolution":{"observed_at":"2026-08-06T10:18:27.996788Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2107.03374","last_updated":"2021-07-14T17:16:02Z","snapshot_observed_at":"2026-08-08T04:13:14.044604Z","submitted_at":"2021-07-07T17:41:24Z","title":"Evaluating Large Language Models Trained on Code","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2107.03374","snapshot_observed_at":"2026-08-06T10:18:24.296351Z","title":null,"venue":null,"work_id":null,"year":2021},"citing_paper":{"arxiv_id":"2508.00324","last_updated":"2025-08-01T05:14:13Z","snapshot_observed_at":"2026-08-07T21:57:02.982200Z","submitted_at":"2025-08-01T05:14:13Z","title":"R1-ACT: Efficient Reasoning Model Safety Alignment by Activating Safety Knowledge","version":1},"reference_index":2,"source":"arxiv_source","source_observed_at":"2026-08-06T10:18:24.296351Z"},"links":{"cited_paper":"/paper/2107.03374","citing_paper":"/paper/2508.00324"},"observation_digest":"sha256:0b74d51065d2876d5646d44c86d1f2a394a1eb177bd90edb039fa0d9f3a3cede","observation_id":"25df598f-45d3-4427-bdfd-463adfc71eab","resolution":{"observed_at":"2026-08-06T10:18:24.296351Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2110.14168","last_updated":"2021-11-18T00:23:45Z","snapshot_observed_at":"2026-08-07T01:45:38.840969Z","submitted_at":"2021-10-27T04:49:45Z","title":"Training Verifiers to Solve Math Word Problems","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2110.14168","snapshot_observed_at":"2026-08-06T10:18:24.367710Z","title":null,"venue":null,"work_id":null,"year":2021},"citing_paper":{"arxiv_id":"2508.00324","last_updated":"2025-08-01T05:14:13Z","snapshot_observed_at":"2026-08-07T21:57:02.982200Z","submitted_at":"2025-08-01T05:14:13Z","title":"R1-ACT: Efficient Reasoning Model Safety Alignment by Activating Safety Knowledge","version":1},"reference_index":3,"source":"arxiv_source","source_observed_at":"2026-08-06T10:18:24.367710Z"},"links":{"cited_paper":"/paper/2110.14168","citing_paper":"/paper/2508.00324"},"observation_digest":"sha256:a47dcbfad68b7e30ab1a439d977ac5dc894fb6a70b62c93027e20cd60a82d2b5","observation_id":"99928bbd-a201-4d73-8f94-a9e0ae89fa79","resolution":{"observed_at":"2026-08-06T10:18:24.367710Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T10:18:24.405850Z","title":null,"venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2508.00324","last_updated":"2025-08-01T05:14:13Z","snapshot_observed_at":"2026-08-07T21:57:02.982200Z","submitted_at":"2025-08-01T05:14:13Z","title":"R1-ACT: Efficient Reasoning Model Safety Alignment by Activating Safety Knowledge","version":1},"reference_index":4,"source":"arxiv_source","source_observed_at":"2026-08-06T10:18:24.405850Z"},"links":{"citing_paper":"/paper/2508.00324"},"observation_digest":"sha256:0b237241e42709621658546bc7374ef4dc11c1553fcfaaab638eb8ecfc8bec1b","observation_id":"63ce70aa-11b6-4875-8f83-7d9c8e29362a","resolution":{"observed_at":"2026-08-06T10:18:24.405850Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2412.16339","last_updated":"2025-01-08T20:11:59Z","snapshot_observed_at":"2026-08-07T23:13:03.652587Z","submitted_at":"2024-12-20T21:00:11Z","title":"Deliberative Alignment: Reasoning Enables Safer Language Models","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2412.16339","snapshot_observed_at":"2026-08-06T10:18:24.466989Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2508.00324","last_updated":"2025-08-01T05:14:13Z","snapshot_observed_at":"2026-08-07T21:57:02.982200Z","submitted_at":"2025-08-01T05:14:13Z","title":"R1-ACT: Efficient Reasoning Model Safety Alignment by Activating Safety Knowledge","version":1},"reference_index":5,"source":"arxiv_source","source_observed_at":"2026-08-06T10:18:24.466989Z"},"links":{"cited_paper":"/paper/2412.16339","citing_paper":"/paper/2508.00324"},"observation_digest":"sha256:b92d2b8c18739d63b0c78ec4df76337ef92af133980d210ce4eda25dd5d73d96","observation_id":"a5a3d952-e6d6-49f3-a72c-8cf9204b488b","resolution":{"observed_at":"2026-08-06T10:18:24.466989Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2501.12948","last_updated":"2026-01-04T03:57:36Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2025-01-22T15:19:35Z","title":"DeepSeek-R1: Incentivizing Reasoning Capability in LLMs via Reinforcement Learning","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2501.12948","snapshot_observed_at":"2026-08-06T10:18:24.526526Z","title":null,"venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2508.00324","last_updated":"2025-08-01T05:14:13Z","snapshot_observed_at":"2026-08-07T21:57:02.982200Z","submitted_at":"2025-08-01T05:14:13Z","title":"R1-ACT: Efficient Reasoning Model Safety Alignment by Activating Safety Knowledge","version":1},"reference_index":6,"source":"arxiv_source","source_observed_at":"2026-08-06T10:18:24.526526Z"},"links":{"cited_paper":"/paper/2501.12948","citing_paper":"/paper/2508.00324"},"observation_digest":"sha256:ad68b7218d426f4ae55e4ca083358f7e988f58c9b8485e008ab8018bae17e737","observation_id":"169f18f3-c7f8-4951-acc6-b197dd78dc57","resolution":{"observed_at":"2026-08-06T10:18:24.526526Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2502.19361","last_updated":"2025-03-30T14:48:59Z","snapshot_observed_at":"2026-08-07T17:45:25.804804Z","submitted_at":"2025-02-26T17:59:27Z","title":"Can Large Language Models Detect Errors in Long Chain-of-Thought Reasoning?","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2502.19361","snapshot_observed_at":"2026-08-06T10:18:24.586763Z","title":null,"venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2508.00324","last_updated":"2025-08-01T05:14:13Z","snapshot_observed_at":"2026-08-07T21:57:02.982200Z","submitted_at":"2025-08-01T05:14:13Z","title":"R1-ACT: Efficient Reasoning Model Safety Alignment by Activating Safety Knowledge","version":1},"reference_index":7,"source":"arxiv_source","source_observed_at":"2026-08-06T10:18:24.586763Z"},"links":{"cited_paper":"/paper/2502.19361","citing_paper":"/paper/2508.00324"},"observation_digest":"sha256:d9978edc9ac08356da66b817e5f3eaf74f4ab0a12041ee88e0692dbffc4b9b64","observation_id":"a352b9f5-de06-469f-b40d-07e19cea8e24","resolution":{"observed_at":"2026-08-06T10:18:24.586763Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T10:18:27.728591Z","title":null,"venue":null,"work_id":"f7861d64-0358-4b0a-93f0-07946e44a3e5","year":1996},"citing_paper":{"arxiv_id":"2508.00324","last_updated":"2025-08-01T05:14:13Z","snapshot_observed_at":"2026-08-07T21:57:02.982200Z","submitted_at":"2025-08-01T05:14:13Z","title":"R1-ACT: Efficient Reasoning Model Safety Alignment by Activating Safety Knowledge","version":1},"reference_index":8,"source":"arxiv_source","source_observed_at":"2026-08-06T10:18:24.674776Z"},"links":{"citing_paper":"/paper/2508.00324"},"observation_digest":"sha256:3575d43d40b914a0aeebe691e433d2377114eeca87fb2779b05077d6f5c9188c","observation_id":"282c23b5-9a8f-4599-acc0-f377b7f77596","resolution":{"observed_at":"2026-08-06T10:18:27.785881Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2503.00555","last_updated":"2025-06-05T03:20:54Z","snapshot_observed_at":"2026-08-07T17:37:05.037079Z","submitted_at":"2025-03-01T16:42:01Z","title":"Safety Tax: Safety Alignment Makes Your Large Reasoning Models Less Reasonable","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2503.00555","snapshot_observed_at":"2026-08-06T10:18:24.777601Z","title":null,"venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2508.00324","last_updated":"2025-08-01T05:14:13Z","snapshot_observed_at":"2026-08-07T21:57:02.982200Z","submitted_at":"2025-08-01T05:14:13Z","title":"R1-ACT: Efficient Reasoning Model Safety Alignment by Activating Safety Knowledge","version":1},"reference_index":9,"source":"arxiv_source","source_observed_at":"2026-08-06T10:18:24.777601Z"},"links":{"cited_paper":"/paper/2503.00555","citing_paper":"/paper/2508.00324"},"observation_digest":"sha256:37f6a8344d5336a91764e14cca310dba27356272ac391ab4adef19c6d608941c","observation_id":"56d45a70-54e2-4893-8624-885b8fd276dc","resolution":{"observed_at":"2026-08-06T10:18:24.777601Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T10:18:24.829088Z","title":null,"venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2508.00324","last_updated":"2025-08-01T05:14:13Z","snapshot_observed_at":"2026-08-07T21:57:02.982200Z","submitted_at":"2025-08-01T05:14:13Z","title":"R1-ACT: Efficient Reasoning Model Safety Alignment by Activating Safety Knowledge","version":1},"reference_index":10,"source":"arxiv_source","source_observed_at":"2026-08-06T10:18:24.829088Z"},"links":{"citing_paper":"/paper/2508.00324"},"observation_digest":"sha256:c2a5d6e8b014653cf539fa6a2597e3122228920e4b530271533f3756e5b139e4","observation_id":"3491aa49-5990-4352-97d0-2b0d20a00cf0","resolution":{"observed_at":"2026-08-06T10:18:24.829088Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2412.16720","last_updated":"2026-04-30T02:46:40Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2024-12-21T18:04:31Z","title":"OpenAI o1 System Card","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2412.16720","snapshot_observed_at":"2026-08-06T10:18:24.887219Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2508.00324","last_updated":"2025-08-01T05:14:13Z","snapshot_observed_at":"2026-08-07T21:57:02.982200Z","submitted_at":"2025-08-01T05:14:13Z","title":"R1-ACT: Efficient Reasoning Model Safety Alignment by Activating Safety Knowledge","version":1},"reference_index":11,"source":"arxiv_source","source_observed_at":"2026-08-06T10:18:24.887219Z"},"links":{"cited_paper":"/paper/2412.16720","citing_paper":"/paper/2508.00324"},"observation_digest":"sha256:c4e8aa5c6cf6006bad7bd529721e12368eff059f11bfae2562500a11f65bf46d","observation_id":"3813ebaf-d21a-480e-955a-a77586ea0aec","resolution":{"observed_at":"2026-08-06T10:18:24.887219Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2502.12025","last_updated":"2025-02-17T16:57:56Z","snapshot_observed_at":"2026-08-07T18:11:43.608954Z","submitted_at":"2025-02-17T16:57:56Z","title":"SafeChain: Safety of Language Models with Long Chain-of-Thought Reasoning Capabilities","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2502.12025","snapshot_observed_at":"2026-08-06T10:18:24.965589Z","title":null,"venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2508.00324","last_updated":"2025-08-01T05:14:13Z","snapshot_observed_at":"2026-08-07T21:57:02.982200Z","submitted_at":"2025-08-01T05:14:13Z","title":"R1-ACT: Efficient Reasoning Model Safety Alignment by Activating Safety Knowledge","version":1},"reference_index":12,"source":"arxiv_source","source_observed_at":"2026-08-06T10:18:24.965589Z"},"links":{"cited_paper":"/paper/2502.12025","citing_paper":"/paper/2508.00324"},"observation_digest":"sha256:8a92d0e192980df250a256ef7e4a29cf92e18f85694f4d0680cac372e1854dc6","observation_id":"95da808b-9f6f-439c-b054-0b6db5c73715","resolution":{"observed_at":"2026-08-06T10:18:24.965589Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T10:18:27.582189Z","title":null,"venue":null,"work_id":"a20bec5f-bf86-46d0-990c-59b252fbebe5","year":2024},"citing_paper":{"arxiv_id":"2508.00324","last_updated":"2025-08-01T05:14:13Z","snapshot_observed_at":"2026-08-07T21:57:02.982200Z","submitted_at":"2025-08-01T05:14:13Z","title":"R1-ACT: Efficient Reasoning Model Safety Alignment by Activating Safety Knowledge","version":1},"reference_index":13,"source":"arxiv_source","source_observed_at":"2026-08-06T10:18:25.102197Z"},"links":{"citing_paper":"/paper/2508.00324"},"observation_digest":"sha256:1e19c7436649d65c60722fcacb51bb92bdbbdbc318d0967d74846d7be7b657ea","observation_id":"8516d96d-07f7-4324-a998-88a420aeb5dc","resolution":{"observed_at":"2026-08-06T10:18:27.661356Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T10:18:25.202097Z","title":null,"venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2508.00324","last_updated":"2025-08-01T05:14:13Z","snapshot_observed_at":"2026-08-07T21:57:02.982200Z","submitted_at":"2025-08-01T05:14:13Z","title":"R1-ACT: Efficient Reasoning Model Safety Alignment by Activating Safety Knowledge","version":1},"reference_index":14,"source":"arxiv_source","source_observed_at":"2026-08-06T10:18:25.202097Z"},"links":{"citing_paper":"/paper/2508.00324"},"observation_digest":"sha256:460b3ada328038cb257e9ab987818f9f42d5b9ca496a6bf357294da39b3449ca","observation_id":"4d6d1703-27a6-4005-ba2c-bd32f3a00ca1","resolution":{"observed_at":"2026-08-06T10:18:25.202097Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2501.19393","last_updated":"2025-03-01T06:07:39Z","snapshot_observed_at":"2026-07-06T20:29:11.710285Z","submitted_at":"2025-01-31T18:48:08Z","title":"s1: Simple test-time scaling","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2501.19393","snapshot_observed_at":"2026-08-06T10:18:25.282728Z","title":null,"venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2508.00324","last_updated":"2025-08-01T05:14:13Z","snapshot_observed_at":"2026-08-07T21:57:02.982200Z","submitted_at":"2025-08-01T05:14:13Z","title":"R1-ACT: Efficient Reasoning Model Safety Alignment by Activating Safety Knowledge","version":1},"reference_index":15,"source":"arxiv_source","source_observed_at":"2026-08-06T10:18:25.282728Z"},"links":{"cited_paper":"/paper/2501.19393","citing_paper":"/paper/2508.00324"},"observation_digest":"sha256:ffb4fb020ac1efdb9574883a8bed70009500377e7642512352ed9b7716ad0191","observation_id":"a5ca5ce9-3da0-471c-b728-7a9702899547","resolution":{"observed_at":"2026-08-06T10:18:25.282728Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2308.01263","last_updated":"2024-04-01T11:50:35Z","snapshot_observed_at":"2026-08-03T00:58:55.865010Z","submitted_at":"2023-08-02T16:30:40Z","title":"XSTest: A Test Suite for Identifying Exaggerated Safety Behaviours in Large Language Models","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2308.01263","snapshot_observed_at":"2026-08-06T10:18:25.385420Z","title":null,"venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2508.00324","last_updated":"2025-08-01T05:14:13Z","snapshot_observed_at":"2026-08-07T21:57:02.982200Z","submitted_at":"2025-08-01T05:14:13Z","title":"R1-ACT: Efficient Reasoning Model Safety Alignment by Activating Safety Knowledge","version":1},"reference_index":16,"source":"arxiv_source","source_observed_at":"2026-08-06T10:18:25.385420Z"},"links":{"cited_paper":"/paper/2308.01263","citing_paper":"/paper/2508.00324"},"observation_digest":"sha256:f8677698b834253fcc14103d8eaa6cae7300fc40a66d2378cbabee7894781d8b","observation_id":"fee9320a-6f73-4392-88dd-e22aaae91eb3","resolution":{"observed_at":"2026-08-06T10:18:25.385420Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2402.03300","last_updated":"2024-04-27T15:25:53Z","snapshot_observed_at":"2026-08-06T14:58:42.911363Z","submitted_at":"2024-02-05T18:55:32Z","title":"DeepSeekMath: Pushing the Limits of Mathematical Reasoning in Open Language Models","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2402.03300","snapshot_observed_at":"2026-08-06T10:18:25.473811Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2508.00324","last_updated":"2025-08-01T05:14:13Z","snapshot_observed_at":"2026-08-07T21:57:02.982200Z","submitted_at":"2025-08-01T05:14:13Z","title":"R1-ACT: Efficient Reasoning Model Safety Alignment by Activating Safety Knowledge","version":1},"reference_index":17,"source":"arxiv_source","source_observed_at":"2026-08-06T10:18:25.473811Z"},"links":{"cited_paper":"/paper/2402.03300","citing_paper":"/paper/2508.00324"},"observation_digest":"sha256:af15cf5f1be2c715481daa456199133c4f159c8ec63ad3a5ac5c778719e30592","observation_id":"18f5d5f6-e434-480f-9c36-f15af4156978","resolution":{"observed_at":"2026-08-06T10:18:25.473811Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2402.10260","last_updated":"2024-08-27T03:32:47Z","snapshot_observed_at":"2026-08-04T21:32:35.483431Z","submitted_at":"2024-02-15T18:58:09Z","title":"A StrongREJECT for Empty Jailbreaks","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2402.10260","snapshot_observed_at":"2026-08-06T10:18:25.592722Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2508.00324","last_updated":"2025-08-01T05:14:13Z","snapshot_observed_at":"2026-08-07T21:57:02.982200Z","submitted_at":"2025-08-01T05:14:13Z","title":"R1-ACT: Efficient Reasoning Model Safety Alignment by Activating Safety Knowledge","version":1},"reference_index":18,"source":"arxiv_source","source_observed_at":"2026-08-06T10:18:25.592722Z"},"links":{"cited_paper":"/paper/2402.10260","citing_paper":"/paper/2508.00324"},"observation_digest":"sha256:dbe04c8315895b90ea0d24da20625681557d3c950fc55789087806420e04fa21","observation_id":"077607d6-89a4-4943-b341-4d1d0cd3ef9c","resolution":{"observed_at":"2026-08-06T10:18:25.592722Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T10:18:25.690409Z","title":null,"venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2508.00324","last_updated":"2025-08-01T05:14:13Z","snapshot_observed_at":"2026-08-07T21:57:02.982200Z","submitted_at":"2025-08-01T05:14:13Z","title":"R1-ACT: Efficient Reasoning Model Safety Alignment by Activating Safety Knowledge","version":1},"reference_index":19,"source":"arxiv_source","source_observed_at":"2026-08-06T10:18:25.690409Z"},"links":{"citing_paper":"/paper/2508.00324"},"observation_digest":"sha256:5f040bea1caac39375f7fe680e8b62a84c860990fe13c7535782d2b02db6e119","observation_id":"8f062020-bac7-4c7c-8d06-74fc588d23f5","resolution":{"observed_at":"2026-08-06T10:18:25.690409Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T10:18:25.798387Z","title":null,"venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2508.00324","last_updated":"2025-08-01T05:14:13Z","snapshot_observed_at":"2026-08-07T21:57:02.982200Z","submitted_at":"2025-08-01T05:14:13Z","title":"R1-ACT: Efficient Reasoning Model Safety Alignment by Activating Safety Knowledge","version":1},"reference_index":20,"source":"arxiv_source","source_observed_at":"2026-08-06T10:18:25.798387Z"},"links":{"citing_paper":"/paper/2508.00324"},"observation_digest":"sha256:52efe858ec395af35a56d3cf4783265a852a5a9458f566917742c4227aafee74","observation_id":"a5d58cb8-17e3-48b1-933f-5a9d56d56b24","resolution":{"observed_at":"2026-08-06T10:18:25.798387Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T10:18:25.883185Z","title":null,"venue":null,"work_id":null,"year":2022},"citing_paper":{"arxiv_id":"2508.00324","last_updated":"2025-08-01T05:14:13Z","snapshot_observed_at":"2026-08-07T21:57:02.982200Z","submitted_at":"2025-08-01T05:14:13Z","title":"R1-ACT: Efficient Reasoning Model Safety Alignment by Activating Safety Knowledge","version":1},"reference_index":21,"source":"arxiv_source","source_observed_at":"2026-08-06T10:18:25.883185Z"},"links":{"citing_paper":"/paper/2508.00324"},"observation_digest":"sha256:d45b3a283b733aa3cc28d5b06871f1129a028ffe8afb302331987ecb66a1c6b5","observation_id":"6f4c8f8e-809e-438b-8126-b3bc6754e37e","resolution":{"observed_at":"2026-08-06T10:18:25.883185Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2406.14598","last_updated":"2025-03-01T21:45:36Z","snapshot_observed_at":"2026-07-06T18:34:29.513732Z","submitted_at":"2024-06-20T17:56:07Z","title":"SORRY-Bench: Systematically Evaluating Large Language Model Safety Refusal","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2406.14598","snapshot_observed_at":"2026-08-06T10:18:25.973754Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2508.00324","last_updated":"2025-08-01T05:14:13Z","snapshot_observed_at":"2026-08-07T21:57:02.982200Z","submitted_at":"2025-08-01T05:14:13Z","title":"R1-ACT: Efficient Reasoning Model Safety Alignment by Activating Safety Knowledge","version":1},"reference_index":22,"source":"arxiv_source","source_observed_at":"2026-08-06T10:18:25.973754Z"},"links":{"cited_paper":"/paper/2406.14598","citing_paper":"/paper/2508.00324"},"observation_digest":"sha256:253a996910fc28b9ba11cefdf546616ffe8c39c0c4cd4d3919d9569d719060e0","observation_id":"3a452ff4-aec0-4560-81e1-dc39de49abfd","resolution":{"observed_at":"2026-08-06T10:18:25.973754Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2505.15214","last_updated":"2025-05-27T07:37:47Z","snapshot_observed_at":"2026-08-07T21:55:27.920047Z","submitted_at":"2025-05-21T07:44:30Z","title":"R-TOFU: Unlearning in Large Reasoning Models","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2505.15214","snapshot_observed_at":"2026-08-06T10:18:26.115945Z","title":null,"venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2508.00324","last_updated":"2025-08-01T05:14:13Z","snapshot_observed_at":"2026-08-07T21:57:02.982200Z","submitted_at":"2025-08-01T05:14:13Z","title":"R1-ACT: Efficient Reasoning Model Safety Alignment by Activating Safety Knowledge","version":1},"reference_index":23,"source":"arxiv_source","source_observed_at":"2026-08-06T10:18:26.115945Z"},"links":{"cited_paper":"/paper/2505.15214","citing_paper":"/paper/2508.00324"},"observation_digest":"sha256:1fbcc8e04d0b36ba65254b5e91e9824c4341cb9f833057d52d0bf36f0a76fc2e","observation_id":"9baf45b4-6b44-44e2-84f2-b7243c4134f7","resolution":{"observed_at":"2026-08-06T10:18:26.115945Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2504.10081","last_updated":"2025-04-14T10:26:37Z","snapshot_observed_at":"2026-08-07T16:05:57.274080Z","submitted_at":"2025-04-14T10:26:37Z","title":"RealSafe-R1: Safety-Aligned DeepSeek-R1 without Compromising Reasoning Capability","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2504.10081","snapshot_observed_at":"2026-08-06T10:18:26.207413Z","title":null,"venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2508.00324","last_updated":"2025-08-01T05:14:13Z","snapshot_observed_at":"2026-08-07T21:57:02.982200Z","submitted_at":"2025-08-01T05:14:13Z","title":"R1-ACT: Efficient Reasoning Model Safety Alignment by Activating Safety Knowledge","version":1},"reference_index":24,"source":"arxiv_source","source_observed_at":"2026-08-06T10:18:26.207413Z"},"links":{"cited_paper":"/paper/2504.10081","citing_paper":"/paper/2508.00324"},"observation_digest":"sha256:90812094a18b7e9eeb5a6ddb4dd34959f5baf6289b4d7dfd6c38c23310d3364f","observation_id":"01dce2b2-92b8-4db2-ba16-40aefb5788d8","resolution":{"observed_at":"2026-08-06T10:18:26.207413Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2505.15404","last_updated":"2026-04-20T03:42:03Z","snapshot_observed_at":"2026-08-08T05:48:26.464887Z","submitted_at":"2025-05-21T11:45:29Z","title":"How Should We Enhance the Safety of Large Reasoning Models: An Empirical Study","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2505.15404","snapshot_observed_at":"2026-08-06T10:18:26.329974Z","title":null,"venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2508.00324","last_updated":"2025-08-01T05:14:13Z","snapshot_observed_at":"2026-08-07T21:57:02.982200Z","submitted_at":"2025-08-01T05:14:13Z","title":"R1-ACT: Efficient Reasoning Model Safety Alignment by Activating Safety Knowledge","version":1},"reference_index":25,"source":"arxiv_source","source_observed_at":"2026-08-06T10:18:26.329974Z"},"links":{"cited_paper":"/paper/2505.15404","citing_paper":"/paper/2508.00324"},"observation_digest":"sha256:f5d099097859f90c8f68cc5256b4e8ec340f4fce31bd7cc231af6f4f259ee5a6","observation_id":"78132a7e-f53b-4d0c-b0f6-0771465e260a","resolution":{"observed_at":"2026-08-06T10:18:26.329974Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T10:18:26.443904Z","title":null,"venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2508.00324","last_updated":"2025-08-01T05:14:13Z","snapshot_observed_at":"2026-08-07T21:57:02.982200Z","submitted_at":"2025-08-01T05:14:13Z","title":"R1-ACT: Efficient Reasoning Model Safety Alignment by Activating Safety Knowledge","version":1},"reference_index":26,"source":"arxiv_source","source_observed_at":"2026-08-06T10:18:26.443904Z"},"links":{"citing_paper":"/paper/2508.00324"},"observation_digest":"sha256:295169f2b4f50b8808bef9027576a3ba49d865393032b751b82ae2a4314d2022","observation_id":"7e7a4f2e-d2c6-4ecc-908a-b3159cc50641","resolution":{"observed_at":"2026-08-06T10:18:26.443904Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T10:18:26.512082Z","title":null,"venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2508.00324","last_updated":"2025-08-01T05:14:13Z","snapshot_observed_at":"2026-08-07T21:57:02.982200Z","submitted_at":"2025-08-01T05:14:13Z","title":"R1-ACT: Efficient Reasoning Model Safety Alignment by Activating Safety Knowledge","version":1},"reference_index":27,"source":"arxiv_source","source_observed_at":"2026-08-06T10:18:26.512082Z"},"links":{"citing_paper":"/paper/2508.00324"},"observation_digest":"sha256:323663e39988d8dbf225b51e36e93a33997cf5ea551dd459bd0cf23c49c85ef2","observation_id":"61e4de22-82a8-4f38-85f4-f4d2e50d5d3f","resolution":{"observed_at":"2026-08-06T10:18:26.512082Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T10:18:26.602669Z","title":"URL: \" 'urlintro :=","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2508.00324","last_updated":"2025-08-01T05:14:13Z","snapshot_observed_at":"2026-08-07T21:57:02.982200Z","submitted_at":"2025-08-01T05:14:13Z","title":"R1-ACT: Efficient Reasoning Model Safety Alignment by Activating Safety Knowledge","version":1},"reference_index":28,"source":"arxiv_source","source_observed_at":"2026-08-06T10:18:26.602669Z"},"links":{"citing_paper":"/paper/2508.00324"},"observation_digest":"sha256:c1160636e98775629df327f5f61b283ba444b4cfe1b60c2fb762437cc9cf8322","observation_id":"6df207fd-d5dc-40cf-a0ec-116e38c703b2","resolution":{"observed_at":"2026-08-06T10:18:26.602669Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T10:18:26.686397Z","title":"write newline","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2508.00324","last_updated":"2025-08-01T05:14:13Z","snapshot_observed_at":"2026-08-07T21:57:02.982200Z","submitted_at":"2025-08-01T05:14:13Z","title":"R1-ACT: Efficient Reasoning Model Safety Alignment by Activating Safety Knowledge","version":1},"reference_index":29,"source":"arxiv_source","source_observed_at":"2026-08-06T10:18:26.686397Z"},"links":{"citing_paper":"/paper/2508.00324"},"observation_digest":"sha256:c7e4042635f21df36d2287d27295f509e73b7c4fc049dde6a68e069be90d2fcb","observation_id":"053cfe66-1469-4c3b-b4e7-0047f45465df","resolution":{"observed_at":"2026-08-06T10:18:26.686397Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}}],"paper":{"arxiv_id":"2508.00324","last_updated":"2025-08-01T05:14:13Z","latest_version":1,"primary_category":"cs.AI","snapshot_observed_at":"2026-08-07T21:57:02.982200Z","submitted_at":"2025-08-01T05:14:13Z","title":"R1-ACT: Efficient Reasoning Model Safety Alignment by Activating Safety Knowledge"},"reference_resolution":{"displayed":29,"state_counts":{"malformed_identifier":0,"metadata_mismatch":0,"parse_uncertain":0,"unresolved":29,"verified_exact":0,"verified_fuzzy":0},"total_outbound_references":29},"refusal":"A citation records a reference. It does not transfer a finding from one paper to another.","schema":"pith.paper-citation-record.v1","standing_sources":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"thesis":"As of 8 August 2026, this Paper Citation Record lists 29 of 29 outbound references and 1 inbound Pith citation observation for arXiv:2508.00324."}