{"as_of":"2026-08-19T10:04:00Z","caps":{"database_statements":6,"inbound":100,"outbound":100},"context_digest":"sha256:af7c9ea4aa5fe920a864cc5336bb6b619c5a3a60e3cffabff56005f7c0c191cd","coverage":[{"denominator":25,"lane":"reference_resolution","note":"Typed states for the displayed outbound observations.","records_observed":25,"source":"paper_references, paper_reference_links","source_observed_at":"2026-08-10T04:37:17.038006Z","state":"measured"},{"denominator":27,"lane":"standing_notices","note":"One-hop event checks from named stored sources.","records_observed":27,"source":"scholarly_work_events, retraction_status_cache","source_observed_at":"2026-08-19T06:32:44.657259+00:00","state":"measured"},{"denominator":2,"lane":"inbound_itemization","note":"Pith citing papers itemized under the disclosed page cap.","records_observed":2,"source":"paper_references, paper_reference_links","source_observed_at":"2026-08-07T14:13:50.508020Z","state":"measured"},{"denominator":1,"lane":"external_citation_measurements","note":"A source-named dated measurement, never combined with another source.","records_observed":0,"source":"arxiv_reference","source_observed_at":"2026-05-12T08:40:41.350151Z","state":"measured"}],"external_citation_measurements":[],"inbound":[{"citation":{"cited_paper":{"arxiv_id":"2501.17749","last_updated":"2025-01-29T16:36:53Z","snapshot_observed_at":"2026-08-19T02:28:54.940509Z","submitted_at":"2025-01-29T16:36:53Z","title":"Early External Safety Testing of OpenAI's o3-mini: Insights from the Pre-Deployment Evaluation","version":1},"cited_work":{"arxiv_id":"2501.17749","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2501.17749","snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Early external safety testing of openai’s o3-mini: Insights from the pre-deployment evaluation","venue":null,"work_id":"42dd7709-4567-4062-a49f-0f4b83e104f7","year":2025},"citing_paper":{"arxiv_id":"2503.09567","last_updated":"2025-07-18T15:57:54Z","snapshot_observed_at":"2026-08-08T22:33:20.124926Z","submitted_at":"2025-03-12T17:35:03Z","title":"Towards Reasoning Era: A Survey of Long Chain-of-Thought for Reasoning Large Language Models","version":5},"reference_index":23,"source":"pdf_text","source_observed_at":"2026-05-12T08:40:40.910461Z"},"links":{"cited_paper":"/paper/2501.17749","citing_paper":"/paper/2503.09567"},"observation_digest":"sha256:6136316c9b41961bbd41ca3f9995bc93e2ff1ee4f82f707d559d2104b22d09ea","observation_id":"f91ec8fc-2a41-4e0c-9663-6682edc638e3","resolution":{"observed_at":"2026-05-12T08:40:41.353112Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-19T06:32:44.657259+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-19T06:32:44.657259+00:00","source":"crossref"},{"observed_at":"2026-08-19T06:32:39.956319+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2501.17749","last_updated":"2025-01-29T16:36:53Z","snapshot_observed_at":"2026-08-19T02:28:54.940509Z","submitted_at":"2025-01-29T16:36:53Z","title":"Early External Safety Testing of OpenAI's o3-mini: Insights from the Pre-Deployment Evaluation","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2501.17749","snapshot_observed_at":"2026-08-07T14:13:50.508020Z","title":"Early external safety testing of openai’s o3-mini: Insights from the pre-deployment evaluation,","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2505.19690","last_updated":"2025-05-26T08:49:19Z","snapshot_observed_at":"2026-08-17T01:37:19.231497Z","submitted_at":"2025-05-26T08:49:19Z","title":"Beyond Safe Answers: A Benchmark for Evaluating True Risk Awareness in Large Reasoning Models","version":1},"reference_index":12,"source":"pdf_text","source_observed_at":"2026-08-07T14:13:50.508020Z"},"links":{"cited_paper":"/paper/2501.17749","citing_paper":"/paper/2505.19690"},"observation_digest":"sha256:c38da65004adc6085fa41f239cf54ff156a6eb162d5d1c5fd2f86b23b0efdab5","observation_id":"15ecdde5-f430-4616-8d8d-7541c6213c5d","resolution":{"observed_at":"2026-08-07T14:13:50.508020Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}}],"links":{"evidence":"/evidence","html":"/paper/2501.17749/citation-record","integrity":"/paper/2501.17749/integrity","json":"/paper/2501.17749/citation-record.json","paper":"/paper/2501.17749"},"outbound":[{"citation":{"cited_paper":{"arxiv_id":"2406.14598","last_updated":"2025-03-01T21:45:36Z","snapshot_observed_at":"2026-08-17T11:35:24.473105Z","submitted_at":"2024-06-20T17:56:07Z","title":"SORRY-Bench: Systematically Evaluating Large Language Model Safety Refusal","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2406.14598","snapshot_observed_at":"2026-08-10T04:37:16.943712Z","title":"Sorry-bench: Systematically evaluating large language model safety ref usal behaviors,","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2501.17749","last_updated":"2025-01-29T16:36:53Z","snapshot_observed_at":"2026-08-19T02:28:54.940509Z","submitted_at":"2025-01-29T16:36:53Z","title":"Early External Safety Testing of OpenAI's o3-mini: Insights from the Pre-Deployment Evaluation","version":1},"reference_index":1,"source":"pdf_text","source_observed_at":"2026-08-10T04:37:16.943712Z"},"links":{"cited_paper":"/paper/2406.14598","citing_paper":"/paper/2501.17749"},"observation_digest":"sha256:8a39a055b0d215126bfcaf8aa6b4290c31413218ea0f44b6764e891b30ba29cf","observation_id":"00626506-bf78-4424-b696-1cffda6c7dca","resolution":{"observed_at":"2026-08-10T04:37:16.943712Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2405.14191","last_updated":"2025-04-07T07:52:28Z","snapshot_observed_at":"2026-08-18T14:32:35.674645Z","submitted_at":"2024-05-23T05:34:31Z","title":"S-Eval: Towards Automated and Comprehensive Safety Evaluation for Large Language Models","version":4},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2405.14191","snapshot_observed_at":"2026-08-10T04:37:16.948537Z","title":"S-eval: Auto- matic and adaptive test generation for benchmarking safety evaluation of large language models,","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2501.17749","last_updated":"2025-01-29T16:36:53Z","snapshot_observed_at":"2026-08-19T02:28:54.940509Z","submitted_at":"2025-01-29T16:36:53Z","title":"Early External Safety Testing of OpenAI's o3-mini: Insights from the Pre-Deployment Evaluation","version":1},"reference_index":2,"source":"pdf_text","source_observed_at":"2026-08-10T04:37:16.948537Z"},"links":{"cited_paper":"/paper/2405.14191","citing_paper":"/paper/2501.17749"},"observation_digest":"sha256:c9248e7785c46e2b77553fdc6198759989aca982e16920ef7f9986898c1a72b1","observation_id":"2b96bf4d-ed93-4bed-b67b-095f5ad09327","resolution":{"observed_at":"2026-08-10T04:37:16.948537Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2309.07045","last_updated":"2024-06-24T04:04:21Z","snapshot_observed_at":"2026-08-18T14:35:02.737196Z","submitted_at":"2023-09-13T15:56:50Z","title":"SafetyBench: Evaluating the Safety of Large Language Models","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2309.07045","snapshot_observed_at":"2026-08-10T04:37:16.952819Z","title":"Safetybench: Evaluating the safety of large language models with multiple choice questions,","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2501.17749","last_updated":"2025-01-29T16:36:53Z","snapshot_observed_at":"2026-08-19T02:28:54.940509Z","submitted_at":"2025-01-29T16:36:53Z","title":"Early External Safety Testing of OpenAI's o3-mini: Insights from the Pre-Deployment Evaluation","version":1},"reference_index":3,"source":"pdf_text","source_observed_at":"2026-08-10T04:37:16.952819Z"},"links":{"cited_paper":"/paper/2309.07045","citing_paper":"/paper/2501.17749"},"observation_digest":"sha256:b65d8d78141c134069f6c4a9720583ad5f6b8445296c5fc08377136901fad2c0","observation_id":"10b0be7e-f833-4089-a919-4ac79cdbfcff","resolution":{"observed_at":"2026-08-10T04:37:16.952819Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2406.10311","last_updated":"2024-09-02T03:37:35Z","snapshot_observed_at":"2026-08-16T13:43:05.211611Z","submitted_at":"2024-06-14T06:47:40Z","title":"CHiSafetyBench: A Chinese Hierarchical Safety Benchmark for Large Language Models","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2406.10311","snapshot_observed_at":"2026-08-10T04:37:16.956518Z","title":"Chisafetybench: A chinese hierar- chical safety benchmark for large language models,","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2501.17749","last_updated":"2025-01-29T16:36:53Z","snapshot_observed_at":"2026-08-19T02:28:54.940509Z","submitted_at":"2025-01-29T16:36:53Z","title":"Early External Safety Testing of OpenAI's o3-mini: Insights from the Pre-Deployment Evaluation","version":1},"reference_index":4,"source":"pdf_text","source_observed_at":"2026-08-10T04:37:16.956518Z"},"links":{"cited_paper":"/paper/2406.10311","citing_paper":"/paper/2501.17749"},"observation_digest":"sha256:dc5dad103ea1d03b989a45c5708e6968464a3a2c265854a80fc7b4645c4b55b6","observation_id":"bb7698cd-626c-41ce-8617-d4761030ec04","resolution":{"observed_at":"2026-08-10T04:37:16.956518Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2410.18927","last_updated":"2024-10-24T17:14:40Z","snapshot_observed_at":"2026-08-16T13:06:11.757117Z","submitted_at":"2024-10-24T17:14:40Z","title":"SafeBench: A Safety Evaluation Framework for Multimodal Large Language Models","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2410.18927","snapshot_observed_at":"2026-08-10T04:37:16.960430Z","title":"Safebench: A safety evaluation framework for multimodal large language models,","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2501.17749","last_updated":"2025-01-29T16:36:53Z","snapshot_observed_at":"2026-08-19T02:28:54.940509Z","submitted_at":"2025-01-29T16:36:53Z","title":"Early External Safety Testing of OpenAI's o3-mini: Insights from the Pre-Deployment Evaluation","version":1},"reference_index":5,"source":"pdf_text","source_observed_at":"2026-08-10T04:37:16.960430Z"},"links":{"cited_paper":"/paper/2410.18927","citing_paper":"/paper/2501.17749"},"observation_digest":"sha256:c64343ec61a08b192057634ba4ec0cc5c53264e7e0f9b46552dde49702dee73a","observation_id":"842d4427-4229-4d9f-b061-38f3e911ef4c","resolution":{"observed_at":"2026-08-10T04:37:16.960430Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2411.06899","last_updated":"2025-02-27T13:08:46Z","snapshot_observed_at":"2026-08-16T13:01:12.867974Z","submitted_at":"2024-11-11T11:57:37Z","title":"LongSafety: Enhance Safety for Long-Context LLMs","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2411.06899","snapshot_observed_at":"2026-08-10T04:37:16.964203Z","title":"Longsafetybench: Long-context llms struggle with safety issues,","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2501.17749","last_updated":"2025-01-29T16:36:53Z","snapshot_observed_at":"2026-08-19T02:28:54.940509Z","submitted_at":"2025-01-29T16:36:53Z","title":"Early External Safety Testing of OpenAI's o3-mini: Insights from the Pre-Deployment Evaluation","version":1},"reference_index":6,"source":"pdf_text","source_observed_at":"2026-08-10T04:37:16.964203Z"},"links":{"cited_paper":"/paper/2411.06899","citing_paper":"/paper/2501.17749"},"observation_digest":"sha256:1450dcffa73d1ea8b305ed9a66b4bab6b2a489751a7449833f13d3e4be4c4a0c","observation_id":"d0ff1b51-4537-40f7-b0e4-8ba34959e8ea","resolution":{"observed_at":"2026-08-10T04:37:16.964203Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2402.05044","last_updated":"2024-06-07T12:05:46Z","snapshot_observed_at":"2026-08-16T14:20:31.299446Z","submitted_at":"2024-02-07T17:33:54Z","title":"SALAD-Bench: A Hierarchical and Comprehensive Safety Benchmark for Large Language Models","version":4},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2402.05044","snapshot_observed_at":"2026-08-10T04:37:16.968365Z","title":"Salad-bench: A hierarchical and comprehensive safety benchmark for large language models,","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2501.17749","last_updated":"2025-01-29T16:36:53Z","snapshot_observed_at":"2026-08-19T02:28:54.940509Z","submitted_at":"2025-01-29T16:36:53Z","title":"Early External Safety Testing of OpenAI's o3-mini: Insights from the Pre-Deployment Evaluation","version":1},"reference_index":7,"source":"pdf_text","source_observed_at":"2026-08-10T04:37:16.968365Z"},"links":{"cited_paper":"/paper/2402.05044","citing_paper":"/paper/2501.17749"},"observation_digest":"sha256:e66ff752e1cdbb540dc38f19a173be3095ee9ef0a673de3de52a6e8ab85d4b54","observation_id":"bb969384-9c6c-4882-b7fe-83ea091a5669","resolution":{"observed_at":"2026-08-10T04:37:16.968365Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-10T04:37:17.310841Z","title":"Beavertails: Towards improved safety alignment of LLM via a human-preference dataset,","venue":null,"work_id":"d70bf0c2-8ecc-492b-a7ec-6ac5707a2722","year":2024},"citing_paper":{"arxiv_id":"2501.17749","last_updated":"2025-01-29T16:36:53Z","snapshot_observed_at":"2026-08-19T02:28:54.940509Z","submitted_at":"2025-01-29T16:36:53Z","title":"Early External Safety Testing of OpenAI's o3-mini: Insights from the Pre-Deployment Evaluation","version":1},"reference_index":8,"source":"pdf_text","source_observed_at":"2026-08-10T04:37:16.971913Z"},"links":{"citing_paper":"/paper/2501.17749"},"observation_digest":"sha256:7c872c75835bfe992608fd60c7b057fd3e2303b1060c052f724ef88eb0b21363","observation_id":"a893c6c9-40be-435a-b83f-d26c4f08386d","resolution":{"observed_at":"2026-08-10T04:37:17.314896Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-19T06:32:44.657259+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-19T06:32:44.657259+00:00","source":"crossref"},{"observed_at":"2026-08-19T06:32:39.956319+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2311.08370","last_updated":"2024-02-16T09:42:19Z","snapshot_observed_at":"2026-08-19T06:38:07.871486Z","submitted_at":"2023-11-14T18:33:43Z","title":"SimpleSafetyTests: a Test Suite for Identifying Critical Safety Risks in Large Language Models","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2311.08370","snapshot_observed_at":"2026-08-10T04:37:16.975584Z","title":"Simplesafetytests: a test suite for identifying critical safety risks in large langua ge models,","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2501.17749","last_updated":"2025-01-29T16:36:53Z","snapshot_observed_at":"2026-08-19T02:28:54.940509Z","submitted_at":"2025-01-29T16:36:53Z","title":"Early External Safety Testing of OpenAI's o3-mini: Insights from the Pre-Deployment Evaluation","version":1},"reference_index":9,"source":"pdf_text","source_observed_at":"2026-08-10T04:37:16.975584Z"},"links":{"cited_paper":"/paper/2311.08370","citing_paper":"/paper/2501.17749"},"observation_digest":"sha256:dc7fac40f4144d9fa975105d90eed6f6dce6150ea3c1db72749064d15213d81e","observation_id":"47aba5ba-9f30-4c3b-9b95-6928a2e9bb0b","resolution":{"observed_at":"2026-08-10T04:37:16.975584Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-10T04:37:17.298677Z","title":"Astral: Automated safety testing of large language models,","venue":null,"work_id":"554e9eb4-c969-449e-801c-688852a0085c","year":2025},"citing_paper":{"arxiv_id":"2501.17749","last_updated":"2025-01-29T16:36:53Z","snapshot_observed_at":"2026-08-19T02:28:54.940509Z","submitted_at":"2025-01-29T16:36:53Z","title":"Early External Safety Testing of OpenAI's o3-mini: Insights from the Pre-Deployment Evaluation","version":1},"reference_index":10,"source":"pdf_text","source_observed_at":"2026-08-10T04:37:16.979482Z"},"links":{"citing_paper":"/paper/2501.17749"},"observation_digest":"sha256:c3d61bc8618f3f8b9e847c71be826811fe71ab6ca1d7ae447b665b7f9e205dc6","observation_id":"cccad398-d7c5-4490-8bf7-dd5ffd2b1e11","resolution":{"observed_at":"2026-08-10T04:37:17.302892Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-19T06:32:44.657259+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-19T06:32:44.657259+00:00","source":"crossref"},{"observed_at":"2026-08-19T06:32:39.956319+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-10T04:37:17.287574Z","title":"A survey on metamorphic testing,","venue":null,"work_id":"b9960d50-d12d-42b8-a42b-03f4c270ec11","year":2016},"citing_paper":{"arxiv_id":"2501.17749","last_updated":"2025-01-29T16:36:53Z","snapshot_observed_at":"2026-08-19T02:28:54.940509Z","submitted_at":"2025-01-29T16:36:53Z","title":"Early External Safety Testing of OpenAI's o3-mini: Insights from the Pre-Deployment Evaluation","version":1},"reference_index":11,"source":"pdf_text","source_observed_at":"2026-08-10T04:37:16.983004Z"},"links":{"citing_paper":"/paper/2501.17749"},"observation_digest":"sha256:65ceb6516018b7933c88a30c2fba4d9c3868f062c26dab756403887424f03cbd","observation_id":"d79e7cbe-e727-45df-bafb-bce88417f312","resolution":{"observed_at":"2026-08-10T04:37:17.291051Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-19T06:32:44.657259+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-19T06:32:44.657259+00:00","source":"crossref"},{"observed_at":"2026-08-19T06:32:39.956319+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-10T04:37:17.275830Z","title":"Guardrails for trust, safet y, and ethical development and deployment of large language models (llm),","venue":null,"work_id":"ca60be87-ec82-416e-a1b9-d0b7deee29fb","year":2023},"citing_paper":{"arxiv_id":"2501.17749","last_updated":"2025-01-29T16:36:53Z","snapshot_observed_at":"2026-08-19T02:28:54.940509Z","submitted_at":"2025-01-29T16:36:53Z","title":"Early External Safety Testing of OpenAI's o3-mini: Insights from the Pre-Deployment Evaluation","version":1},"reference_index":12,"source":"pdf_text","source_observed_at":"2026-08-10T04:37:16.986274Z"},"links":{"citing_paper":"/paper/2501.17749"},"observation_digest":"sha256:885074bf05d8645906c306866aef484d583d7bbeaca13ce0c76cef68a4b5a7f3","observation_id":"04dee86e-138f-4109-a9fe-b5931220627e","resolution":{"observed_at":"2026-08-10T04:37:17.280008Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-19T06:32:44.657259+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-19T06:32:44.657259+00:00","source":"crossref"},{"observed_at":"2026-08-19T06:32:39.956319+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-10T04:37:17.262829Z","title":"European Commission AI Act","venue":null,"work_id":"f4ef093a-e06f-43ce-b962-e1f78f7c5f9f","year":null},"citing_paper":{"arxiv_id":"2501.17749","last_updated":"2025-01-29T16:36:53Z","snapshot_observed_at":"2026-08-19T02:28:54.940509Z","submitted_at":"2025-01-29T16:36:53Z","title":"Early External Safety Testing of OpenAI's o3-mini: Insights from the Pre-Deployment Evaluation","version":1},"reference_index":13,"source":"pdf_text","source_observed_at":"2026-08-10T04:37:16.990397Z"},"links":{"citing_paper":"/paper/2501.17749"},"observation_digest":"sha256:5ba882154565e5ca2df1f562676cd656a6f144790ed85ab36cc7591d0bd95916","observation_id":"d8295da2-1dd8-43a7-b7fa-3d9105a14fa2","resolution":{"observed_at":"2026-08-10T04:37:17.267036Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-19T06:32:44.657259+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-19T06:32:44.657259+00:00","source":"crossref"},{"observed_at":"2026-08-19T06:32:39.956319+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-10T04:37:17.251337Z","title":"Artiﬁcial Intelligence Act (Regulation (EU) 2024/16 89), Ofﬁcial Journal version of 13 June 2024","venue":null,"work_id":"83f01c9e-5ed4-401c-a1c9-f47d3735121d","year":2024},"citing_paper":{"arxiv_id":"2501.17749","last_updated":"2025-01-29T16:36:53Z","snapshot_observed_at":"2026-08-19T02:28:54.940509Z","submitted_at":"2025-01-29T16:36:53Z","title":"Early External Safety Testing of OpenAI's o3-mini: Insights from the Pre-Deployment Evaluation","version":1},"reference_index":14,"source":"pdf_text","source_observed_at":"2026-08-10T04:37:16.993873Z"},"links":{"citing_paper":"/paper/2501.17749"},"observation_digest":"sha256:f25dfaf7ed1875171dcc62dc1f937669d714ac6682ef93f918fc4265beea1247","observation_id":"f86c79da-9169-4d0b-8b5b-d75b3355f507","resolution":{"observed_at":"2026-08-10T04:37:17.255058Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-19T06:32:44.657259+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-19T06:32:44.657259+00:00","source":"crossref"},{"observed_at":"2026-08-19T06:32:39.956319+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2312.06674","last_updated":"2023-12-07T19:40:50Z","snapshot_observed_at":"2026-08-14T15:42:19.849118Z","submitted_at":"2023-12-07T19:40:50Z","title":"Llama Guard: LLM-based Input-Output Safeguard for Human-AI Conversations","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2312.06674","snapshot_observed_at":"2026-08-10T04:37:16.998077Z","title":"Llama guard: Llm-based input-output safeguard for human- ai conversations,","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2501.17749","last_updated":"2025-01-29T16:36:53Z","snapshot_observed_at":"2026-08-19T02:28:54.940509Z","submitted_at":"2025-01-29T16:36:53Z","title":"Early External Safety Testing of OpenAI's o3-mini: Insights from the Pre-Deployment Evaluation","version":1},"reference_index":15,"source":"pdf_text","source_observed_at":"2026-08-10T04:37:16.998077Z"},"links":{"cited_paper":"/paper/2312.06674","citing_paper":"/paper/2501.17749"},"observation_digest":"sha256:7dcfad751d63d5d75607ee1c1e497e1746c05bbb5b7874037c313e564402f3c6","observation_id":"1cb77ed6-af48-4b3f-a363-3dafa0f2183a","resolution":{"observed_at":"2026-08-10T04:37:16.998077Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2402.16444","last_updated":"2024-11-05T02:13:59Z","snapshot_observed_at":"2026-08-18T03:58:02.605905Z","submitted_at":"2024-02-26T09:43:02Z","title":"ShieldLM: Empowering LLMs as Aligned, Customizable and Explainable Safety Detectors","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2402.16444","snapshot_observed_at":"2026-08-10T04:37:17.002476Z","title":"ShieldLM: Empowering LLMs as aligned, customizable and explainable safety detec tors,","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2501.17749","last_updated":"2025-01-29T16:36:53Z","snapshot_observed_at":"2026-08-19T02:28:54.940509Z","submitted_at":"2025-01-29T16:36:53Z","title":"Early External Safety Testing of OpenAI's o3-mini: Insights from the Pre-Deployment Evaluation","version":1},"reference_index":16,"source":"pdf_text","source_observed_at":"2026-08-10T04:37:17.002476Z"},"links":{"cited_paper":"/paper/2402.16444","citing_paper":"/paper/2501.17749"},"observation_digest":"sha256:c343c4147c3c0e49b2c6d229c70028686b33ba92f38a902d4464b1b556c7b6d8","observation_id":"5ab876b4-fd61-4419-aeff-ba967d05e0b7","resolution":{"observed_at":"2026-08-10T04:37:17.002476Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2402.10260","last_updated":"2024-08-27T03:32:47Z","snapshot_observed_at":"2026-08-04T21:32:35.483431Z","submitted_at":"2024-02-15T18:58:09Z","title":"A StrongREJECT for Empty Jailbreaks","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2402.10260","snapshot_observed_at":"2026-08-10T04:37:17.006505Z","title":"A strongreject for empty jailbreaks,","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2501.17749","last_updated":"2025-01-29T16:36:53Z","snapshot_observed_at":"2026-08-19T02:28:54.940509Z","submitted_at":"2025-01-29T16:36:53Z","title":"Early External Safety Testing of OpenAI's o3-mini: Insights from the Pre-Deployment Evaluation","version":1},"reference_index":17,"source":"pdf_text","source_observed_at":"2026-08-10T04:37:17.006505Z"},"links":{"cited_paper":"/paper/2402.10260","citing_paper":"/paper/2501.17749"},"observation_digest":"sha256:a5024135f7083f1f434101a5f3a95598890d0bbdedc5324a25380fedea3607ab","observation_id":"32549227-2140-4849-890a-5b7dd32a6123","resolution":{"observed_at":"2026-08-10T04:37:17.006505Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2209.07858","last_updated":"2022-11-22T19:12:57Z","snapshot_observed_at":"2026-08-17T08:35:42.452149Z","submitted_at":"2022-08-23T23:37:14Z","title":"Red Teaming Language Models to Reduce Harms: Methods, Scaling Behaviors, and Lessons Learned","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2209.07858","snapshot_observed_at":"2026-08-10T04:37:17.010084Z","title":"Red teaming language models to reduce harms: Methods, sca ling behaviors, and lessons learned,","venue":null,"work_id":null,"year":2022},"citing_paper":{"arxiv_id":"2501.17749","last_updated":"2025-01-29T16:36:53Z","snapshot_observed_at":"2026-08-19T02:28:54.940509Z","submitted_at":"2025-01-29T16:36:53Z","title":"Early External Safety Testing of OpenAI's o3-mini: Insights from the Pre-Deployment Evaluation","version":1},"reference_index":18,"source":"pdf_text","source_observed_at":"2026-08-10T04:37:17.010084Z"},"links":{"cited_paper":"/paper/2209.07858","citing_paper":"/paper/2501.17749"},"observation_digest":"sha256:aea2005ee38cf7f6e40473d0bc4cfa3e5140ed7be4be1f912dafdff3692b84ac","observation_id":"0797e6cc-8eec-43f3-af8c-bedab0117b73","resolution":{"observed_at":"2026-08-10T04:37:17.010084Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2310.06987","last_updated":"2023-10-10T20:15:54Z","snapshot_observed_at":"2026-08-15T18:47:46.763989Z","submitted_at":"2023-10-10T20:15:54Z","title":"Catastrophic Jailbreak of Open-source LLMs via Exploiting Generation","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2310.06987","snapshot_observed_at":"2026-08-10T04:37:17.014284Z","title":"Catastrop hic jailbreak of open-source llms via exploiting generation,","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2501.17749","last_updated":"2025-01-29T16:36:53Z","snapshot_observed_at":"2026-08-19T02:28:54.940509Z","submitted_at":"2025-01-29T16:36:53Z","title":"Early External Safety Testing of OpenAI's o3-mini: Insights from the Pre-Deployment Evaluation","version":1},"reference_index":19,"source":"pdf_text","source_observed_at":"2026-08-10T04:37:17.014284Z"},"links":{"cited_paper":"/paper/2310.06987","citing_paper":"/paper/2501.17749"},"observation_digest":"sha256:59bbd8322833cf09083c686e652fc394fc98362179eeb8185b4807ab1779d375","observation_id":"90df0369-a45a-42dc-a089-0ad59a44f7b3","resolution":{"observed_at":"2026-08-10T04:37:17.014284Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2307.15043","last_updated":"2023-12-20T20:48:57Z","snapshot_observed_at":"2026-08-12T09:06:50.363435Z","submitted_at":"2023-07-27T17:49:12Z","title":"Universal and Transferable Adversarial Attacks on Aligned Language Models","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2307.15043","snapshot_observed_at":"2026-08-10T04:37:17.018055Z","title":"Universal and transferable adversarial attacks on aligned language models,","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2501.17749","last_updated":"2025-01-29T16:36:53Z","snapshot_observed_at":"2026-08-19T02:28:54.940509Z","submitted_at":"2025-01-29T16:36:53Z","title":"Early External Safety Testing of OpenAI's o3-mini: Insights from the Pre-Deployment Evaluation","version":1},"reference_index":20,"source":"pdf_text","source_observed_at":"2026-08-10T04:37:17.018055Z"},"links":{"cited_paper":"/paper/2307.15043","citing_paper":"/paper/2501.17749"},"observation_digest":"sha256:541cf82cbe8121848a34da4c3ee6da8a629def285b3d2b5c253ddf99dd07b36d","observation_id":"afba7695-fc6d-448e-bfff-2b153789b5f7","resolution":{"observed_at":"2026-08-10T04:37:17.018055Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2402.04249","last_updated":"2024-02-27T04:43:08Z","snapshot_observed_at":"2026-08-16T09:07:20.265665Z","submitted_at":"2024-02-06T18:59:08Z","title":"HarmBench: A Standardized Evaluation Framework for Automated Red Teaming and Robust Refusal","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2402.04249","snapshot_observed_at":"2026-08-10T04:37:17.021691Z","title":"Harm- bench: A standardized evaluation framework for automated r ed teaming and robust refusal,","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2501.17749","last_updated":"2025-01-29T16:36:53Z","snapshot_observed_at":"2026-08-19T02:28:54.940509Z","submitted_at":"2025-01-29T16:36:53Z","title":"Early External Safety Testing of OpenAI's o3-mini: Insights from the Pre-Deployment Evaluation","version":1},"reference_index":21,"source":"pdf_text","source_observed_at":"2026-08-10T04:37:17.021691Z"},"links":{"cited_paper":"/paper/2402.04249","citing_paper":"/paper/2501.17749"},"observation_digest":"sha256:9084aa6e8fe3b420d0cee69f45ad15abec363cbf05eaf71963fdd8f34241d86c","observation_id":"38e85f31-01f8-42c7-bad5-766efcf9985a","resolution":{"observed_at":"2026-08-10T04:37:17.021691Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2308.03825","last_updated":"2024-05-15T12:06:31Z","snapshot_observed_at":"2026-08-16T09:32:10.049300Z","submitted_at":"2023-08-07T16:55:20Z","title":"\"Do Anything Now\": Characterizing and Evaluating In-The-Wild Jailbreak Prompts on Large Language Models","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2308.03825","snapshot_observed_at":"2026-08-10T04:37:17.025612Z","title":"\" do an ything now","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2501.17749","last_updated":"2025-01-29T16:36:53Z","snapshot_observed_at":"2026-08-19T02:28:54.940509Z","submitted_at":"2025-01-29T16:36:53Z","title":"Early External Safety Testing of OpenAI's o3-mini: Insights from the Pre-Deployment Evaluation","version":1},"reference_index":22,"source":"pdf_text","source_observed_at":"2026-08-10T04:37:17.025612Z"},"links":{"cited_paper":"/paper/2308.03825","citing_paper":"/paper/2501.17749"},"observation_digest":"sha256:3bf31b1ded4db3b8d9927fc0146ff092796ff2bce3195b6a402992583accb75b","observation_id":"1fd88780-7f4e-4dc4-a69b-db7ef0a9d757","resolution":{"observed_at":"2026-08-10T04:37:17.025612Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-10T04:37:17.237698Z","title":"Jailbroken: H ow does llm safety training fail?,","venue":null,"work_id":"e4230fea-74fb-480d-b9c8-77fdb1eb322f","year":2024},"citing_paper":{"arxiv_id":"2501.17749","last_updated":"2025-01-29T16:36:53Z","snapshot_observed_at":"2026-08-19T02:28:54.940509Z","submitted_at":"2025-01-29T16:36:53Z","title":"Early External Safety Testing of OpenAI's o3-mini: Insights from the Pre-Deployment Evaluation","version":1},"reference_index":23,"source":"pdf_text","source_observed_at":"2026-08-10T04:37:17.029662Z"},"links":{"citing_paper":"/paper/2501.17749"},"observation_digest":"sha256:f8eb80380dcd593b5bbbc2ae751919a7db81bd4634eb3ca8c601df1fd7245cb3","observation_id":"befa1f1a-7b65-4894-9e48-d125dd22d8da","resolution":{"observed_at":"2026-08-10T04:37:17.243309Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-19T06:32:44.657259+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-19T06:32:44.657259+00:00","source":"crossref"},{"observed_at":"2026-08-19T06:32:39.956319+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2408.03837","last_updated":"2024-08-19T21:57:24Z","snapshot_observed_at":"2026-08-16T13:27:56.095333Z","submitted_at":"2024-08-07T15:22:44Z","title":"WalledEval: A Comprehensive Safety Evaluation Toolkit for Large Language Models","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2408.03837","snapshot_observed_at":"2026-08-10T04:37:17.033434Z","title":"Walledeval: A comprehensive safety evaluation toolkit f or large language models,","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2501.17749","last_updated":"2025-01-29T16:36:53Z","snapshot_observed_at":"2026-08-19T02:28:54.940509Z","submitted_at":"2025-01-29T16:36:53Z","title":"Early External Safety Testing of OpenAI's o3-mini: Insights from the Pre-Deployment Evaluation","version":1},"reference_index":24,"source":"pdf_text","source_observed_at":"2026-08-10T04:37:17.033434Z"},"links":{"cited_paper":"/paper/2408.03837","citing_paper":"/paper/2501.17749"},"observation_digest":"sha256:d01f3e7b81d3faea7c64a7d0ed8a90f01e8ef20f0627f91d75abe71a47dad2e0","observation_id":"f68c3631-350a-45ac-90f4-496ba9779108","resolution":{"observed_at":"2026-08-10T04:37:17.033434Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2408.15313","last_updated":"2025-04-08T11:04:33Z","snapshot_observed_at":"2026-08-16T13:23:25.714111Z","submitted_at":"2024-08-27T17:31:21Z","title":"Bi-Factorial Preference Optimization: Balancing Safety-Helpfulness in Language Models","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2408.15313","snapshot_observed_at":"2026-08-10T04:37:17.038006Z","title":"Bi-fact orial preference optimization: Balancing safety- helpfulness in language models,","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2501.17749","last_updated":"2025-01-29T16:36:53Z","snapshot_observed_at":"2026-08-19T02:28:54.940509Z","submitted_at":"2025-01-29T16:36:53Z","title":"Early External Safety Testing of OpenAI's o3-mini: Insights from the Pre-Deployment Evaluation","version":1},"reference_index":25,"source":"pdf_text","source_observed_at":"2026-08-10T04:37:17.038006Z"},"links":{"cited_paper":"/paper/2408.15313","citing_paper":"/paper/2501.17749"},"observation_digest":"sha256:35a1ee13ed378f038701b976f5a7cab5d2bd790f6f3fbe250cc907524d4a3cf9","observation_id":"7d4fb587-da85-4b7b-ac9f-b386a5e3ed18","resolution":{"observed_at":"2026-08-10T04:37:17.038006Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}}],"paper":{"arxiv_id":"2501.17749","last_updated":"2025-01-29T16:36:53Z","latest_version":1,"primary_category":"cs.SE","snapshot_observed_at":"2026-08-19T02:28:54.940509Z","submitted_at":"2025-01-29T16:36:53Z","title":"Early External Safety Testing of OpenAI's o3-mini: Insights from the Pre-Deployment Evaluation"},"reference_resolution":{"displayed":25,"state_counts":{"malformed_identifier":0,"metadata_mismatch":0,"parse_uncertain":0,"unresolved":18,"verified_exact":0,"verified_fuzzy":7},"total_outbound_references":25},"refusal":"A citation records a reference. It does not transfer a finding from one paper to another.","schema":"pith.paper-citation-record.v1","standing_sources":[{"observed_at":"2026-08-19T06:32:44.657259+00:00","source":"crossref"},{"observed_at":"2026-08-19T06:32:39.956319+00:00","source":"retraction_watch"}],"thesis":"As of 19 August 2026, this Paper Citation Record lists 25 of 25 outbound references and 2 inbound Pith citation observations for arXiv:2501.17749."}