{"as_of":"2026-07-29T06:54:00Z","caps":{"database_statements":6,"inbound":100,"outbound":100},"context_digest":"sha256:41aa502c1ea7e5ba63259286408f9c336592b982af2633c75381752b983e8002","coverage":[{"denominator":0,"lane":"reference_resolution","note":"Typed states for the displayed outbound observations.","records_observed":0,"source":"paper_references, paper_reference_links","source_observed_at":null,"state":"measured"},{"denominator":8,"lane":"standing_notices","note":"One-hop event checks from named stored sources.","records_observed":8,"source":"scholarly_work_events, retraction_status_cache","source_observed_at":"2026-07-28T06:31:03.373048+00:00","state":"measured"},{"denominator":8,"lane":"inbound_itemization","note":"Pith citing papers itemized under the disclosed page cap.","records_observed":8,"source":"paper_references, paper_reference_links","source_observed_at":"2026-06-29T17:34:22.676341Z","state":"measured"},{"denominator":1,"lane":"external_citation_measurements","note":"A source-named dated measurement, never combined with another source.","records_observed":0,"source":"arxiv_reference","source_observed_at":"2026-07-04T19:40:06.723481Z","state":"measured"}],"external_citation_measurements":[],"inbound":[{"citation":{"cited_paper":{"arxiv_id":"2507.19672","last_updated":"2025-07-25T20:52:58Z","snapshot_observed_at":"2026-07-06T22:03:05.943334Z","submitted_at":"2025-07-25T20:52:58Z","title":"Alignment and Safety in Large Language Models: Safety Mechanisms, Training Paradigms, and Emerging Challenges","version":1},"cited_work":{"arxiv_id":"2507.19672","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2507.19672","snapshot_observed_at":"2026-07-04T19:40:06.723481Z","title":"Alignment and safety in large language models: Safety mechanisms, training paradigms, and emerging challenges.arXiv preprint arXiv:2507.19672","venue":null,"work_id":"517eb84e-95d6-41e8-98e4-900bab886e19","year":2025},"citing_paper":{"arxiv_id":"2509.22510","last_updated":"2026-05-15T12:53:06Z","snapshot_observed_at":"2026-07-06T22:30:55.313733Z","submitted_at":"2025-09-26T15:52:21Z","title":"We Think, Therefore We Align LLMs to Helpful, Harmless and Honest Before They Go Wrong","version":3},"reference_index":22,"source":"arxiv_source","source_observed_at":"2026-05-21T22:11:20.651761Z"},"links":{"cited_paper":"/paper/2507.19672","citing_paper":"/paper/2509.22510"},"observation_digest":"sha256:c1d813313dea2bdf5749b9432af68c8d09fc20acdc77167c954799aa7b9c286c","observation_id":"f4e43ee4-2ac9-4299-aba7-52d612727faa","resolution":{"observed_at":"2026-05-21T22:14:23.539591Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-07-28T06:31:03.373048+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-07-28T06:31:03.373048+00:00","source":"crossref"},{"observed_at":"2026-07-28T06:30:57.601408+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2507.19672","last_updated":"2025-07-25T20:52:58Z","snapshot_observed_at":"2026-07-06T22:03:05.943334Z","submitted_at":"2025-07-25T20:52:58Z","title":"Alignment and Safety in Large Language Models: Safety Mechanisms, Training Paradigms, and Emerging Challenges","version":1},"cited_work":{"arxiv_id":"2507.19672","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2507.19672","snapshot_observed_at":"2026-07-04T19:40:06.723481Z","title":"Alignment and safety in large language models: Safety mechanisms, training paradigms, and emerging challenges.arXiv preprint arXiv:2507.19672","venue":null,"work_id":"517eb84e-95d6-41e8-98e4-900bab886e19","year":2025},"citing_paper":{"arxiv_id":"2602.07340","last_updated":"2026-05-21T07:39:41Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2026-02-07T03:46:33Z","title":"Revisiting Robustness for LLM Safety Alignment via Selective Geometry Control","version":2},"reference_index":16,"source":"pdf_text","source_observed_at":"2026-05-22T11:17:03.104902Z"},"links":{"cited_paper":"/paper/2507.19672","citing_paper":"/paper/2602.07340"},"observation_digest":"sha256:1f087b0c2d20924d4267d5d7cf6bc7f9082667b826a7818e4722780a4328dedb","observation_id":"5f5d3065-24a6-4d14-a956-fc6c4b12fee4","resolution":{"observed_at":"2026-05-22T11:21:29.042697Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-07-28T06:31:03.373048+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-07-28T06:31:03.373048+00:00","source":"crossref"},{"observed_at":"2026-07-28T06:30:57.601408+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2507.19672","last_updated":"2025-07-25T20:52:58Z","snapshot_observed_at":"2026-07-06T22:03:05.943334Z","submitted_at":"2025-07-25T20:52:58Z","title":"Alignment and Safety in Large Language Models: Safety Mechanisms, Training Paradigms, and Emerging Challenges","version":1},"cited_work":{"arxiv_id":"2507.19672","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2507.19672","snapshot_observed_at":"2026-07-04T19:40:06.723481Z","title":"Alignment and safety in large language models: Safety mechanisms, training paradigms, and emerging challenges.arXiv preprint arXiv:2507.19672","venue":null,"work_id":"517eb84e-95d6-41e8-98e4-900bab886e19","year":2025},"citing_paper":{"arxiv_id":"2604.01346","last_updated":"2026-04-06T19:07:02Z","snapshot_observed_at":"2026-07-06T22:51:31.209896Z","submitted_at":"2026-04-01T19:57:33Z","title":"Safety, Security, and Cognitive Risks in World Models","version":2},"reference_index":66,"source":"pdf_text","source_observed_at":"2026-05-13T22:35:46.126714Z"},"links":{"cited_paper":"/paper/2507.19672","citing_paper":"/paper/2604.01346"},"observation_digest":"sha256:ad22d32f530d95f2bda5c11d8062b2faab890bd04495c691c4588a9234c24307","observation_id":"3d1ce0aa-2a3e-49cc-b456-ca80140266be","resolution":{"observed_at":"2026-05-13T22:38:22.232235Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-07-28T06:31:03.373048+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-07-28T06:31:03.373048+00:00","source":"crossref"},{"observed_at":"2026-07-28T06:30:57.601408+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2507.19672","last_updated":"2025-07-25T20:52:58Z","snapshot_observed_at":"2026-07-06T22:03:05.943334Z","submitted_at":"2025-07-25T20:52:58Z","title":"Alignment and Safety in Large Language Models: Safety Mechanisms, Training Paradigms, and Emerging Challenges","version":1},"cited_work":{"arxiv_id":"2507.19672","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2507.19672","snapshot_observed_at":"2026-07-04T19:40:06.723481Z","title":"Alignment and safety in large language models: Safety mechanisms, training paradigms, and emerging challenges.arXiv preprint arXiv:2507.19672","venue":null,"work_id":"517eb84e-95d6-41e8-98e4-900bab886e19","year":2025},"citing_paper":{"arxiv_id":"2604.01444","last_updated":"2026-04-03T15:46:12Z","snapshot_observed_at":"2026-07-06T22:51:35.522923Z","submitted_at":"2026-04-01T22:38:38Z","title":"Cooking Up Risks: Benchmarking and Reducing Food Safety Risks in Large Language Models","version":2},"reference_index":16,"source":"pdf_text","source_observed_at":"2026-05-13T22:02:12.555644Z"},"links":{"cited_paper":"/paper/2507.19672","citing_paper":"/paper/2604.01444"},"observation_digest":"sha256:c75831db72a75e23bd54195d0225d247ccbd0d4173d27d50841af527762954a1","observation_id":"226f53af-b87c-4feb-b09d-03645da07d83","resolution":{"observed_at":"2026-05-13T22:03:20.355783Z","resolver_source":"arxiv_id","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-07-28T06:31:03.373048+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-07-28T06:31:03.373048+00:00","source":"crossref"},{"observed_at":"2026-07-28T06:30:57.601408+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2507.19672","last_updated":"2025-07-25T20:52:58Z","snapshot_observed_at":"2026-07-06T22:03:05.943334Z","submitted_at":"2025-07-25T20:52:58Z","title":"Alignment and Safety in Large Language Models: Safety Mechanisms, Training Paradigms, and Emerging Challenges","version":1},"cited_work":{"arxiv_id":"2507.19672","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2507.19672","snapshot_observed_at":"2026-07-04T19:40:06.723481Z","title":"Alignment and safety in large language models: Safety mechanisms, training paradigms, and emerging challenges.arXiv preprint arXiv:2507.19672","venue":null,"work_id":"517eb84e-95d6-41e8-98e4-900bab886e19","year":2025},"citing_paper":{"arxiv_id":"2605.27110","last_updated":"2026-05-26T14:51:13Z","snapshot_observed_at":"2026-07-06T23:36:53.330986Z","submitted_at":"2026-05-26T14:51:13Z","title":"BAIT: Boundary-Guided Disclosure Escalation via Self-Conditioned Reasoning","version":1},"reference_index":2,"source":"pdf_text","source_observed_at":"2026-06-29T17:34:22.676341Z"},"links":{"cited_paper":"/paper/2507.19672","citing_paper":"/paper/2605.27110"},"observation_digest":"sha256:2a52399b224f7c48a5169f8342f62973b788d5cede936a0ce03d91a1a3df5bae","observation_id":"2fcd960c-d01d-4acd-bf48-ae0a30c6c98e","resolution":{"observed_at":"2026-06-29T18:03:48.662941Z","resolver_source":"arxiv_id","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-07-28T06:31:03.373048+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-07-28T06:31:03.373048+00:00","source":"crossref"},{"observed_at":"2026-07-28T06:30:57.601408+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2507.19672","last_updated":"2025-07-25T20:52:58Z","snapshot_observed_at":"2026-07-06T22:03:05.943334Z","submitted_at":"2025-07-25T20:52:58Z","title":"Alignment and Safety in Large Language Models: Safety Mechanisms, Training Paradigms, and Emerging Challenges","version":1},"cited_work":{"arxiv_id":"2507.19672","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2507.19672","snapshot_observed_at":"2026-07-04T19:40:06.723481Z","title":"Alignment and safety in large language models: Safety mechanisms, training paradigms, and emerging challenges.arXiv preprint arXiv:2507.19672","venue":null,"work_id":"517eb84e-95d6-41e8-98e4-900bab886e19","year":2025},"citing_paper":{"arxiv_id":"2606.07678","last_updated":"2026-06-04T20:23:23Z","snapshot_observed_at":"2026-07-06T23:47:19.773490Z","submitted_at":"2026-06-04T20:23:23Z","title":"DOG-DPO:Dynamic Optimization in Geometry for Safety Alignment","version":1},"reference_index":5,"source":"pdf_text","source_observed_at":"2026-06-28T02:28:54.243385Z"},"links":{"cited_paper":"/paper/2507.19672","citing_paper":"/paper/2606.07678"},"observation_digest":"sha256:38cfeed58b84cfa39b388bcab76646f831308a89040158cbe9e6e09472ec6718","observation_id":"1d6b9f48-f207-4170-af72-b14f4c133553","resolution":{"observed_at":"2026-07-02T12:06:56.036195Z","resolver_source":"arxiv_id","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-07-28T06:31:03.373048+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-07-28T06:31:03.373048+00:00","source":"crossref"},{"observed_at":"2026-07-28T06:30:57.601408+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2507.19672","last_updated":"2025-07-25T20:52:58Z","snapshot_observed_at":"2026-07-06T22:03:05.943334Z","submitted_at":"2025-07-25T20:52:58Z","title":"Alignment and Safety in Large Language Models: Safety Mechanisms, Training Paradigms, and Emerging Challenges","version":1},"cited_work":{"arxiv_id":"2507.19672","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2507.19672","snapshot_observed_at":"2026-07-04T19:40:06.723481Z","title":"Alignment and safety in large language models: Safety mechanisms, training paradigms, and emerging challenges.arXiv preprint arXiv:2507.19672","venue":null,"work_id":"517eb84e-95d6-41e8-98e4-900bab886e19","year":2025},"citing_paper":{"arxiv_id":"2606.08044","last_updated":"2026-06-06T08:10:56Z","snapshot_observed_at":"2026-07-06T23:47:33.680573Z","submitted_at":"2026-06-06T08:10:56Z","title":"When Behavioral Safety Evaluation Fails: A Representation-Level Perspective","version":1},"reference_index":38,"source":"arxiv_source","source_observed_at":"2026-06-27T20:04:17.744876Z"},"links":{"cited_paper":"/paper/2507.19672","citing_paper":"/paper/2606.08044"},"observation_digest":"sha256:f247920fc5c9766b9581d2110b7de3feac86505dd02bc9e123f2f442d966977d","observation_id":"ff803701-5de4-4afb-8251-2b5470c70872","resolution":{"observed_at":"2026-07-02T20:57:23.060356Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-07-28T06:31:03.373048+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-07-28T06:31:03.373048+00:00","source":"crossref"},{"observed_at":"2026-07-28T06:30:57.601408+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2507.19672","last_updated":"2025-07-25T20:52:58Z","snapshot_observed_at":"2026-07-06T22:03:05.943334Z","submitted_at":"2025-07-25T20:52:58Z","title":"Alignment and Safety in Large Language Models: Safety Mechanisms, Training Paradigms, and Emerging Challenges","version":1},"cited_work":{"arxiv_id":"2507.19672","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2507.19672","snapshot_observed_at":"2026-07-04T19:40:06.723481Z","title":"Alignment and safety in large language models: Safety mechanisms, training paradigms, and emerging challenges.arXiv preprint arXiv:2507.19672","venue":null,"work_id":"517eb84e-95d6-41e8-98e4-900bab886e19","year":2025},"citing_paper":{"arxiv_id":"2606.25442","last_updated":"2026-06-24T06:10:33Z","snapshot_observed_at":"2026-07-06T23:59:52.373272Z","submitted_at":"2026-06-24T06:10:33Z","title":"PolicyAlign: Direct Policy-Based Safety Alignment for Large Language Models","version":1},"reference_index":2,"source":"arxiv_source","source_observed_at":"2026-06-25T21:09:19.727723Z"},"links":{"cited_paper":"/paper/2507.19672","citing_paper":"/paper/2606.25442"},"observation_digest":"sha256:df43bb7be9a742405e6869783280109df5fbc7e4be19dc3f5aa7c704422ce36e","observation_id":"38dfe6e5-ba1c-4977-afac-63d06655746c","resolution":{"observed_at":"2026-07-04T19:40:06.724811Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-07-28T06:31:03.373048+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-07-28T06:31:03.373048+00:00","source":"crossref"},{"observed_at":"2026-07-28T06:30:57.601408+00:00","source":"retraction_watch"}],"state":"measured"}}],"links":{"evidence":"/evidence","html":"/paper/2507.19672/citation-record","integrity":"/paper/2507.19672/integrity","json":"/paper/2507.19672/citation-record.json","paper":"/paper/2507.19672"},"outbound":[],"paper":{"arxiv_id":"2507.19672","last_updated":"2025-07-25T20:52:58Z","latest_version":1,"primary_category":"cs.AI","snapshot_observed_at":"2026-07-06T22:03:05.943334Z","submitted_at":"2025-07-25T20:52:58Z","title":"Alignment and Safety in Large Language Models: Safety Mechanisms, Training Paradigms, and Emerging Challenges"},"reference_resolution":{"displayed":0,"state_counts":{"malformed_identifier":0,"metadata_mismatch":0,"parse_uncertain":0,"unresolved":0,"verified_exact":0,"verified_fuzzy":0},"total_outbound_references":0},"refusal":"A citation records a reference. It does not transfer a finding from one paper to another.","schema":"pith.paper-citation-record.v1","standing_sources":[{"observed_at":"2026-07-28T06:31:03.373048+00:00","source":"crossref"},{"observed_at":"2026-07-28T06:30:57.601408+00:00","source":"retraction_watch"}],"thesis":"As of 29 July 2026, this Paper Citation Record lists 0 of 0 outbound references and 8 inbound Pith citation observations for arXiv:2507.19672."}