{"as_of":"2026-08-16T11:16:00Z","caps":{"database_statements":6,"inbound":100,"outbound":100},"context_digest":"sha256:e318b2add34d5656bcf6872e2a22c37fd237d19d71ed451fcb05015b65be89cc","coverage":[{"denominator":60,"lane":"reference_resolution","note":"Typed states for the displayed outbound observations.","records_observed":60,"source":"paper_references, paper_reference_links","source_observed_at":"2026-08-15T18:13:58.320855Z","state":"measured"},{"denominator":60,"lane":"standing_notices","note":"One-hop event checks from named stored sources.","records_observed":60,"source":"scholarly_work_events, retraction_status_cache","source_observed_at":"2026-08-16T06:30:59.297886+00:00","state":"measured"},{"denominator":0,"lane":"inbound_itemization","note":"Pith citing papers itemized under the disclosed page cap.","records_observed":0,"source":"paper_references, paper_reference_links","source_observed_at":null,"state":"measured"},{"denominator":1,"lane":"external_citation_measurements","note":"A source-named dated measurement, never combined with another source.","records_observed":0,"source":"cited_works","source_observed_at":null,"state":"measured"}],"external_citation_measurements":[],"inbound":[],"links":{"evidence":"/evidence","html":"/paper/2507.18631/citation-record","integrity":"/paper/2507.18631/integrity","json":"/paper/2507.18631/citation-record.json","paper":"/paper/2507.18631"},"outbound":[{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-15T18:13:58.063350Z","title":"online\" 'onlinestring :=","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2507.18631","last_updated":"2025-07-25T07:20:24Z","snapshot_observed_at":"2026-08-15T18:07:02.667074Z","submitted_at":"2025-07-24T17:59:24Z","title":"Layer-Aware Representation Filtering: Purifying Finetuning Data to Preserve LLM Safety Alignment","version":2},"reference_index":1,"source":"arxiv_source","source_observed_at":"2026-08-15T18:13:58.063350Z"},"links":{"citing_paper":"/paper/2507.18631"},"observation_digest":"sha256:ddc06d20840b135c1c3afd3322f012970be32d34408463953b60bd76d2db4949","observation_id":"853d5b88-5b37-4838-8b8d-3928eb6c2755","resolution":{"observed_at":"2026-08-15T18:13:58.063350Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-15T18:13:58.068887Z","title":"write newline","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2507.18631","last_updated":"2025-07-25T07:20:24Z","snapshot_observed_at":"2026-08-15T18:07:02.667074Z","submitted_at":"2025-07-24T17:59:24Z","title":"Layer-Aware Representation Filtering: Purifying Finetuning Data to Preserve LLM Safety Alignment","version":2},"reference_index":2,"source":"arxiv_source","source_observed_at":"2026-08-15T18:13:58.068887Z"},"links":{"citing_paper":"/paper/2507.18631"},"observation_digest":"sha256:4bca881a8c1e7919432f304a95875b591bf2a1ed8102093226d63ba3c573763c","observation_id":"c4ea7ce4-9aba-42c9-b50b-17e999d18e22","resolution":{"observed_at":"2026-08-15T18:13:58.068887Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2404.14219","last_updated":"2024-08-30T21:17:17Z","snapshot_observed_at":"2026-08-10T14:07:02.234322Z","submitted_at":"2024-04-22T14:32:33Z","title":"Phi-3 Technical Report: A Highly Capable Language Model Locally on Your Phone","version":4},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2404.14219","snapshot_observed_at":"2026-08-15T18:13:58.073741Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2507.18631","last_updated":"2025-07-25T07:20:24Z","snapshot_observed_at":"2026-08-15T18:07:02.667074Z","submitted_at":"2025-07-24T17:59:24Z","title":"Layer-Aware Representation Filtering: Purifying Finetuning Data to Preserve LLM Safety Alignment","version":2},"reference_index":3,"source":"arxiv_source","source_observed_at":"2026-08-15T18:13:58.073741Z"},"links":{"cited_paper":"/paper/2404.14219","citing_paper":"/paper/2507.18631"},"observation_digest":"sha256:58d165d37a700b987ec2fa8367801b922343d09bc51adee327d1fa89041ce17b","observation_id":"f8963a37-f4d4-4fe0-a1fe-642bb77782a8","resolution":{"observed_at":"2026-08-15T18:13:58.073741Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-15T18:13:58.079059Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2507.18631","last_updated":"2025-07-25T07:20:24Z","snapshot_observed_at":"2026-08-15T18:07:02.667074Z","submitted_at":"2025-07-24T17:59:24Z","title":"Layer-Aware Representation Filtering: Purifying Finetuning Data to Preserve LLM Safety Alignment","version":2},"reference_index":4,"source":"arxiv_source","source_observed_at":"2026-08-15T18:13:58.079059Z"},"links":{"citing_paper":"/paper/2507.18631"},"observation_digest":"sha256:4cfbad2e82513dca9f83fce12e3b9375b9625f7ad2339d0022a05fb1f93458aa","observation_id":"0feb4913-2f92-4f62-b2a4-f21cdc343e30","resolution":{"observed_at":"2026-08-15T18:13:58.079059Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2107.03374","last_updated":"2021-07-14T17:16:02Z","snapshot_observed_at":"2026-08-08T11:58:24.516369Z","submitted_at":"2021-07-07T17:41:24Z","title":"Evaluating Large Language Models Trained on Code","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2107.03374","snapshot_observed_at":"2026-08-15T18:13:58.083669Z","title":null,"venue":null,"work_id":null,"year":2021},"citing_paper":{"arxiv_id":"2507.18631","last_updated":"2025-07-25T07:20:24Z","snapshot_observed_at":"2026-08-15T18:07:02.667074Z","submitted_at":"2025-07-24T17:59:24Z","title":"Layer-Aware Representation Filtering: Purifying Finetuning Data to Preserve LLM Safety Alignment","version":2},"reference_index":5,"source":"arxiv_source","source_observed_at":"2026-08-15T18:13:58.083669Z"},"links":{"cited_paper":"/paper/2107.03374","citing_paper":"/paper/2507.18631"},"observation_digest":"sha256:a352544b4c2c7ec5dd8306fc4a0222b63c882dd98ceb413ba555698de825c9be","observation_id":"34d717fb-f455-4513-98da-761d3975516e","resolution":{"observed_at":"2026-08-15T18:13:58.083669Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-15T18:13:58.087943Z","title":null,"venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2507.18631","last_updated":"2025-07-25T07:20:24Z","snapshot_observed_at":"2026-08-15T18:07:02.667074Z","submitted_at":"2025-07-24T17:59:24Z","title":"Layer-Aware Representation Filtering: Purifying Finetuning Data to Preserve LLM Safety Alignment","version":2},"reference_index":6,"source":"arxiv_source","source_observed_at":"2026-08-15T18:13:58.087943Z"},"links":{"citing_paper":"/paper/2507.18631"},"observation_digest":"sha256:247c6ef154dfb92520d1ea2dc35cb445e1d953a037945fdfe66d40d7bcbfd981","observation_id":"95f71dd0-a009-47e1-9b09-0a6ea0771e2b","resolution":{"observed_at":"2026-08-15T18:13:58.087943Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-15T18:13:59.347554Z","title":null,"venue":null,"work_id":"b15cb002-dc5c-4520-b4e8-43981505d246","year":2025},"citing_paper":{"arxiv_id":"2507.18631","last_updated":"2025-07-25T07:20:24Z","snapshot_observed_at":"2026-08-15T18:07:02.667074Z","submitted_at":"2025-07-24T17:59:24Z","title":"Layer-Aware Representation Filtering: Purifying Finetuning Data to Preserve LLM Safety Alignment","version":2},"reference_index":7,"source":"arxiv_source","source_observed_at":"2026-08-15T18:13:58.092309Z"},"links":{"citing_paper":"/paper/2507.18631"},"observation_digest":"sha256:847a3120d917ca03ca76f75ff103ad74c1b07bac79edffb1b6ec8292eabd9008","observation_id":"eb39f237-eb76-4b3e-9b9e-f641630cc130","resolution":{"observed_at":"2026-08-15T18:13:59.351858Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-16T06:30:59.297886+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-16T06:30:59.297886+00:00","source":"crossref"},{"observed_at":"2026-08-16T06:30:54.164669+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2410.04524","last_updated":"2025-02-17T02:32:33Z","snapshot_observed_at":"2026-08-13T16:17:41.474017Z","submitted_at":"2024-10-06T15:34:04Z","title":"Toward Secure Tuning: Mitigating Security Risks from Instruction Fine-Tuning","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2410.04524","snapshot_observed_at":"2026-08-15T18:13:58.096730Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2507.18631","last_updated":"2025-07-25T07:20:24Z","snapshot_observed_at":"2026-08-15T18:07:02.667074Z","submitted_at":"2025-07-24T17:59:24Z","title":"Layer-Aware Representation Filtering: Purifying Finetuning Data to Preserve LLM Safety Alignment","version":2},"reference_index":8,"source":"arxiv_source","source_observed_at":"2026-08-15T18:13:58.096730Z"},"links":{"cited_paper":"/paper/2410.04524","citing_paper":"/paper/2507.18631"},"observation_digest":"sha256:54acc6a89f2030edec777a9eae29cc2f911276b37e1a22ec54c3b2d838f5274c","observation_id":"d108daec-9c67-4b2a-95f9-5488ec0d0216","resolution":{"observed_at":"2026-08-15T18:13:58.096730Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2412.19512","last_updated":"2025-08-28T01:13:45Z","snapshot_observed_at":"2026-08-14T04:53:46.476562Z","submitted_at":"2024-12-27T08:03:22Z","title":"Safeguard Fine-Tuned LLMs Through Pre- and Post-Tuning Model Merging","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2412.19512","snapshot_observed_at":"2026-08-15T18:13:58.102071Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2507.18631","last_updated":"2025-07-25T07:20:24Z","snapshot_observed_at":"2026-08-15T18:07:02.667074Z","submitted_at":"2025-07-24T17:59:24Z","title":"Layer-Aware Representation Filtering: Purifying Finetuning Data to Preserve LLM Safety Alignment","version":2},"reference_index":9,"source":"arxiv_source","source_observed_at":"2026-08-15T18:13:58.102071Z"},"links":{"cited_paper":"/paper/2412.19512","citing_paper":"/paper/2507.18631"},"observation_digest":"sha256:59f48f36d8ef537c74f0a95576719577061b551174f2c5ba2a60354526712cd3","observation_id":"b0cb59f1-cc2b-48e8-9ece-a4a3dd1880a4","resolution":{"observed_at":"2026-08-15T18:13:58.102071Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-15T18:13:58.107373Z","title":null,"venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2507.18631","last_updated":"2025-07-25T07:20:24Z","snapshot_observed_at":"2026-08-15T18:07:02.667074Z","submitted_at":"2025-07-24T17:59:24Z","title":"Layer-Aware Representation Filtering: Purifying Finetuning Data to Preserve LLM Safety Alignment","version":2},"reference_index":10,"source":"arxiv_source","source_observed_at":"2026-08-15T18:13:58.107373Z"},"links":{"citing_paper":"/paper/2507.18631"},"observation_digest":"sha256:c9604746fbb9068b4746778477449eb2d2249e56bc8ae0dfff52a58563d64da2","observation_id":"eab81e58-646d-4830-9ef8-64a26358ae24","resolution":{"observed_at":"2026-08-15T18:13:58.107373Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-15T18:13:58.113830Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2507.18631","last_updated":"2025-07-25T07:20:24Z","snapshot_observed_at":"2026-08-15T18:07:02.667074Z","submitted_at":"2025-07-24T17:59:24Z","title":"Layer-Aware Representation Filtering: Purifying Finetuning Data to Preserve LLM Safety Alignment","version":2},"reference_index":11,"source":"arxiv_source","source_observed_at":"2026-08-15T18:13:58.113830Z"},"links":{"citing_paper":"/paper/2507.18631"},"observation_digest":"sha256:be0b98ee1ea002dcc71366a51a219a0f48623dab2ea3acc748fe7655f75095ca","observation_id":"3ec54ef0-dfa6-463a-a861-6c37894113b3","resolution":{"observed_at":"2026-08-15T18:13:58.113830Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-15T18:13:59.326426Z","title":null,"venue":null,"work_id":"b3aa3042-f6c1-4138-a6ae-0ec77698a833","year":2024},"citing_paper":{"arxiv_id":"2507.18631","last_updated":"2025-07-25T07:20:24Z","snapshot_observed_at":"2026-08-15T18:07:02.667074Z","submitted_at":"2025-07-24T17:59:24Z","title":"Layer-Aware Representation Filtering: Purifying Finetuning Data to Preserve LLM Safety Alignment","version":2},"reference_index":12,"source":"arxiv_source","source_observed_at":"2026-08-15T18:13:58.119017Z"},"links":{"citing_paper":"/paper/2507.18631"},"observation_digest":"sha256:319786664ce3a3d9be50098412bc87125057d53f28079d2436ef1a86a8a787e9","observation_id":"21ffb058-7435-4eee-b72a-03113c0aa761","resolution":{"observed_at":"2026-08-15T18:13:59.330442Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-16T06:30:59.297886+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-16T06:30:59.297886+00:00","source":"crossref"},{"observed_at":"2026-08-16T06:30:54.164669+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-15T18:13:59.313092Z","title":null,"venue":null,"work_id":"04d43e6e-6e06-4b6b-a50c-c8a916034b85","year":2024},"citing_paper":{"arxiv_id":"2507.18631","last_updated":"2025-07-25T07:20:24Z","snapshot_observed_at":"2026-08-15T18:07:02.667074Z","submitted_at":"2025-07-24T17:59:24Z","title":"Layer-Aware Representation Filtering: Purifying Finetuning Data to Preserve LLM Safety Alignment","version":2},"reference_index":13,"source":"arxiv_source","source_observed_at":"2026-08-15T18:13:58.122946Z"},"links":{"citing_paper":"/paper/2507.18631"},"observation_digest":"sha256:35bfc2f63153079697c333ecaa2722a74d497c695d227787ddc36e489b7b8460","observation_id":"3e89adbe-4c88-49a1-a2a2-dbf903111736","resolution":{"observed_at":"2026-08-15T18:13:59.317624Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-16T06:30:59.297886+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-16T06:30:59.297886+00:00","source":"crossref"},{"observed_at":"2026-08-16T06:30:54.164669+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-15T18:13:58.126797Z","title":null,"venue":null,"work_id":null,"year":2021},"citing_paper":{"arxiv_id":"2507.18631","last_updated":"2025-07-25T07:20:24Z","snapshot_observed_at":"2026-08-15T18:07:02.667074Z","submitted_at":"2025-07-24T17:59:24Z","title":"Layer-Aware Representation Filtering: Purifying Finetuning Data to Preserve LLM Safety Alignment","version":2},"reference_index":14,"source":"arxiv_source","source_observed_at":"2026-08-15T18:13:58.126797Z"},"links":{"citing_paper":"/paper/2507.18631"},"observation_digest":"sha256:d404c9860ac7b0112f87406d7db57e63c7c54039d030e41d93b3cfee5a3e9117","observation_id":"6f82fae8-639f-4c48-80c1-85495d458c34","resolution":{"observed_at":"2026-08-15T18:13:58.126797Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-15T18:13:58.130907Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2507.18631","last_updated":"2025-07-25T07:20:24Z","snapshot_observed_at":"2026-08-15T18:07:02.667074Z","submitted_at":"2025-07-24T17:59:24Z","title":"Layer-Aware Representation Filtering: Purifying Finetuning Data to Preserve LLM Safety Alignment","version":2},"reference_index":15,"source":"arxiv_source","source_observed_at":"2026-08-15T18:13:58.130907Z"},"links":{"citing_paper":"/paper/2507.18631"},"observation_digest":"sha256:392c9a16e806740b8d243e53edc62232d10b74e3c5be249f7b246bc72d9015ab","observation_id":"31ccae20-a47e-4daa-b67b-a85a597e9050","resolution":{"observed_at":"2026-08-15T18:13:58.130907Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2411.19939","last_updated":"2025-05-17T15:14:14Z","snapshot_observed_at":"2026-08-15T10:46:50.125849Z","submitted_at":"2024-11-29T18:56:37Z","title":"VLSBench: Unveiling Visual Leakage in Multimodal Safety","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2411.19939","snapshot_observed_at":"2026-08-15T18:13:58.134740Z","title":null,"venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2507.18631","last_updated":"2025-07-25T07:20:24Z","snapshot_observed_at":"2026-08-15T18:07:02.667074Z","submitted_at":"2025-07-24T17:59:24Z","title":"Layer-Aware Representation Filtering: Purifying Finetuning Data to Preserve LLM Safety Alignment","version":2},"reference_index":16,"source":"arxiv_source","source_observed_at":"2026-08-15T18:13:58.134740Z"},"links":{"cited_paper":"/paper/2411.19939","citing_paper":"/paper/2507.18631"},"observation_digest":"sha256:dd7ba6b1aa22402451d12fb85d994accd99b9fa9915773e57401763f71c405db","observation_id":"184995cf-ca86-4440-a72a-559a7f391ea4","resolution":{"observed_at":"2026-08-15T18:13:58.134740Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2409.18169","last_updated":"2026-04-23T18:48:49Z","snapshot_observed_at":"2026-08-14T15:55:00.676298Z","submitted_at":"2024-09-26T17:55:22Z","title":"Harmful Fine-tuning Attacks and Defenses for Large Language Models: A Survey","version":6},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2409.18169","snapshot_observed_at":"2026-08-15T18:13:58.138769Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2507.18631","last_updated":"2025-07-25T07:20:24Z","snapshot_observed_at":"2026-08-15T18:07:02.667074Z","submitted_at":"2025-07-24T17:59:24Z","title":"Layer-Aware Representation Filtering: Purifying Finetuning Data to Preserve LLM Safety Alignment","version":2},"reference_index":17,"source":"arxiv_source","source_observed_at":"2026-08-15T18:13:58.138769Z"},"links":{"cited_paper":"/paper/2409.18169","citing_paper":"/paper/2507.18631"},"observation_digest":"sha256:562b3f7db56cb27174dac0d76b38793ac34ceb15c624f82f0de0a4283ddc7f7d","observation_id":"bd18d47d-7abf-4508-898c-8623a8408cbe","resolution":{"observed_at":"2026-08-15T18:13:58.138769Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-15T18:13:59.283806Z","title":null,"venue":null,"work_id":"ed6aeaaf-613f-4fd8-a954-c22e23ae1af7","year":2024},"citing_paper":{"arxiv_id":"2507.18631","last_updated":"2025-07-25T07:20:24Z","snapshot_observed_at":"2026-08-15T18:07:02.667074Z","submitted_at":"2025-07-24T17:59:24Z","title":"Layer-Aware Representation Filtering: Purifying Finetuning Data to Preserve LLM Safety Alignment","version":2},"reference_index":18,"source":"arxiv_source","source_observed_at":"2026-08-15T18:13:58.143005Z"},"links":{"citing_paper":"/paper/2507.18631"},"observation_digest":"sha256:7fef3c05b9b0e0cb8f82a9101cb9d09670a44799ed225f32e54b587373a18ff3","observation_id":"73f76105-b10d-4c98-9208-e362b0c12fcd","resolution":{"observed_at":"2026-08-15T18:13:59.287956Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-16T06:30:59.297886+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-16T06:30:59.297886+00:00","source":"crossref"},{"observed_at":"2026-08-16T06:30:54.164669+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2310.06825","last_updated":"2023-10-10T17:54:58Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2023-10-10T17:54:58Z","title":"Mistral 7B","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2310.06825","snapshot_observed_at":"2026-08-15T18:13:58.147434Z","title":null,"venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2507.18631","last_updated":"2025-07-25T07:20:24Z","snapshot_observed_at":"2026-08-15T18:07:02.667074Z","submitted_at":"2025-07-24T17:59:24Z","title":"Layer-Aware Representation Filtering: Purifying Finetuning Data to Preserve LLM Safety Alignment","version":2},"reference_index":19,"source":"arxiv_source","source_observed_at":"2026-08-15T18:13:58.147434Z"},"links":{"cited_paper":"/paper/2310.06825","citing_paper":"/paper/2507.18631"},"observation_digest":"sha256:3cef1fb38ad117ea83c9ea0d180c91c88b54df1cca18a9ecce1d971c351c049b","observation_id":"3de25048-435c-4f7c-883c-9b4a34220732","resolution":{"observed_at":"2026-08-15T18:13:58.147434Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"1909.06146","last_updated":"2019-09-13T11:18:20Z","snapshot_observed_at":"2026-08-13T18:54:55.815938Z","submitted_at":"2019-09-13T11:18:20Z","title":"PubMedQA: A Dataset for Biomedical Research Question Answering","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"1909.06146","snapshot_observed_at":"2026-08-15T18:13:58.151662Z","title":null,"venue":null,"work_id":null,"year":2019},"citing_paper":{"arxiv_id":"2507.18631","last_updated":"2025-07-25T07:20:24Z","snapshot_observed_at":"2026-08-15T18:07:02.667074Z","submitted_at":"2025-07-24T17:59:24Z","title":"Layer-Aware Representation Filtering: Purifying Finetuning Data to Preserve LLM Safety Alignment","version":2},"reference_index":20,"source":"arxiv_source","source_observed_at":"2026-08-15T18:13:58.151662Z"},"links":{"cited_paper":"/paper/1909.06146","citing_paper":"/paper/2507.18631"},"observation_digest":"sha256:ac4492892dacbb9edbe9fafdc0c31095b1aad0a01143137b9409a87d646e5d7f","observation_id":"e0600772-ae63-4681-a634-df59de2a1b78","resolution":{"observed_at":"2026-08-15T18:13:58.151662Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-15T18:13:59.270450Z","title":null,"venue":null,"work_id":"5ea9105a-4ca1-418c-a655-416cd39bdb2e","year":2024},"citing_paper":{"arxiv_id":"2507.18631","last_updated":"2025-07-25T07:20:24Z","snapshot_observed_at":"2026-08-15T18:07:02.667074Z","submitted_at":"2025-07-24T17:59:24Z","title":"Layer-Aware Representation Filtering: Purifying Finetuning Data to Preserve LLM Safety Alignment","version":2},"reference_index":21,"source":"arxiv_source","source_observed_at":"2026-08-15T18:13:58.156078Z"},"links":{"citing_paper":"/paper/2507.18631"},"observation_digest":"sha256:58448b8bf87fe6cde4a2c4fbecc2711429c03a3a589e4096d53e51f238a6f334","observation_id":"0525ec50-605a-4d1a-a707-e5adc41d6dce","resolution":{"observed_at":"2026-08-15T18:13:59.274830Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-16T06:30:59.297886+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-16T06:30:59.297886+00:00","source":"crossref"},{"observed_at":"2026-08-16T06:30:54.164669+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2402.05044","last_updated":"2024-06-07T12:05:46Z","snapshot_observed_at":"2026-08-13T04:24:04.479852Z","submitted_at":"2024-02-07T17:33:54Z","title":"SALAD-Bench: A Hierarchical and Comprehensive Safety Benchmark for Large Language Models","version":4},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2402.05044","snapshot_observed_at":"2026-08-15T18:13:58.160012Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2507.18631","last_updated":"2025-07-25T07:20:24Z","snapshot_observed_at":"2026-08-15T18:07:02.667074Z","submitted_at":"2025-07-24T17:59:24Z","title":"Layer-Aware Representation Filtering: Purifying Finetuning Data to Preserve LLM Safety Alignment","version":2},"reference_index":22,"source":"arxiv_source","source_observed_at":"2026-08-15T18:13:58.160012Z"},"links":{"cited_paper":"/paper/2402.05044","citing_paper":"/paper/2507.18631"},"observation_digest":"sha256:cbe7c7cdf997fd953a459dbf94ad3bc039738a0d3f3f210dc5f20e137c2ffe58","observation_id":"1f89ed07-990b-4cfc-8f93-ba7549c2b795","resolution":{"observed_at":"2026-08-15T18:13:58.160012Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-15T18:13:59.256905Z","title":null,"venue":null,"work_id":"74c133e2-93a0-4c49-b33a-f0fcb2130d0c","year":2025},"citing_paper":{"arxiv_id":"2507.18631","last_updated":"2025-07-25T07:20:24Z","snapshot_observed_at":"2026-08-15T18:07:02.667074Z","submitted_at":"2025-07-24T17:59:24Z","title":"Layer-Aware Representation Filtering: Purifying Finetuning Data to Preserve LLM Safety Alignment","version":2},"reference_index":23,"source":"arxiv_source","source_observed_at":"2026-08-15T18:13:58.164377Z"},"links":{"citing_paper":"/paper/2507.18631"},"observation_digest":"sha256:4353f00209c897f0347d3fe7dbe2986e15c532ec10ec02cc4770016a03219052","observation_id":"155e2823-2219-4059-b2ed-1ce90deb75d9","resolution":{"observed_at":"2026-08-15T18:13:59.261019Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-16T06:30:59.297886+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-16T06:30:59.297886+00:00","source":"crossref"},{"observed_at":"2026-08-16T06:30:54.164669+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-15T18:13:59.243190Z","title":null,"venue":null,"work_id":"cb9fb9e3-7abc-4d56-a6d6-7ca68744a983","year":2025},"citing_paper":{"arxiv_id":"2507.18631","last_updated":"2025-07-25T07:20:24Z","snapshot_observed_at":"2026-08-15T18:07:02.667074Z","submitted_at":"2025-07-24T17:59:24Z","title":"Layer-Aware Representation Filtering: Purifying Finetuning Data to Preserve LLM Safety Alignment","version":2},"reference_index":24,"source":"arxiv_source","source_observed_at":"2026-08-15T18:13:58.168244Z"},"links":{"citing_paper":"/paper/2507.18631"},"observation_digest":"sha256:336a277ff0e516c9cb4b8790047ddf2e765e8fb7fa1380a1511e0456e8d80eac","observation_id":"1827fb46-584a-4792-9e48-80aa1a195127","resolution":{"observed_at":"2026-08-15T18:13:59.248134Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-16T06:30:59.297886+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-16T06:30:59.297886+00:00","source":"crossref"},{"observed_at":"2026-08-16T06:30:54.164669+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-15T18:13:59.230372Z","title":null,"venue":null,"work_id":"8bbc5993-e893-43c3-9b6c-37ee8f410e55","year":2025},"citing_paper":{"arxiv_id":"2507.18631","last_updated":"2025-07-25T07:20:24Z","snapshot_observed_at":"2026-08-15T18:07:02.667074Z","submitted_at":"2025-07-24T17:59:24Z","title":"Layer-Aware Representation Filtering: Purifying Finetuning Data to Preserve LLM Safety Alignment","version":2},"reference_index":25,"source":"arxiv_source","source_observed_at":"2026-08-15T18:13:58.172127Z"},"links":{"citing_paper":"/paper/2507.18631"},"observation_digest":"sha256:47c22c1f0f5341c800f11b8bc7c769278b9bdaecf25df22a8f71a8bc55893bc0","observation_id":"67c4c9a9-cb02-4d1c-ac31-595be8e96163","resolution":{"observed_at":"2026-08-15T18:13:59.234462Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-16T06:30:59.297886+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-16T06:30:59.297886+00:00","source":"crossref"},{"observed_at":"2026-08-16T06:30:54.164669+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2407.21783","last_updated":"2024-11-23T23:27:33Z","snapshot_observed_at":"2026-08-13T17:20:44.002518Z","submitted_at":"2024-07-31T17:54:27Z","title":"The Llama 3 Herd of Models","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2407.21783","snapshot_observed_at":"2026-08-15T18:13:58.176161Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2507.18631","last_updated":"2025-07-25T07:20:24Z","snapshot_observed_at":"2026-08-15T18:07:02.667074Z","submitted_at":"2025-07-24T17:59:24Z","title":"Layer-Aware Representation Filtering: Purifying Finetuning Data to Preserve LLM Safety Alignment","version":2},"reference_index":26,"source":"arxiv_source","source_observed_at":"2026-08-15T18:13:58.176161Z"},"links":{"cited_paper":"/paper/2407.21783","citing_paper":"/paper/2507.18631"},"observation_digest":"sha256:748fc94e89e43218fd350cbe19267467c438c2cad339024e3725168536dec8fa","observation_id":"ef3e8b43-1d09-417c-92bd-386a0cc1f3b2","resolution":{"observed_at":"2026-08-15T18:13:58.176161Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-15T18:13:59.216359Z","title":"Le, Barret Zoph, Jason Wei, and Adam Roberts","venue":null,"work_id":"f6f30eec-3a0b-4fb4-82e7-184ece175709","year":2023},"citing_paper":{"arxiv_id":"2507.18631","last_updated":"2025-07-25T07:20:24Z","snapshot_observed_at":"2026-08-15T18:07:02.667074Z","submitted_at":"2025-07-24T17:59:24Z","title":"Layer-Aware Representation Filtering: Purifying Finetuning Data to Preserve LLM Safety Alignment","version":2},"reference_index":27,"source":"arxiv_source","source_observed_at":"2026-08-15T18:13:58.181169Z"},"links":{"citing_paper":"/paper/2507.18631"},"observation_digest":"sha256:bfc74cef6e2cc67cad840c576dab08ea1b1fd2b5f37d773d6f0279c7d9be3305","observation_id":"0ccec216-7161-42f3-8662-6d52a06b04c7","resolution":{"observed_at":"2026-08-15T18:13:59.221218Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-16T06:30:59.297886+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-16T06:30:59.297886+00:00","source":"crossref"},{"observed_at":"2026-08-16T06:30:54.164669+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-15T18:13:58.185761Z","title":null,"venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2507.18631","last_updated":"2025-07-25T07:20:24Z","snapshot_observed_at":"2026-08-15T18:07:02.667074Z","submitted_at":"2025-07-24T17:59:24Z","title":"Layer-Aware Representation Filtering: Purifying Finetuning Data to Preserve LLM Safety Alignment","version":2},"reference_index":28,"source":"arxiv_source","source_observed_at":"2026-08-15T18:13:58.185761Z"},"links":{"citing_paper":"/paper/2507.18631"},"observation_digest":"sha256:a27ef9ac1b6fdcf001b9dde14e6030416a407dd79902acf29c1014fb2e4c09a3","observation_id":"63b3335f-14ac-4fd6-b0a2-446ee4ea9f79","resolution":{"observed_at":"2026-08-15T18:13:58.185761Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-15T18:13:59.202364Z","title":null,"venue":null,"work_id":"620cc394-6ff4-44d8-9cc0-1b41a6c223ac","year":2024},"citing_paper":{"arxiv_id":"2507.18631","last_updated":"2025-07-25T07:20:24Z","snapshot_observed_at":"2026-08-15T18:07:02.667074Z","submitted_at":"2025-07-24T17:59:24Z","title":"Layer-Aware Representation Filtering: Purifying Finetuning Data to Preserve LLM Safety Alignment","version":2},"reference_index":29,"source":"arxiv_source","source_observed_at":"2026-08-15T18:13:58.189793Z"},"links":{"citing_paper":"/paper/2507.18631"},"observation_digest":"sha256:2b9d7f67eee11002c99cb3b119097ecaf5a0e5c285a9930e70211c56148f8030","observation_id":"92a0e61a-e8a1-4e22-a4f0-894b0f6f7b7c","resolution":{"observed_at":"2026-08-15T18:13:59.207039Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-16T06:30:59.297886+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-16T06:30:59.297886+00:00","source":"crossref"},{"observed_at":"2026-08-16T06:30:54.164669+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-15T18:13:58.193740Z","title":null,"venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2507.18631","last_updated":"2025-07-25T07:20:24Z","snapshot_observed_at":"2026-08-15T18:07:02.667074Z","submitted_at":"2025-07-24T17:59:24Z","title":"Layer-Aware Representation Filtering: Purifying Finetuning Data to Preserve LLM Safety Alignment","version":2},"reference_index":30,"source":"arxiv_source","source_observed_at":"2026-08-15T18:13:58.193740Z"},"links":{"citing_paper":"/paper/2507.18631"},"observation_digest":"sha256:7f742086397f288f90f73c75e8f28f5a7fc198aef77cfd3ae89cc2672f2b09d0","observation_id":"f901c20c-96aa-45b4-b63b-0e2b2411d0a8","resolution":{"observed_at":"2026-08-15T18:13:58.193740Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2402.04249","last_updated":"2024-02-27T04:43:08Z","snapshot_observed_at":"2026-08-16T09:07:20.265665Z","submitted_at":"2024-02-06T18:59:08Z","title":"HarmBench: A Standardized Evaluation Framework for Automated Red Teaming and Robust Refusal","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2402.04249","snapshot_observed_at":"2026-08-15T18:13:58.197941Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2507.18631","last_updated":"2025-07-25T07:20:24Z","snapshot_observed_at":"2026-08-15T18:07:02.667074Z","submitted_at":"2025-07-24T17:59:24Z","title":"Layer-Aware Representation Filtering: Purifying Finetuning Data to Preserve LLM Safety Alignment","version":2},"reference_index":31,"source":"arxiv_source","source_observed_at":"2026-08-15T18:13:58.197941Z"},"links":{"cited_paper":"/paper/2402.04249","citing_paper":"/paper/2507.18631"},"observation_digest":"sha256:483dd3bd5a70dc8a267d015ccacd204bb4509dc2f3a0580d99a1ad373450e228","observation_id":"14f144ce-bfac-4c26-a57e-291a99e6f821","resolution":{"observed_at":"2026-08-15T18:13:58.197941Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-15T18:13:58.202174Z","title":null,"venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2507.18631","last_updated":"2025-07-25T07:20:24Z","snapshot_observed_at":"2026-08-15T18:07:02.667074Z","submitted_at":"2025-07-24T17:59:24Z","title":"Layer-Aware Representation Filtering: Purifying Finetuning Data to Preserve LLM Safety Alignment","version":2},"reference_index":32,"source":"arxiv_source","source_observed_at":"2026-08-15T18:13:58.202174Z"},"links":{"citing_paper":"/paper/2507.18631"},"observation_digest":"sha256:af4d58e43f14e8970f0d22018010c951d5ce319512385b618bd64734e3946437","observation_id":"4a5a84a6-a340-46a4-8c79-755e0cf8f018","resolution":{"observed_at":"2026-08-15T18:13:58.202174Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2306.02707","last_updated":"2023-06-05T08:58:39Z","snapshot_observed_at":"2026-08-09T03:29:22.849759Z","submitted_at":"2023-06-05T08:58:39Z","title":"Orca: Progressive Learning from Complex Explanation Traces of GPT-4","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2306.02707","snapshot_observed_at":"2026-08-15T18:13:58.206597Z","title":null,"venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2507.18631","last_updated":"2025-07-25T07:20:24Z","snapshot_observed_at":"2026-08-15T18:07:02.667074Z","submitted_at":"2025-07-24T17:59:24Z","title":"Layer-Aware Representation Filtering: Purifying Finetuning Data to Preserve LLM Safety Alignment","version":2},"reference_index":33,"source":"arxiv_source","source_observed_at":"2026-08-15T18:13:58.206597Z"},"links":{"cited_paper":"/paper/2306.02707","citing_paper":"/paper/2507.18631"},"observation_digest":"sha256:78d01403d82b5262dc645df317553bff57d095c1f52e71d1a39e11dc5e3ba09a","observation_id":"bc8e3bf2-20bd-495b-b762-f8539d6ebb80","resolution":{"observed_at":"2026-08-15T18:13:58.206597Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2502.09674","last_updated":"2025-05-27T08:40:42Z","snapshot_observed_at":"2026-08-13T22:50:26.223810Z","submitted_at":"2025-02-13T06:39:22Z","title":"The Hidden Dimensions of LLM Alignment: A Multi-Dimensional Analysis of Orthogonal Safety Directions","version":4},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2502.09674","snapshot_observed_at":"2026-08-15T18:13:58.211236Z","title":null,"venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2507.18631","last_updated":"2025-07-25T07:20:24Z","snapshot_observed_at":"2026-08-15T18:07:02.667074Z","submitted_at":"2025-07-24T17:59:24Z","title":"Layer-Aware Representation Filtering: Purifying Finetuning Data to Preserve LLM Safety Alignment","version":2},"reference_index":34,"source":"arxiv_source","source_observed_at":"2026-08-15T18:13:58.211236Z"},"links":{"cited_paper":"/paper/2502.09674","citing_paper":"/paper/2507.18631"},"observation_digest":"sha256:73d0931b3596ff138c9fad809470067e71975bfd7bcd2223b4f3f7fa5effe656","observation_id":"7ec92b30-77f1-4820-b4c2-2cdb605dc7f3","resolution":{"observed_at":"2026-08-15T18:13:58.211236Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-15T18:13:58.216544Z","title":null,"venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2507.18631","last_updated":"2025-07-25T07:20:24Z","snapshot_observed_at":"2026-08-15T18:07:02.667074Z","submitted_at":"2025-07-24T17:59:24Z","title":"Layer-Aware Representation Filtering: Purifying Finetuning Data to Preserve LLM Safety Alignment","version":2},"reference_index":35,"source":"arxiv_source","source_observed_at":"2026-08-15T18:13:58.216544Z"},"links":{"citing_paper":"/paper/2507.18631"},"observation_digest":"sha256:37bd70d8af638fd8f917272e9f0ac6e02a25c5ac8fbda4cef172c249ea87bbf8","observation_id":"963cd09d-a9a1-409b-a56d-39695c76b58e","resolution":{"observed_at":"2026-08-15T18:13:58.216544Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2304.03277","last_updated":"2023-04-06T17:58:09Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2023-04-06T17:58:09Z","title":"Instruction Tuning with GPT-4","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2304.03277","snapshot_observed_at":"2026-08-15T18:13:58.220450Z","title":null,"venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2507.18631","last_updated":"2025-07-25T07:20:24Z","snapshot_observed_at":"2026-08-15T18:07:02.667074Z","submitted_at":"2025-07-24T17:59:24Z","title":"Layer-Aware Representation Filtering: Purifying Finetuning Data to Preserve LLM Safety Alignment","version":2},"reference_index":36,"source":"arxiv_source","source_observed_at":"2026-08-15T18:13:58.220450Z"},"links":{"cited_paper":"/paper/2304.03277","citing_paper":"/paper/2507.18631"},"observation_digest":"sha256:1f764d0be83ddd13de9fb7d323ed1cd1ebbc107ad56f7e713ba4792be5af36e4","observation_id":"171f5cbf-5b79-4229-aa8b-5396ee8de8bf","resolution":{"observed_at":"2026-08-15T18:13:58.220450Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-15T18:13:58.224735Z","title":null,"venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2507.18631","last_updated":"2025-07-25T07:20:24Z","snapshot_observed_at":"2026-08-15T18:07:02.667074Z","submitted_at":"2025-07-24T17:59:24Z","title":"Layer-Aware Representation Filtering: Purifying Finetuning Data to Preserve LLM Safety Alignment","version":2},"reference_index":37,"source":"arxiv_source","source_observed_at":"2026-08-15T18:13:58.224735Z"},"links":{"citing_paper":"/paper/2507.18631"},"observation_digest":"sha256:404a5644fa21c4545b59c8155f7749b90484f47caf437c47ed4f7da0f5865111","observation_id":"96781a24-793d-4aca-b356-404b7ec5096d","resolution":{"observed_at":"2026-08-15T18:13:58.224735Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-15T18:13:59.173106Z","title":null,"venue":null,"work_id":"856f1ce1-7976-41da-b4f9-956f14b81a4f","year":2024},"citing_paper":{"arxiv_id":"2507.18631","last_updated":"2025-07-25T07:20:24Z","snapshot_observed_at":"2026-08-15T18:07:02.667074Z","submitted_at":"2025-07-24T17:59:24Z","title":"Layer-Aware Representation Filtering: Purifying Finetuning Data to Preserve LLM Safety Alignment","version":2},"reference_index":38,"source":"arxiv_source","source_observed_at":"2026-08-15T18:13:58.228636Z"},"links":{"citing_paper":"/paper/2507.18631"},"observation_digest":"sha256:9e9fceda7d8492abecac6be19b394084b61391cb711351410ebb6c77466b221b","observation_id":"ab883154-8edd-4b91-a4b1-73af5c88e90a","resolution":{"observed_at":"2026-08-15T18:13:59.177364Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-16T06:30:59.297886+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-16T06:30:59.297886+00:00","source":"crossref"},{"observed_at":"2026-08-16T06:30:54.164669+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2412.15115","last_updated":"2025-01-03T02:18:21Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2024-12-19T17:56:09Z","title":"Qwen2.5 Technical Report","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2412.15115","snapshot_observed_at":"2026-08-15T18:13:58.232479Z","title":null,"venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2507.18631","last_updated":"2025-07-25T07:20:24Z","snapshot_observed_at":"2026-08-15T18:07:02.667074Z","submitted_at":"2025-07-24T17:59:24Z","title":"Layer-Aware Representation Filtering: Purifying Finetuning Data to Preserve LLM Safety Alignment","version":2},"reference_index":39,"source":"arxiv_source","source_observed_at":"2026-08-15T18:13:58.232479Z"},"links":{"cited_paper":"/paper/2412.15115","citing_paper":"/paper/2507.18631"},"observation_digest":"sha256:d0a81ad7f63469a9244bdf7e793a942b38d94ac75da2595b66560d4a2ba40192","observation_id":"451dedea-7b0b-4cc8-9c59-e863dfd9534a","resolution":{"observed_at":"2026-08-15T18:13:58.232479Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-15T18:13:58.236321Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2507.18631","last_updated":"2025-07-25T07:20:24Z","snapshot_observed_at":"2026-08-15T18:07:02.667074Z","submitted_at":"2025-07-24T17:59:24Z","title":"Layer-Aware Representation Filtering: Purifying Finetuning Data to Preserve LLM Safety Alignment","version":2},"reference_index":40,"source":"arxiv_source","source_observed_at":"2026-08-15T18:13:58.236321Z"},"links":{"citing_paper":"/paper/2507.18631"},"observation_digest":"sha256:c12a2d7609842d0a4535c521eaf54b3c3932c712f3d099d8e4e40f33336a7feb","observation_id":"d49d37f8-03e6-4018-8feb-6f349ab08c85","resolution":{"observed_at":"2026-08-15T18:13:58.236321Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-15T18:13:59.159786Z","title":null,"venue":null,"work_id":"e632171c-c8aa-46a7-a67f-e3d3fd1f964b","year":2024},"citing_paper":{"arxiv_id":"2507.18631","last_updated":"2025-07-25T07:20:24Z","snapshot_observed_at":"2026-08-15T18:07:02.667074Z","submitted_at":"2025-07-24T17:59:24Z","title":"Layer-Aware Representation Filtering: Purifying Finetuning Data to Preserve LLM Safety Alignment","version":2},"reference_index":41,"source":"arxiv_source","source_observed_at":"2026-08-15T18:13:58.240260Z"},"links":{"citing_paper":"/paper/2507.18631"},"observation_digest":"sha256:ecaa10a461e342bbe67ff46abb63adb9dd527e1af7ab678700913b797ddb1204","observation_id":"2888ba6a-04df-4405-b813-0250bd762f94","resolution":{"observed_at":"2026-08-15T18:13:59.164425Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-16T06:30:59.297886+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-16T06:30:59.297886+00:00","source":"crossref"},{"observed_at":"2026-08-16T06:30:54.164669+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-15T18:13:59.146395Z","title":null,"venue":null,"work_id":"d8badd77-36f2-429c-a485-260d0b4c8cac","year":2007},"citing_paper":{"arxiv_id":"2507.18631","last_updated":"2025-07-25T07:20:24Z","snapshot_observed_at":"2026-08-15T18:07:02.667074Z","submitted_at":"2025-07-24T17:59:24Z","title":"Layer-Aware Representation Filtering: Purifying Finetuning Data to Preserve LLM Safety Alignment","version":2},"reference_index":42,"source":"arxiv_source","source_observed_at":"2026-08-15T18:13:58.244219Z"},"links":{"citing_paper":"/paper/2507.18631"},"observation_digest":"sha256:45411c34e89abbf327e3f90a2c48a40e8b2db771fbf87fdbe99c65bd8dc0f297","observation_id":"564797de-c99d-4454-b729-85aa6e7e69ba","resolution":{"observed_at":"2026-08-15T18:13:59.150561Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-16T06:30:59.297886+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-16T06:30:59.297886+00:00","source":"crossref"},{"observed_at":"2026-08-16T06:30:54.164669+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-15T18:13:59.133056Z","title":null,"venue":null,"work_id":"6c9272a6-2737-431b-8402-9dee9c15caa4","year":2025},"citing_paper":{"arxiv_id":"2507.18631","last_updated":"2025-07-25T07:20:24Z","snapshot_observed_at":"2026-08-15T18:07:02.667074Z","submitted_at":"2025-07-24T17:59:24Z","title":"Layer-Aware Representation Filtering: Purifying Finetuning Data to Preserve LLM Safety Alignment","version":2},"reference_index":43,"source":"arxiv_source","source_observed_at":"2026-08-15T18:13:58.248560Z"},"links":{"citing_paper":"/paper/2507.18631"},"observation_digest":"sha256:48fcf2ab455a4d1c00713f998362671b07f53e083368195684612ca9bdc8fbfc","observation_id":"5e4944a2-8146-4fa9-bc7e-422ed32d7d06","resolution":{"observed_at":"2026-08-15T18:13:59.137665Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-16T06:30:59.297886+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-16T06:30:59.297886+00:00","source":"crossref"},{"observed_at":"2026-08-16T06:30:54.164669+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-15T18:13:58.252634Z","title":"Hashimoto","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2507.18631","last_updated":"2025-07-25T07:20:24Z","snapshot_observed_at":"2026-08-15T18:07:02.667074Z","submitted_at":"2025-07-24T17:59:24Z","title":"Layer-Aware Representation Filtering: Purifying Finetuning Data to Preserve LLM Safety Alignment","version":2},"reference_index":44,"source":"arxiv_source","source_observed_at":"2026-08-15T18:13:58.252634Z"},"links":{"citing_paper":"/paper/2507.18631"},"observation_digest":"sha256:48464fb4bcd69756bca1b84d6100f4f4bd2b7be746804490fe7be3ae8c5ad8a5","observation_id":"fb90d4b0-da16-44bc-bb0f-d2074f7446a6","resolution":{"observed_at":"2026-08-15T18:13:58.252634Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-15T18:13:58.256725Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2507.18631","last_updated":"2025-07-25T07:20:24Z","snapshot_observed_at":"2026-08-15T18:07:02.667074Z","submitted_at":"2025-07-24T17:59:24Z","title":"Layer-Aware Representation Filtering: Purifying Finetuning Data to Preserve LLM Safety Alignment","version":2},"reference_index":45,"source":"arxiv_source","source_observed_at":"2026-08-15T18:13:58.256725Z"},"links":{"citing_paper":"/paper/2507.18631"},"observation_digest":"sha256:3bf1396462532a678f1a360a97bf0ffa63dbdcb43704ff346813517d6a610729","observation_id":"72a653c1-e72f-4547-8f1a-fbec0a5d1d4f","resolution":{"observed_at":"2026-08-15T18:13:58.256725Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2412.17522","last_updated":"2025-01-05T04:44:32Z","snapshot_observed_at":"2026-08-15T15:38:59.741982Z","submitted_at":"2024-12-23T12:44:54Z","title":"DiffusionAttacker: Diffusion-Driven Prompt Manipulation for LLM Jailbreak","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2412.17522","snapshot_observed_at":"2026-08-15T18:13:58.260498Z","title":null,"venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2507.18631","last_updated":"2025-07-25T07:20:24Z","snapshot_observed_at":"2026-08-15T18:07:02.667074Z","submitted_at":"2025-07-24T17:59:24Z","title":"Layer-Aware Representation Filtering: Purifying Finetuning Data to Preserve LLM Safety Alignment","version":2},"reference_index":46,"source":"arxiv_source","source_observed_at":"2026-08-15T18:13:58.260498Z"},"links":{"cited_paper":"/paper/2412.17522","citing_paper":"/paper/2507.18631"},"observation_digest":"sha256:ad055edf1304cdc8d68094e92600889dcd9afcacd73a071352c143d712927cea","observation_id":"53369d06-f6a7-46af-a47e-8343964a5297","resolution":{"observed_at":"2026-08-15T18:13:58.260498Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-15T18:13:59.103518Z","title":null,"venue":null,"work_id":"f6855e68-1736-4fd5-add9-3463d417e1d7","year":2024},"citing_paper":{"arxiv_id":"2507.18631","last_updated":"2025-07-25T07:20:24Z","snapshot_observed_at":"2026-08-15T18:07:02.667074Z","submitted_at":"2025-07-24T17:59:24Z","title":"Layer-Aware Representation Filtering: Purifying Finetuning Data to Preserve LLM Safety Alignment","version":2},"reference_index":47,"source":"arxiv_source","source_observed_at":"2026-08-15T18:13:58.265181Z"},"links":{"citing_paper":"/paper/2507.18631"},"observation_digest":"sha256:9d9c3e6c3b9091f1e638e74dadf42d9ea9d2bd77576658ed69cfef897868ebde","observation_id":"59a330a3-7630-4926-a78e-eb2d59ded940","resolution":{"observed_at":"2026-08-15T18:13:59.107927Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-16T06:30:59.297886+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-16T06:30:59.297886+00:00","source":"crossref"},{"observed_at":"2026-08-16T06:30:54.164669+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2303.17564","last_updated":"2023-12-21T06:21:11Z","snapshot_observed_at":"2026-08-14T12:54:48.492396Z","submitted_at":"2023-03-30T17:30:36Z","title":"BloombergGPT: A Large Language Model for Finance","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2303.17564","snapshot_observed_at":"2026-08-15T18:13:58.269384Z","title":null,"venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2507.18631","last_updated":"2025-07-25T07:20:24Z","snapshot_observed_at":"2026-08-15T18:07:02.667074Z","submitted_at":"2025-07-24T17:59:24Z","title":"Layer-Aware Representation Filtering: Purifying Finetuning Data to Preserve LLM Safety Alignment","version":2},"reference_index":48,"source":"arxiv_source","source_observed_at":"2026-08-15T18:13:58.269384Z"},"links":{"cited_paper":"/paper/2303.17564","citing_paper":"/paper/2507.18631"},"observation_digest":"sha256:a21c58baca95f9974432094eb71ab2e67eeee7c3d42f47d062d3a7edf17a0d5c","observation_id":"f04d4964-7648-41d6-8fcc-161f0321a0db","resolution":{"observed_at":"2026-08-15T18:13:58.269384Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-15T18:13:58.273453Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2507.18631","last_updated":"2025-07-25T07:20:24Z","snapshot_observed_at":"2026-08-15T18:07:02.667074Z","submitted_at":"2025-07-24T17:59:24Z","title":"Layer-Aware Representation Filtering: Purifying Finetuning Data to Preserve LLM Safety Alignment","version":2},"reference_index":49,"source":"arxiv_source","source_observed_at":"2026-08-15T18:13:58.273453Z"},"links":{"citing_paper":"/paper/2507.18631"},"observation_digest":"sha256:ab8b896fdec73ea842b126f70e7ca2bae0d252e22ade9c68848e9fb6f8e60fdc","observation_id":"f7aec56e-9a6d-4a00-a148-0ec82b4c2fd5","resolution":{"observed_at":"2026-08-15T18:13:58.273453Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-15T18:13:59.082257Z","title":null,"venue":null,"work_id":"4cc4552f-bd1e-4738-bd32-34b89e1ae204","year":2024},"citing_paper":{"arxiv_id":"2507.18631","last_updated":"2025-07-25T07:20:24Z","snapshot_observed_at":"2026-08-15T18:07:02.667074Z","submitted_at":"2025-07-24T17:59:24Z","title":"Layer-Aware Representation Filtering: Purifying Finetuning Data to Preserve LLM Safety Alignment","version":2},"reference_index":50,"source":"arxiv_source","source_observed_at":"2026-08-15T18:13:58.277662Z"},"links":{"citing_paper":"/paper/2507.18631"},"observation_digest":"sha256:645d62170d25addf8493a581f8fc72a85eff8e55fc3a34cd7a11e9b51ffc8a56","observation_id":"06b08e16-65e4-416d-9cca-039eeb1a402e","resolution":{"observed_at":"2026-08-15T18:13:59.086526Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-16T06:30:59.297886+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-16T06:30:59.297886+00:00","source":"crossref"},{"observed_at":"2026-08-16T06:30:54.164669+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2309.12284","last_updated":"2024-05-03T17:36:07Z","snapshot_observed_at":"2026-08-13T10:57:13.012119Z","submitted_at":"2023-09-21T17:45:42Z","title":"MetaMath: Bootstrap Your Own Mathematical Questions for Large Language Models","version":4},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2309.12284","snapshot_observed_at":"2026-08-15T18:13:58.282824Z","title":null,"venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2507.18631","last_updated":"2025-07-25T07:20:24Z","snapshot_observed_at":"2026-08-15T18:07:02.667074Z","submitted_at":"2025-07-24T17:59:24Z","title":"Layer-Aware Representation Filtering: Purifying Finetuning Data to Preserve LLM Safety Alignment","version":2},"reference_index":51,"source":"arxiv_source","source_observed_at":"2026-08-15T18:13:58.282824Z"},"links":{"cited_paper":"/paper/2309.12284","citing_paper":"/paper/2507.18631"},"observation_digest":"sha256:1ed6db28d07ff9bb1e8986db3be016ee8bb8725941b365c3db65de96bbb78cea","observation_id":"6ad2df5b-c7ae-4824-9c5c-3e0207162186","resolution":{"observed_at":"2026-08-15T18:13:58.282824Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-15T18:13:58.287060Z","title":null,"venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2507.18631","last_updated":"2025-07-25T07:20:24Z","snapshot_observed_at":"2026-08-15T18:07:02.667074Z","submitted_at":"2025-07-24T17:59:24Z","title":"Layer-Aware Representation Filtering: Purifying Finetuning Data to Preserve LLM Safety Alignment","version":2},"reference_index":52,"source":"arxiv_source","source_observed_at":"2026-08-15T18:13:58.287060Z"},"links":{"citing_paper":"/paper/2507.18631"},"observation_digest":"sha256:523f9b013d43d3cdd66e8678ca554d7f59f855a7c96ff27f09b2f0e76f45c33b","observation_id":"c22a5ecc-d184-4848-8f80-e37a9ff148b5","resolution":{"observed_at":"2026-08-15T18:13:58.287060Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-15T18:13:59.068627Z","title":null,"venue":null,"work_id":"b9ce506a-a511-4215-a026-4ddf863195ae","year":2024},"citing_paper":{"arxiv_id":"2507.18631","last_updated":"2025-07-25T07:20:24Z","snapshot_observed_at":"2026-08-15T18:07:02.667074Z","submitted_at":"2025-07-24T17:59:24Z","title":"Layer-Aware Representation Filtering: Purifying Finetuning Data to Preserve LLM Safety Alignment","version":2},"reference_index":53,"source":"arxiv_source","source_observed_at":"2026-08-15T18:13:58.291006Z"},"links":{"citing_paper":"/paper/2507.18631"},"observation_digest":"sha256:8e6e1f02169f6c1ed5be2e3344b6c5914c0347f1970fed91c564232226eb780c","observation_id":"500277c0-7578-402f-91d5-72864343f2d5","resolution":{"observed_at":"2026-08-15T18:13:59.073501Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-16T06:30:59.297886+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-16T06:30:59.297886+00:00","source":"crossref"},{"observed_at":"2026-08-16T06:30:54.164669+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-15T18:13:59.054736Z","title":null,"venue":null,"work_id":"9e5b7cce-c112-45c8-8a0a-3d6b3176028d","year":2025},"citing_paper":{"arxiv_id":"2507.18631","last_updated":"2025-07-25T07:20:24Z","snapshot_observed_at":"2026-08-15T18:07:02.667074Z","submitted_at":"2025-07-24T17:59:24Z","title":"Layer-Aware Representation Filtering: Purifying Finetuning Data to Preserve LLM Safety Alignment","version":2},"reference_index":54,"source":"arxiv_source","source_observed_at":"2026-08-15T18:13:58.295289Z"},"links":{"citing_paper":"/paper/2507.18631"},"observation_digest":"sha256:b1ba5447a22f5de31b7ee173e9bd25d47c97d495d0a56d927938b734d1d6d079","observation_id":"448351e1-35eb-436c-8526-e03f6041d412","resolution":{"observed_at":"2026-08-15T18:13:59.059765Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-16T06:30:59.297886+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-16T06:30:59.297886+00:00","source":"crossref"},{"observed_at":"2026-08-16T06:30:54.164669+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-15T18:13:59.038976Z","title":null,"venue":null,"work_id":"69ff480c-e7e5-4380-88b2-fd0119c656c5","year":2025},"citing_paper":{"arxiv_id":"2507.18631","last_updated":"2025-07-25T07:20:24Z","snapshot_observed_at":"2026-08-15T18:07:02.667074Z","submitted_at":"2025-07-24T17:59:24Z","title":"Layer-Aware Representation Filtering: Purifying Finetuning Data to Preserve LLM Safety Alignment","version":2},"reference_index":55,"source":"arxiv_source","source_observed_at":"2026-08-15T18:13:58.299423Z"},"links":{"citing_paper":"/paper/2507.18631"},"observation_digest":"sha256:d00bbbeee85e7b7b2dfb65a6e850b7af8b27f53bad224ff9b1c8b9ef3ad4c7cd","observation_id":"9801f925-faa7-4bbf-8d7b-df8ab9294fd2","resolution":{"observed_at":"2026-08-15T18:13:59.045166Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-16T06:30:59.297886+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-16T06:30:59.297886+00:00","source":"crossref"},{"observed_at":"2026-08-16T06:30:54.164669+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2403.12171","last_updated":"2024-03-18T18:39:53Z","snapshot_observed_at":"2026-08-15T12:12:25.322622Z","submitted_at":"2024-03-18T18:39:53Z","title":"EasyJailbreak: A Unified Framework for Jailbreaking Large Language Models","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2403.12171","snapshot_observed_at":"2026-08-15T18:13:58.303164Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2507.18631","last_updated":"2025-07-25T07:20:24Z","snapshot_observed_at":"2026-08-15T18:07:02.667074Z","submitted_at":"2025-07-24T17:59:24Z","title":"Layer-Aware Representation Filtering: Purifying Finetuning Data to Preserve LLM Safety Alignment","version":2},"reference_index":56,"source":"arxiv_source","source_observed_at":"2026-08-15T18:13:58.303164Z"},"links":{"cited_paper":"/paper/2403.12171","citing_paper":"/paper/2507.18631"},"observation_digest":"sha256:b906588dac63b1717104d693d903e67265be3c44d54d6881d02a4f21b1f8efca","observation_id":"2f1efd34-6d5b-4aef-998c-1c6f65924b32","resolution":{"observed_at":"2026-08-15T18:13:58.303164Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2402.02207","last_updated":"2024-06-17T22:26:32Z","snapshot_observed_at":"2026-08-13T04:27:42.124177Z","submitted_at":"2024-02-03T16:43:42Z","title":"Safety Fine-Tuning at (Almost) No Cost: A Baseline for Vision Large Language Models","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2402.02207","snapshot_observed_at":"2026-08-15T18:13:58.308125Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2507.18631","last_updated":"2025-07-25T07:20:24Z","snapshot_observed_at":"2026-08-15T18:07:02.667074Z","submitted_at":"2025-07-24T17:59:24Z","title":"Layer-Aware Representation Filtering: Purifying Finetuning Data to Preserve LLM Safety Alignment","version":2},"reference_index":57,"source":"arxiv_source","source_observed_at":"2026-08-15T18:13:58.308125Z"},"links":{"cited_paper":"/paper/2402.02207","citing_paper":"/paper/2507.18631"},"observation_digest":"sha256:7ae140bd171c7fced114269254c0ac115f600438757140c8b499ad87532edabb","observation_id":"4c70829f-df76-4c64-aee7-e9700c6203d7","resolution":{"observed_at":"2026-08-15T18:13:58.308125Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2310.01405","last_updated":"2025-03-03T06:14:14Z","snapshot_observed_at":"2026-07-06T16:26:38.284922Z","submitted_at":"2023-10-02T17:59:07Z","title":"Representation Engineering: A Top-Down Approach to AI Transparency","version":4},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2310.01405","snapshot_observed_at":"2026-08-15T18:13:58.312212Z","title":null,"venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2507.18631","last_updated":"2025-07-25T07:20:24Z","snapshot_observed_at":"2026-08-15T18:07:02.667074Z","submitted_at":"2025-07-24T17:59:24Z","title":"Layer-Aware Representation Filtering: Purifying Finetuning Data to Preserve LLM Safety Alignment","version":2},"reference_index":58,"source":"arxiv_source","source_observed_at":"2026-08-15T18:13:58.312212Z"},"links":{"cited_paper":"/paper/2310.01405","citing_paper":"/paper/2507.18631"},"observation_digest":"sha256:61433c1fca6d40559dd75b5fcb5d652fdaa01e99aff76ee3021dbc3a685daee1","observation_id":"172b5635-ec66-430d-811f-7942b5821db1","resolution":{"observed_at":"2026-08-15T18:13:58.312212Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2406.04313","last_updated":"2024-07-12T16:51:07Z","snapshot_observed_at":"2026-08-12T23:48:30.125202Z","submitted_at":"2024-06-06T17:57:04Z","title":"Improving Alignment and Robustness with Circuit Breakers","version":4},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2406.04313","snapshot_observed_at":"2026-08-15T18:13:58.316618Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2507.18631","last_updated":"2025-07-25T07:20:24Z","snapshot_observed_at":"2026-08-15T18:07:02.667074Z","submitted_at":"2025-07-24T17:59:24Z","title":"Layer-Aware Representation Filtering: Purifying Finetuning Data to Preserve LLM Safety Alignment","version":2},"reference_index":59,"source":"arxiv_source","source_observed_at":"2026-08-15T18:13:58.316618Z"},"links":{"cited_paper":"/paper/2406.04313","citing_paper":"/paper/2507.18631"},"observation_digest":"sha256:52a47196925521caa8818b3bed1ae3cf4faf507ffabfc11e6a48fe2c6e44a34d","observation_id":"123f864d-d679-420f-b0a0-ce10dd35ac5d","resolution":{"observed_at":"2026-08-15T18:13:58.316618Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2307.15043","last_updated":"2023-12-20T20:48:57Z","snapshot_observed_at":"2026-08-12T09:06:50.363435Z","submitted_at":"2023-07-27T17:49:12Z","title":"Universal and Transferable Adversarial Attacks on Aligned Language Models","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2307.15043","snapshot_observed_at":"2026-08-15T18:13:58.320855Z","title":null,"venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2507.18631","last_updated":"2025-07-25T07:20:24Z","snapshot_observed_at":"2026-08-15T18:07:02.667074Z","submitted_at":"2025-07-24T17:59:24Z","title":"Layer-Aware Representation Filtering: Purifying Finetuning Data to Preserve LLM Safety Alignment","version":2},"reference_index":60,"source":"arxiv_source","source_observed_at":"2026-08-15T18:13:58.320855Z"},"links":{"cited_paper":"/paper/2307.15043","citing_paper":"/paper/2507.18631"},"observation_digest":"sha256:bcbab27ea482b0c96e62fc34ab5b41885426d8b8927976ee673af3e4cea09058","observation_id":"0fb29fb7-df3b-42eb-953d-146c7be5fd32","resolution":{"observed_at":"2026-08-15T18:13:58.320855Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}}],"paper":{"arxiv_id":"2507.18631","last_updated":"2025-07-25T07:20:24Z","latest_version":2,"primary_category":"cs.CR","snapshot_observed_at":"2026-08-15T18:07:02.667074Z","submitted_at":"2025-07-24T17:59:24Z","title":"Layer-Aware Representation Filtering: Purifying Finetuning Data to Preserve LLM Safety Alignment"},"reference_resolution":{"displayed":60,"state_counts":{"malformed_identifier":0,"metadata_mismatch":0,"parse_uncertain":0,"unresolved":59,"verified_exact":0,"verified_fuzzy":1},"total_outbound_references":60},"refusal":"A citation records a reference. It does not transfer a finding from one paper to another.","schema":"pith.paper-citation-record.v1","standing_sources":[{"observed_at":"2026-08-16T06:30:59.297886+00:00","source":"crossref"},{"observed_at":"2026-08-16T06:30:54.164669+00:00","source":"retraction_watch"}],"thesis":"As of 16 August 2026, this Paper Citation Record lists 60 of 60 outbound references and 0 inbound Pith citation observations for arXiv:2507.18631."}