{"as_of":"2026-08-12T15:05:00Z","caps":{"database_statements":6,"inbound":100,"outbound":100},"context_digest":"sha256:fd0f231a9886be7317fbac211b068de2e883728f20f45f883742169fe3869656","coverage":[{"denominator":26,"lane":"reference_resolution","note":"Typed states for the displayed outbound observations.","records_observed":26,"source":"paper_references, paper_reference_links","source_observed_at":"2026-08-10T22:53:20.963955Z","state":"measured"},{"denominator":26,"lane":"standing_notices","note":"One-hop event checks from named stored sources.","records_observed":26,"source":"scholarly_work_events, retraction_status_cache","source_observed_at":"2026-08-12T06:34:41.77262+00:00","state":"measured"},{"denominator":0,"lane":"inbound_itemization","note":"Pith citing papers itemized under the disclosed page cap.","records_observed":0,"source":"paper_references, paper_reference_links","source_observed_at":null,"state":"measured"},{"denominator":1,"lane":"external_citation_measurements","note":"A source-named dated measurement, never combined with another source.","records_observed":0,"source":"cited_works","source_observed_at":null,"state":"measured"}],"external_citation_measurements":[],"inbound":[],"links":{"evidence":"/evidence","html":"/paper/2501.00517/citation-record","integrity":"/paper/2501.00517/integrity","json":"/paper/2501.00517/citation-record.json","paper":"/paper/2501.00517"},"outbound":[{"citation":{"cited_paper":{"arxiv_id":"2310.17884","last_updated":"2024-06-28T23:27:10Z","snapshot_observed_at":"2026-08-07T22:36:46.691875Z","submitted_at":"2023-10-27T04:15:30Z","title":"Can LLMs Keep a Secret? Testing Privacy Implications of Language Models via Contextual Integrity Theory","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2310.17884","snapshot_observed_at":"2026-08-10T22:53:20.344092Z","title":"Can LLMs Keep a Secret? Testing Privacy Implications of Language Models via Contextual Integrity Theory","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2501.00517","last_updated":"2024-12-31T16:01:25Z","snapshot_observed_at":"2026-08-12T01:14:06.357738Z","submitted_at":"2024-12-31T16:01:25Z","title":"A Method for Enhancing the Safety of Large Model Generation Based on Multi-dimensional Attack and Defense","version":1},"reference_index":1,"source":"pdf_text","source_observed_at":"2026-08-10T22:53:20.344092Z"},"links":{"cited_paper":"/paper/2310.17884","citing_paper":"/paper/2501.00517"},"observation_digest":"sha256:9f3a1a5a02c7d3f33294a7b8e9eaab7aee80a9cc74c1c973ffb933f4951f6cad","observation_id":"31561fc9-05c1-44cb-bacf-fbb7550e4640","resolution":{"observed_at":"2026-08-10T22:53:20.344092Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2309.13788","last_updated":"2024-04-23T22:59:13Z","snapshot_observed_at":"2026-08-12T01:13:43.935792Z","submitted_at":"2023-09-25T00:45:07Z","title":"Can LLM-Generated Misinformation Be Detected?","version":5},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2309.13788","snapshot_observed_at":"2026-08-10T22:53:20.372871Z","title":"Can LLM -Generated Misinformation Be Detected?","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2501.00517","last_updated":"2024-12-31T16:01:25Z","snapshot_observed_at":"2026-08-12T01:14:06.357738Z","submitted_at":"2024-12-31T16:01:25Z","title":"A Method for Enhancing the Safety of Large Model Generation Based on Multi-dimensional Attack and Defense","version":1},"reference_index":2,"source":"pdf_text","source_observed_at":"2026-08-10T22:53:20.372871Z"},"links":{"cited_paper":"/paper/2309.13788","citing_paper":"/paper/2501.00517"},"observation_digest":"sha256:10877f0ca7e719002c5cd517920c0979041507d439eceacd2328508d6327818b","observation_id":"087a0014-2e97-4466-9c8b-600f6881ab33","resolution":{"observed_at":"2026-08-10T22:53:20.372871Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2404.09932","last_updated":"2024-09-06T00:46:40Z","snapshot_observed_at":"2026-07-06T18:00:30.424554Z","submitted_at":"2024-04-15T16:58:28Z","title":"Foundational Challenges in Assuring Alignment and Safety of Large Language Models","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2404.09932","snapshot_observed_at":"2026-08-10T22:53:20.428520Z","title":"Foundational Challenges in Assuring Alignment and Safety of Large Language Models","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2501.00517","last_updated":"2024-12-31T16:01:25Z","snapshot_observed_at":"2026-08-12T01:14:06.357738Z","submitted_at":"2024-12-31T16:01:25Z","title":"A Method for Enhancing the Safety of Large Model Generation Based on Multi-dimensional Attack and Defense","version":1},"reference_index":3,"source":"pdf_text","source_observed_at":"2026-08-10T22:53:20.428520Z"},"links":{"cited_paper":"/paper/2404.09932","citing_paper":"/paper/2501.00517"},"observation_digest":"sha256:b966696695ddd0f13dd47318eb3440dee09ebd2e6a14c455643c6647fe73faee","observation_id":"9c9485fd-3d9b-48a2-a535-f7d52dd85d64","resolution":{"observed_at":"2026-08-10T22:53:20.428520Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2210.11416","last_updated":"2022-12-06T21:39:48Z","snapshot_observed_at":"2026-07-06T14:08:18.855958Z","submitted_at":"2022-10-20T16:58:32Z","title":"Scaling Instruction-Finetuned Language Models","version":5},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2210.11416","snapshot_observed_at":"2026-08-10T22:53:20.480547Z","title":"Scaling Instruction -Finetuned Language Models","venue":null,"work_id":null,"year":2022},"citing_paper":{"arxiv_id":"2501.00517","last_updated":"2024-12-31T16:01:25Z","snapshot_observed_at":"2026-08-12T01:14:06.357738Z","submitted_at":"2024-12-31T16:01:25Z","title":"A Method for Enhancing the Safety of Large Model Generation Based on Multi-dimensional Attack and Defense","version":1},"reference_index":4,"source":"pdf_text","source_observed_at":"2026-08-10T22:53:20.480547Z"},"links":{"cited_paper":"/paper/2210.11416","citing_paper":"/paper/2501.00517"},"observation_digest":"sha256:52b0afccbd87e026cb22afa79fa8cb1eeb1d798932c181dbf29eff18fd9cb806","observation_id":"bdc615e5-ee8f-44f9-88ff-58dd2205d023","resolution":{"observed_at":"2026-08-10T22:53:20.480547Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2203.02155","last_updated":"2022-03-04T07:04:42Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2022-03-04T07:04:42Z","title":"Training language models to follow instructions with human feedback","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2203.02155","snapshot_observed_at":"2026-08-10T22:53:20.557650Z","title":"Training language models to follow instructions with human feedback","venue":null,"work_id":null,"year":2022},"citing_paper":{"arxiv_id":"2501.00517","last_updated":"2024-12-31T16:01:25Z","snapshot_observed_at":"2026-08-12T01:14:06.357738Z","submitted_at":"2024-12-31T16:01:25Z","title":"A Method for Enhancing the Safety of Large Model Generation Based on Multi-dimensional Attack and Defense","version":1},"reference_index":5,"source":"pdf_text","source_observed_at":"2026-08-10T22:53:20.557650Z"},"links":{"cited_paper":"/paper/2203.02155","citing_paper":"/paper/2501.00517"},"observation_digest":"sha256:affc0d6abf6807a53277d5b5db5c4defd1de2c25e2d06c35ec4f6834132f8329","observation_id":"2d85b8b3-0060-4a14-8675-58054c37313f","resolution":{"observed_at":"2026-08-10T22:53:20.557650Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2309.07875","last_updated":"2024-03-19T16:50:50Z","snapshot_observed_at":"2026-08-06T23:45:57.910675Z","submitted_at":"2023-09-14T17:23:37Z","title":"Safety-Tuned LLaMAs: Lessons From Improving the Safety of Large Language Models that Follow Instructions","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2309.07875","snapshot_observed_at":"2026-08-10T22:53:20.585262Z","title":"Safety-Tuned LLaMAs: Lessons From Improving the Safety of Large Language Models that Follow Instructions","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2501.00517","last_updated":"2024-12-31T16:01:25Z","snapshot_observed_at":"2026-08-12T01:14:06.357738Z","submitted_at":"2024-12-31T16:01:25Z","title":"A Method for Enhancing the Safety of Large Model Generation Based on Multi-dimensional Attack and Defense","version":1},"reference_index":6,"source":"pdf_text","source_observed_at":"2026-08-10T22:53:20.585262Z"},"links":{"cited_paper":"/paper/2309.07875","citing_paper":"/paper/2501.00517"},"observation_digest":"sha256:766b7442c071907c1cdbe553fcbb256c46ae966a0293fd6db2dc263dff13667a","observation_id":"fe6fb275-3db4-472c-97dd-84e84362f4b5","resolution":{"observed_at":"2026-08-10T22:53:20.585262Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2305.18290","last_updated":"2024-07-29T22:26:36Z","snapshot_observed_at":"2026-08-01T16:34:38.795326Z","submitted_at":"2023-05-29T17:57:46Z","title":"Direct Preference Optimization: Your Language Model is Secretly a Reward Model","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2305.18290","snapshot_observed_at":"2026-08-10T22:53:20.610290Z","title":"Direct Preference Optimization: Your Language Model is Secretly a Reward Model","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2501.00517","last_updated":"2024-12-31T16:01:25Z","snapshot_observed_at":"2026-08-12T01:14:06.357738Z","submitted_at":"2024-12-31T16:01:25Z","title":"A Method for Enhancing the Safety of Large Model Generation Based on Multi-dimensional Attack and Defense","version":1},"reference_index":7,"source":"pdf_text","source_observed_at":"2026-08-10T22:53:20.610290Z"},"links":{"cited_paper":"/paper/2305.18290","citing_paper":"/paper/2501.00517"},"observation_digest":"sha256:e3db8ec7d9b99bb49ac9a2f156d0a5a2e2511312ff332f468a0f4e5388a2a85f","observation_id":"375f6167-abdc-4ad2-a10b-031127dcf45c","resolution":{"observed_at":"2026-08-10T22:53:20.610290Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2402.01306","last_updated":"2024-11-19T18:12:45Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2024-02-02T10:53:36Z","title":"KTO: Model Alignment as Prospect Theoretic Optimization","version":4},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2402.01306","snapshot_observed_at":"2026-08-10T22:53:20.621497Z","title":"KTO: Model Alignment as Prospect Theoretic Optimization","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2501.00517","last_updated":"2024-12-31T16:01:25Z","snapshot_observed_at":"2026-08-12T01:14:06.357738Z","submitted_at":"2024-12-31T16:01:25Z","title":"A Method for Enhancing the Safety of Large Model Generation Based on Multi-dimensional Attack and Defense","version":1},"reference_index":8,"source":"pdf_text","source_observed_at":"2026-08-10T22:53:20.621497Z"},"links":{"cited_paper":"/paper/2402.01306","citing_paper":"/paper/2501.00517"},"observation_digest":"sha256:e63e2b7a405acac65b794189d9e66bb018f934c5d6059d6cf04ed8e01274232d","observation_id":"77e9c73c-6acf-4760-b655-cab31d36e36a","resolution":{"observed_at":"2026-08-10T22:53:20.621497Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-10T22:53:21.675728Z","title":"Beavertails: Towards improved safety alignment of llm via a human - preference dataset","venue":null,"work_id":"f3bdafea-5022-40a6-847f-bfcbd4404f77","year":2024},"citing_paper":{"arxiv_id":"2501.00517","last_updated":"2024-12-31T16:01:25Z","snapshot_observed_at":"2026-08-12T01:14:06.357738Z","submitted_at":"2024-12-31T16:01:25Z","title":"A Method for Enhancing the Safety of Large Model Generation Based on Multi-dimensional Attack and Defense","version":1},"reference_index":9,"source":"pdf_text","source_observed_at":"2026-08-10T22:53:20.633584Z"},"links":{"citing_paper":"/paper/2501.00517"},"observation_digest":"sha256:dcea287a737b22a8ba4a20f68cf603ea08208376e67a82dec18e64af9bd7da9c","observation_id":"fb1d6030-b4c2-4c56-9435-800cc824a30b","resolution":{"observed_at":"2026-08-10T22:53:21.679484Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-12T06:34:41.77262+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-12T06:34:41.77262+00:00","source":"crossref"},{"observed_at":"2026-08-12T06:34:36.333875+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2310.12773","last_updated":"2023-10-19T14:22:03Z","snapshot_observed_at":"2026-08-02T16:56:38.535065Z","submitted_at":"2023-10-19T14:22:03Z","title":"Safe RLHF: Safe Reinforcement Learning from Human Feedback","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2310.12773","snapshot_observed_at":"2026-08-10T22:53:20.636788Z","title":"Safe RLHF: Safe Reinforcement Learning from Human Feedback","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2501.00517","last_updated":"2024-12-31T16:01:25Z","snapshot_observed_at":"2026-08-12T01:14:06.357738Z","submitted_at":"2024-12-31T16:01:25Z","title":"A Method for Enhancing the Safety of Large Model Generation Based on Multi-dimensional Attack and Defense","version":1},"reference_index":10,"source":"pdf_text","source_observed_at":"2026-08-10T22:53:20.636788Z"},"links":{"cited_paper":"/paper/2310.12773","citing_paper":"/paper/2501.00517"},"observation_digest":"sha256:21be0195bfcdb403eda5d4600f980b921d274ecabd3423952536b340a7c9ffed","observation_id":"aa68c2c4-eda7-44fe-a71c-452ae2bd883d","resolution":{"observed_at":"2026-08-10T22:53:20.636788Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2410.05451","last_updated":"2025-07-03T05:45:41Z","snapshot_observed_at":"2026-08-09T06:23:57.503977Z","submitted_at":"2024-10-07T19:34:35Z","title":"SecAlign: Defending Against Prompt Injection with Preference Optimization","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2410.05451","snapshot_observed_at":"2026-08-10T22:53:20.639962Z","title":"Aligning LLMs to Be Robust Against Prompt Injection","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2501.00517","last_updated":"2024-12-31T16:01:25Z","snapshot_observed_at":"2026-08-12T01:14:06.357738Z","submitted_at":"2024-12-31T16:01:25Z","title":"A Method for Enhancing the Safety of Large Model Generation Based on Multi-dimensional Attack and Defense","version":1},"reference_index":11,"source":"pdf_text","source_observed_at":"2026-08-10T22:53:20.639962Z"},"links":{"cited_paper":"/paper/2410.05451","citing_paper":"/paper/2501.00517"},"observation_digest":"sha256:c85b8fef54183a47c61630140ba0f7955de90972b9f0687dede82b288b1732cb","observation_id":"f0db6692-c76b-48b3-b8b8-14609638632d","resolution":{"observed_at":"2026-08-10T22:53:20.639962Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2309.03883","last_updated":"2024-03-11T02:01:09Z","snapshot_observed_at":"2026-08-06T11:43:44.726916Z","submitted_at":"2023-09-07T17:45:31Z","title":"DoLa: Decoding by Contrasting Layers Improves Factuality in Large Language Models","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2309.03883","snapshot_observed_at":"2026-08-10T22:53:20.644003Z","title":"DoLa: Decoding by Contrasting Layers Improves Factuality in Large Language Models","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2501.00517","last_updated":"2024-12-31T16:01:25Z","snapshot_observed_at":"2026-08-12T01:14:06.357738Z","submitted_at":"2024-12-31T16:01:25Z","title":"A Method for Enhancing the Safety of Large Model Generation Based on Multi-dimensional Attack and Defense","version":1},"reference_index":12,"source":"pdf_text","source_observed_at":"2026-08-10T22:53:20.644003Z"},"links":{"cited_paper":"/paper/2309.03883","citing_paper":"/paper/2501.00517"},"observation_digest":"sha256:7cb49cf45823204788df21fc112600d2ad9bfb699cbca787c843221bc21598cb","observation_id":"c5d680cc-0042-4fe3-b989-cce3c0ef94c0","resolution":{"observed_at":"2026-08-10T22:53:20.644003Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-10T22:53:20.650695Z","title":"Superficial Safety Alignment Hypothesis","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2501.00517","last_updated":"2024-12-31T16:01:25Z","snapshot_observed_at":"2026-08-12T01:14:06.357738Z","submitted_at":"2024-12-31T16:01:25Z","title":"A Method for Enhancing the Safety of Large Model Generation Based on Multi-dimensional Attack and Defense","version":1},"reference_index":13,"source":"pdf_text","source_observed_at":"2026-08-10T22:53:20.650695Z"},"links":{"citing_paper":"/paper/2501.00517"},"observation_digest":"sha256:de05d8764280398df463278c59d2c053dc1ebf532a0beba00d2fe766960b529b","observation_id":"d671b371-d266-4ce3-87f8-5be94a9224b0","resolution":{"observed_at":"2026-08-10T22:53:20.650695Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-10T22:53:21.639389Z","title":"Safety Layers in Aligned Large Language Models: The Key to LLM Security","venue":null,"work_id":"bc6329ac-c87e-40bf-8427-43ada07f6fda","year":2024},"citing_paper":{"arxiv_id":"2501.00517","last_updated":"2024-12-31T16:01:25Z","snapshot_observed_at":"2026-08-12T01:14:06.357738Z","submitted_at":"2024-12-31T16:01:25Z","title":"A Method for Enhancing the Safety of Large Model Generation Based on Multi-dimensional Attack and Defense","version":1},"reference_index":14,"source":"pdf_text","source_observed_at":"2026-08-10T22:53:20.655282Z"},"links":{"citing_paper":"/paper/2501.00517"},"observation_digest":"sha256:7004652746351c059d0d476da9d0ede19f9466a590faff6adc5e154a3a9f0a36","observation_id":"c3bbda7f-1a48-422f-97ce-e35bce0d84c7","resolution":{"observed_at":"2026-08-10T22:53:21.668623Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-12T06:34:41.77262+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-12T06:34:41.77262+00:00","source":"crossref"},{"observed_at":"2026-08-12T06:34:36.333875+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2310.06474","last_updated":"2024-03-04T04:03:54Z","snapshot_observed_at":"2026-08-08T01:21:39.731592Z","submitted_at":"2023-10-10T09:44:06Z","title":"Multilingual Jailbreak Challenges in Large Language Models","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2310.06474","snapshot_observed_at":"2026-08-10T22:53:20.662431Z","title":"Multilingual Jailbreak Challenges in Large Language Models","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2501.00517","last_updated":"2024-12-31T16:01:25Z","snapshot_observed_at":"2026-08-12T01:14:06.357738Z","submitted_at":"2024-12-31T16:01:25Z","title":"A Method for Enhancing the Safety of Large Model Generation Based on Multi-dimensional Attack and Defense","version":1},"reference_index":15,"source":"pdf_text","source_observed_at":"2026-08-10T22:53:20.662431Z"},"links":{"cited_paper":"/paper/2310.06474","citing_paper":"/paper/2501.00517"},"observation_digest":"sha256:5913d69ce2b65b16eebb882589722d43baf392b2a5ec28621106918ad7e121a9","observation_id":"e796599c-3564-4d6a-96ca-4edc5fbbd1cd","resolution":{"observed_at":"2026-08-10T22:53:20.662431Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2309.07045","last_updated":"2024-06-24T04:04:21Z","snapshot_observed_at":"2026-08-09T10:42:31.794884Z","submitted_at":"2023-09-13T15:56:50Z","title":"SafetyBench: Evaluating the Safety of Large Language Models","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2309.07045","snapshot_observed_at":"2026-08-10T22:53:20.710337Z","title":"SafetyBench: Evaluating the Safety of Large Language Models with Multiple Choice Questions","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2501.00517","last_updated":"2024-12-31T16:01:25Z","snapshot_observed_at":"2026-08-12T01:14:06.357738Z","submitted_at":"2024-12-31T16:01:25Z","title":"A Method for Enhancing the Safety of Large Model Generation Based on Multi-dimensional Attack and Defense","version":1},"reference_index":16,"source":"pdf_text","source_observed_at":"2026-08-10T22:53:20.710337Z"},"links":{"cited_paper":"/paper/2309.07045","citing_paper":"/paper/2501.00517"},"observation_digest":"sha256:ebdcf4876f7a2bab1ae974a929ba1b475aaaffac9b19fbf63f5cb883e268fdf1","observation_id":"a82aabb7-4b40-4daa-a556-e8e0d0903a6b","resolution":{"observed_at":"2026-08-10T22:53:20.710337Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2307.09705","last_updated":"2023-07-19T01:22:40Z","snapshot_observed_at":"2026-08-12T10:57:13.152138Z","submitted_at":"2023-07-19T01:22:40Z","title":"CValues: Measuring the Values of Chinese Large Language Models from Safety to Responsibility","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2307.09705","snapshot_observed_at":"2026-08-10T22:53:20.778379Z","title":"CValues: Measuring the Values of Chinese Large Language Models from Safety to Responsibility","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2501.00517","last_updated":"2024-12-31T16:01:25Z","snapshot_observed_at":"2026-08-12T01:14:06.357738Z","submitted_at":"2024-12-31T16:01:25Z","title":"A Method for Enhancing the Safety of Large Model Generation Based on Multi-dimensional Attack and Defense","version":1},"reference_index":17,"source":"pdf_text","source_observed_at":"2026-08-10T22:53:20.778379Z"},"links":{"cited_paper":"/paper/2307.09705","citing_paper":"/paper/2501.00517"},"observation_digest":"sha256:bdfa5b151adf76a0d449d1435b6d4e43969c21532b88595d7a528abcb46c68ab","observation_id":"502968ec-c692-4812-b17d-d6c4964ee2d8","resolution":{"observed_at":"2026-08-10T22:53:20.778379Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2405.14191","last_updated":"2025-04-07T07:52:28Z","snapshot_observed_at":"2026-08-10T20:03:49.769456Z","submitted_at":"2024-05-23T05:34:31Z","title":"S-Eval: Towards Automated and Comprehensive Safety Evaluation for Large Language Models","version":4},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2405.14191","snapshot_observed_at":"2026-08-10T22:53:20.842188Z","title":"S-Eval: Automatic and Adaptive Test Generation for Benchmarking Safety Evaluation of Large Language Models","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2501.00517","last_updated":"2024-12-31T16:01:25Z","snapshot_observed_at":"2026-08-12T01:14:06.357738Z","submitted_at":"2024-12-31T16:01:25Z","title":"A Method for Enhancing the Safety of Large Model Generation Based on Multi-dimensional Attack and Defense","version":1},"reference_index":18,"source":"pdf_text","source_observed_at":"2026-08-10T22:53:20.842188Z"},"links":{"cited_paper":"/paper/2405.14191","citing_paper":"/paper/2501.00517"},"observation_digest":"sha256:36a864c6f744811d2dde84a64042f76d73a3ea71290bdda9843abe4b0397989d","observation_id":"3c2d03b5-39a7-4daf-9977-694f7bab023b","resolution":{"observed_at":"2026-08-10T22:53:20.842188Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-10T22:53:21.606093Z","title":"A Post-Training Enhanced Optimization Approach for Small Language Models","venue":null,"work_id":"d0d57f95-46bc-4733-af73-e86458f1a197","year":2024},"citing_paper":{"arxiv_id":"2501.00517","last_updated":"2024-12-31T16:01:25Z","snapshot_observed_at":"2026-08-12T01:14:06.357738Z","submitted_at":"2024-12-31T16:01:25Z","title":"A Method for Enhancing the Safety of Large Model Generation Based on Multi-dimensional Attack and Defense","version":1},"reference_index":19,"source":"pdf_text","source_observed_at":"2026-08-10T22:53:20.867904Z"},"links":{"citing_paper":"/paper/2501.00517"},"observation_digest":"sha256:a0968c581c818c430a6cdce040466bd0600d461363a4d321b3307cfc06db7c16","observation_id":"4e73d1ec-f0c0-4127-af19-9d94cf0b1394","resolution":{"observed_at":"2026-08-10T22:53:21.614718Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-12T06:34:41.77262+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-12T06:34:41.77262+00:00","source":"crossref"},{"observed_at":"2026-08-12T06:34:36.333875+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2304.10436","last_updated":"2023-04-20T16:27:35Z","snapshot_observed_at":"2026-07-06T15:17:59.640907Z","submitted_at":"2023-04-20T16:27:35Z","title":"Safety Assessment of Chinese Large Language Models","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2304.10436","snapshot_observed_at":"2026-08-10T22:53:20.901818Z","title":"Safety Assessment of Chinese Large Language Models","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2501.00517","last_updated":"2024-12-31T16:01:25Z","snapshot_observed_at":"2026-08-12T01:14:06.357738Z","submitted_at":"2024-12-31T16:01:25Z","title":"A Method for Enhancing the Safety of Large Model Generation Based on Multi-dimensional Attack and Defense","version":1},"reference_index":20,"source":"pdf_text","source_observed_at":"2026-08-10T22:53:20.901818Z"},"links":{"cited_paper":"/paper/2304.10436","citing_paper":"/paper/2501.00517"},"observation_digest":"sha256:940fe3adfab1125e410196a6019eb92254009759e6bf60bef8e49df44173db7f","observation_id":"4a17f408-296d-49d9-a27a-5a28bccd50cc","resolution":{"observed_at":"2026-08-10T22:53:20.901818Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-10T22:53:21.460457Z","title":null,"venue":null,"work_id":"0e670fb9-c565-4c37-b98d-94554ce0b6d9","year":null},"citing_paper":{"arxiv_id":"2501.00517","last_updated":"2024-12-31T16:01:25Z","snapshot_observed_at":"2026-08-12T01:14:06.357738Z","submitted_at":"2024-12-31T16:01:25Z","title":"A Method for Enhancing the Safety of Large Model Generation Based on Multi-dimensional Attack and Defense","version":1},"reference_index":21,"source":"pdf_text","source_observed_at":"2026-08-10T22:53:20.922825Z"},"links":{"citing_paper":"/paper/2501.00517"},"observation_digest":"sha256:3f69548677e4e759ac9dd8817e47f5be71c447a2979c9055bab9a3cdcf987bce","observation_id":"b273681b-e41e-46a8-9ac5-55a50aa3bb35","resolution":{"observed_at":"2026-08-10T22:53:21.546684Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-12T06:34:41.77262+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-12T06:34:41.77262+00:00","source":"crossref"},{"observed_at":"2026-08-12T06:34:36.333875+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-10T22:53:21.274270Z","title":null,"venue":null,"work_id":"d1bd44f7-56f2-4569-893f-314b6e42de2c","year":null},"citing_paper":{"arxiv_id":"2501.00517","last_updated":"2024-12-31T16:01:25Z","snapshot_observed_at":"2026-08-12T01:14:06.357738Z","submitted_at":"2024-12-31T16:01:25Z","title":"A Method for Enhancing the Safety of Large Model Generation Based on Multi-dimensional Attack and Defense","version":1},"reference_index":22,"source":"pdf_text","source_observed_at":"2026-08-10T22:53:20.940426Z"},"links":{"citing_paper":"/paper/2501.00517"},"observation_digest":"sha256:056ffd3de7ee2c3909ea49adf57221baa658b2255fe9e9c7a934769971f6cbdb","observation_id":"3e6070bb-a32a-4651-8914-a3630ceefd27","resolution":{"observed_at":"2026-08-10T22:53:21.344711Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-12T06:34:41.77262+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-12T06:34:41.77262+00:00","source":"crossref"},{"observed_at":"2026-08-12T06:34:36.333875+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-10T22:53:21.263713Z","title":null,"venue":null,"work_id":"c188e34b-52de-4721-8cc0-9a6c74cfa78a","year":null},"citing_paper":{"arxiv_id":"2501.00517","last_updated":"2024-12-31T16:01:25Z","snapshot_observed_at":"2026-08-12T01:14:06.357738Z","submitted_at":"2024-12-31T16:01:25Z","title":"A Method for Enhancing the Safety of Large Model Generation Based on Multi-dimensional Attack and Defense","version":1},"reference_index":23,"source":"pdf_text","source_observed_at":"2026-08-10T22:53:20.953353Z"},"links":{"citing_paper":"/paper/2501.00517"},"observation_digest":"sha256:b41e0b2af133b3e3c91667e74940b5f61e4945758e07a26d23722f18c17643f2","observation_id":"d97dc54b-3a49-48ad-9948-f7eb9dc9cae8","resolution":{"observed_at":"2026-08-10T22:53:21.267410Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-12T06:34:41.77262+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-12T06:34:41.77262+00:00","source":"crossref"},{"observed_at":"2026-08-12T06:34:36.333875+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-10T22:53:21.253090Z","title":null,"venue":null,"work_id":"68bea2b0-8771-44fa-82d0-bb7049fe0f7b","year":null},"citing_paper":{"arxiv_id":"2501.00517","last_updated":"2024-12-31T16:01:25Z","snapshot_observed_at":"2026-08-12T01:14:06.357738Z","submitted_at":"2024-12-31T16:01:25Z","title":"A Method for Enhancing the Safety of Large Model Generation Based on Multi-dimensional Attack and Defense","version":1},"reference_index":24,"source":"pdf_text","source_observed_at":"2026-08-10T22:53:20.956873Z"},"links":{"citing_paper":"/paper/2501.00517"},"observation_digest":"sha256:468fa855b2ead3bd00d24447780a0282d4766bc185f42622b4a8685da6c6de94","observation_id":"e3800634-b248-4896-93cb-531948b027f4","resolution":{"observed_at":"2026-08-10T22:53:21.256475Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-12T06:34:41.77262+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-12T06:34:41.77262+00:00","source":"crossref"},{"observed_at":"2026-08-12T06:34:36.333875+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-10T22:53:21.243320Z","title":null,"venue":null,"work_id":"7106a9d2-87f6-460a-b9bb-c6a0c6b4e04b","year":null},"citing_paper":{"arxiv_id":"2501.00517","last_updated":"2024-12-31T16:01:25Z","snapshot_observed_at":"2026-08-12T01:14:06.357738Z","submitted_at":"2024-12-31T16:01:25Z","title":"A Method for Enhancing the Safety of Large Model Generation Based on Multi-dimensional Attack and Defense","version":1},"reference_index":25,"source":"pdf_text","source_observed_at":"2026-08-10T22:53:20.960308Z"},"links":{"citing_paper":"/paper/2501.00517"},"observation_digest":"sha256:e898802b75ddceca7d707f83271ce1be3e063139b5122b637e4bb88f55300266","observation_id":"c0877385-0913-4394-9198-bce4417cd12c","resolution":{"observed_at":"2026-08-10T22:53:21.246408Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-12T06:34:41.77262+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-12T06:34:41.77262+00:00","source":"crossref"},{"observed_at":"2026-08-12T06:34:36.333875+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-10T22:53:21.230818Z","title":null,"venue":null,"work_id":"ee2da42b-57c5-493e-9b21-d787c4f41f84","year":null},"citing_paper":{"arxiv_id":"2501.00517","last_updated":"2024-12-31T16:01:25Z","snapshot_observed_at":"2026-08-12T01:14:06.357738Z","submitted_at":"2024-12-31T16:01:25Z","title":"A Method for Enhancing the Safety of Large Model Generation Based on Multi-dimensional Attack and Defense","version":1},"reference_index":26,"source":"pdf_text","source_observed_at":"2026-08-10T22:53:20.963955Z"},"links":{"citing_paper":"/paper/2501.00517"},"observation_digest":"sha256:f1f90ae047c28c0dfcff2e2d4b33499eade833a4ff763d7b567aa963b2da000d","observation_id":"29c8aaf3-22fd-4f92-ac75-196ad80bdcca","resolution":{"observed_at":"2026-08-10T22:53:21.236077Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-12T06:34:41.77262+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-12T06:34:41.77262+00:00","source":"crossref"},{"observed_at":"2026-08-12T06:34:36.333875+00:00","source":"retraction_watch"}],"state":"measured"}}],"paper":{"arxiv_id":"2501.00517","last_updated":"2024-12-31T16:01:25Z","latest_version":1,"primary_category":"cs.CR","snapshot_observed_at":"2026-08-12T01:14:06.357738Z","submitted_at":"2024-12-31T16:01:25Z","title":"A Method for Enhancing the Safety of Large Model Generation Based on Multi-dimensional Attack and Defense"},"reference_resolution":{"displayed":26,"state_counts":{"malformed_identifier":0,"metadata_mismatch":0,"parse_uncertain":0,"unresolved":23,"verified_exact":0,"verified_fuzzy":3},"total_outbound_references":26},"refusal":"A citation records a reference. It does not transfer a finding from one paper to another.","schema":"pith.paper-citation-record.v1","standing_sources":[{"observed_at":"2026-08-12T06:34:41.77262+00:00","source":"crossref"},{"observed_at":"2026-08-12T06:34:36.333875+00:00","source":"retraction_watch"}],"thesis":"As of 12 August 2026, this Paper Citation Record lists 26 of 26 outbound references and 0 inbound Pith citation observations for arXiv:2501.00517."}