{"as_of":"2026-08-07T10:24:00Z","caps":{"database_statements":6,"inbound":100,"outbound":100},"context_digest":"sha256:12c3bf8b62bf886f0a11a3922bc0c97e352533947ba6c4e0ab79face74ca370b","coverage":[{"denominator":0,"lane":"reference_resolution","note":"Typed states for the displayed outbound observations.","records_observed":0,"source":"paper_references, paper_reference_links","source_observed_at":null,"state":"measured"},{"denominator":36,"lane":"standing_notices","note":"One-hop event checks from named stored sources.","records_observed":36,"source":"scholarly_work_events, retraction_status_cache","source_observed_at":"2026-08-07T06:34:17.273281+00:00","state":"measured"},{"denominator":36,"lane":"inbound_itemization","note":"Pith citing papers itemized under the disclosed page cap.","records_observed":36,"source":"paper_references, paper_reference_links","source_observed_at":"2026-08-07T05:42:31.720393Z","state":"measured"},{"denominator":1,"lane":"external_citation_measurements","note":"A source-named dated measurement, never combined with another source.","records_observed":1,"source":"arxiv_reference","source_observed_at":"2026-08-05T02:28:24.338817Z","state":"measured"}],"external_citation_measurements":[{"count":50,"observed_at":"2026-08-05T02:28:24.338817Z","source":"arxiv_reference"}],"inbound":[{"citation":{"cited_paper":{"arxiv_id":"2306.04528","last_updated":"2024-07-16T07:29:49Z","snapshot_observed_at":"2026-07-06T15:39:49.273453Z","submitted_at":"2023-06-07T15:37:00Z","title":"PromptRobust: Towards Evaluating the Robustness of Large Language Models on Adversarial Prompts","version":5},"cited_work":{"arxiv_id":"2306.04528","doi":"10.48550/arxiv.2306.04528","metadata_source":"arxiv_reference","pith_arxiv_id":"2306.04528","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Promptbench: Towards evaluating the robustness of large language models on adversarial prompts","venue":"arXiv (Cornell University)","work_id":"e65e5731-8f09-44e1-94d5-1674f9e08988","year":2024},"citing_paper":{"arxiv_id":"2307.15043","last_updated":"2023-12-20T20:48:57Z","snapshot_observed_at":"2026-07-06T15:59:23.019044Z","submitted_at":"2023-07-27T17:49:12Z","title":"Universal and Transferable Adversarial Attacks on Aligned Language Models","version":2},"reference_index":29,"source":"pdf_text","source_observed_at":"2026-05-24T07:42:09.112946Z"},"links":{"cited_paper":"/paper/2306.04528","citing_paper":"/paper/2307.15043"},"observation_digest":"sha256:5d33d012f364e2fea20cb9af0704cb83fe5131197a5b9301088b702d5add6075","observation_id":"2b2e6116-13c7-4016-8e76-f738ee64014d","resolution":{"observed_at":"2026-05-24T07:44:08.483210Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2306.04528","last_updated":"2024-07-16T07:29:49Z","snapshot_observed_at":"2026-07-06T15:39:49.273453Z","submitted_at":"2023-06-07T15:37:00Z","title":"PromptRobust: Towards Evaluating the Robustness of Large Language Models on Adversarial Prompts","version":5},"cited_work":{"arxiv_id":"2306.04528","doi":"10.48550/arxiv.2306.04528","metadata_source":"arxiv_reference","pith_arxiv_id":"2306.04528","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Promptbench: Towards evaluating the robustness of large language models on adversarial prompts","venue":"arXiv (Cornell University)","work_id":"e65e5731-8f09-44e1-94d5-1674f9e08988","year":2024},"citing_paper":{"arxiv_id":"2309.00614","last_updated":"2023-09-04T17:47:36Z","snapshot_observed_at":"2026-07-06T16:13:23.343694Z","submitted_at":"2023-09-01T17:59:44Z","title":"Baseline Defenses for Adversarial Attacks Against Aligned Language Models","version":2},"reference_index":65,"source":"arxiv_source","source_observed_at":"2026-05-13T23:24:39.835347Z"},"links":{"cited_paper":"/paper/2306.04528","citing_paper":"/paper/2309.00614"},"observation_digest":"sha256:f50ffba53f6530bb7d631550312c6585f0d2b92c98730b3aed409f952b0c12a6","observation_id":"d4a95b09-3832-4a2f-bc31-9fac787648ee","resolution":{"observed_at":"2026-05-13T23:24:40.082946Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2306.04528","last_updated":"2024-07-16T07:29:49Z","snapshot_observed_at":"2026-07-06T15:39:49.273453Z","submitted_at":"2023-06-07T15:37:00Z","title":"PromptRobust: Towards Evaluating the Robustness of Large Language Models on Adversarial Prompts","version":5},"cited_work":{"arxiv_id":"2306.04528","doi":"10.48550/arxiv.2306.04528","metadata_source":"arxiv_reference","pith_arxiv_id":"2306.04528","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Promptbench: Towards evaluating the robustness of large language models on adversarial prompts","venue":"arXiv (Cornell University)","work_id":"e65e5731-8f09-44e1-94d5-1674f9e08988","year":2024},"citing_paper":{"arxiv_id":"2309.01219","last_updated":"2025-09-14T09:34:46Z","snapshot_observed_at":"2026-07-06T16:13:46.112815Z","submitted_at":"2023-09-03T16:56:48Z","title":"Siren's Song in the AI Ocean: A Survey on Hallucination in Large Language Models","version":3},"reference_index":30,"source":"pdf_text","source_observed_at":"2026-05-12T14:21:16.453610Z"},"links":{"cited_paper":"/paper/2306.04528","citing_paper":"/paper/2309.01219"},"observation_digest":"sha256:8751d121cea0ac95efb10947c0c881ee5d97ee9f0eaaf5c4f4fdbc610bc7d560","observation_id":"c781a3b8-e332-4632-8ca8-f7ec654ec293","resolution":{"observed_at":"2026-05-12T14:21:16.477948Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2306.04528","last_updated":"2024-07-16T07:29:49Z","snapshot_observed_at":"2026-07-06T15:39:49.273453Z","submitted_at":"2023-06-07T15:37:00Z","title":"PromptRobust: Towards Evaluating the Robustness of Large Language Models on Adversarial Prompts","version":5},"cited_work":{"arxiv_id":"2306.04528","doi":"10.48550/arxiv.2306.04528","metadata_source":"arxiv_reference","pith_arxiv_id":"2306.04528","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Promptbench: Towards evaluating the robustness of large language models on adversarial prompts","venue":"arXiv (Cornell University)","work_id":"e65e5731-8f09-44e1-94d5-1674f9e08988","year":2024},"citing_paper":{"arxiv_id":"2309.08532","last_updated":"2025-05-01T11:56:52Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2023-09-15T16:50:09Z","title":"EvoPrompt: Connecting LLMs with Evolutionary Algorithms Yields Powerful Prompt Optimizers","version":3},"reference_index":137,"source":"arxiv_source","source_observed_at":"2026-05-16T06:11:49.475825Z"},"links":{"cited_paper":"/paper/2306.04528","citing_paper":"/paper/2309.08532"},"observation_digest":"sha256:e8791e278e8ef0d8b153ebace670bd76af3a9bcd35d81f621f6922557ffbec5a","observation_id":"f2ed5e7c-03b1-4f84-a0fd-c4e85a3b4c55","resolution":{"observed_at":"2026-05-16T06:11:49.585490Z","resolver_source":"arxiv_id","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2306.04528","last_updated":"2024-07-16T07:29:49Z","snapshot_observed_at":"2026-07-06T15:39:49.273453Z","submitted_at":"2023-06-07T15:37:00Z","title":"PromptRobust: Towards Evaluating the Robustness of Large Language Models on Adversarial Prompts","version":5},"cited_work":{"arxiv_id":"2306.04528","doi":"10.48550/arxiv.2306.04528","metadata_source":"arxiv_reference","pith_arxiv_id":"2306.04528","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Promptbench: Towards evaluating the robustness of large language models on adversarial prompts","venue":"arXiv (Cornell University)","work_id":"e65e5731-8f09-44e1-94d5-1674f9e08988","year":2024},"citing_paper":{"arxiv_id":"2401.05561","last_updated":"2024-09-30T10:17:12Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2024-01-10T22:07:21Z","title":"TrustLLM: Trustworthiness in Large Language Models","version":6},"reference_index":172,"source":"pdf_text","source_observed_at":"2026-05-18T11:17:08.108565Z"},"links":{"cited_paper":"/paper/2306.04528","citing_paper":"/paper/2401.05561"},"observation_digest":"sha256:0b580b7d0e60d86c70515926b68901cba41e81e90d1a4c23008835f22ad6fcc7","observation_id":"6503384a-e3c4-4163-9717-5493a50b0b00","resolution":{"observed_at":"2026-05-18T11:17:08.543894Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2306.04528","last_updated":"2024-07-16T07:29:49Z","snapshot_observed_at":"2026-07-06T15:39:49.273453Z","submitted_at":"2023-06-07T15:37:00Z","title":"PromptRobust: Towards Evaluating the Robustness of Large Language Models on Adversarial Prompts","version":5},"cited_work":{"arxiv_id":"2306.04528","doi":"10.48550/arxiv.2306.04528","metadata_source":"arxiv_reference","pith_arxiv_id":"2306.04528","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Promptbench: Towards evaluating the robustness of large language models on adversarial prompts","venue":"arXiv (Cornell University)","work_id":"e65e5731-8f09-44e1-94d5-1674f9e08988","year":2024},"citing_paper":{"arxiv_id":"2402.06922","last_updated":"2026-04-21T14:06:34Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2024-02-10T11:07:24Z","title":"Whispers in the Machine: Confidentiality in Agentic Systems","version":5},"reference_index":30,"source":"pdf_text","source_observed_at":"2026-05-24T03:59:03.972043Z"},"links":{"cited_paper":"/paper/2306.04528","citing_paper":"/paper/2402.06922"},"observation_digest":"sha256:c51139ca5f29885d235910c579819a0a443e4f4c9752b3569973d6ee9d470010","observation_id":"f5db751b-c80a-4b12-a3b9-929b9718e4c5","resolution":{"observed_at":"2026-05-24T04:03:53.865582Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2306.04528","last_updated":"2024-07-16T07:29:49Z","snapshot_observed_at":"2026-07-06T15:39:49.273453Z","submitted_at":"2023-06-07T15:37:00Z","title":"PromptRobust: Towards Evaluating the Robustness of Large Language Models on Adversarial Prompts","version":5},"cited_work":{"arxiv_id":"2306.04528","doi":"10.48550/arxiv.2306.04528","metadata_source":"arxiv_reference","pith_arxiv_id":"2306.04528","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Promptbench: Towards evaluating the robustness of large language models on adversarial prompts","venue":"arXiv (Cornell University)","work_id":"e65e5731-8f09-44e1-94d5-1674f9e08988","year":2024},"citing_paper":{"arxiv_id":"2406.04244","last_updated":"2024-06-06T16:41:39Z","snapshot_observed_at":"2026-07-30T15:43:06.151242Z","submitted_at":"2024-06-06T16:41:39Z","title":"Benchmark Data Contamination of Large Language Models: A Survey","version":1},"reference_index":191,"source":"pdf_text","source_observed_at":"2026-05-22T23:10:40.420241Z"},"links":{"cited_paper":"/paper/2306.04528","citing_paper":"/paper/2406.04244"},"observation_digest":"sha256:9615162639dcb39f24a739308837a2cf8003750dcab377607913a39c50940286","observation_id":"6f1e8eaf-774d-4692-a572-11f1d10e44ce","resolution":{"observed_at":"2026-05-22T23:10:41.169831Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2306.04528","last_updated":"2024-07-16T07:29:49Z","snapshot_observed_at":"2026-07-06T15:39:49.273453Z","submitted_at":"2023-06-07T15:37:00Z","title":"PromptRobust: Towards Evaluating the Robustness of Large Language Models on Adversarial Prompts","version":5},"cited_work":{"arxiv_id":"2306.04528","doi":"10.48550/arxiv.2306.04528","metadata_source":"arxiv_reference","pith_arxiv_id":"2306.04528","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Promptbench: Towards evaluating the robustness of large language models on adversarial prompts","venue":"arXiv (Cornell University)","work_id":"e65e5731-8f09-44e1-94d5-1674f9e08988","year":2024},"citing_paper":{"arxiv_id":"2409.10102","last_updated":"2026-05-16T07:15:07Z","snapshot_observed_at":"2026-08-04T19:12:49.508911Z","submitted_at":"2024-09-16T09:06:44Z","title":"Trustworthiness in Retrieval-Augmented Generation Systems: A Survey","version":2},"reference_index":92,"source":"pdf_text","source_observed_at":"2026-05-23T21:08:11.787013Z"},"links":{"cited_paper":"/paper/2306.04528","citing_paper":"/paper/2409.10102"},"observation_digest":"sha256:9fbd18926faf8e00026d625467bc689da382c7a4450c95e5a67afc09aa7c0c5a","observation_id":"b7e0043e-1fed-4f3a-b51b-9971c85280c1","resolution":{"observed_at":"2026-05-23T21:08:25.916992Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2306.04528","last_updated":"2024-07-16T07:29:49Z","snapshot_observed_at":"2026-07-06T15:39:49.273453Z","submitted_at":"2023-06-07T15:37:00Z","title":"PromptRobust: Towards Evaluating the Robustness of Large Language Models on Adversarial Prompts","version":5},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2306.04528","snapshot_observed_at":"2026-08-07T05:31:52.719510Z","title":"Promptbench: Towards evaluating the robustness of large language models on adversarial prompts.arXiv preprint arXiv:2306.04528, 2023","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2506.07645","last_updated":"2025-06-09T11:09:39Z","snapshot_observed_at":"2026-08-07T05:26:31.212953Z","submitted_at":"2025-06-09T11:09:39Z","title":"Evaluating LLMs Robustness in Less Resourced Languages with Proxy Models","version":1},"reference_index":31,"source":"pdf_text","source_observed_at":"2026-08-07T05:31:52.719510Z"},"links":{"cited_paper":"/paper/2306.04528","citing_paper":"/paper/2506.07645"},"observation_digest":"sha256:5951f0f0107d4c0efbb82159308d164016148a4d4bfabf13fda008ed90e832fd","observation_id":"a07116a8-6252-46a8-9d49-de7c94e931c5","resolution":{"observed_at":"2026-08-07T05:31:52.719510Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2306.04528","last_updated":"2024-07-16T07:29:49Z","snapshot_observed_at":"2026-07-06T15:39:49.273453Z","submitted_at":"2023-06-07T15:37:00Z","title":"PromptRobust: Towards Evaluating the Robustness of Large Language Models on Adversarial Prompts","version":5},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2306.04528","snapshot_observed_at":"2026-08-07T05:42:31.720393Z","title":null,"venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2506.11111","last_updated":"2025-07-09T06:18:33Z","snapshot_observed_at":"2026-08-07T05:37:03.363226Z","submitted_at":"2025-06-08T16:20:12Z","title":"Evaluating and Improving Robustness in Large Language Models: A Survey and Future Directions","version":2},"reference_index":245,"source":"pdf_text","source_observed_at":"2026-08-07T05:42:31.720393Z"},"links":{"cited_paper":"/paper/2306.04528","citing_paper":"/paper/2506.11111"},"observation_digest":"sha256:29c2ca82918d8295ffb83d5b947089dfa379a9abe6db280b0eb6eb3033debfbd","observation_id":"aa8fdce4-2789-41cc-a44b-cb4d9432a2dc","resolution":{"observed_at":"2026-08-07T05:42:31.720393Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2306.04528","last_updated":"2024-07-16T07:29:49Z","snapshot_observed_at":"2026-07-06T15:39:49.273453Z","submitted_at":"2023-06-07T15:37:00Z","title":"PromptRobust: Towards Evaluating the Robustness of Large Language Models on Adversarial Prompts","version":5},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2306.04528","snapshot_observed_at":"2026-08-06T18:55:55.041443Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2507.06956","last_updated":"2025-07-09T15:39:17Z","snapshot_observed_at":"2026-08-06T18:48:24.090150Z","submitted_at":"2025-07-09T15:39:17Z","title":"Investigating the Robustness of Retrieval-Augmented Generation at the Query Level","version":1},"reference_index":44,"source":"arxiv_source","source_observed_at":"2026-08-06T18:55:55.041443Z"},"links":{"cited_paper":"/paper/2306.04528","citing_paper":"/paper/2507.06956"},"observation_digest":"sha256:b3d17ed4a6733c1fc55a548582f8df7fab9b821e30bccede780b77a09582bc74","observation_id":"92ea6742-49b2-4302-9880-1e77754532cb","resolution":{"observed_at":"2026-08-06T18:55:55.041443Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2306.04528","last_updated":"2024-07-16T07:29:49Z","snapshot_observed_at":"2026-07-06T15:39:49.273453Z","submitted_at":"2023-06-07T15:37:00Z","title":"PromptRobust: Towards Evaluating the Robustness of Large Language Models on Adversarial Prompts","version":5},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2306.04528","snapshot_observed_at":"2026-08-06T17:45:58.463121Z","title":"Promptbench: Towards evaluating the robustness of large language models on adversarial prompts,","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2507.10054","last_updated":"2025-07-23T10:43:29Z","snapshot_observed_at":"2026-08-07T08:22:10.123810Z","submitted_at":"2025-07-14T08:36:26Z","title":"Explicit Vulnerability Generation with LLMs: An Investigation Beyond Adversarial Attacks","version":2},"reference_index":8,"source":"pdf_text","source_observed_at":"2026-08-06T17:45:58.463121Z"},"links":{"cited_paper":"/paper/2306.04528","citing_paper":"/paper/2507.10054"},"observation_digest":"sha256:2e25a5aef7ecd4a1ffabb363f0588a85afbfc45327204916148b66807f67f3a9","observation_id":"47108bc9-e6b0-48df-a5cc-ffe314623e52","resolution":{"observed_at":"2026-08-06T17:45:58.463121Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2306.04528","last_updated":"2024-07-16T07:29:49Z","snapshot_observed_at":"2026-07-06T15:39:49.273453Z","submitted_at":"2023-06-07T15:37:00Z","title":"PromptRobust: Towards Evaluating the Robustness of Large Language Models on Adversarial Prompts","version":5},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2306.04528","snapshot_observed_at":"2026-08-06T14:57:06.882060Z","title":"PromptBench: Towards evaluating the robustness of Large Language Models on adversarial prompts.arXiv: 2306.04528 [cs.CL], June 2023","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2507.17257","last_updated":"2025-07-23T06:56:15Z","snapshot_observed_at":"2026-08-06T18:04:24.113657Z","submitted_at":"2025-07-23T06:56:15Z","title":"Agent Identity Evals: Measuring Agentic Identity","version":1},"reference_index":77,"source":"pdf_text","source_observed_at":"2026-08-06T14:57:06.882060Z"},"links":{"cited_paper":"/paper/2306.04528","citing_paper":"/paper/2507.17257"},"observation_digest":"sha256:ad197adb354482d8b57ab5babc2939f8a506743df0f1ec4fd36482f4ea0ef337","observation_id":"15fb529e-3caf-4a08-bb0b-0efca3b5e99f","resolution":{"observed_at":"2026-08-06T14:57:06.882060Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2306.04528","last_updated":"2024-07-16T07:29:49Z","snapshot_observed_at":"2026-07-06T15:39:49.273453Z","submitted_at":"2023-06-07T15:37:00Z","title":"PromptRobust: Towards Evaluating the Robustness of Large Language Models on Adversarial Prompts","version":5},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2306.04528","snapshot_observed_at":"2026-08-05T05:59:28.243698Z","title":"Promptrobust: Towards evaluating the robustness of large language models on adversarial prompts, 2024","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2509.04615","last_updated":"2025-09-04T18:59:07Z","snapshot_observed_at":"2026-08-05T05:59:27.514485Z","submitted_at":"2025-09-04T18:59:07Z","title":"Breaking to Build: A Threat Model of Prompt-Based Attacks for Securing LLMs","version":1},"reference_index":6,"source":"pdf_text","source_observed_at":"2026-08-05T05:59:28.243698Z"},"links":{"cited_paper":"/paper/2306.04528","citing_paper":"/paper/2509.04615"},"observation_digest":"sha256:449e9e99776876ed0651aec6a55f12f653c8f159d3b8b41056b266305677bff5","observation_id":"7b63df69-81fe-40e7-a062-bae3a724ff7f","resolution":{"observed_at":"2026-08-05T05:59:28.243698Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2306.04528","last_updated":"2024-07-16T07:29:49Z","snapshot_observed_at":"2026-07-06T15:39:49.273453Z","submitted_at":"2023-06-07T15:37:00Z","title":"PromptRobust: Towards Evaluating the Robustness of Large Language Models on Adversarial Prompts","version":5},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2306.04528","snapshot_observed_at":"2026-08-03T06:46:26.373904Z","title":"Promptbench: Towards eval- uating the robustness of large language models on adversarial prompts.arXiv preprint arXiv:2306.04528, 2023","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2601.22025","last_updated":"2026-06-09T23:57:32Z","snapshot_observed_at":"2026-08-05T13:58:45.670760Z","submitted_at":"2026-01-29T17:32:34Z","title":"When Generic Prompt Improvements Hurt: Evaluation-Driven Iteration for LLM Applications","version":2},"reference_index":28,"source":"pdf_text","source_observed_at":"2026-08-03T06:46:26.373904Z"},"links":{"cited_paper":"/paper/2306.04528","citing_paper":"/paper/2601.22025"},"observation_digest":"sha256:ee6bb5b09be21c91e6a592c55ba311785429c138dbd7987e70b4697d7c087590","observation_id":"0495afe7-d37d-4a4b-8fe3-b7e7c29526c7","resolution":{"observed_at":"2026-08-03T06:46:26.373904Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2306.04528","last_updated":"2024-07-16T07:29:49Z","snapshot_observed_at":"2026-07-06T15:39:49.273453Z","submitted_at":"2023-06-07T15:37:00Z","title":"PromptRobust: Towards Evaluating the Robustness of Large Language Models on Adversarial Prompts","version":5},"cited_work":{"arxiv_id":"2306.04528","doi":"10.48550/arxiv.2306.04528","metadata_source":"arxiv_reference","pith_arxiv_id":"2306.04528","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Promptbench: Towards evaluating the robustness of large language models on adversarial prompts","venue":"arXiv (Cornell University)","work_id":"e65e5731-8f09-44e1-94d5-1674f9e08988","year":2024},"citing_paper":{"arxiv_id":"2603.03332","last_updated":"2026-04-16T23:41:41Z","snapshot_observed_at":"2026-08-04T02:21:24.392691Z","submitted_at":"2026-02-11T03:11:30Z","title":"Fragile Thoughts: How Large Language Models Handle Chain-of-Thought Perturbations","version":3},"reference_index":12,"source":"pdf_text","source_observed_at":"2026-05-16T03:43:18.987241Z"},"links":{"cited_paper":"/paper/2306.04528","citing_paper":"/paper/2603.03332"},"observation_digest":"sha256:968330fc70b8648c4634a3f268c32aebc22cbdaf87fce400ab8df252343a0763","observation_id":"3bf2096c-140c-4017-84eb-0dd8394f2a8b","resolution":{"observed_at":"2026-05-16T03:47:15.338790Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2306.04528","last_updated":"2024-07-16T07:29:49Z","snapshot_observed_at":"2026-07-06T15:39:49.273453Z","submitted_at":"2023-06-07T15:37:00Z","title":"PromptRobust: Towards Evaluating the Robustness of Large Language Models on Adversarial Prompts","version":5},"cited_work":{"arxiv_id":"2306.04528","doi":"10.48550/arxiv.2306.04528","metadata_source":"arxiv_reference","pith_arxiv_id":"2306.04528","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Promptbench: Towards evaluating the robustness of large language models on adversarial prompts","venue":"arXiv (Cornell University)","work_id":"e65e5731-8f09-44e1-94d5-1674f9e08988","year":2024},"citing_paper":{"arxiv_id":"2603.10477","last_updated":"2026-04-08T07:03:59Z","snapshot_observed_at":"2026-08-03T13:46:05.416163Z","submitted_at":"2026-03-11T07:00:59Z","title":"PEEM: Prompt Engineering Evaluation Metrics for Interpretable Joint Evaluation of Prompts and Responses","version":2},"reference_index":43,"source":"pdf_text","source_observed_at":"2026-05-15T13:57:41.428695Z"},"links":{"cited_paper":"/paper/2306.04528","citing_paper":"/paper/2603.10477"},"observation_digest":"sha256:b58284a95f48bd07d070d5e2510bb324060b3d868622ea17c9059df7d9ef7915","observation_id":"abd144a7-b5e5-4d2b-8c87-04c8cdef2ffd","resolution":{"observed_at":"2026-05-15T14:00:03.073494Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2306.04528","last_updated":"2024-07-16T07:29:49Z","snapshot_observed_at":"2026-07-06T15:39:49.273453Z","submitted_at":"2023-06-07T15:37:00Z","title":"PromptRobust: Towards Evaluating the Robustness of Large Language Models on Adversarial Prompts","version":5},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2306.04528","snapshot_observed_at":"2026-07-13T14:38:58.879673Z","title":"Promptbench: Towards evaluating the robustness of large language models on adversarial prompts,","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2604.01039","last_updated":"2026-06-05T19:49:24Z","snapshot_observed_at":"2026-08-02T23:40:30.095065Z","submitted_at":"2026-04-01T15:45:56Z","title":"Automated Framework to Evaluate and Harden LLM System Instructions against Encoding Attacks","version":2},"reference_index":21,"source":"pdf_text","source_observed_at":"2026-07-13T14:38:58.879673Z"},"links":{"cited_paper":"/paper/2306.04528","citing_paper":"/paper/2604.01039"},"observation_digest":"sha256:74f57ffcd2a5f13bdc8cdce03887ef693f089a72e5c175fa31f8d1013d0e3f4d","observation_id":"282e714c-1f2b-4b3a-ac98-aedba4179bcb","resolution":{"observed_at":"2026-07-13T14:38:58.879673Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2306.04528","last_updated":"2024-07-16T07:29:49Z","snapshot_observed_at":"2026-07-06T15:39:49.273453Z","submitted_at":"2023-06-07T15:37:00Z","title":"PromptRobust: Towards Evaluating the Robustness of Large Language Models on Adversarial Prompts","version":5},"cited_work":{"arxiv_id":"2306.04528","doi":"10.48550/arxiv.2306.04528","metadata_source":"arxiv_reference","pith_arxiv_id":"2306.04528","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Promptbench: Towards evaluating the robustness of large language models on adversarial prompts","venue":"arXiv (Cornell University)","work_id":"e65e5731-8f09-44e1-94d5-1674f9e08988","year":2024},"citing_paper":{"arxiv_id":"2604.16421","last_updated":"2026-04-03T11:36:49Z","snapshot_observed_at":"2026-07-06T23:03:43.854612Z","submitted_at":"2026-04-03T11:36:49Z","title":"Measuring Representation Robustness in Large Language Models for Geometry","version":1},"reference_index":40,"source":"pdf_text","source_observed_at":"2026-05-13T19:35:32.531660Z"},"links":{"cited_paper":"/paper/2306.04528","citing_paper":"/paper/2604.16421"},"observation_digest":"sha256:44be8abd4ee197313aef2b0ea81fe9cc7ae042cd63600d333991d658a9df1fe8","observation_id":"ed3491bb-86d1-4f3e-aa00-134f2634b233","resolution":{"observed_at":"2026-05-13T19:38:10.484726Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2306.04528","last_updated":"2024-07-16T07:29:49Z","snapshot_observed_at":"2026-07-06T15:39:49.273453Z","submitted_at":"2023-06-07T15:37:00Z","title":"PromptRobust: Towards Evaluating the Robustness of Large Language Models on Adversarial Prompts","version":5},"cited_work":{"arxiv_id":"2306.04528","doi":"10.48550/arxiv.2306.04528","metadata_source":"arxiv_reference","pith_arxiv_id":"2306.04528","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Promptbench: Towards evaluating the robustness of large language models on adversarial prompts","venue":"arXiv (Cornell University)","work_id":"e65e5731-8f09-44e1-94d5-1674f9e08988","year":2024},"citing_paper":{"arxiv_id":"2604.23135","last_updated":"2026-05-17T15:51:14Z","snapshot_observed_at":"2026-07-06T23:09:23.591544Z","submitted_at":"2026-04-25T04:26:19Z","title":"Characterizing Paraphrase-Induced Failures in Lean 4 Autoformalization","version":1},"reference_index":11,"source":"pdf_text","source_observed_at":"2026-05-08T08:33:34.423179Z"},"links":{"cited_paper":"/paper/2306.04528","citing_paper":"/paper/2604.23135"},"observation_digest":"sha256:60489f12277d3278140a17871a261431000fcced3444f2d6d357e5d55d97fbce","observation_id":"11f307d6-7263-46c5-8027-d796552a41b7","resolution":{"observed_at":"2026-05-11T20:36:08.464012Z","resolver_source":"arxiv_id","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2306.04528","last_updated":"2024-07-16T07:29:49Z","snapshot_observed_at":"2026-07-06T15:39:49.273453Z","submitted_at":"2023-06-07T15:37:00Z","title":"PromptRobust: Towards Evaluating the Robustness of Large Language Models on Adversarial Prompts","version":5},"cited_work":{"arxiv_id":"2306.04528","doi":"10.48550/arxiv.2306.04528","metadata_source":"arxiv_reference","pith_arxiv_id":"2306.04528","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Promptbench: Towards evaluating the robustness of large language models on adversarial prompts","venue":"arXiv (Cornell University)","work_id":"e65e5731-8f09-44e1-94d5-1674f9e08988","year":2024},"citing_paper":{"arxiv_id":"2604.23135","last_updated":"2026-05-17T15:51:14Z","snapshot_observed_at":"2026-07-06T23:09:23.591544Z","submitted_at":"2026-04-25T04:26:19Z","title":"Characterizing Paraphrase-Induced Failures in Lean 4 Autoformalization","version":2},"reference_index":11,"source":"pdf_text","source_observed_at":"2026-05-21T00:49:58.959331Z"},"links":{"cited_paper":"/paper/2306.04528","citing_paper":"/paper/2604.23135"},"observation_digest":"sha256:859271512e6973d4ac7eeaae8e3fd2309323c5ec98bbb9a9b6042e89b751b9bc","observation_id":"be275719-33b2-4a79-a3bc-77875de9734d","resolution":{"observed_at":"2026-05-21T00:53:53.714254Z","resolver_source":"arxiv_id","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2306.04528","last_updated":"2024-07-16T07:29:49Z","snapshot_observed_at":"2026-07-06T15:39:49.273453Z","submitted_at":"2023-06-07T15:37:00Z","title":"PromptRobust: Towards Evaluating the Robustness of Large Language Models on Adversarial Prompts","version":5},"cited_work":{"arxiv_id":"2306.04528","doi":"10.48550/arxiv.2306.04528","metadata_source":"arxiv_reference","pith_arxiv_id":"2306.04528","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Promptbench: Towards evaluating the robustness of large language models on adversarial prompts","venue":"arXiv (Cornell University)","work_id":"e65e5731-8f09-44e1-94d5-1674f9e08988","year":2024},"citing_paper":{"arxiv_id":"2604.23338","last_updated":"2026-05-06T17:17:02Z","snapshot_observed_at":"2026-08-06T19:46:39.219000Z","submitted_at":"2026-04-25T14:57:15Z","title":"A Systematic Survey of Security Threats and Defenses in LLM-Based AI Agents: A Layered Attack Surface Framework","version":2},"reference_index":77,"source":"pdf_text","source_observed_at":"2026-05-08T07:53:13.746141Z"},"links":{"cited_paper":"/paper/2306.04528","citing_paper":"/paper/2604.23338"},"observation_digest":"sha256:1b2f4a9933ac8b18ab851644a3cb6e09a759488d4e703539e4029105988f75e3","observation_id":"615c03fa-9e15-4524-9310-144f87348904","resolution":{"observed_at":"2026-05-11T20:51:09.098195Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2306.04528","last_updated":"2024-07-16T07:29:49Z","snapshot_observed_at":"2026-07-06T15:39:49.273453Z","submitted_at":"2023-06-07T15:37:00Z","title":"PromptRobust: Towards Evaluating the Robustness of Large Language Models on Adversarial Prompts","version":5},"cited_work":{"arxiv_id":"2306.04528","doi":"10.48550/arxiv.2306.04528","metadata_source":"arxiv_reference","pith_arxiv_id":"2306.04528","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Promptbench: Towards evaluating the robustness of large language models on adversarial prompts","venue":"arXiv (Cornell University)","work_id":"e65e5731-8f09-44e1-94d5-1674f9e08988","year":2024},"citing_paper":{"arxiv_id":"2604.24712","last_updated":"2026-04-27T17:21:09Z","snapshot_observed_at":"2026-07-06T23:10:42.926677Z","submitted_at":"2026-04-27T17:21:09Z","title":"When Prompt Under-Specification Improves Code Correctness: An Exploratory Study of Prompt Wording and Structure Effects on LLM-Based Code Generation","version":1},"reference_index":46,"source":"pdf_text","source_observed_at":"2026-05-08T03:00:26.137401Z"},"links":{"cited_paper":"/paper/2306.04528","citing_paper":"/paper/2604.24712"},"observation_digest":"sha256:8864c11479aec1c73fdb9eb3c5ad3067fee6d24bf0c89381195b448ec8583c9d","observation_id":"91f507b1-67ba-48d8-bccc-2c084d245b36","resolution":{"observed_at":"2026-05-11T22:17:07.253922Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2306.04528","last_updated":"2024-07-16T07:29:49Z","snapshot_observed_at":"2026-07-06T15:39:49.273453Z","submitted_at":"2023-06-07T15:37:00Z","title":"PromptRobust: Towards Evaluating the Robustness of Large Language Models on Adversarial Prompts","version":5},"cited_work":{"arxiv_id":"2306.04528","doi":"10.48550/arxiv.2306.04528","metadata_source":"arxiv_reference","pith_arxiv_id":"2306.04528","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Promptbench: Towards evaluating the robustness of large language models on adversarial prompts","venue":"arXiv (Cornell University)","work_id":"e65e5731-8f09-44e1-94d5-1674f9e08988","year":2024},"citing_paper":{"arxiv_id":"2605.04665","last_updated":"2026-05-09T22:09:59Z","snapshot_observed_at":"2026-07-06T23:17:23.641984Z","submitted_at":"2026-05-06T09:11:10Z","title":"Paraphrase-Induced Output-Mode Collapse: When LLMs Break Character Under Semantically Equivalent Inputs","version":1},"reference_index":5,"source":"pdf_text","source_observed_at":"2026-05-08T16:26:54.068997Z"},"links":{"cited_paper":"/paper/2306.04528","citing_paper":"/paper/2605.04665"},"observation_digest":"sha256:e6143118a976ee7679454fa7a405c9469dcf91a498943a5486e7c9528571f8f2","observation_id":"b9a6b961-0c06-4f58-8018-8c947129bdfd","resolution":{"observed_at":"2026-05-11T18:16:07.328226Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2306.04528","last_updated":"2024-07-16T07:29:49Z","snapshot_observed_at":"2026-07-06T15:39:49.273453Z","submitted_at":"2023-06-07T15:37:00Z","title":"PromptRobust: Towards Evaluating the Robustness of Large Language Models on Adversarial Prompts","version":5},"cited_work":{"arxiv_id":"2306.04528","doi":"10.48550/arxiv.2306.04528","metadata_source":"arxiv_reference","pith_arxiv_id":"2306.04528","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Promptbench: Towards evaluating the robustness of large language models on adversarial prompts","venue":"arXiv (Cornell University)","work_id":"e65e5731-8f09-44e1-94d5-1674f9e08988","year":2024},"citing_paper":{"arxiv_id":"2605.04665","last_updated":"2026-05-09T22:09:59Z","snapshot_observed_at":"2026-07-06T23:17:23.641984Z","submitted_at":"2026-05-06T09:11:10Z","title":"Paraphrase-Induced Output-Mode Collapse: When LLMs Break Character Under Semantically Equivalent Inputs","version":2},"reference_index":5,"source":"pdf_text","source_observed_at":"2026-05-12T02:13:48.672711Z"},"links":{"cited_paper":"/paper/2306.04528","citing_paper":"/paper/2605.04665"},"observation_digest":"sha256:9b051502c34efefd0f04dd6959fd1c496087c8ed84abfbaaea86181e59718957","observation_id":"0810c81a-95cc-4e59-8a2b-0761e6fcc17e","resolution":{"observed_at":"2026-05-12T02:16:16.124945Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2306.04528","last_updated":"2024-07-16T07:29:49Z","snapshot_observed_at":"2026-07-06T15:39:49.273453Z","submitted_at":"2023-06-07T15:37:00Z","title":"PromptRobust: Towards Evaluating the Robustness of Large Language Models on Adversarial Prompts","version":5},"cited_work":{"arxiv_id":"2306.04528","doi":"10.48550/arxiv.2306.04528","metadata_source":"arxiv_reference","pith_arxiv_id":"2306.04528","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Promptbench: Towards evaluating the robustness of large language models on adversarial prompts","venue":"arXiv (Cornell University)","work_id":"e65e5731-8f09-44e1-94d5-1674f9e08988","year":2024},"citing_paper":{"arxiv_id":"2605.09041","last_updated":"2026-05-09T16:26:49Z","snapshot_observed_at":"2026-08-03T06:54:18.779987Z","submitted_at":"2026-05-09T16:26:49Z","title":"BiAxisAudit: A Novel Framework to Evaluate LLM Bias Across Prompt Sensitivity and Response-Layer Divergence","version":1},"reference_index":34,"source":"pdf_text","source_observed_at":"2026-05-12T03:26:18.375974Z"},"links":{"cited_paper":"/paper/2306.04528","citing_paper":"/paper/2605.09041"},"observation_digest":"sha256:2e00f09ab9742b3ccb93b221c3d493caa9518773e83002e337bc91769a0f33e7","observation_id":"8f2c0e70-523f-4127-bf25-7274bc3202c0","resolution":{"observed_at":"2026-05-12T07:21:26.819643Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2306.04528","last_updated":"2024-07-16T07:29:49Z","snapshot_observed_at":"2026-07-06T15:39:49.273453Z","submitted_at":"2023-06-07T15:37:00Z","title":"PromptRobust: Towards Evaluating the Robustness of Large Language Models on Adversarial Prompts","version":5},"cited_work":{"arxiv_id":"2306.04528","doi":"10.48550/arxiv.2306.04528","metadata_source":"arxiv_reference","pith_arxiv_id":"2306.04528","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Promptbench: Towards evaluating the robustness of large language models on adversarial prompts","venue":"arXiv (Cornell University)","work_id":"e65e5731-8f09-44e1-94d5-1674f9e08988","year":2024},"citing_paper":{"arxiv_id":"2605.10516","last_updated":"2026-05-11T13:06:24Z","snapshot_observed_at":"2026-07-06T23:22:28.318343Z","submitted_at":"2026-05-11T13:06:24Z","title":"Consistency as a Testable Property: Statistical Methods to Evaluate AI Agent Reliability","version":1},"reference_index":35,"source":"pdf_text","source_observed_at":"2026-05-12T04:41:15.286881Z"},"links":{"cited_paper":"/paper/2306.04528","citing_paper":"/paper/2605.10516"},"observation_digest":"sha256:a78be57d9ad7de51c846ccf794717ce2bb2ec5fabab542027f469ddd3ea42244","observation_id":"cb4f7e11-7e6d-44ef-a76b-25a5eccf2e66","resolution":{"observed_at":"2026-05-12T04:41:21.819932Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2306.04528","last_updated":"2024-07-16T07:29:49Z","snapshot_observed_at":"2026-07-06T15:39:49.273453Z","submitted_at":"2023-06-07T15:37:00Z","title":"PromptRobust: Towards Evaluating the Robustness of Large Language Models on Adversarial Prompts","version":5},"cited_work":{"arxiv_id":"2306.04528","doi":"10.48550/arxiv.2306.04528","metadata_source":"arxiv_reference","pith_arxiv_id":"2306.04528","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Promptbench: Towards evaluating the robustness of large language models on adversarial prompts","venue":"arXiv (Cornell University)","work_id":"e65e5731-8f09-44e1-94d5-1674f9e08988","year":2024},"citing_paper":{"arxiv_id":"2606.01210","last_updated":"2026-05-31T13:00:57Z","snapshot_observed_at":"2026-08-06T16:10:14.400447Z","submitted_at":"2026-05-31T13:00:57Z","title":"Can we trust LLM Self-Explanations for Entity Resolution?","version":1},"reference_index":73,"source":"pdf_text","source_observed_at":"2026-06-28T16:13:14.653064Z"},"links":{"cited_paper":"/paper/2306.04528","citing_paper":"/paper/2606.01210"},"observation_digest":"sha256:dc5be28e22ccd9828b0e1280bb9c3f086327f005e6e2096fc16e5d3217b61c51","observation_id":"ee06c8c2-73dc-445b-95b2-7422e7f31eb1","resolution":{"observed_at":"2026-07-01T21:56:15.179442Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2306.04528","last_updated":"2024-07-16T07:29:49Z","snapshot_observed_at":"2026-07-06T15:39:49.273453Z","submitted_at":"2023-06-07T15:37:00Z","title":"PromptRobust: Towards Evaluating the Robustness of Large Language Models on Adversarial Prompts","version":5},"cited_work":{"arxiv_id":"2306.04528","doi":"10.48550/arxiv.2306.04528","metadata_source":"arxiv_reference","pith_arxiv_id":"2306.04528","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Promptbench: Towards evaluating the robustness of large language models on adversarial prompts","venue":"arXiv (Cornell University)","work_id":"e65e5731-8f09-44e1-94d5-1674f9e08988","year":2024},"citing_paper":{"arxiv_id":"2606.01441","last_updated":"2026-05-31T20:20:53Z","snapshot_observed_at":"2026-08-05T08:36:53.588712Z","submitted_at":"2026-05-31T20:20:53Z","title":"Dive into Ambiguity: A*-Inspired Multi-Agents Commonsense Obfuscation Attack on LLM Prompts","version":1},"reference_index":43,"source":"pdf_text","source_observed_at":"2026-06-28T16:54:12.354178Z"},"links":{"cited_paper":"/paper/2306.04528","citing_paper":"/paper/2606.01441"},"observation_digest":"sha256:d5d6d3800708ce3142ad1b1e44708d4f38a51bb70f213b4eca64f6fcc94bc59b","observation_id":"f4b554b0-5b1b-4c6d-bf9d-daca796d5c9e","resolution":{"observed_at":"2026-07-01T21:26:16.387346Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2306.04528","last_updated":"2024-07-16T07:29:49Z","snapshot_observed_at":"2026-07-06T15:39:49.273453Z","submitted_at":"2023-06-07T15:37:00Z","title":"PromptRobust: Towards Evaluating the Robustness of Large Language Models on Adversarial Prompts","version":5},"cited_work":{"arxiv_id":"2306.04528","doi":"10.48550/arxiv.2306.04528","metadata_source":"arxiv_reference","pith_arxiv_id":"2306.04528","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Promptbench: Towards evaluating the robustness of large language models on adversarial prompts","venue":"arXiv (Cornell University)","work_id":"e65e5731-8f09-44e1-94d5-1674f9e08988","year":2024},"citing_paper":{"arxiv_id":"2606.06924","last_updated":"2026-06-05T05:42:00Z","snapshot_observed_at":"2026-08-06T11:49:12.989321Z","submitted_at":"2026-06-05T05:42:00Z","title":"From Sampled Outcomes to Capability Distributions: Rethinking Supervision for LLM Routing","version":1},"reference_index":158,"source":"arxiv_source","source_observed_at":"2026-06-27T22:54:28.452796Z"},"links":{"cited_paper":"/paper/2306.04528","citing_paper":"/paper/2606.06924"},"observation_digest":"sha256:a83611f5c0cfd9c81f502a4e676e536a08aed52fde479af7da11ad3d086095dd","observation_id":"5e033c3d-c72d-4347-a826-3d07a3b969b3","resolution":{"observed_at":"2026-07-02T16:17:08.772986Z","resolver_source":"arxiv_id","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2306.04528","last_updated":"2024-07-16T07:29:49Z","snapshot_observed_at":"2026-07-06T15:39:49.273453Z","submitted_at":"2023-06-07T15:37:00Z","title":"PromptRobust: Towards Evaluating the Robustness of Large Language Models on Adversarial Prompts","version":5},"cited_work":{"arxiv_id":"2306.04528","doi":"10.48550/arxiv.2306.04528","metadata_source":"arxiv_reference","pith_arxiv_id":"2306.04528","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Promptbench: Towards evaluating the robustness of large language models on adversarial prompts","venue":"arXiv (Cornell University)","work_id":"e65e5731-8f09-44e1-94d5-1674f9e08988","year":2024},"citing_paper":{"arxiv_id":"2606.12978","last_updated":"2026-06-11T07:12:17Z","snapshot_observed_at":"2026-08-03T00:41:34.363368Z","submitted_at":"2026-06-11T07:12:17Z","title":"Trajectory-Level Redirection Attacks on Vision-Language-Action Models","version":1},"reference_index":43,"source":"pdf_text","source_observed_at":"2026-06-27T06:33:53.013076Z"},"links":{"cited_paper":"/paper/2306.04528","citing_paper":"/paper/2606.12978"},"observation_digest":"sha256:be21380096bed45adf4954367305f249eefafa84a76cce193085e8f6afae1982","observation_id":"2ec8c2be-f86e-4035-be37-84907a53bfb9","resolution":{"observed_at":"2026-07-03T15:18:33.763432Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2306.04528","last_updated":"2024-07-16T07:29:49Z","snapshot_observed_at":"2026-07-06T15:39:49.273453Z","submitted_at":"2023-06-07T15:37:00Z","title":"PromptRobust: Towards Evaluating the Robustness of Large Language Models on Adversarial Prompts","version":5},"cited_work":{"arxiv_id":"2306.04528","doi":"10.48550/arxiv.2306.04528","metadata_source":"arxiv_reference","pith_arxiv_id":"2306.04528","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Promptbench: Towards evaluating the robustness of large language models on adversarial prompts","venue":"arXiv (Cornell University)","work_id":"e65e5731-8f09-44e1-94d5-1674f9e08988","year":2024},"citing_paper":{"arxiv_id":"2606.23716","last_updated":"2026-06-16T14:19:04Z","snapshot_observed_at":"2026-08-03T11:45:50.371507Z","submitted_at":"2026-06-16T14:19:04Z","title":"Legal Reasoning Is Not Lawyering: Rethinking Legal Benchmarks for Pro Se Access to Justice","version":1},"reference_index":17,"source":"arxiv_source","source_observed_at":"2026-06-26T22:26:54.239505Z"},"links":{"cited_paper":"/paper/2306.04528","citing_paper":"/paper/2606.23716"},"observation_digest":"sha256:d74d4f2cbc3fe62f1c0d953b08d5a35f61a1404ca2516d7811350a2c6770fb63","observation_id":"25f70ffd-2c35-4b17-b7af-c765deb85791","resolution":{"observed_at":"2026-07-03T23:19:03.950183Z","resolver_source":"arxiv_id","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2306.04528","last_updated":"2024-07-16T07:29:49Z","snapshot_observed_at":"2026-07-06T15:39:49.273453Z","submitted_at":"2023-06-07T15:37:00Z","title":"PromptRobust: Towards Evaluating the Robustness of Large Language Models on Adversarial Prompts","version":5},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2306.04528","snapshot_observed_at":"2026-07-14T19:17:45.652076Z","title":"Zhu, K., Wang, J., Zhou, J., Wang, Z., Chen, H., Wang, Y., Yang, L., Ye, W., Gong, N","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2607.09665","last_updated":"2026-05-02T22:51:41Z","snapshot_observed_at":"2026-08-04T00:20:04.568648Z","submitted_at":"2026-05-02T22:51:41Z","title":"Format Sensitivity Index: Token-Controlled Prompt Wrapper Robustness and Schema Compliance in LLM Benchmarking","version":1},"reference_index":16,"source":"pdf_text","source_observed_at":"2026-07-14T19:17:45.652076Z"},"links":{"cited_paper":"/paper/2306.04528","citing_paper":"/paper/2607.09665"},"observation_digest":"sha256:77e8b32a9784a449845a7664f4063fb22f70cb93062c691e206aaa941d55622f","observation_id":"4448a36d-8194-49fc-9258-dbf5724d335e","resolution":{"observed_at":"2026-07-14T19:17:45.652076Z","resolver_source":null,"status":"malformed_identifier"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2306.04528","last_updated":"2024-07-16T07:29:49Z","snapshot_observed_at":"2026-07-06T15:39:49.273453Z","submitted_at":"2023-06-07T15:37:00Z","title":"PromptRobust: Towards Evaluating the Robustness of Large Language Models on Adversarial Prompts","version":5},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2306.04528","snapshot_observed_at":"2026-07-14T19:17:45.652076Z","title":"https://arxiv.org/abs/2306.04528","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2607.09665","last_updated":"2026-05-02T22:51:41Z","snapshot_observed_at":"2026-08-04T00:20:04.568648Z","submitted_at":"2026-05-02T22:51:41Z","title":"Format Sensitivity Index: Token-Controlled Prompt Wrapper Robustness and Schema Compliance in LLM Benchmarking","version":1},"reference_index":17,"source":"pdf_text","source_observed_at":"2026-07-14T19:17:45.652076Z"},"links":{"cited_paper":"/paper/2306.04528","citing_paper":"/paper/2607.09665"},"observation_digest":"sha256:dd50eb57ee0311566de2343c654abba12f6534d257a494967c182f05b3aaeae2","observation_id":"3e67b377-fb5f-4839-9e62-3a2f3fa3b155","resolution":{"observed_at":"2026-07-14T19:17:45.652076Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2306.04528","last_updated":"2024-07-16T07:29:49Z","snapshot_observed_at":"2026-07-06T15:39:49.273453Z","submitted_at":"2023-06-07T15:37:00Z","title":"PromptRobust: Towards Evaluating the Robustness of Large Language Models on Adversarial Prompts","version":5},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2306.04528","snapshot_observed_at":"2026-08-01T14:35:01.351523Z","title":"Promptbench: Towards evaluating the robustness of large language models on adversarial prompts,","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2607.18725","last_updated":"2026-07-21T05:32:40Z","snapshot_observed_at":"2026-08-06T06:33:08.110077Z","submitted_at":"2026-07-21T05:32:40Z","title":"Find Before You Fine-Tune: A Diagnostic Study of Small LLMs for Cybersecurity QA","version":1},"reference_index":16,"source":"pdf_text","source_observed_at":"2026-08-01T14:35:01.351523Z"},"links":{"cited_paper":"/paper/2306.04528","citing_paper":"/paper/2607.18725"},"observation_digest":"sha256:1136ffbbf8c5ddf201026e809a765266ddb1960fd6f9a614ac15a5a0a7edfce6","observation_id":"0fb91a91-a8a6-4c45-b562-b96172b8ce3a","resolution":{"observed_at":"2026-08-01T14:35:01.351523Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2306.04528","last_updated":"2024-07-16T07:29:49Z","snapshot_observed_at":"2026-07-06T15:39:49.273453Z","submitted_at":"2023-06-07T15:37:00Z","title":"PromptRobust: Towards Evaluating the Robustness of Large Language Models on Adversarial Prompts","version":5},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2306.04528","snapshot_observed_at":"2026-08-02T07:07:45.246508Z","title":null,"venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2607.22683","last_updated":"2026-07-13T05:02:43Z","snapshot_observed_at":"2026-08-07T06:32:39.192476Z","submitted_at":"2026-07-13T05:02:43Z","title":"Imprompt: A Language Framework for Prompt Programming","version":1},"reference_index":66,"source":"pdf_text","source_observed_at":"2026-08-02T07:07:45.246508Z"},"links":{"cited_paper":"/paper/2306.04528","citing_paper":"/paper/2607.22683"},"observation_digest":"sha256:8d0b68f483cfea924ffeb3b7eea3384097e6d773f8ebc090909aaca251dc8e63","observation_id":"d1c5e8b1-a155-4fd8-817c-65fea8b70be4","resolution":{"observed_at":"2026-08-02T07:07:45.246508Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}}],"links":{"evidence":"/evidence","html":"/paper/2306.04528/citation-record","integrity":"/paper/2306.04528/integrity","json":"/paper/2306.04528/citation-record.json","paper":"/paper/2306.04528"},"outbound":[],"paper":{"arxiv_id":"2306.04528","last_updated":"2024-07-16T07:29:49Z","latest_version":5,"primary_category":"cs.CL","snapshot_observed_at":"2026-07-06T15:39:49.273453Z","submitted_at":"2023-06-07T15:37:00Z","title":"PromptRobust: Towards Evaluating the Robustness of Large Language Models on Adversarial Prompts"},"reference_resolution":{"displayed":0,"state_counts":{"malformed_identifier":0,"metadata_mismatch":0,"parse_uncertain":0,"unresolved":0,"verified_exact":0,"verified_fuzzy":0},"total_outbound_references":0},"refusal":"A citation records a reference. It does not transfer a finding from one paper to another.","schema":"pith.paper-citation-record.v1","standing_sources":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"thesis":"As of 7 August 2026, this Paper Citation Record lists 0 of 0 outbound references and 36 inbound Pith citation observations for arXiv:2306.04528."}