{"as_of":"2026-08-20T17:47:00Z","caps":{"database_statements":6,"inbound":100,"outbound":100},"context_digest":"sha256:703c93062b541cec7582102d6bd820c5ad9354355ab32c4984edb708b4c4db79","coverage":[{"denominator":0,"lane":"reference_resolution","note":"Typed states for the displayed outbound observations.","records_observed":0,"source":"paper_references, paper_reference_links","source_observed_at":null,"state":"measured"},{"denominator":19,"lane":"standing_notices","note":"One-hop event checks from named stored sources.","records_observed":19,"source":"scholarly_work_events, retraction_status_cache","source_observed_at":"2026-08-20T06:33:59.587034+00:00","state":"measured"},{"denominator":19,"lane":"inbound_itemization","note":"Pith citing papers itemized under the disclosed page cap.","records_observed":19,"source":"paper_references, paper_reference_links","source_observed_at":"2026-08-16T12:32:36.863642Z","state":"measured"},{"denominator":1,"lane":"external_citation_measurements","note":"A source-named dated measurement, never combined with another source.","records_observed":0,"source":"pith","source_observed_at":"2026-07-10T15:27:20.049648Z","state":"measured"}],"external_citation_measurements":[],"inbound":[{"citation":{"cited_paper":{"arxiv_id":"2401.03729","last_updated":"2024-04-01T20:56:11Z","snapshot_observed_at":"2026-08-20T05:18:39.545724Z","submitted_at":"2024-01-08T08:28:08Z","title":"The Butterfly Effect of Altering Prompts: How Small Changes and Jailbreaks Affect Large Language Model Performance","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2401.03729","snapshot_observed_at":"2026-08-12T17:39:08.412464Z","title":"The butterfly effect of altering prompts: How small changes and jailbreaks affect large language model perfor- mance,","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2411.12395","last_updated":"2024-11-19T10:27:26Z","snapshot_observed_at":"2026-08-20T01:48:54.739502Z","submitted_at":"2024-11-19T10:27:26Z","title":"Do LLMs Understand Ambiguity in Text? A Case Study in Open-world Question Answering","version":1},"reference_index":14,"source":"pdf_text","source_observed_at":"2026-08-12T17:39:08.412464Z"},"links":{"cited_paper":"/paper/2401.03729","citing_paper":"/paper/2411.12395"},"observation_digest":"sha256:cc0cabbad49ab05d296929d697246e8295aff054764789792dd9b82287625438","observation_id":"3285a012-fc59-4c10-93a6-87c5269e170f","resolution":{"observed_at":"2026-08-12T17:39:08.412464Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2401.03729","last_updated":"2024-04-01T20:56:11Z","snapshot_observed_at":"2026-08-20T05:18:39.545724Z","submitted_at":"2024-01-08T08:28:08Z","title":"The Butterfly Effect of Altering Prompts: How Small Changes and Jailbreaks Affect Large Language Model Performance","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2401.03729","snapshot_observed_at":"2026-08-11T15:07:59.832911Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2412.11318","last_updated":"2024-12-15T21:30:21Z","snapshot_observed_at":"2026-08-19T12:24:38.926664Z","submitted_at":"2024-12-15T21:30:21Z","title":"Generics are puzzling. Can language models find the missing piece?","version":1},"reference_index":52,"source":"arxiv_source","source_observed_at":"2026-08-11T15:07:59.832911Z"},"links":{"cited_paper":"/paper/2401.03729","citing_paper":"/paper/2412.11318"},"observation_digest":"sha256:906e931c0a32eb0c5857ab7510c4887e8eced8cb02b4be657c46b7d0122dcbdc","observation_id":"3db360a3-e6a6-45e8-81ab-a0736da0b52f","resolution":{"observed_at":"2026-08-11T15:07:59.832911Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2401.03729","last_updated":"2024-04-01T20:56:11Z","snapshot_observed_at":"2026-08-20T05:18:39.545724Z","submitted_at":"2024-01-08T08:28:08Z","title":"The Butterfly Effect of Altering Prompts: How Small Changes and Jailbreaks Affect Large Language Model Performance","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2401.03729","snapshot_observed_at":"2026-08-11T11:31:33.131012Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2412.15404","last_updated":"2024-12-19T21:14:54Z","snapshot_observed_at":"2026-08-18T06:12:01.389771Z","submitted_at":"2024-12-19T21:14:54Z","title":"A Retrieval-Augmented Generation Framework for Academic Literature Navigation in Data Science","version":1},"reference_index":30,"source":"pdf_text","source_observed_at":"2026-08-11T11:31:33.131012Z"},"links":{"cited_paper":"/paper/2401.03729","citing_paper":"/paper/2412.15404"},"observation_digest":"sha256:bc0c0603700ded3467b5b8db9e1f3d12b1dcb4ddaa25f1a306f57fd0db123381","observation_id":"d3c367d7-fb86-4925-80ed-5d9e242e0645","resolution":{"observed_at":"2026-08-11T11:31:33.131012Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2401.03729","last_updated":"2024-04-01T20:56:11Z","snapshot_observed_at":"2026-08-20T05:18:39.545724Z","submitted_at":"2024-01-08T08:28:08Z","title":"The Butterfly Effect of Altering Prompts: How Small Changes and Jailbreaks Affect Large Language Model Performance","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2401.03729","snapshot_observed_at":"2026-08-10T23:27:17.135931Z","title":"arXiv:2401.03729","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2412.20440","last_updated":"2024-12-29T11:33:51Z","snapshot_observed_at":"2026-08-17T01:21:34.781628Z","submitted_at":"2024-12-29T11:33:51Z","title":"Enhancing Entertainment Translation for Indian Languages using Adaptive Context, Style and LLMs","version":1},"reference_index":20,"source":"pdf_text","source_observed_at":"2026-08-10T23:27:17.135931Z"},"links":{"cited_paper":"/paper/2401.03729","citing_paper":"/paper/2412.20440"},"observation_digest":"sha256:27e90d060f0b39e64700ded859eefcd10dff7193fbb47b78a45a7a71863cc76f","observation_id":"5efc4130-5073-45dd-946b-66d52469b952","resolution":{"observed_at":"2026-08-10T23:27:17.135931Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2401.03729","last_updated":"2024-04-01T20:56:11Z","snapshot_observed_at":"2026-08-20T05:18:39.545724Z","submitted_at":"2024-01-08T08:28:08Z","title":"The Butterfly Effect of Altering Prompts: How Small Changes and Jailbreaks Affect Large Language Model Performance","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2401.03729","snapshot_observed_at":"2026-08-08T13:53:58.420252Z","title":"and Morstatter, F","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2502.07087","last_updated":"2025-02-10T22:27:02Z","snapshot_observed_at":"2026-08-12T08:05:20.556490Z","submitted_at":"2025-02-10T22:27:02Z","title":"Evaluating the Systematic Reasoning Abilities of Large Language Models through Graph Coloring","version":1},"reference_index":17,"source":"arxiv_source","source_observed_at":"2026-08-08T13:53:58.420252Z"},"links":{"cited_paper":"/paper/2401.03729","citing_paper":"/paper/2502.07087"},"observation_digest":"sha256:ffc6f92768a073787c5b2763c2edeba0a74cde882b2c9e4df96f16aa88962b94","observation_id":"308db8f5-f1ff-4f17-9dcf-044cde671c0e","resolution":{"observed_at":"2026-08-08T13:53:58.420252Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2401.03729","last_updated":"2024-04-01T20:56:11Z","snapshot_observed_at":"2026-08-20T05:18:39.545724Z","submitted_at":"2024-01-08T08:28:08Z","title":"The Butterfly Effect of Altering Prompts: How Small Changes and Jailbreaks Affect Large Language Model Performance","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2401.03729","snapshot_observed_at":"2026-08-16T12:32:36.863642Z","title":"The butterfly effect of altering prompts: How small changes and jailbreaks affect large language model performance","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2504.12585","last_updated":"2025-04-17T02:00:53Z","snapshot_observed_at":"2026-08-18T10:31:42.044859Z","submitted_at":"2025-04-17T02:00:53Z","title":"Identifying and Mitigating the Influence of the Prior Distribution in Large Language Models","version":1},"reference_index":22,"source":"pdf_text","source_observed_at":"2026-08-16T12:32:36.863642Z"},"links":{"cited_paper":"/paper/2401.03729","citing_paper":"/paper/2504.12585"},"observation_digest":"sha256:ba17b079c9f025b6f828ae64288ba6932de84aa17e23cba63749ec63d3f52bbc","observation_id":"03dbe367-2f0d-4faf-b97d-0be9f20dd802","resolution":{"observed_at":"2026-08-16T12:32:36.863642Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2401.03729","last_updated":"2024-04-01T20:56:11Z","snapshot_observed_at":"2026-08-20T05:18:39.545724Z","submitted_at":"2024-01-08T08:28:08Z","title":"The Butterfly Effect of Altering Prompts: How Small Changes and Jailbreaks Affect Large Language Model Performance","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2401.03729","snapshot_observed_at":"2026-08-16T11:28:46.321815Z","title":"The Butterfly Effect of Altering Prompts: How Small Changes and Jailbreaks Affect Large Language Model Performance, January 9, 2024","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2504.15499","last_updated":"2025-04-22T00:29:18Z","snapshot_observed_at":"2026-08-20T16:45:37.285794Z","submitted_at":"2025-04-22T00:29:18Z","title":"Guillotine: Hypervisors for Isolating Malicious AIs","version":1},"reference_index":56,"source":"pdf_text","source_observed_at":"2026-08-16T11:28:46.321815Z"},"links":{"cited_paper":"/paper/2401.03729","citing_paper":"/paper/2504.15499"},"observation_digest":"sha256:9091b39250d55b7894fce68c5e71fdbaa39156d24c02e89d699c5c52b961bf42","observation_id":"31484b16-63e3-4747-82a6-b11e12f9c682","resolution":{"observed_at":"2026-08-16T11:28:46.321815Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2401.03729","last_updated":"2024-04-01T20:56:11Z","snapshot_observed_at":"2026-08-20T05:18:39.545724Z","submitted_at":"2024-01-08T08:28:08Z","title":"The Butterfly Effect of Altering Prompts: How Small Changes and Jailbreaks Affect Large Language Model Performance","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2401.03729","snapshot_observed_at":"2026-08-16T01:03:46.977332Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2505.02172","last_updated":"2025-05-24T19:09:19Z","snapshot_observed_at":"2026-08-19T05:24:21.375512Z","submitted_at":"2025-05-04T16:24:12Z","title":"Identifying Legal Holdings with LLMs: A Systematic Study of Performance, Scale, and Memorization","version":3},"reference_index":16,"source":"pdf_text","source_observed_at":"2026-08-16T01:03:46.977332Z"},"links":{"cited_paper":"/paper/2401.03729","citing_paper":"/paper/2505.02172"},"observation_digest":"sha256:0b5157e1767a8ab1e8ea20a69b50f5c40deedb2ad1a73e5662f312f5c2839756","observation_id":"0442d9d6-3cfa-4a0e-8560-2f13e029f4d3","resolution":{"observed_at":"2026-08-16T01:03:46.977332Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2401.03729","last_updated":"2024-04-01T20:56:11Z","snapshot_observed_at":"2026-08-20T05:18:39.545724Z","submitted_at":"2024-01-08T08:28:08Z","title":"The Butterfly Effect of Altering Prompts: How Small Changes and Jailbreaks Affect Large Language Model Performance","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2401.03729","snapshot_observed_at":"2026-08-06T21:34:41.685451Z","title":"The butterfly effect of altering prompts: How small changes and jailbreaks affect large language model perfor- mance,","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2506.23930","last_updated":"2025-06-30T14:59:25Z","snapshot_observed_at":"2026-08-17T05:06:24.099153Z","submitted_at":"2025-06-30T14:59:25Z","title":"Leveraging the Potential of Prompt Engineering for Hate Speech Detection in Low-Resource Languages","version":1},"reference_index":64,"source":"pdf_text","source_observed_at":"2026-08-06T21:34:41.685451Z"},"links":{"cited_paper":"/paper/2401.03729","citing_paper":"/paper/2506.23930"},"observation_digest":"sha256:fd29c4d6e21d06f399142539cf52071b579807e6b1bc241e2975006badf7a870","observation_id":"82edd1f4-63aa-4f5d-b684-a8351c01cf05","resolution":{"observed_at":"2026-08-06T21:34:41.685451Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2401.03729","last_updated":"2024-04-01T20:56:11Z","snapshot_observed_at":"2026-08-20T05:18:39.545724Z","submitted_at":"2024-01-08T08:28:08Z","title":"The Butterfly Effect of Altering Prompts: How Small Changes and Jailbreaks Affect Large Language Model Performance","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2401.03729","snapshot_observed_at":"2026-08-06T15:46:55.535021Z","title":"The butterfly effect of altering prompts: How small changes and jailbreaks affect large lan- guage model performance,","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2507.15104","last_updated":"2026-06-15T19:35:17Z","snapshot_observed_at":"2026-08-13T10:11:17.212034Z","submitted_at":"2025-07-20T19:57:07Z","title":"AnalogFed: Privacy-Preserving Discovery of Analog Circuits at Scale with Federated Generative AI","version":2},"reference_index":49,"source":"pdf_text","source_observed_at":"2026-08-06T15:46:55.535021Z"},"links":{"cited_paper":"/paper/2401.03729","citing_paper":"/paper/2507.15104"},"observation_digest":"sha256:e5cefd6276c81dd3d4d6cf20a0f1b444c49ad3b222c2bb7810c5f90266e2043a","observation_id":"6d9e59a4-05b4-4e49-8fed-f6cc9a53dd5f","resolution":{"observed_at":"2026-08-06T15:46:55.535021Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2401.03729","last_updated":"2024-04-01T20:56:11Z","snapshot_observed_at":"2026-08-20T05:18:39.545724Z","submitted_at":"2024-01-08T08:28:08Z","title":"The Butterfly Effect of Altering Prompts: How Small Changes and Jailbreaks Affect Large Language Model Performance","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2401.03729","snapshot_observed_at":"2026-08-05T18:53:54.428372Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2508.13948","last_updated":"2025-08-19T15:37:29Z","snapshot_observed_at":"2026-08-16T14:10:23.470524Z","submitted_at":"2025-08-19T15:37:29Z","title":"Prompt Orchestration Markup Language","version":1},"reference_index":73,"source":"pdf_text","source_observed_at":"2026-08-05T18:53:54.428372Z"},"links":{"cited_paper":"/paper/2401.03729","citing_paper":"/paper/2508.13948"},"observation_digest":"sha256:3f18fae9a62734963d177d19d8c01b26bda398493982ffafb6ee1736fbc9c8d5","observation_id":"acc19eb2-f10b-4509-834c-1bb77eb616fa","resolution":{"observed_at":"2026-08-05T18:53:54.428372Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2401.03729","last_updated":"2024-04-01T20:56:11Z","snapshot_observed_at":"2026-08-20T05:18:39.545724Z","submitted_at":"2024-01-08T08:28:08Z","title":"The Butterfly Effect of Altering Prompts: How Small Changes and Jailbreaks Affect Large Language Model Performance","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2401.03729","snapshot_observed_at":"2026-08-05T14:23:37.545935Z","title":"& Morstatter F","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2508.21389","last_updated":"2025-08-29T08:05:00Z","snapshot_observed_at":"2026-08-13T03:10:57.103348Z","submitted_at":"2025-08-29T08:05:00Z","title":"AllSummedUp: un framework open-source pour comparer les metriques d'evaluation de resume","version":1},"reference_index":13,"source":"arxiv_source","source_observed_at":"2026-08-05T14:23:37.545935Z"},"links":{"cited_paper":"/paper/2401.03729","citing_paper":"/paper/2508.21389"},"observation_digest":"sha256:ff0d44e34424905a7dc4da68da986d3d187ca896b8954787413b6956bcc74f3f","observation_id":"269b5749-f9dc-4495-b10e-9f21b5199d3c","resolution":{"observed_at":"2026-08-05T14:23:37.545935Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2401.03729","last_updated":"2024-04-01T20:56:11Z","snapshot_observed_at":"2026-08-20T05:18:39.545724Z","submitted_at":"2024-01-08T08:28:08Z","title":"The Butterfly Effect of Altering Prompts: How Small Changes and Jailbreaks Affect Large Language Model Performance","version":3},"cited_work":{"arxiv_id":"2401.03729","doi":null,"metadata_source":"pith","pith_arxiv_id":"2401.03729","snapshot_observed_at":"2026-07-10T15:27:20.049648Z","title":"The butterfly effect of altering prompts: How small changes and jailbreaks affect large language model performance.arXiv preprint arXiv:2401.03729","venue":"cs.CL","work_id":"150a45e5-3469-496b-b2f3-50ce7822e435","year":2024},"citing_paper":{"arxiv_id":"2605.09271","last_updated":"2026-05-10T02:42:29Z","snapshot_observed_at":"2026-08-12T15:02:13.436236Z","submitted_at":"2026-05-10T02:42:29Z","title":"Shaping Schema via Language Representation as the Next Frontier for LLM Intelligence Expanding","version":1},"reference_index":64,"source":"pdf_text","source_observed_at":"2026-05-12T05:01:38.118237Z"},"links":{"cited_paper":"/paper/2401.03729","citing_paper":"/paper/2605.09271"},"observation_digest":"sha256:198d7b28b5ee3254afdaf25d5ed2110ca464329b7ef4d402399a37e3b11b8b7a","observation_id":"dd558f8d-8b20-4b78-bc61-9fd9b2eabf5d","resolution":{"observed_at":"2026-05-12T05:41:27.189901Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-20T06:33:59.587034+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-20T06:33:59.587034+00:00","source":"crossref"},{"observed_at":"2026-08-20T06:33:54.927442+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2401.03729","last_updated":"2024-04-01T20:56:11Z","snapshot_observed_at":"2026-08-20T05:18:39.545724Z","submitted_at":"2024-01-08T08:28:08Z","title":"The Butterfly Effect of Altering Prompts: How Small Changes and Jailbreaks Affect Large Language Model Performance","version":3},"cited_work":{"arxiv_id":"2401.03729","doi":null,"metadata_source":"pith","pith_arxiv_id":"2401.03729","snapshot_observed_at":"2026-07-10T15:27:20.049648Z","title":"The butterfly effect of altering prompts: How small changes and jailbreaks affect large language model performance.arXiv preprint arXiv:2401.03729","venue":"cs.CL","work_id":"150a45e5-3469-496b-b2f3-50ce7822e435","year":2024},"citing_paper":{"arxiv_id":"2605.17436","last_updated":"2026-05-17T13:11:38Z","snapshot_observed_at":"2026-08-16T14:30:40.491700Z","submitted_at":"2026-05-17T13:11:38Z","title":"Medical Context Distorts Decisions in Clinical Vision Language Models","version":1},"reference_index":14,"source":"pdf_text","source_observed_at":"2026-05-20T14:59:11.496757Z"},"links":{"cited_paper":"/paper/2401.03729","citing_paper":"/paper/2605.17436"},"observation_digest":"sha256:fa1b4cd0cfc8fccf989571966e9b2302224bf779e4f3300526b0f356794268cc","observation_id":"83fbe9da-285b-4947-8221-5771043a3d43","resolution":{"observed_at":"2026-05-20T15:03:24.788659Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-20T06:33:59.587034+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-20T06:33:59.587034+00:00","source":"crossref"},{"observed_at":"2026-08-20T06:33:54.927442+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2401.03729","last_updated":"2024-04-01T20:56:11Z","snapshot_observed_at":"2026-08-20T05:18:39.545724Z","submitted_at":"2024-01-08T08:28:08Z","title":"The Butterfly Effect of Altering Prompts: How Small Changes and Jailbreaks Affect Large Language Model Performance","version":3},"cited_work":{"arxiv_id":"2401.03729","doi":null,"metadata_source":"pith","pith_arxiv_id":"2401.03729","snapshot_observed_at":"2026-07-10T15:27:20.049648Z","title":"The butterfly effect of altering prompts: How small changes and jailbreaks affect large language model performance.arXiv preprint arXiv:2401.03729","venue":"cs.CL","work_id":"150a45e5-3469-496b-b2f3-50ce7822e435","year":2024},"citing_paper":{"arxiv_id":"2605.23040","last_updated":"2026-05-21T21:13:14Z","snapshot_observed_at":"2026-08-19T07:16:42.923325Z","submitted_at":"2026-05-21T21:13:14Z","title":"Steered Generation via Gradient-Based Optimization on Sparse Query Features","version":1},"reference_index":36,"source":"pdf_text","source_observed_at":"2026-05-25T05:31:29.510639Z"},"links":{"cited_paper":"/paper/2401.03729","citing_paper":"/paper/2605.23040"},"observation_digest":"sha256:df7c8505bb2d18647228db881bcc9fba28d8f072e3a37f56c8c2e9a0a93ace29","observation_id":"0695012b-c7b1-4bc7-8797-a76f055ca876","resolution":{"observed_at":"2026-05-25T05:36:40.337827Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-20T06:33:59.587034+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-20T06:33:59.587034+00:00","source":"crossref"},{"observed_at":"2026-08-20T06:33:54.927442+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2401.03729","last_updated":"2024-04-01T20:56:11Z","snapshot_observed_at":"2026-08-20T05:18:39.545724Z","submitted_at":"2024-01-08T08:28:08Z","title":"The Butterfly Effect of Altering Prompts: How Small Changes and Jailbreaks Affect Large Language Model Performance","version":3},"cited_work":{"arxiv_id":"2401.03729","doi":null,"metadata_source":"pith","pith_arxiv_id":"2401.03729","snapshot_observed_at":"2026-07-10T15:27:20.049648Z","title":"The butterfly effect of altering prompts: How small changes and jailbreaks affect large language model performance.arXiv preprint arXiv:2401.03729","venue":"cs.CL","work_id":"150a45e5-3469-496b-b2f3-50ce7822e435","year":2024},"citing_paper":{"arxiv_id":"2605.25151","last_updated":"2026-05-24T16:07:34Z","snapshot_observed_at":"2026-08-11T14:14:39.274356Z","submitted_at":"2026-05-24T16:07:34Z","title":"Representation Without Control: Testing the Realization Effect in Language Models","version":1},"reference_index":12,"source":"pdf_text","source_observed_at":"2026-06-30T11:02:37.361704Z"},"links":{"cited_paper":"/paper/2401.03729","citing_paper":"/paper/2605.25151"},"observation_digest":"sha256:e49aec0ac653283cf7d847d7e780aecfe4de3395c1bc1536d4e1cf643f7dfb6d","observation_id":"2028af79-bfeb-48e9-bff5-536091de88df","resolution":{"observed_at":"2026-06-30T11:04:37.539468Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-20T06:33:59.587034+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-20T06:33:59.587034+00:00","source":"crossref"},{"observed_at":"2026-08-20T06:33:54.927442+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2401.03729","last_updated":"2024-04-01T20:56:11Z","snapshot_observed_at":"2026-08-20T05:18:39.545724Z","submitted_at":"2024-01-08T08:28:08Z","title":"The Butterfly Effect of Altering Prompts: How Small Changes and Jailbreaks Affect Large Language Model Performance","version":3},"cited_work":{"arxiv_id":"2401.03729","doi":null,"metadata_source":"pith","pith_arxiv_id":"2401.03729","snapshot_observed_at":"2026-07-10T15:27:20.049648Z","title":"The butterfly effect of altering prompts: How small changes and jailbreaks affect large language model performance.arXiv preprint arXiv:2401.03729","venue":"cs.CL","work_id":"150a45e5-3469-496b-b2f3-50ce7822e435","year":2024},"citing_paper":{"arxiv_id":"2606.20603","last_updated":"2026-05-20T14:54:09Z","snapshot_observed_at":"2026-08-03T02:28:08.992682Z","submitted_at":"2026-05-20T14:54:09Z","title":"A Survey of Large Language Models for Perception and Measurement of Human Psychology","version":1},"reference_index":196,"source":"pdf_text","source_observed_at":"2026-06-30T16:59:25.825681Z"},"links":{"cited_paper":"/paper/2401.03729","citing_paper":"/paper/2606.20603"},"observation_digest":"sha256:631fe3bab85e55e935481d6d55d27c179716ae2c033eeb1f976305c990096cca","observation_id":"a615d5ba-9f72-4f48-8ed0-fb98fb70fbb3","resolution":{"observed_at":"2026-06-30T17:04:57.570895Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-20T06:33:59.587034+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-20T06:33:59.587034+00:00","source":"crossref"},{"observed_at":"2026-08-20T06:33:54.927442+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2401.03729","last_updated":"2024-04-01T20:56:11Z","snapshot_observed_at":"2026-08-20T05:18:39.545724Z","submitted_at":"2024-01-08T08:28:08Z","title":"The Butterfly Effect of Altering Prompts: How Small Changes and Jailbreaks Affect Large Language Model Performance","version":3},"cited_work":{"arxiv_id":"2401.03729","doi":null,"metadata_source":"pith","pith_arxiv_id":"2401.03729","snapshot_observed_at":"2026-07-10T15:27:20.049648Z","title":"The butterfly effect of altering prompts: How small changes and jailbreaks affect large language model performance.arXiv preprint arXiv:2401.03729","venue":"cs.CL","work_id":"150a45e5-3469-496b-b2f3-50ce7822e435","year":2024},"citing_paper":{"arxiv_id":"2607.07918","last_updated":"2026-08-14T18:25:59Z","snapshot_observed_at":"2026-08-20T17:09:47.116748Z","submitted_at":"2026-07-08T21:03:27Z","title":"Efficient Safety Alignment of Language Models via Latent Personality Traits","version":1},"reference_index":37,"source":"arxiv_source","source_observed_at":"2026-07-10T15:26:23.290009Z"},"links":{"cited_paper":"/paper/2401.03729","citing_paper":"/paper/2607.07918"},"observation_digest":"sha256:96f8574211124291a3bef6ba54e1d87fe7e807a5af16975b0d9a218a5d38e336","observation_id":"022ce466-ebc4-422d-833a-d5511697f8f2","resolution":{"observed_at":"2026-07-10T15:27:20.050919Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-20T06:33:59.587034+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-20T06:33:59.587034+00:00","source":"crossref"},{"observed_at":"2026-08-20T06:33:54.927442+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2401.03729","last_updated":"2024-04-01T20:56:11Z","snapshot_observed_at":"2026-08-20T05:18:39.545724Z","submitted_at":"2024-01-08T08:28:08Z","title":"The Butterfly Effect of Altering Prompts: How Small Changes and Jailbreaks Affect Large Language Model Performance","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2401.03729","snapshot_observed_at":"2026-08-06T00:31:32.404917Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2608.01193","last_updated":"2026-08-02T12:18:00Z","snapshot_observed_at":"2026-08-17T18:26:14.803657Z","submitted_at":"2026-08-02T12:18:00Z","title":"Humans Are More Diverse: Frontier LLMs Show Extreme Policies in Idealised AI Development Races","version":1},"reference_index":30,"source":"pdf_text","source_observed_at":"2026-08-06T00:31:32.404917Z"},"links":{"cited_paper":"/paper/2401.03729","citing_paper":"/paper/2608.01193"},"observation_digest":"sha256:580f55e88893456612d56b7abc12e859cd1201e6d5faf8db626879529be7f872","observation_id":"14c74172-9201-4e24-8e74-faa72f1dc988","resolution":{"observed_at":"2026-08-06T00:31:32.404917Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}}],"links":{"evidence":"/evidence","html":"/paper/2401.03729/citation-record","integrity":"/paper/2401.03729/integrity","json":"/paper/2401.03729/citation-record.json","paper":"/paper/2401.03729"},"outbound":[],"paper":{"arxiv_id":"2401.03729","last_updated":"2024-04-01T20:56:11Z","latest_version":3,"primary_category":"cs.CL","snapshot_observed_at":"2026-08-20T05:18:39.545724Z","submitted_at":"2024-01-08T08:28:08Z","title":"The Butterfly Effect of Altering Prompts: How Small Changes and Jailbreaks Affect Large Language Model Performance"},"reference_resolution":{"displayed":0,"state_counts":{"malformed_identifier":0,"metadata_mismatch":0,"parse_uncertain":0,"unresolved":0,"verified_exact":0,"verified_fuzzy":0},"total_outbound_references":0},"refusal":"A citation records a reference. It does not transfer a finding from one paper to another.","schema":"pith.paper-citation-record.v1","standing_sources":[{"observed_at":"2026-08-20T06:33:59.587034+00:00","source":"crossref"},{"observed_at":"2026-08-20T06:33:54.927442+00:00","source":"retraction_watch"}],"thesis":"As of 20 August 2026, this Paper Citation Record lists 0 of 0 outbound references and 19 inbound Pith citation observations for arXiv:2401.03729."}