{"as_of":"2026-08-08T21:15:00Z","caps":{"database_statements":6,"inbound":100,"outbound":100},"context_digest":"sha256:59642280d34aca4ede9a785894216e5bc730c8d8cee83be954c8d53cbf399695","coverage":[{"denominator":0,"lane":"reference_resolution","note":"Typed states for the displayed outbound observations.","records_observed":0,"source":"paper_references, paper_reference_links","source_observed_at":null,"state":"measured"},{"denominator":32,"lane":"standing_notices","note":"One-hop event checks from named stored sources.","records_observed":32,"source":"scholarly_work_events, retraction_status_cache","source_observed_at":"2026-08-08T06:32:00.761636+00:00","state":"measured"},{"denominator":32,"lane":"inbound_itemization","note":"Pith citing papers itemized under the disclosed page cap.","records_observed":32,"source":"paper_references, paper_reference_links","source_observed_at":"2026-08-07T11:59:22.616939Z","state":"measured"},{"denominator":1,"lane":"external_citation_measurements","note":"A source-named dated measurement, never combined with another source.","records_observed":1,"source":"arxiv_reference","source_observed_at":"2026-08-05T02:28:24.338817Z","state":"measured"}],"external_citation_measurements":[{"count":54,"observed_at":"2026-08-05T02:28:24.338817Z","source":"arxiv_reference"}],"inbound":[{"citation":{"cited_paper":{"arxiv_id":"2305.14387","last_updated":"2024-01-08T04:46:56Z","snapshot_observed_at":"2026-07-06T15:31:51.249701Z","submitted_at":"2023-05-22T17:55:50Z","title":"AlpacaFarm: A Simulation Framework for Methods that Learn from Human Feedback","version":4},"cited_work":{"arxiv_id":"2305.14387","doi":"10.48550/arxiv.2305.14387","metadata_source":"arxiv_reference","pith_arxiv_id":"2305.14387","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Yann Dubois, Xuechen Li, Rohan Taori, Tianyi Zhang, Ishaan Gulrajani, Jimmy Ba, Carlos Guestrin, Percy Liang, and Tatsunori B Hashimoto","venue":"arXiv (Cornell University)","work_id":"a875adf2-8826-466d-bf52-896ee15632ea","year":2023},"citing_paper":{"arxiv_id":"2305.17926","last_updated":"2023-08-30T13:22:35Z","snapshot_observed_at":"2026-08-08T03:32:23.667881Z","submitted_at":"2023-05-29T07:41:03Z","title":"Large Language Models are not Fair Evaluators","version":2},"reference_index":78,"source":"arxiv_source","source_observed_at":"2026-05-17T12:10:42.248005Z"},"links":{"cited_paper":"/paper/2305.14387","citing_paper":"/paper/2305.17926"},"observation_digest":"sha256:176b3e716222ed6d8b796a94d246aab2c12a33c673b5807e1dc1673229f12767","observation_id":"c633b906-40bb-48c1-b2ad-7cb858c7dcb6","resolution":{"observed_at":"2026-05-17T12:10:42.405853Z","resolver_source":"arxiv_id","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2305.14387","last_updated":"2024-01-08T04:46:56Z","snapshot_observed_at":"2026-07-06T15:31:51.249701Z","submitted_at":"2023-05-22T17:55:50Z","title":"AlpacaFarm: A Simulation Framework for Methods that Learn from Human Feedback","version":4},"cited_work":{"arxiv_id":"2305.14387","doi":"10.48550/arxiv.2305.14387","metadata_source":"arxiv_reference","pith_arxiv_id":"2305.14387","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Yann Dubois, Xuechen Li, Rohan Taori, Tianyi Zhang, Ishaan Gulrajani, Jimmy Ba, Carlos Guestrin, Percy Liang, and Tatsunori B Hashimoto","venue":"arXiv (Cornell University)","work_id":"a875adf2-8826-466d-bf52-896ee15632ea","year":2023},"citing_paper":{"arxiv_id":"2306.05685","last_updated":"2023-12-24T02:01:34Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2023-06-09T05:55:52Z","title":"Judging LLM-as-a-Judge with MT-Bench and Chatbot Arena","version":4},"reference_index":14,"source":"pdf_text","source_observed_at":"2026-05-10T18:52:59.033645Z"},"links":{"cited_paper":"/paper/2305.14387","citing_paper":"/paper/2306.05685"},"observation_digest":"sha256:fa48f10926994f5574fad8990adf7b9ed4dfe177967ccc6f3ad20b8112324001","observation_id":"41c776e8-d9d2-4e35-a3a2-e329d193b2fa","resolution":{"observed_at":"2026-05-10T18:52:59.178783Z","resolver_source":"arxiv_id","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2305.14387","last_updated":"2024-01-08T04:46:56Z","snapshot_observed_at":"2026-07-06T15:31:51.249701Z","submitted_at":"2023-05-22T17:55:50Z","title":"AlpacaFarm: A Simulation Framework for Methods that Learn from Human Feedback","version":4},"cited_work":{"arxiv_id":"2305.14387","doi":"10.48550/arxiv.2305.14387","metadata_source":"arxiv_reference","pith_arxiv_id":"2305.14387","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Yann Dubois, Xuechen Li, Rohan Taori, Tianyi Zhang, Ishaan Gulrajani, Jimmy Ba, Carlos Guestrin, Percy Liang, and Tatsunori B Hashimoto","venue":"arXiv (Cornell University)","work_id":"a875adf2-8826-466d-bf52-896ee15632ea","year":2023},"citing_paper":{"arxiv_id":"2306.11644","last_updated":"2023-10-02T06:12:30Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2023-06-20T16:14:25Z","title":"Textbooks Are All You Need","version":2},"reference_index":11,"source":"pdf_text","source_observed_at":"2026-05-13T04:44:03.148223Z"},"links":{"cited_paper":"/paper/2305.14387","citing_paper":"/paper/2306.11644"},"observation_digest":"sha256:477aa4150a17d0f0bbfa6e5ff666e7429f79d91d43dea1a5aeff3240c853056f","observation_id":"d04362a3-2474-402b-83e2-bd6d5380fe24","resolution":{"observed_at":"2026-05-13T04:44:03.204657Z","resolver_source":"arxiv_id","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2305.14387","last_updated":"2024-01-08T04:46:56Z","snapshot_observed_at":"2026-07-06T15:31:51.249701Z","submitted_at":"2023-05-22T17:55:50Z","title":"AlpacaFarm: A Simulation Framework for Methods that Learn from Human Feedback","version":4},"cited_work":{"arxiv_id":"2305.14387","doi":"10.48550/arxiv.2305.14387","metadata_source":"arxiv_reference","pith_arxiv_id":"2305.14387","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Yann Dubois, Xuechen Li, Rohan Taori, Tianyi Zhang, Ishaan Gulrajani, Jimmy Ba, Carlos Guestrin, Percy Liang, and Tatsunori B Hashimoto","venue":"arXiv (Cornell University)","work_id":"a875adf2-8826-466d-bf52-896ee15632ea","year":2023},"citing_paper":{"arxiv_id":"2307.06435","last_updated":"2024-10-17T01:10:40Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2023-07-12T20:01:52Z","title":"A Comprehensive Overview of Large Language Models","version":10},"reference_index":174,"source":"pdf_text","source_observed_at":"2026-05-19T20:28:38.900026Z"},"links":{"cited_paper":"/paper/2305.14387","citing_paper":"/paper/2307.06435"},"observation_digest":"sha256:921aa67069b16df0122df1345ea12170233536ddb0ed416e34f5843f50a2c583","observation_id":"0404f66a-62af-479f-a6d8-65ffa17db7b1","resolution":{"observed_at":"2026-05-19T20:28:39.557301Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2305.14387","last_updated":"2024-01-08T04:46:56Z","snapshot_observed_at":"2026-07-06T15:31:51.249701Z","submitted_at":"2023-05-22T17:55:50Z","title":"AlpacaFarm: A Simulation Framework for Methods that Learn from Human Feedback","version":4},"cited_work":{"arxiv_id":"2305.14387","doi":"10.48550/arxiv.2305.14387","metadata_source":"arxiv_reference","pith_arxiv_id":"2305.14387","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Yann Dubois, Xuechen Li, Rohan Taori, Tianyi Zhang, Ishaan Gulrajani, Jimmy Ba, Carlos Guestrin, Percy Liang, and Tatsunori B Hashimoto","venue":"arXiv (Cornell University)","work_id":"a875adf2-8826-466d-bf52-896ee15632ea","year":2023},"citing_paper":{"arxiv_id":"2309.00614","last_updated":"2023-09-04T17:47:36Z","snapshot_observed_at":"2026-07-06T16:13:23.343694Z","submitted_at":"2023-09-01T17:59:44Z","title":"Baseline Defenses for Adversarial Attacks Against Aligned Language Models","version":2},"reference_index":11,"source":"arxiv_source","source_observed_at":"2026-05-13T23:24:39.835347Z"},"links":{"cited_paper":"/paper/2305.14387","citing_paper":"/paper/2309.00614"},"observation_digest":"sha256:8cc8d8de1a6b2410b813114873d8c1896509c8e190c363f84ca1ef2ef8518c35","observation_id":"cd184d77-cc36-4485-9d2c-7d16d87aa969","resolution":{"observed_at":"2026-05-13T23:24:40.111516Z","resolver_source":"arxiv_id","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2305.14387","last_updated":"2024-01-08T04:46:56Z","snapshot_observed_at":"2026-07-06T15:31:51.249701Z","submitted_at":"2023-05-22T17:55:50Z","title":"AlpacaFarm: A Simulation Framework for Methods that Learn from Human Feedback","version":4},"cited_work":{"arxiv_id":"2305.14387","doi":"10.48550/arxiv.2305.14387","metadata_source":"arxiv_reference","pith_arxiv_id":"2305.14387","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Yann Dubois, Xuechen Li, Rohan Taori, Tianyi Zhang, Ishaan Gulrajani, Jimmy Ba, Carlos Guestrin, Percy Liang, and Tatsunori B Hashimoto","venue":"arXiv (Cornell University)","work_id":"a875adf2-8826-466d-bf52-896ee15632ea","year":2023},"citing_paper":{"arxiv_id":"2309.11495","last_updated":"2023-09-25T15:25:49Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2023-09-20T17:50:55Z","title":"Chain-of-Verification Reduces Hallucination in Large Language Models","version":2},"reference_index":138,"source":"arxiv_source","source_observed_at":"2026-05-18T01:06:49.811982Z"},"links":{"cited_paper":"/paper/2305.14387","citing_paper":"/paper/2309.11495"},"observation_digest":"sha256:329b96ff7c6fa65814a60215029289325336f4faf433f79e5764a16ce319b2b3","observation_id":"fbf82ab5-4dae-4e5f-ae5a-b0b1bec77a67","resolution":{"observed_at":"2026-05-18T01:06:50.331746Z","resolver_source":"arxiv_id","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2305.14387","last_updated":"2024-01-08T04:46:56Z","snapshot_observed_at":"2026-07-06T15:31:51.249701Z","submitted_at":"2023-05-22T17:55:50Z","title":"AlpacaFarm: A Simulation Framework for Methods that Learn from Human Feedback","version":4},"cited_work":{"arxiv_id":"2305.14387","doi":"10.48550/arxiv.2305.14387","metadata_source":"arxiv_reference","pith_arxiv_id":"2305.14387","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Yann Dubois, Xuechen Li, Rohan Taori, Tianyi Zhang, Ishaan Gulrajani, Jimmy Ba, Carlos Guestrin, Percy Liang, and Tatsunori B Hashimoto","venue":"arXiv (Cornell University)","work_id":"a875adf2-8826-466d-bf52-896ee15632ea","year":2023},"citing_paper":{"arxiv_id":"2309.14525","last_updated":"2023-09-25T20:59:33Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2023-09-25T20:59:33Z","title":"Aligning Large Multimodal Models with Factually Augmented RLHF","version":1},"reference_index":12,"source":"pdf_text","source_observed_at":"2026-05-15T17:58:17.699042Z"},"links":{"cited_paper":"/paper/2305.14387","citing_paper":"/paper/2309.14525"},"observation_digest":"sha256:fb70dda63bb2c9381d3322d15a20916c837fc7848e2931522c812e0f1fd2831f","observation_id":"59a8c904-4c23-47ce-8f37-ff774f2cd88e","resolution":{"observed_at":"2026-05-15T17:58:17.843448Z","resolver_source":"arxiv_id","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2305.14387","last_updated":"2024-01-08T04:46:56Z","snapshot_observed_at":"2026-07-06T15:31:51.249701Z","submitted_at":"2023-05-22T17:55:50Z","title":"AlpacaFarm: A Simulation Framework for Methods that Learn from Human Feedback","version":4},"cited_work":{"arxiv_id":"2305.14387","doi":"10.48550/arxiv.2305.14387","metadata_source":"arxiv_reference","pith_arxiv_id":"2305.14387","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Yann Dubois, Xuechen Li, Rohan Taori, Tianyi Zhang, Ishaan Gulrajani, Jimmy Ba, Carlos Guestrin, Percy Liang, and Tatsunori B Hashimoto","venue":"arXiv (Cornell University)","work_id":"a875adf2-8826-466d-bf52-896ee15632ea","year":2023},"citing_paper":{"arxiv_id":"2310.11511","last_updated":"2023-10-17T18:18:32Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2023-10-17T18:18:32Z","title":"Self-RAG: Learning to Retrieve, Generate, and Critique through Self-Reflection","version":1},"reference_index":137,"source":"arxiv_source","source_observed_at":"2026-05-12T14:15:10.907921Z"},"links":{"cited_paper":"/paper/2305.14387","citing_paper":"/paper/2310.11511"},"observation_digest":"sha256:722abeda36b0cd226ef8b5a574826b77679f58193289ad79529a34df4339d9cb","observation_id":"d214e8c3-495a-4794-b92a-6fdfe98f9429","resolution":{"observed_at":"2026-05-12T14:15:11.113391Z","resolver_source":"arxiv_id","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2305.14387","last_updated":"2024-01-08T04:46:56Z","snapshot_observed_at":"2026-07-06T15:31:51.249701Z","submitted_at":"2023-05-22T17:55:50Z","title":"AlpacaFarm: A Simulation Framework for Methods that Learn from Human Feedback","version":4},"cited_work":{"arxiv_id":"2305.14387","doi":"10.48550/arxiv.2305.14387","metadata_source":"arxiv_reference","pith_arxiv_id":"2305.14387","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Yann Dubois, Xuechen Li, Rohan Taori, Tianyi Zhang, Ishaan Gulrajani, Jimmy Ba, Carlos Guestrin, Percy Liang, and Tatsunori B Hashimoto","venue":"arXiv (Cornell University)","work_id":"a875adf2-8826-466d-bf52-896ee15632ea","year":2023},"citing_paper":{"arxiv_id":"2311.16867","last_updated":"2023-11-29T19:45:10Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2023-11-28T15:12:47Z","title":"The Falcon Series of Open Language Models","version":2},"reference_index":276,"source":"arxiv_source","source_observed_at":"2026-05-16T09:46:09.701440Z"},"links":{"cited_paper":"/paper/2305.14387","citing_paper":"/paper/2311.16867"},"observation_digest":"sha256:bdb97735b3b1ffd9b6fcd1fce960287fb1ed35eee22d5b2ca591ef954bb16b90","observation_id":"6182808e-2263-4bc4-ae30-791c538dcd7c","resolution":{"observed_at":"2026-05-16T09:46:10.083261Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2305.14387","last_updated":"2024-01-08T04:46:56Z","snapshot_observed_at":"2026-07-06T15:31:51.249701Z","submitted_at":"2023-05-22T17:55:50Z","title":"AlpacaFarm: A Simulation Framework for Methods that Learn from Human Feedback","version":4},"cited_work":{"arxiv_id":"2305.14387","doi":"10.48550/arxiv.2305.14387","metadata_source":"arxiv_reference","pith_arxiv_id":"2305.14387","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Yann Dubois, Xuechen Li, Rohan Taori, Tianyi Zhang, Ishaan Gulrajani, Jimmy Ba, Carlos Guestrin, Percy Liang, and Tatsunori B Hashimoto","venue":"arXiv (Cornell University)","work_id":"a875adf2-8826-466d-bf52-896ee15632ea","year":2023},"citing_paper":{"arxiv_id":"2401.10020","last_updated":"2025-03-28T00:06:51Z","snapshot_observed_at":"2026-08-07T08:02:34.857823Z","submitted_at":"2024-01-18T14:43:47Z","title":"Self-Rewarding Language Models","version":3},"reference_index":95,"source":"arxiv_source","source_observed_at":"2026-05-13T12:01:42.290502Z"},"links":{"cited_paper":"/paper/2305.14387","citing_paper":"/paper/2401.10020"},"observation_digest":"sha256:51bf751f7603ad7a8520f467d486daf84fdfa697cdb26aef4feff3dc96d2a66c","observation_id":"650f1700-58a3-4bc1-8196-94d98ab8091e","resolution":{"observed_at":"2026-05-13T12:01:42.411055Z","resolver_source":"arxiv_id","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2305.14387","last_updated":"2024-01-08T04:46:56Z","snapshot_observed_at":"2026-07-06T15:31:51.249701Z","submitted_at":"2023-05-22T17:55:50Z","title":"AlpacaFarm: A Simulation Framework for Methods that Learn from Human Feedback","version":4},"cited_work":{"arxiv_id":"2305.14387","doi":"10.48550/arxiv.2305.14387","metadata_source":"arxiv_reference","pith_arxiv_id":"2305.14387","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Yann Dubois, Xuechen Li, Rohan Taori, Tianyi Zhang, Ishaan Gulrajani, Jimmy Ba, Carlos Guestrin, Percy Liang, and Tatsunori B Hashimoto","venue":"arXiv (Cornell University)","work_id":"a875adf2-8826-466d-bf52-896ee15632ea","year":2023},"citing_paper":{"arxiv_id":"2401.15884","last_updated":"2024-10-07T02:19:21Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2024-01-29T04:36:39Z","title":"Corrective Retrieval Augmented Generation","version":3},"reference_index":8,"source":"arxiv_source","source_observed_at":"2026-05-12T11:19:17.120464Z"},"links":{"cited_paper":"/paper/2305.14387","citing_paper":"/paper/2401.15884"},"observation_digest":"sha256:c4668ce73dd8dcc8d08240a1e445d12eeb4088e9d933f8a2b16b8a43809bb09a","observation_id":"c0008a37-ac01-4e72-8eb4-452843efdb90","resolution":{"observed_at":"2026-05-12T11:19:17.301138Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2305.14387","last_updated":"2024-01-08T04:46:56Z","snapshot_observed_at":"2026-07-06T15:31:51.249701Z","submitted_at":"2023-05-22T17:55:50Z","title":"AlpacaFarm: A Simulation Framework for Methods that Learn from Human Feedback","version":4},"cited_work":{"arxiv_id":"2305.14387","doi":"10.48550/arxiv.2305.14387","metadata_source":"arxiv_reference","pith_arxiv_id":"2305.14387","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Yann Dubois, Xuechen Li, Rohan Taori, Tianyi Zhang, Ishaan Gulrajani, Jimmy Ba, Carlos Guestrin, Percy Liang, and Tatsunori B Hashimoto","venue":"arXiv (Cornell University)","work_id":"a875adf2-8826-466d-bf52-896ee15632ea","year":2023},"citing_paper":{"arxiv_id":"2406.04244","last_updated":"2024-06-06T16:41:39Z","snapshot_observed_at":"2026-07-30T15:43:06.151242Z","submitted_at":"2024-06-06T16:41:39Z","title":"Benchmark Data Contamination of Large Language Models: A Survey","version":1},"reference_index":40,"source":"pdf_text","source_observed_at":"2026-05-22T23:10:40.420241Z"},"links":{"cited_paper":"/paper/2305.14387","citing_paper":"/paper/2406.04244"},"observation_digest":"sha256:39c7c9661a6eb20c8f346568b3972a53cfb21cba7198664ff8b6accf829f4f88","observation_id":"98eb55fd-94b7-446f-b2ff-518e5b4fab1b","resolution":{"observed_at":"2026-05-22T23:10:40.926678Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2305.14387","last_updated":"2024-01-08T04:46:56Z","snapshot_observed_at":"2026-07-06T15:31:51.249701Z","submitted_at":"2023-05-22T17:55:50Z","title":"AlpacaFarm: A Simulation Framework for Methods that Learn from Human Feedback","version":4},"cited_work":{"arxiv_id":"2305.14387","doi":"10.48550/arxiv.2305.14387","metadata_source":"arxiv_reference","pith_arxiv_id":"2305.14387","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Yann Dubois, Xuechen Li, Rohan Taori, Tianyi Zhang, Ishaan Gulrajani, Jimmy Ba, Carlos Guestrin, Percy Liang, and Tatsunori B Hashimoto","venue":"arXiv (Cornell University)","work_id":"a875adf2-8826-466d-bf52-896ee15632ea","year":2023},"citing_paper":{"arxiv_id":"2408.00724","last_updated":"2025-03-03T07:53:32Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2024-08-01T17:16:04Z","title":"Inference Scaling Laws: An Empirical Analysis of Compute-Optimal Inference for Problem-Solving with Language Models","version":3},"reference_index":288,"source":"arxiv_source","source_observed_at":"2026-05-18T06:38:36.517935Z"},"links":{"cited_paper":"/paper/2305.14387","citing_paper":"/paper/2408.00724"},"observation_digest":"sha256:6d16e3cee44dad847ec17f0d6de59dca3ccb7249237985a4da76bb472423db28","observation_id":"a4007580-dd99-4d7b-99b3-570889290493","resolution":{"observed_at":"2026-05-18T06:38:37.165372Z","resolver_source":"arxiv_id","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2305.14387","last_updated":"2024-01-08T04:46:56Z","snapshot_observed_at":"2026-07-06T15:31:51.249701Z","submitted_at":"2023-05-22T17:55:50Z","title":"AlpacaFarm: A Simulation Framework for Methods that Learn from Human Feedback","version":4},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2305.14387","snapshot_observed_at":"2026-08-07T11:59:22.616939Z","title":"Alpacafarm: A simulation framework for methods that learn from human feedback","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2506.01055","last_updated":"2025-06-01T15:48:06Z","snapshot_observed_at":"2026-08-08T14:39:19.171040Z","submitted_at":"2025-06-01T15:48:06Z","title":"Simple Prompt Injection Attacks Can Leak Personal Data Observed by LLM Agents During Task Execution","version":1},"reference_index":12,"source":"pdf_text","source_observed_at":"2026-08-07T11:59:22.616939Z"},"links":{"cited_paper":"/paper/2305.14387","citing_paper":"/paper/2506.01055"},"observation_digest":"sha256:b6aa90ae41335562d0a6573716cc85dab1161b4f5b1ac06a805dc6b9dddc87e2","observation_id":"662f2b23-c9d8-4462-b276-496aa0e30648","resolution":{"observed_at":"2026-08-07T11:59:22.616939Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2305.14387","last_updated":"2024-01-08T04:46:56Z","snapshot_observed_at":"2026-07-06T15:31:51.249701Z","submitted_at":"2023-05-22T17:55:50Z","title":"AlpacaFarm: A Simulation Framework for Methods that Learn from Human Feedback","version":4},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2305.14387","snapshot_observed_at":"2026-08-07T05:21:58.639987Z","title":"Hashimoto","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2506.08266","last_updated":"2025-06-09T22:03:56Z","snapshot_observed_at":"2026-08-07T05:12:42.307191Z","submitted_at":"2025-06-09T22:03:56Z","title":"Reinforcement Learning from Human Feedback with High-Confidence Safety Constraints","version":1},"reference_index":11,"source":"arxiv_source","source_observed_at":"2026-08-07T05:21:58.639987Z"},"links":{"cited_paper":"/paper/2305.14387","citing_paper":"/paper/2506.08266"},"observation_digest":"sha256:37873435f72e90d3fca776452bce0d32f018108bc9e7213c2549d0bc610684df","observation_id":"e8c042f1-7573-493f-aac2-c9a443c337f4","resolution":{"observed_at":"2026-08-07T05:21:58.639987Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2305.14387","last_updated":"2024-01-08T04:46:56Z","snapshot_observed_at":"2026-07-06T15:31:51.249701Z","submitted_at":"2023-05-22T17:55:50Z","title":"AlpacaFarm: A Simulation Framework for Methods that Learn from Human Feedback","version":4},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2305.14387","snapshot_observed_at":"2026-08-06T14:51:53.887056Z","title":"arXiv:2305.14387","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2507.17477","last_updated":"2025-07-23T13:00:00Z","snapshot_observed_at":"2026-08-07T11:10:51.753612Z","submitted_at":"2025-07-23T13:00:00Z","title":"An Uncertainty-Driven Adaptive Self-Alignment Framework for Large Language Models","version":1},"reference_index":3,"source":"pdf_text","source_observed_at":"2026-08-06T14:51:53.887056Z"},"links":{"cited_paper":"/paper/2305.14387","citing_paper":"/paper/2507.17477"},"observation_digest":"sha256:9bae0f1119c3116199c94c39a4b8715802c5b395641248fff05a91fa403e1bf6","observation_id":"7ec1677e-173c-4452-ab99-ce10bfeee767","resolution":{"observed_at":"2026-08-06T14:51:53.887056Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2305.14387","last_updated":"2024-01-08T04:46:56Z","snapshot_observed_at":"2026-07-06T15:31:51.249701Z","submitted_at":"2023-05-22T17:55:50Z","title":"AlpacaFarm: A Simulation Framework for Methods that Learn from Human Feedback","version":4},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2305.14387","snapshot_observed_at":"2026-08-04T17:56:38.763239Z","title":"Hashimoto","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2509.10397","last_updated":"2026-07-25T20:09:07Z","snapshot_observed_at":"2026-08-07T01:01:23.391389Z","submitted_at":"2025-09-12T16:44:34Z","title":"RecoWorld: Building Simulated Environments for Agentic Recommender Systems","version":3},"reference_index":13,"source":"arxiv_source","source_observed_at":"2026-08-04T17:56:38.763239Z"},"links":{"cited_paper":"/paper/2305.14387","citing_paper":"/paper/2509.10397"},"observation_digest":"sha256:060849eca78ef7fe4d24ecf4ce63d150b4118fcbdbb5a9f1dab4d3a1050087d5","observation_id":"c0e8ab4e-2af1-4680-84be-b13782bead09","resolution":{"observed_at":"2026-08-04T17:56:38.763239Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2305.14387","last_updated":"2024-01-08T04:46:56Z","snapshot_observed_at":"2026-07-06T15:31:51.249701Z","submitted_at":"2023-05-22T17:55:50Z","title":"AlpacaFarm: A Simulation Framework for Methods that Learn from Human Feedback","version":4},"cited_work":{"arxiv_id":"2305.14387","doi":"10.48550/arxiv.2305.14387","metadata_source":"arxiv_reference","pith_arxiv_id":"2305.14387","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Yann Dubois, Xuechen Li, Rohan Taori, Tianyi Zhang, Ishaan Gulrajani, Jimmy Ba, Carlos Guestrin, Percy Liang, and Tatsunori B Hashimoto","venue":"arXiv (Cornell University)","work_id":"a875adf2-8826-466d-bf52-896ee15632ea","year":2023},"citing_paper":{"arxiv_id":"2509.23542","last_updated":"2026-04-19T04:59:30Z","snapshot_observed_at":"2026-08-03T18:10:00.237919Z","submitted_at":"2025-09-28T00:43:52Z","title":"On the Shelf Life of Fine-Tuned LLM-Judges: Future-Proofing, Backward-Compatibility, and Question Generalization","version":2},"reference_index":9,"source":"arxiv_source","source_observed_at":"2026-05-18T12:53:45.767341Z"},"links":{"cited_paper":"/paper/2305.14387","citing_paper":"/paper/2509.23542"},"observation_digest":"sha256:ff474b444baf95c93678e0983ec475426b9d09e814d62dfdcb17177a9df96074","observation_id":"0954badc-258c-4482-983e-fd21e59143d4","resolution":{"observed_at":"2026-05-18T12:56:24.568395Z","resolver_source":"arxiv_id","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2305.14387","last_updated":"2024-01-08T04:46:56Z","snapshot_observed_at":"2026-07-06T15:31:51.249701Z","submitted_at":"2023-05-22T17:55:50Z","title":"AlpacaFarm: A Simulation Framework for Methods that Learn from Human Feedback","version":4},"cited_work":{"arxiv_id":"2305.14387","doi":"10.48550/arxiv.2305.14387","metadata_source":"arxiv_reference","pith_arxiv_id":"2305.14387","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Yann Dubois, Xuechen Li, Rohan Taori, Tianyi Zhang, Ishaan Gulrajani, Jimmy Ba, Carlos Guestrin, Percy Liang, and Tatsunori B Hashimoto","venue":"arXiv (Cornell University)","work_id":"a875adf2-8826-466d-bf52-896ee15632ea","year":2023},"citing_paper":{"arxiv_id":"2601.14053","last_updated":"2026-04-16T15:26:17Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2026-01-20T15:06:19Z","title":"LLMOrbit: A Circular Taxonomy of Large Language Models -From Scaling Walls to Agentic AI Systems","version":2},"reference_index":49,"source":"pdf_text","source_observed_at":"2026-05-16T12:47:28.248540Z"},"links":{"cited_paper":"/paper/2305.14387","citing_paper":"/paper/2601.14053"},"observation_digest":"sha256:2f0aac9811be0577a02dbd6bb85af799fe03074300b09cf3f1c00b0e613f9d4f","observation_id":"603a3ffd-b380-48b2-b709-f295c5a3a440","resolution":{"observed_at":"2026-05-16T12:47:53.751787Z","resolver_source":"arxiv_id","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2305.14387","last_updated":"2024-01-08T04:46:56Z","snapshot_observed_at":"2026-07-06T15:31:51.249701Z","submitted_at":"2023-05-22T17:55:50Z","title":"AlpacaFarm: A Simulation Framework for Methods that Learn from Human Feedback","version":4},"cited_work":{"arxiv_id":"2305.14387","doi":"10.48550/arxiv.2305.14387","metadata_source":"arxiv_reference","pith_arxiv_id":"2305.14387","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Yann Dubois, Xuechen Li, Rohan Taori, Tianyi Zhang, Ishaan Gulrajani, Jimmy Ba, Carlos Guestrin, Percy Liang, and Tatsunori B Hashimoto","venue":"arXiv (Cornell University)","work_id":"a875adf2-8826-466d-bf52-896ee15632ea","year":2023},"citing_paper":{"arxiv_id":"2605.00195","last_updated":"2026-05-11T12:48:16Z","snapshot_observed_at":"2026-07-31T18:51:54.041588Z","submitted_at":"2026-04-30T20:20:59Z","title":"Diversity in Large Language Models under Supervised Fine-Tuning","version":1},"reference_index":46,"source":"arxiv_source","source_observed_at":"2026-05-09T20:32:37.788283Z"},"links":{"cited_paper":"/paper/2305.14387","citing_paper":"/paper/2605.00195"},"observation_digest":"sha256:12f1088f652a1a5cdb155168906ee7a392485c15704eb01c268f3ca7a3948bef","observation_id":"9fb779b0-903a-4a0c-898c-e1821bbcfcf5","resolution":{"observed_at":"2026-05-09T20:37:32.112040Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2305.14387","last_updated":"2024-01-08T04:46:56Z","snapshot_observed_at":"2026-07-06T15:31:51.249701Z","submitted_at":"2023-05-22T17:55:50Z","title":"AlpacaFarm: A Simulation Framework for Methods that Learn from Human Feedback","version":4},"cited_work":{"arxiv_id":"2305.14387","doi":"10.48550/arxiv.2305.14387","metadata_source":"arxiv_reference","pith_arxiv_id":"2305.14387","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Yann Dubois, Xuechen Li, Rohan Taori, Tianyi Zhang, Ishaan Gulrajani, Jimmy Ba, Carlos Guestrin, Percy Liang, and Tatsunori B Hashimoto","venue":"arXiv (Cornell University)","work_id":"a875adf2-8826-466d-bf52-896ee15632ea","year":2023},"citing_paper":{"arxiv_id":"2605.00195","last_updated":"2026-05-11T12:48:16Z","snapshot_observed_at":"2026-07-31T18:51:54.041588Z","submitted_at":"2026-04-30T20:20:59Z","title":"Diversity in Large Language Models under Supervised Fine-Tuning","version":2},"reference_index":46,"source":"arxiv_source","source_observed_at":"2026-05-12T03:10:22.314719Z"},"links":{"cited_paper":"/paper/2305.14387","citing_paper":"/paper/2605.00195"},"observation_digest":"sha256:c371e653f194e8860bbd4e168365c5913a62587f9346b3e441d00c5797cb6ee1","observation_id":"42f50f90-736f-45cc-a9cd-a95063fe62fc","resolution":{"observed_at":"2026-05-12T03:11:18.198968Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2305.14387","last_updated":"2024-01-08T04:46:56Z","snapshot_observed_at":"2026-07-06T15:31:51.249701Z","submitted_at":"2023-05-22T17:55:50Z","title":"AlpacaFarm: A Simulation Framework for Methods that Learn from Human Feedback","version":4},"cited_work":{"arxiv_id":"2305.14387","doi":"10.48550/arxiv.2305.14387","metadata_source":"arxiv_reference","pith_arxiv_id":"2305.14387","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Yann Dubois, Xuechen Li, Rohan Taori, Tianyi Zhang, Ishaan Gulrajani, Jimmy Ba, Carlos Guestrin, Percy Liang, and Tatsunori B Hashimoto","venue":"arXiv (Cornell University)","work_id":"a875adf2-8826-466d-bf52-896ee15632ea","year":2023},"citing_paper":{"arxiv_id":"2605.08904","last_updated":"2026-05-09T11:51:34Z","snapshot_observed_at":"2026-07-06T23:21:02.177557Z","submitted_at":"2026-05-09T11:51:34Z","title":"OPT-BENCH: Evaluating the Iterative Self-Optimization of LLM Agents in Large-Scale Search Spaces","version":1},"reference_index":64,"source":"arxiv_source","source_observed_at":"2026-05-12T02:57:15.521594Z"},"links":{"cited_paper":"/paper/2305.14387","citing_paper":"/paper/2605.08904"},"observation_digest":"sha256:b2ffecca5c9390122fc6d11f08339f5a0ca2f19bee8dd8c120b03838684adea2","observation_id":"642a627a-3497-4036-8e18-857282d8e03d","resolution":{"observed_at":"2026-05-12T03:01:18.800129Z","resolver_source":"arxiv_id","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2305.14387","last_updated":"2024-01-08T04:46:56Z","snapshot_observed_at":"2026-07-06T15:31:51.249701Z","submitted_at":"2023-05-22T17:55:50Z","title":"AlpacaFarm: A Simulation Framework for Methods that Learn from Human Feedback","version":4},"cited_work":{"arxiv_id":"2305.14387","doi":"10.48550/arxiv.2305.14387","metadata_source":"arxiv_reference","pith_arxiv_id":"2305.14387","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Yann Dubois, Xuechen Li, Rohan Taori, Tianyi Zhang, Ishaan Gulrajani, Jimmy Ba, Carlos Guestrin, Percy Liang, and Tatsunori B Hashimoto","venue":"arXiv (Cornell University)","work_id":"a875adf2-8826-466d-bf52-896ee15632ea","year":2023},"citing_paper":{"arxiv_id":"2605.23171","last_updated":"2026-05-22T02:43:19Z","snapshot_observed_at":"2026-08-01T19:48:27.330086Z","submitted_at":"2026-05-22T02:43:19Z","title":"Understanding and Improving Noisy Embedding Techniques in Instruction Finetuning","version":1},"reference_index":7,"source":"pdf_text","source_observed_at":"2026-05-25T04:51:05.359636Z"},"links":{"cited_paper":"/paper/2305.14387","citing_paper":"/paper/2605.23171"},"observation_digest":"sha256:e289606e7266dd8097e42d047a9fb11aef7b6a0eea6a966fd99fca42008b9433","observation_id":"24387372-a44b-410c-8017-87796548bf21","resolution":{"observed_at":"2026-05-25T04:55:24.092587Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2305.14387","last_updated":"2024-01-08T04:46:56Z","snapshot_observed_at":"2026-07-06T15:31:51.249701Z","submitted_at":"2023-05-22T17:55:50Z","title":"AlpacaFarm: A Simulation Framework for Methods that Learn from Human Feedback","version":4},"cited_work":{"arxiv_id":"2305.14387","doi":"10.48550/arxiv.2305.14387","metadata_source":"arxiv_reference","pith_arxiv_id":"2305.14387","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Yann Dubois, Xuechen Li, Rohan Taori, Tianyi Zhang, Ishaan Gulrajani, Jimmy Ba, Carlos Guestrin, Percy Liang, and Tatsunori B Hashimoto","venue":"arXiv (Cornell University)","work_id":"a875adf2-8826-466d-bf52-896ee15632ea","year":2023},"citing_paper":{"arxiv_id":"2606.01811","last_updated":"2026-06-01T07:27:43Z","snapshot_observed_at":"2026-08-01T23:50:35.081089Z","submitted_at":"2026-06-01T07:27:43Z","title":"\"I've Seen How This Goes\": Characterizing Diversity via Progressive Conditional Surprise","version":1},"reference_index":7,"source":"arxiv_source","source_observed_at":"2026-06-28T14:51:06.448678Z"},"links":{"cited_paper":"/paper/2305.14387","citing_paper":"/paper/2606.01811"},"observation_digest":"sha256:426926b148049bc3de40492ddd1ea1dd98100c61430d86dea000e7b23cadfed7","observation_id":"d819dc1b-c8ce-47a8-913b-7bdd4b79b5c3","resolution":{"observed_at":"2026-07-01T22:56:20.797161Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2305.14387","last_updated":"2024-01-08T04:46:56Z","snapshot_observed_at":"2026-07-06T15:31:51.249701Z","submitted_at":"2023-05-22T17:55:50Z","title":"AlpacaFarm: A Simulation Framework for Methods that Learn from Human Feedback","version":4},"cited_work":{"arxiv_id":"2305.14387","doi":"10.48550/arxiv.2305.14387","metadata_source":"arxiv_reference","pith_arxiv_id":"2305.14387","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Yann Dubois, Xuechen Li, Rohan Taori, Tianyi Zhang, Ishaan Gulrajani, Jimmy Ba, Carlos Guestrin, Percy Liang, and Tatsunori B Hashimoto","venue":"arXiv (Cornell University)","work_id":"a875adf2-8826-466d-bf52-896ee15632ea","year":2023},"citing_paper":{"arxiv_id":"2606.03137","last_updated":"2026-07-01T08:33:01Z","snapshot_observed_at":"2026-08-05T05:03:21.120024Z","submitted_at":"2026-06-02T04:26:01Z","title":"Think-Before-Speak: From Internal Evaluation to Public Expression in Multi-Agent Social Simulation","version":1},"reference_index":4,"source":"pdf_text","source_observed_at":"2026-06-28T10:24:30.660372Z"},"links":{"cited_paper":"/paper/2305.14387","citing_paper":"/paper/2606.03137"},"observation_digest":"sha256:e874738f0a8e7b4ac544b62c643d5a47a893040e4d2a6e93e611d051e55d225a","observation_id":"87967254-7898-4ef0-9b2e-114ae8fcd293","resolution":{"observed_at":"2026-07-02T03:06:29.378053Z","resolver_source":"arxiv_id","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2305.14387","last_updated":"2024-01-08T04:46:56Z","snapshot_observed_at":"2026-07-06T15:31:51.249701Z","submitted_at":"2023-05-22T17:55:50Z","title":"AlpacaFarm: A Simulation Framework for Methods that Learn from Human Feedback","version":4},"cited_work":{"arxiv_id":"2305.14387","doi":"10.48550/arxiv.2305.14387","metadata_source":"arxiv_reference","pith_arxiv_id":"2305.14387","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Yann Dubois, Xuechen Li, Rohan Taori, Tianyi Zhang, Ishaan Gulrajani, Jimmy Ba, Carlos Guestrin, Percy Liang, and Tatsunori B Hashimoto","venue":"arXiv (Cornell University)","work_id":"a875adf2-8826-466d-bf52-896ee15632ea","year":2023},"citing_paper":{"arxiv_id":"2606.03137","last_updated":"2026-07-01T08:33:01Z","snapshot_observed_at":"2026-08-05T05:03:21.120024Z","submitted_at":"2026-06-02T04:26:01Z","title":"Think-Before-Speak: From Internal Evaluation to Public Expression in Multi-Agent Social Simulation","version":2},"reference_index":7,"source":"pdf_text","source_observed_at":"2026-07-02T23:10:03.733636Z"},"links":{"cited_paper":"/paper/2305.14387","citing_paper":"/paper/2606.03137"},"observation_digest":"sha256:0dcc4fbd3e36f7bb2737db74a800bc99fb4750df0caf90e4c4083670bb776acf","observation_id":"1fa36e31-3765-47a6-896b-81e21a2f5453","resolution":{"observed_at":"2026-07-02T23:17:29.021269Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2305.14387","last_updated":"2024-01-08T04:46:56Z","snapshot_observed_at":"2026-07-06T15:31:51.249701Z","submitted_at":"2023-05-22T17:55:50Z","title":"AlpacaFarm: A Simulation Framework for Methods that Learn from Human Feedback","version":4},"cited_work":{"arxiv_id":"2305.14387","doi":"10.48550/arxiv.2305.14387","metadata_source":"arxiv_reference","pith_arxiv_id":"2305.14387","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Yann Dubois, Xuechen Li, Rohan Taori, Tianyi Zhang, Ishaan Gulrajani, Jimmy Ba, Carlos Guestrin, Percy Liang, and Tatsunori B Hashimoto","venue":"arXiv (Cornell University)","work_id":"a875adf2-8826-466d-bf52-896ee15632ea","year":2023},"citing_paper":{"arxiv_id":"2606.19714","last_updated":"2026-06-18T02:26:05Z","snapshot_observed_at":"2026-08-01T18:27:22.876701Z","submitted_at":"2026-06-18T02:26:05Z","title":"AURA: Adaptive Uncertainty-aware Refinement for LLM-as-a-Judge Auditing","version":1},"reference_index":7,"source":"arxiv_source","source_observed_at":"2026-06-26T15:48:26.303462Z"},"links":{"cited_paper":"/paper/2305.14387","citing_paper":"/paper/2606.19714"},"observation_digest":"sha256:532daa255b6e5e0ab4f03d154f4a6fe11411b1a8975db066e3d4871d9dd68e41","observation_id":"5b97ba62-2938-4c33-adcf-a8c4fc23d29d","resolution":{"observed_at":"2026-07-04T05:39:40.156723Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2305.14387","last_updated":"2024-01-08T04:46:56Z","snapshot_observed_at":"2026-07-06T15:31:51.249701Z","submitted_at":"2023-05-22T17:55:50Z","title":"AlpacaFarm: A Simulation Framework for Methods that Learn from Human Feedback","version":4},"cited_work":{"arxiv_id":"2305.14387","doi":"10.48550/arxiv.2305.14387","metadata_source":"arxiv_reference","pith_arxiv_id":"2305.14387","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Yann Dubois, Xuechen Li, Rohan Taori, Tianyi Zhang, Ishaan Gulrajani, Jimmy Ba, Carlos Guestrin, Percy Liang, and Tatsunori B Hashimoto","venue":"arXiv (Cornell University)","work_id":"a875adf2-8826-466d-bf52-896ee15632ea","year":2023},"citing_paper":{"arxiv_id":"2606.19993","last_updated":"2026-06-18T09:31:31Z","snapshot_observed_at":"2026-08-07T06:04:20.969024Z","submitted_at":"2026-06-18T09:31:31Z","title":"Activation- and Influence-Aware Ranks (AIR): Function-Preserving SVD Compression for LLMs","version":1},"reference_index":174,"source":"arxiv_source","source_observed_at":"2026-06-26T18:09:17.414031Z"},"links":{"cited_paper":"/paper/2305.14387","citing_paper":"/paper/2606.19993"},"observation_digest":"sha256:af5c1cc78b4fdaaff9c1c4224d04d33fecff4ee87fe29a9c4ddc6ec6f72deaeb","observation_id":"c68d1aae-9a53-4eb8-8a41-dcb39cc5f058","resolution":{"observed_at":"2026-07-04T03:19:31.548571Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2305.14387","last_updated":"2024-01-08T04:46:56Z","snapshot_observed_at":"2026-07-06T15:31:51.249701Z","submitted_at":"2023-05-22T17:55:50Z","title":"AlpacaFarm: A Simulation Framework for Methods that Learn from Human Feedback","version":4},"cited_work":{"arxiv_id":"2305.14387","doi":"10.48550/arxiv.2305.14387","metadata_source":"arxiv_reference","pith_arxiv_id":"2305.14387","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Yann Dubois, Xuechen Li, Rohan Taori, Tianyi Zhang, Ishaan Gulrajani, Jimmy Ba, Carlos Guestrin, Percy Liang, and Tatsunori B Hashimoto","venue":"arXiv (Cornell University)","work_id":"a875adf2-8826-466d-bf52-896ee15632ea","year":2023},"citing_paper":{"arxiv_id":"2606.25445","last_updated":"2026-06-24T06:15:24Z","snapshot_observed_at":"2026-07-06T23:59:52.373272Z","submitted_at":"2026-06-24T06:15:24Z","title":"C3-Bench: A Context-Aware Change Captioning Benchmark","version":1},"reference_index":26,"source":"pdf_text","source_observed_at":"2026-06-25T21:02:52.529391Z"},"links":{"cited_paper":"/paper/2305.14387","citing_paper":"/paper/2606.25445"},"observation_digest":"sha256:c75d61d8e2bfb70d64483478cd72a118c15d5d78dcfdc499ead7dae6215aa02d","observation_id":"004c6577-3963-4e08-ba4f-a124e17a89a7","resolution":{"observed_at":"2026-07-04T19:50:10.108934Z","resolver_source":"arxiv_id","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2305.14387","last_updated":"2024-01-08T04:46:56Z","snapshot_observed_at":"2026-07-06T15:31:51.249701Z","submitted_at":"2023-05-22T17:55:50Z","title":"AlpacaFarm: A Simulation Framework for Methods that Learn from Human Feedback","version":4},"cited_work":{"arxiv_id":"2305.14387","doi":"10.48550/arxiv.2305.14387","metadata_source":"arxiv_reference","pith_arxiv_id":"2305.14387","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Yann Dubois, Xuechen Li, Rohan Taori, Tianyi Zhang, Ishaan Gulrajani, Jimmy Ba, Carlos Guestrin, Percy Liang, and Tatsunori B Hashimoto","venue":"arXiv (Cornell University)","work_id":"a875adf2-8826-466d-bf52-896ee15632ea","year":2023},"citing_paper":{"arxiv_id":"2606.26346","last_updated":"2026-06-24T19:38:21Z","snapshot_observed_at":"2026-08-08T08:43:50.725073Z","submitted_at":"2026-06-24T19:38:21Z","title":"How Do Tool-Augmented LLM Agents Perform on Real-World Energy Analytics Tasks?","version":1},"reference_index":46,"source":"pdf_text","source_observed_at":"2026-06-26T01:34:10.103638Z"},"links":{"cited_paper":"/paper/2305.14387","citing_paper":"/paper/2606.26346"},"observation_digest":"sha256:ebf9104a7eeafa7dbe83d7ba0fb9df3294ec881f36482584edc73e272d41d996","observation_id":"4fbfc001-c736-41d4-ac46-390a0ee1ca24","resolution":{"observed_at":"2026-07-04T15:29:56.664257Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2305.14387","last_updated":"2024-01-08T04:46:56Z","snapshot_observed_at":"2026-07-06T15:31:51.249701Z","submitted_at":"2023-05-22T17:55:50Z","title":"AlpacaFarm: A Simulation Framework for Methods that Learn from Human Feedback","version":4},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2305.14387","snapshot_observed_at":"2026-07-11T13:53:36.775836Z","title":"arXiv preprint arXiv:2305.14387 , year=","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2607.04763","last_updated":"2026-07-26T14:17:18Z","snapshot_observed_at":"2026-08-02T10:24:43.977557Z","submitted_at":"2026-07-06T07:56:53Z","title":"Multi-Turn On-Policy Distillation with Prefix Replay","version":1},"reference_index":268,"source":"arxiv_source","source_observed_at":"2026-07-11T13:53:36.775836Z"},"links":{"cited_paper":"/paper/2305.14387","citing_paper":"/paper/2607.04763"},"observation_digest":"sha256:5ca5ddcbc1792b0bddb5b20fa3befed0d85e90e66227eb752b865bf540e3df57","observation_id":"08d8ff0c-5ae4-41e2-b896-c58ddbfc6f99","resolution":{"observed_at":"2026-07-11T13:53:36.775836Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2305.14387","last_updated":"2024-01-08T04:46:56Z","snapshot_observed_at":"2026-07-06T15:31:51.249701Z","submitted_at":"2023-05-22T17:55:50Z","title":"AlpacaFarm: A Simulation Framework for Methods that Learn from Human Feedback","version":4},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2305.14387","snapshot_observed_at":"2026-08-02T08:41:03.696680Z","title":"arXiv preprint arXiv:2305.14387 , year=","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2607.04763","last_updated":"2026-07-26T14:17:18Z","snapshot_observed_at":"2026-08-02T10:24:43.977557Z","submitted_at":"2026-07-06T07:56:53Z","title":"Multi-Turn On-Policy Distillation with Prefix Replay","version":3},"reference_index":269,"source":"arxiv_source","source_observed_at":"2026-08-02T08:41:03.696680Z"},"links":{"cited_paper":"/paper/2305.14387","citing_paper":"/paper/2607.04763"},"observation_digest":"sha256:366aba32dd41cbc195e17fcbe559c6babf153af45ef96b9a5e8bd44ecb76b971","observation_id":"7cc20e6e-6e09-4de3-ace2-c50a812b577f","resolution":{"observed_at":"2026-08-02T08:41:03.696680Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}}],"links":{"evidence":"/evidence","html":"/paper/2305.14387/citation-record","integrity":"/paper/2305.14387/integrity","json":"/paper/2305.14387/citation-record.json","paper":"/paper/2305.14387"},"outbound":[],"paper":{"arxiv_id":"2305.14387","last_updated":"2024-01-08T04:46:56Z","latest_version":4,"primary_category":"cs.LG","snapshot_observed_at":"2026-07-06T15:31:51.249701Z","submitted_at":"2023-05-22T17:55:50Z","title":"AlpacaFarm: A Simulation Framework for Methods that Learn from Human Feedback"},"reference_resolution":{"displayed":0,"state_counts":{"malformed_identifier":0,"metadata_mismatch":0,"parse_uncertain":0,"unresolved":0,"verified_exact":0,"verified_fuzzy":0},"total_outbound_references":0},"refusal":"A citation records a reference. It does not transfer a finding from one paper to another.","schema":"pith.paper-citation-record.v1","standing_sources":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"thesis":"As of 8 August 2026, this Paper Citation Record lists 0 of 0 outbound references and 32 inbound Pith citation observations for arXiv:2305.14387."}