{"as_of":"2026-08-05T23:50:00Z","caps":{"database_statements":6,"inbound":100,"outbound":100},"context_digest":"sha256:b49dafc6f81a5876e68df2882394225c211f1e31695fc85f6daf9c536d92dd32","coverage":[{"denominator":0,"lane":"reference_resolution","note":"Typed states for the displayed outbound observations.","records_observed":0,"source":"paper_references, paper_reference_links","source_observed_at":null,"state":"measured"},{"denominator":9,"lane":"standing_notices","note":"One-hop event checks from named stored sources.","records_observed":9,"source":"scholarly_work_events, retraction_status_cache","source_observed_at":"2026-08-05T06:32:48.257954+00:00","state":"measured"},{"denominator":9,"lane":"inbound_itemization","note":"Pith citing papers itemized under the disclosed page cap.","records_observed":9,"source":"paper_references, paper_reference_links","source_observed_at":"2026-08-02T08:40:48.986418Z","state":"measured"},{"denominator":1,"lane":"external_citation_measurements","note":"A source-named dated measurement, never combined with another source.","records_observed":0,"source":"arxiv_reference","source_observed_at":"2026-07-04T00:29:16.737553Z","state":"measured"}],"external_citation_measurements":[],"inbound":[{"citation":{"cited_paper":{"arxiv_id":"2312.16682","last_updated":"2024-04-22T22:51:32Z","snapshot_observed_at":"2026-07-06T17:09:02.167931Z","submitted_at":"2023-12-27T18:53:09Z","title":"Some things are more CRINGE than others: Iterative Preference Optimization with the Pairwise Cringe Loss","version":2},"cited_work":{"arxiv_id":"2312.16682","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2312.16682","snapshot_observed_at":"2026-07-04T00:29:16.737553Z","title":"Rui Yang, Ruomeng Ding, Yong Lin, Huan Zhang, and Tong Zhang","venue":null,"work_id":"82c78c41-62c3-4eb5-b38d-0263edfb7f2e","year":2024},"citing_paper":{"arxiv_id":"2401.10020","last_updated":"2025-03-28T00:06:51Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2024-01-18T14:43:47Z","title":"Self-Rewarding Language Models","version":3},"reference_index":121,"source":"arxiv_source","source_observed_at":"2026-05-13T12:01:42.290502Z"},"links":{"cited_paper":"/paper/2312.16682","citing_paper":"/paper/2401.10020"},"observation_digest":"sha256:537fea74fb23a66852f31633b44956ce4b9fff4369645d463f1d23c17c1e76bc","observation_id":"bdd04f55-6a75-4367-b275-0d6bb227348b","resolution":{"observed_at":"2026-05-13T12:01:42.518264Z","resolver_source":"arxiv_id","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2312.16682","last_updated":"2024-04-22T22:51:32Z","snapshot_observed_at":"2026-07-06T17:09:02.167931Z","submitted_at":"2023-12-27T18:53:09Z","title":"Some things are more CRINGE than others: Iterative Preference Optimization with the Pairwise Cringe Loss","version":2},"cited_work":{"arxiv_id":"2312.16682","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2312.16682","snapshot_observed_at":"2026-07-04T00:29:16.737553Z","title":"Rui Yang, Ruomeng Ding, Yong Lin, Huan Zhang, and Tong Zhang","venue":null,"work_id":"82c78c41-62c3-4eb5-b38d-0263edfb7f2e","year":2024},"citing_paper":{"arxiv_id":"2408.15339","last_updated":"2026-05-07T20:32:48Z","snapshot_observed_at":"2026-08-02T15:53:33.114051Z","submitted_at":"2024-08-27T18:04:07Z","title":"UNA: A Unified Supervised Framework for Efficient LLM Alignment Across Feedback Types","version":4},"reference_index":22,"source":"pdf_text","source_observed_at":"2026-05-23T21:22:36.970101Z"},"links":{"cited_paper":"/paper/2312.16682","citing_paper":"/paper/2408.15339"},"observation_digest":"sha256:acee1a2a47b93c878a382ab527f022e0bd985e9546c2f470ba56e20ece74d886","observation_id":"3107b263-5fce-43a2-99a7-30228e15d883","resolution":{"observed_at":"2026-05-23T21:23:27.459223Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2312.16682","last_updated":"2024-04-22T22:51:32Z","snapshot_observed_at":"2026-07-06T17:09:02.167931Z","submitted_at":"2023-12-27T18:53:09Z","title":"Some things are more CRINGE than others: Iterative Preference Optimization with the Pairwise Cringe Loss","version":2},"cited_work":{"arxiv_id":"2312.16682","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2312.16682","snapshot_observed_at":"2026-07-04T00:29:16.737553Z","title":"Rui Yang, Ruomeng Ding, Yong Lin, Huan Zhang, and Tong Zhang","venue":null,"work_id":"82c78c41-62c3-4eb5-b38d-0263edfb7f2e","year":2024},"citing_paper":{"arxiv_id":"2509.20265","last_updated":"2026-04-29T12:32:16Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2025-09-24T15:52:36Z","title":"Failure Modes of Maximum Entropy RLHF","version":3},"reference_index":56,"source":"arxiv_source","source_observed_at":"2026-05-18T14:02:11.084514Z"},"links":{"cited_paper":"/paper/2312.16682","citing_paper":"/paper/2509.20265"},"observation_digest":"sha256:e5dea7616835a9cae7a3978274002baf1d0f8219dae36dc9e883e66193c0bdac","observation_id":"499568bc-e163-4943-a102-469d44a34e5d","resolution":{"observed_at":"2026-05-18T14:02:39.774025Z","resolver_source":"arxiv_id","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2312.16682","last_updated":"2024-04-22T22:51:32Z","snapshot_observed_at":"2026-07-06T17:09:02.167931Z","submitted_at":"2023-12-27T18:53:09Z","title":"Some things are more CRINGE than others: Iterative Preference Optimization with the Pairwise Cringe Loss","version":2},"cited_work":{"arxiv_id":"2312.16682","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2312.16682","snapshot_observed_at":"2026-07-04T00:29:16.737553Z","title":"Rui Yang, Ruomeng Ding, Yong Lin, Huan Zhang, and Tong Zhang","venue":null,"work_id":"82c78c41-62c3-4eb5-b38d-0263edfb7f2e","year":2024},"citing_paper":{"arxiv_id":"2604.17543","last_updated":"2026-04-19T17:06:56Z","snapshot_observed_at":"2026-07-06T23:04:37.370465Z","submitted_at":"2026-04-19T17:06:56Z","title":"PoliLegalLM: A Technical Report on a Large Language Model for Political and Legal Affairs","version":1},"reference_index":40,"source":"arxiv_source","source_observed_at":"2026-05-10T06:03:01.460352Z"},"links":{"cited_paper":"/paper/2312.16682","citing_paper":"/paper/2604.17543"},"observation_digest":"sha256:0225fab021472366fae581e5b832c2dd9b00d95123122cac2a1be538c2e67b61","observation_id":"504c69b8-2128-4b3f-84ce-6f097940e55d","resolution":{"observed_at":"2026-05-10T06:06:19.077085Z","resolver_source":"arxiv_id","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2312.16682","last_updated":"2024-04-22T22:51:32Z","snapshot_observed_at":"2026-07-06T17:09:02.167931Z","submitted_at":"2023-12-27T18:53:09Z","title":"Some things are more CRINGE than others: Iterative Preference Optimization with the Pairwise Cringe Loss","version":2},"cited_work":{"arxiv_id":"2312.16682","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2312.16682","snapshot_observed_at":"2026-07-04T00:29:16.737553Z","title":"Rui Yang, Ruomeng Ding, Yong Lin, Huan Zhang, and Tong Zhang","venue":null,"work_id":"82c78c41-62c3-4eb5-b38d-0263edfb7f2e","year":2024},"citing_paper":{"arxiv_id":"2605.15012","last_updated":"2026-05-14T16:12:30Z","snapshot_observed_at":"2026-07-06T23:26:23.179482Z","submitted_at":"2026-05-14T16:12:30Z","title":"Boosting Reinforcement Learning with Verifiable Rewards via Randomly Selected Few-Shot Guidance","version":1},"reference_index":89,"source":"pdf_text","source_observed_at":"2026-05-15T03:18:26.590871Z"},"links":{"cited_paper":"/paper/2312.16682","citing_paper":"/paper/2605.15012"},"observation_digest":"sha256:4e540872d007996b45c966eb34ca4102a73666430be737346f581d2163b76fe2","observation_id":"748042aa-09d4-4c81-afdf-ab8cd8118bff","resolution":{"observed_at":"2026-05-15T03:19:42.952153Z","resolver_source":"arxiv_id","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2312.16682","last_updated":"2024-04-22T22:51:32Z","snapshot_observed_at":"2026-07-06T17:09:02.167931Z","submitted_at":"2023-12-27T18:53:09Z","title":"Some things are more CRINGE than others: Iterative Preference Optimization with the Pairwise Cringe Loss","version":2},"cited_work":{"arxiv_id":"2312.16682","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2312.16682","snapshot_observed_at":"2026-07-04T00:29:16.737553Z","title":"Rui Yang, Ruomeng Ding, Yong Lin, Huan Zhang, and Tong Zhang","venue":null,"work_id":"82c78c41-62c3-4eb5-b38d-0263edfb7f2e","year":2024},"citing_paper":{"arxiv_id":"2606.21090","last_updated":"2026-06-17T18:03:06Z","snapshot_observed_at":"2026-07-30T06:23:05.041940Z","submitted_at":"2026-06-17T18:03:06Z","title":"Self-Improvement Can Self-Regress: The Rise-and-Collapse Failure Mode of LLM Self-Training","version":1},"reference_index":19,"source":"pdf_text","source_observed_at":"2026-06-26T21:07:28.660119Z"},"links":{"cited_paper":"/paper/2312.16682","citing_paper":"/paper/2606.21090"},"observation_digest":"sha256:739b3563f166fddf68e0f7380718baf3a74f3616f2397ad643275cd2a6b5a0c0","observation_id":"b4ada445-124f-4b33-984a-81144f4444a0","resolution":{"observed_at":"2026-07-04T00:29:16.740668Z","resolver_source":"arxiv_id","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2312.16682","last_updated":"2024-04-22T22:51:32Z","snapshot_observed_at":"2026-07-06T17:09:02.167931Z","submitted_at":"2023-12-27T18:53:09Z","title":"Some things are more CRINGE than others: Iterative Preference Optimization with the Pairwise Cringe Loss","version":2},"cited_work":{"arxiv_id":"2312.16682","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2312.16682","snapshot_observed_at":"2026-07-04T00:29:16.737553Z","title":"Rui Yang, Ruomeng Ding, Yong Lin, Huan Zhang, and Tong Zhang","venue":null,"work_id":"82c78c41-62c3-4eb5-b38d-0263edfb7f2e","year":2024},"citing_paper":{"arxiv_id":"2607.01470","last_updated":"2026-07-01T21:02:54Z","snapshot_observed_at":"2026-08-03T01:52:00.202969Z","submitted_at":"2026-07-01T21:02:54Z","title":"World Feedback for Clinical Agents: Diagnosing RL in FHIR Environments","version":1},"reference_index":12,"source":"arxiv_source","source_observed_at":"2026-07-03T20:14:46.147637Z"},"links":{"cited_paper":"/paper/2312.16682","citing_paper":"/paper/2607.01470"},"observation_digest":"sha256:e2e8dca13ab8cacfff239ae94c30aa3e7e9e59a8ff3937a0f0386865b8065ad1","observation_id":"c0cc8293-6a3a-4fef-8c49-44de339ed846","resolution":{"observed_at":"2026-07-03T20:18:55.876079Z","resolver_source":"arxiv_id","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2312.16682","last_updated":"2024-04-22T22:51:32Z","snapshot_observed_at":"2026-07-06T17:09:02.167931Z","submitted_at":"2023-12-27T18:53:09Z","title":"Some things are more CRINGE than others: Iterative Preference Optimization with the Pairwise Cringe Loss","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2312.16682","snapshot_observed_at":"2026-07-11T13:53:36.775836Z","title":"arXiv preprint arXiv:2312.16682 , year=","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2607.04763","last_updated":"2026-07-26T14:17:18Z","snapshot_observed_at":"2026-08-02T10:24:43.977557Z","submitted_at":"2026-07-06T07:56:53Z","title":"Multi-Turn On-Policy Distillation with Prefix Replay","version":1},"reference_index":150,"source":"arxiv_source","source_observed_at":"2026-07-11T13:53:36.775836Z"},"links":{"cited_paper":"/paper/2312.16682","citing_paper":"/paper/2607.04763"},"observation_digest":"sha256:9855d28babb56dc39cae2f562ec6c4ef8a230b3483da59047848e7d26d2b7da6","observation_id":"474341e2-d357-4988-9686-54c90fb4d7c6","resolution":{"observed_at":"2026-07-11T13:53:36.775836Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2312.16682","last_updated":"2024-04-22T22:51:32Z","snapshot_observed_at":"2026-07-06T17:09:02.167931Z","submitted_at":"2023-12-27T18:53:09Z","title":"Some things are more CRINGE than others: Iterative Preference Optimization with the Pairwise Cringe Loss","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2312.16682","snapshot_observed_at":"2026-08-02T08:40:48.986418Z","title":"arXiv preprint arXiv:2312.16682 , year=","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2607.04763","last_updated":"2026-07-26T14:17:18Z","snapshot_observed_at":"2026-08-02T10:24:43.977557Z","submitted_at":"2026-07-06T07:56:53Z","title":"Multi-Turn On-Policy Distillation with Prefix Replay","version":3},"reference_index":151,"source":"arxiv_source","source_observed_at":"2026-08-02T08:40:48.986418Z"},"links":{"cited_paper":"/paper/2312.16682","citing_paper":"/paper/2607.04763"},"observation_digest":"sha256:3db051691f37c78fcf3b4a4dc24832b10cae0690944980f7fb06d151fcc000d1","observation_id":"328acb1c-6e3a-4aee-b3d9-ce4f3de8e73f","resolution":{"observed_at":"2026-08-02T08:40:48.986418Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}}],"links":{"evidence":"/evidence","html":"/paper/2312.16682/citation-record","integrity":"/paper/2312.16682/integrity","json":"/paper/2312.16682/citation-record.json","paper":"/paper/2312.16682"},"outbound":[],"paper":{"arxiv_id":"2312.16682","last_updated":"2024-04-22T22:51:32Z","latest_version":2,"primary_category":"cs.CL","snapshot_observed_at":"2026-07-06T17:09:02.167931Z","submitted_at":"2023-12-27T18:53:09Z","title":"Some things are more CRINGE than others: Iterative Preference Optimization with the Pairwise Cringe Loss"},"reference_resolution":{"displayed":0,"state_counts":{"malformed_identifier":0,"metadata_mismatch":0,"parse_uncertain":0,"unresolved":0,"verified_exact":0,"verified_fuzzy":0},"total_outbound_references":0},"refusal":"A citation records a reference. It does not transfer a finding from one paper to another.","schema":"pith.paper-citation-record.v1","standing_sources":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"thesis":"As of 5 August 2026, this Paper Citation Record lists 0 of 0 outbound references and 9 inbound Pith citation observations for arXiv:2312.16682."}