{"as_of":"2026-08-12T04:41:00Z","caps":{"database_statements":6,"inbound":100,"outbound":100},"context_digest":"sha256:51e3d65ba3e280f59f50db7a88ceba2b2037f12ebc0ce4af6410d5bc635e81ca","coverage":[{"denominator":0,"lane":"reference_resolution","note":"Typed states for the displayed outbound observations.","records_observed":0,"source":"paper_references, paper_reference_links","source_observed_at":null,"state":"measured"},{"denominator":10,"lane":"standing_notices","note":"One-hop event checks from named stored sources.","records_observed":10,"source":"scholarly_work_events, retraction_status_cache","source_observed_at":"2026-08-11T06:34:44.6726+00:00","state":"measured"},{"denominator":10,"lane":"inbound_itemization","note":"Pith citing papers itemized under the disclosed page cap.","records_observed":10,"source":"paper_references, paper_reference_links","source_observed_at":"2026-08-07T14:16:40.186644Z","state":"measured"},{"denominator":1,"lane":"external_citation_measurements","note":"A source-named dated measurement, never combined with another source.","records_observed":0,"source":"arxiv_reference","source_observed_at":"2026-07-02T20:57:23.620580Z","state":"measured"}],"external_citation_measurements":[],"inbound":[{"citation":{"cited_paper":{"arxiv_id":"2405.16455","last_updated":"2025-08-25T01:01:38Z","snapshot_observed_at":"2026-08-09T15:50:26.906021Z","submitted_at":"2024-05-26T07:00:05Z","title":"On the Algorithmic Bias of Aligning Large Language Models with RLHF: Preference Collapse and Matching Regularization","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2405.16455","snapshot_observed_at":"2026-08-07T14:16:40.186644Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2505.20627","last_updated":"2025-05-27T02:07:35Z","snapshot_observed_at":"2026-08-11T02:01:04.493797Z","submitted_at":"2025-05-27T02:07:35Z","title":"Fundamental Limits of Game-Theoretic LLM Alignment: Smith Consistency and Preference Matching","version":1},"reference_index":36,"source":"arxiv_source","source_observed_at":"2026-08-07T14:16:40.186644Z"},"links":{"cited_paper":"/paper/2405.16455","citing_paper":"/paper/2505.20627"},"observation_digest":"sha256:9a4209f60454db4c0d916c45c425a80cb6e192cb62a6952e741092b0c6fca226","observation_id":"3bf6d4d8-7973-42df-833b-da06b03a8c7a","resolution":{"observed_at":"2026-08-07T14:16:40.186644Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2405.16455","last_updated":"2025-08-25T01:01:38Z","snapshot_observed_at":"2026-08-09T15:50:26.906021Z","submitted_at":"2024-05-26T07:00:05Z","title":"On the Algorithmic Bias of Aligning Large Language Models with RLHF: Preference Collapse and Matching Regularization","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2405.16455","snapshot_observed_at":"2026-08-07T13:45:03.350482Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2505.21395","last_updated":"2025-05-27T16:23:24Z","snapshot_observed_at":"2026-08-07T13:26:21.678395Z","submitted_at":"2025-05-27T16:23:24Z","title":"Square$\\chi$PO: Differentially Private and Robust $\\chi^2$-Preference Optimization in Offline Direct Alignment","version":1},"reference_index":86,"source":"arxiv_source","source_observed_at":"2026-08-07T13:45:03.350482Z"},"links":{"cited_paper":"/paper/2405.16455","citing_paper":"/paper/2505.21395"},"observation_digest":"sha256:0798e5e0d01453eed30e8480c2775830b9e32678e1e467b7e176310dee12d5bd","observation_id":"2d05d8e8-2c3e-4b5a-81ae-0e73ee18c877","resolution":{"observed_at":"2026-08-07T13:45:03.350482Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2405.16455","last_updated":"2025-08-25T01:01:38Z","snapshot_observed_at":"2026-08-09T15:50:26.906021Z","submitted_at":"2024-05-26T07:00:05Z","title":"On the Algorithmic Bias of Aligning Large Language Models with RLHF: Preference Collapse and Matching Regularization","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2405.16455","snapshot_observed_at":"2026-08-07T01:04:32.454213Z","title":"On the algorithmic bias of aligning large language models with rlhf: Preference collapse and matching regularization.arXiv preprint arXiv:2405.16455,","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2506.12350","last_updated":"2025-06-14T05:14:49Z","snapshot_observed_at":"2026-08-08T15:24:03.755091Z","submitted_at":"2025-06-14T05:14:49Z","title":"Theoretical Tensions in RLHF: Reconciling Empirical Success with Inconsistencies in Social Choice Theory","version":1},"reference_index":22,"source":"pdf_text","source_observed_at":"2026-08-07T01:04:32.454213Z"},"links":{"cited_paper":"/paper/2405.16455","citing_paper":"/paper/2506.12350"},"observation_digest":"sha256:cb9beb0d7a47dd60122f3df2ddb827c5186c77bdcbb3affe47d3187f83303a5c","observation_id":"5f3c4032-9f7c-458e-964b-d3b0f333e6ba","resolution":{"observed_at":"2026-08-07T01:04:32.454213Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2405.16455","last_updated":"2025-08-25T01:01:38Z","snapshot_observed_at":"2026-08-09T15:50:26.906021Z","submitted_at":"2024-05-26T07:00:05Z","title":"On the Algorithmic Bias of Aligning Large Language Models with RLHF: Preference Collapse and Matching Regularization","version":2},"cited_work":{"arxiv_id":"2405.16455","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2405.16455","snapshot_observed_at":"2026-07-02T20:57:23.620580Z","title":"arXiv preprint arXiv:2405.16455 , year=","venue":null,"work_id":"c03c9e1b-99aa-4e63-bb35-db3cb52394b5","year":2024},"citing_paper":{"arxiv_id":"2507.04005","last_updated":"2026-04-05T08:35:28Z","snapshot_observed_at":"2026-08-02T12:46:52.977974Z","submitted_at":"2025-07-05T11:17:20Z","title":"Exploring a Gamified Personality Assessment Method through Interaction with LLM Agents Embodying Different Personalities","version":4},"reference_index":140,"source":"pdf_text","source_observed_at":"2026-05-19T06:35:06.890058Z"},"links":{"cited_paper":"/paper/2405.16455","citing_paper":"/paper/2507.04005"},"observation_digest":"sha256:1bb8d4eabd053776e2ac6d30da8800c388d1213024441f7ed7c0d0720292306b","observation_id":"87142bb4-165a-438c-894a-4ef8a4ece542","resolution":{"observed_at":"2026-05-19T06:37:07.620062Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-11T06:34:44.6726+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-11T06:34:44.6726+00:00","source":"crossref"},{"observed_at":"2026-08-11T06:34:36.301508+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2405.16455","last_updated":"2025-08-25T01:01:38Z","snapshot_observed_at":"2026-08-09T15:50:26.906021Z","submitted_at":"2024-05-26T07:00:05Z","title":"On the Algorithmic Bias of Aligning Large Language Models with RLHF: Preference Collapse and Matching Regularization","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2405.16455","snapshot_observed_at":"2026-08-06T14:13:07.104772Z","title":"On the algorithmic bias of aligning large language models with rlhf: Preference collapse and matching regularization","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2507.19672","last_updated":"2025-07-25T20:52:58Z","snapshot_observed_at":"2026-08-07T12:02:14.124506Z","submitted_at":"2025-07-25T20:52:58Z","title":"Alignment and Safety in Large Language Models: Safety Mechanisms, Training Paradigms, and Emerging Challenges","version":1},"reference_index":221,"source":"arxiv_source","source_observed_at":"2026-08-06T14:13:07.104772Z"},"links":{"cited_paper":"/paper/2405.16455","citing_paper":"/paper/2507.19672"},"observation_digest":"sha256:8754c4cb90bf665930c314d380ea0a6b69da2b3a53fc2233e449fc2fb9fe42eb","observation_id":"4c8b87d4-5aad-4da1-a322-545f90e4eb97","resolution":{"observed_at":"2026-08-06T14:13:07.104772Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2405.16455","last_updated":"2025-08-25T01:01:38Z","snapshot_observed_at":"2026-08-09T15:50:26.906021Z","submitted_at":"2024-05-26T07:00:05Z","title":"On the Algorithmic Bias of Aligning Large Language Models with RLHF: Preference Collapse and Matching Regularization","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2405.16455","snapshot_observed_at":"2026-08-05T22:45:20.828281Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2508.06443","last_updated":"2025-08-08T16:36:16Z","snapshot_observed_at":"2026-08-09T05:04:25.875616Z","submitted_at":"2025-08-08T16:36:16Z","title":"The Fair Game: Auditing & Debiasing AI Algorithms Over Time","version":1},"reference_index":130,"source":"arxiv_source","source_observed_at":"2026-08-05T22:45:20.828281Z"},"links":{"cited_paper":"/paper/2405.16455","citing_paper":"/paper/2508.06443"},"observation_digest":"sha256:fd1496955b6ffde10c49e06fd407eda4b0d85f97314ceebae53b88646086eedf","observation_id":"005718e4-27e6-48cc-8ea7-4e06d3fcd0e4","resolution":{"observed_at":"2026-08-05T22:45:20.828281Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2405.16455","last_updated":"2025-08-25T01:01:38Z","snapshot_observed_at":"2026-08-09T15:50:26.906021Z","submitted_at":"2024-05-26T07:00:05Z","title":"On the Algorithmic Bias of Aligning Large Language Models with RLHF: Preference Collapse and Matching Regularization","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2405.16455","snapshot_observed_at":"2026-08-05T15:44:38.622610Z","title":"J.: On the Algorithmic Bias of Aligning Large Language Models with RLHF: Preference Collapse and Matching Regularization","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2508.19567","last_updated":"2025-08-27T04:54:33Z","snapshot_observed_at":"2026-08-09T09:55:48.990961Z","submitted_at":"2025-08-27T04:54:33Z","title":"Counterfactual Reward Model Training for Bias Mitigation in Multimodal Reinforcement Learning","version":1},"reference_index":4,"source":"pdf_text","source_observed_at":"2026-08-05T15:44:38.622610Z"},"links":{"cited_paper":"/paper/2405.16455","citing_paper":"/paper/2508.19567"},"observation_digest":"sha256:b024a6e6c55ecc348ff54a6fccf12ca3b9ae7da302f08a1cd29f8f9e8fd8d341","observation_id":"7c1c29ba-f599-4abb-8e1f-7192e8ac808e","resolution":{"observed_at":"2026-08-05T15:44:38.622610Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2405.16455","last_updated":"2025-08-25T01:01:38Z","snapshot_observed_at":"2026-08-09T15:50:26.906021Z","submitted_at":"2024-05-26T07:00:05Z","title":"On the Algorithmic Bias of Aligning Large Language Models with RLHF: Preference Collapse and Matching Regularization","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2405.16455","snapshot_observed_at":"2026-08-04T16:29:42.395542Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2509.13679","last_updated":"2026-06-23T10:12:17Z","snapshot_observed_at":"2026-08-10T20:50:08.963394Z","submitted_at":"2025-09-17T04:13:27Z","title":"\"GenAI Defaults to Bias!\" Gamify AI Literacy Through Reflections on Prompts","version":2},"reference_index":83,"source":"pdf_text","source_observed_at":"2026-08-04T16:29:42.395542Z"},"links":{"cited_paper":"/paper/2405.16455","citing_paper":"/paper/2509.13679"},"observation_digest":"sha256:842b97e43c2439f03ce640feef9da5262d217187b7a19249a1c0d3413e5a364f","observation_id":"a09d48ca-5a53-4333-928e-17f03e4641c8","resolution":{"observed_at":"2026-08-04T16:29:42.395542Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2405.16455","last_updated":"2025-08-25T01:01:38Z","snapshot_observed_at":"2026-08-09T15:50:26.906021Z","submitted_at":"2024-05-26T07:00:05Z","title":"On the Algorithmic Bias of Aligning Large Language Models with RLHF: Preference Collapse and Matching Regularization","version":2},"cited_work":{"arxiv_id":"2405.16455","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2405.16455","snapshot_observed_at":"2026-07-02T20:57:23.620580Z","title":"arXiv preprint arXiv:2405.16455 , year=","venue":null,"work_id":"c03c9e1b-99aa-4e63-bb35-db3cb52394b5","year":2024},"citing_paper":{"arxiv_id":"2605.07724","last_updated":"2026-06-03T12:55:17Z","snapshot_observed_at":"2026-07-06T23:20:06.574823Z","submitted_at":"2026-05-08T13:27:23Z","title":"Curated Synthetic Data Doesn't Have to Collapse: A Theoretical Study of Generative Retraining with Pluralistic Preferences","version":1},"reference_index":101,"source":"arxiv_source","source_observed_at":"2026-05-11T02:30:14.693348Z"},"links":{"cited_paper":"/paper/2405.16455","citing_paper":"/paper/2605.07724"},"observation_digest":"sha256:fd7ba072300cd5f68b3d6d7d673605e05840df58e125112a0359a6a5033642de","observation_id":"38527641-703e-44af-98b8-5e64922fefbe","resolution":{"observed_at":"2026-05-11T02:30:54.396711Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-11T06:34:44.6726+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-11T06:34:44.6726+00:00","source":"crossref"},{"observed_at":"2026-08-11T06:34:36.301508+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2405.16455","last_updated":"2025-08-25T01:01:38Z","snapshot_observed_at":"2026-08-09T15:50:26.906021Z","submitted_at":"2024-05-26T07:00:05Z","title":"On the Algorithmic Bias of Aligning Large Language Models with RLHF: Preference Collapse and Matching Regularization","version":2},"cited_work":{"arxiv_id":"2405.16455","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2405.16455","snapshot_observed_at":"2026-07-02T20:57:23.620580Z","title":"arXiv preprint arXiv:2405.16455 , year=","venue":null,"work_id":"c03c9e1b-99aa-4e63-bb35-db3cb52394b5","year":2024},"citing_paper":{"arxiv_id":"2606.07988","last_updated":"2026-06-06T05:35:31Z","snapshot_observed_at":"2026-07-06T23:47:33.680573Z","submitted_at":"2026-06-06T05:35:31Z","title":"PAFO: Pareto Fairness Optimization for Personalized Reward Modeling","version":1},"reference_index":12,"source":"pdf_text","source_observed_at":"2026-06-27T20:00:05.900814Z"},"links":{"cited_paper":"/paper/2405.16455","citing_paper":"/paper/2606.07988"},"observation_digest":"sha256:b85b9652b78de8f7b8b75191ef7ebf1e43fde1673cd62b3cf0c801b158accb6c","observation_id":"7b1e7592-25e7-48c9-bbd3-bbaaa9cc933b","resolution":{"observed_at":"2026-07-02T20:57:23.622083Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-11T06:34:44.6726+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-11T06:34:44.6726+00:00","source":"crossref"},{"observed_at":"2026-08-11T06:34:36.301508+00:00","source":"retraction_watch"}],"state":"measured"}}],"links":{"evidence":"/evidence","html":"/paper/2405.16455/citation-record","integrity":"/paper/2405.16455/integrity","json":"/paper/2405.16455/citation-record.json","paper":"/paper/2405.16455"},"outbound":[],"paper":{"arxiv_id":"2405.16455","last_updated":"2025-08-25T01:01:38Z","latest_version":2,"primary_category":"stat.ML","snapshot_observed_at":"2026-08-09T15:50:26.906021Z","submitted_at":"2024-05-26T07:00:05Z","title":"On the Algorithmic Bias of Aligning Large Language Models with RLHF: Preference Collapse and Matching Regularization"},"reference_resolution":{"displayed":0,"state_counts":{"malformed_identifier":0,"metadata_mismatch":0,"parse_uncertain":0,"unresolved":0,"verified_exact":0,"verified_fuzzy":0},"total_outbound_references":0},"refusal":"A citation records a reference. It does not transfer a finding from one paper to another.","schema":"pith.paper-citation-record.v1","standing_sources":[{"observed_at":"2026-08-11T06:34:44.6726+00:00","source":"crossref"},{"observed_at":"2026-08-11T06:34:36.301508+00:00","source":"retraction_watch"}],"thesis":"As of 12 August 2026, this Paper Citation Record lists 0 of 0 outbound references and 10 inbound Pith citation observations for arXiv:2405.16455."}