{"as_of":"2026-08-08T15:24:00Z","caps":{"database_statements":6,"inbound":100,"outbound":100},"context_digest":"sha256:a600d005d9f6c334bc68687424bb12ee2f4d2cb5030b95b9e24d1ef821350f5d","coverage":[{"denominator":0,"lane":"reference_resolution","note":"Typed states for the displayed outbound observations.","records_observed":0,"source":"paper_references, paper_reference_links","source_observed_at":null,"state":"measured"},{"denominator":22,"lane":"standing_notices","note":"One-hop event checks from named stored sources.","records_observed":22,"source":"scholarly_work_events, retraction_status_cache","source_observed_at":"2026-08-08T06:32:00.761636+00:00","state":"measured"},{"denominator":22,"lane":"inbound_itemization","note":"Pith citing papers itemized under the disclosed page cap.","records_observed":22,"source":"paper_references, paper_reference_links","source_observed_at":"2026-08-08T05:42:43.505543Z","state":"measured"},{"denominator":1,"lane":"external_citation_measurements","note":"A source-named dated measurement, never combined with another source.","records_observed":1,"source":"arxiv_reference","source_observed_at":"2026-08-05T02:28:24.338817Z","state":"measured"}],"external_citation_measurements":[{"count":9,"observed_at":"2026-08-05T02:28:24.338817Z","source":"arxiv_reference"}],"inbound":[{"citation":{"cited_paper":{"arxiv_id":"2309.06256","last_updated":"2024-10-13T19:27:20Z","snapshot_observed_at":"2026-07-06T16:17:26.924098Z","submitted_at":"2023-09-12T14:16:54Z","title":"Mitigating the Alignment Tax of RLHF","version":4},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2309.06256","snapshot_observed_at":"2026-08-08T05:42:43.505543Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2502.08301","last_updated":"2025-06-23T09:04:32Z","snapshot_observed_at":"2026-08-08T05:36:10.034729Z","submitted_at":"2025-02-12T11:02:59Z","title":"Compromising Honesty and Harmlessness in Language Models via Deception Attacks","version":2},"reference_index":31,"source":"pdf_text","source_observed_at":"2026-08-08T05:42:43.505543Z"},"links":{"cited_paper":"/paper/2309.06256","citing_paper":"/paper/2502.08301"},"observation_digest":"sha256:87f525d086839a931579987988feb79e440e50601fe567257be821ffe2585eff","observation_id":"d1c608ac-143c-4e23-8001-10d4687a0371","resolution":{"observed_at":"2026-08-08T05:42:43.505543Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2309.06256","last_updated":"2024-10-13T19:27:20Z","snapshot_observed_at":"2026-07-06T16:17:26.924098Z","submitted_at":"2023-09-12T14:16:54Z","title":"Mitigating the Alignment Tax of RLHF","version":4},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2309.06256","snapshot_observed_at":"2026-08-07T14:56:03.062616Z","title":"Mitigating the alignment tax of rlhf, 2024 b","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2505.17231","last_updated":"2025-05-22T19:13:34Z","snapshot_observed_at":"2026-08-07T14:48:15.518394Z","submitted_at":"2025-05-22T19:13:34Z","title":"ExeSQL: Self-Taught Text-to-SQL Models with Execution-Driven Bootstrapping for SQL Dialects","version":1},"reference_index":69,"source":"arxiv_source","source_observed_at":"2026-08-07T14:56:03.062616Z"},"links":{"cited_paper":"/paper/2309.06256","citing_paper":"/paper/2505.17231"},"observation_digest":"sha256:a9d97bad13568b0f80b79ca53e678a679b120cec6651b2f7a09007d80a0ff34b","observation_id":"02b76a65-7a54-43de-9bbc-5c7ad1d15192","resolution":{"observed_at":"2026-08-07T14:56:03.062616Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2309.06256","last_updated":"2024-10-13T19:27:20Z","snapshot_observed_at":"2026-07-06T16:17:26.924098Z","submitted_at":"2023-09-12T14:16:54Z","title":"Mitigating the Alignment Tax of RLHF","version":4},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2309.06256","snapshot_observed_at":"2026-08-07T11:40:09.918152Z","title":"Mitigating the alignment tax of RLHF","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2506.01901","last_updated":"2025-06-02T17:23:16Z","snapshot_observed_at":"2026-08-07T11:29:33.060031Z","submitted_at":"2025-06-02T17:23:16Z","title":"Understanding Overadaptation in Supervised Fine-Tuning: The Role of Ensemble Methods","version":1},"reference_index":27,"source":"arxiv_source","source_observed_at":"2026-08-07T11:40:09.918152Z"},"links":{"cited_paper":"/paper/2309.06256","citing_paper":"/paper/2506.01901"},"observation_digest":"sha256:0c612252dfcfe906182f4096d34336e2d611e3ada9280af3ecbe273e2e141cb8","observation_id":"68814b27-f0a7-4b71-bc79-75bc0a926d01","resolution":{"observed_at":"2026-08-07T11:40:09.918152Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2309.06256","last_updated":"2024-10-13T19:27:20Z","snapshot_observed_at":"2026-07-06T16:17:26.924098Z","submitted_at":"2023-09-12T14:16:54Z","title":"Mitigating the Alignment Tax of RLHF","version":4},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2309.06256","snapshot_observed_at":"2026-08-06T18:51:31.112415Z","title":"Mitigating the alignment tax of rlhf","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2507.07375","last_updated":"2025-07-10T01:56:56Z","snapshot_observed_at":"2026-08-08T11:06:36.428118Z","submitted_at":"2025-07-10T01:56:56Z","title":"Bradley-Terry and Multi-Objective Reward Modeling Are Complementary","version":1},"reference_index":60,"source":"pdf_text","source_observed_at":"2026-08-06T18:51:31.112415Z"},"links":{"cited_paper":"/paper/2309.06256","citing_paper":"/paper/2507.07375"},"observation_digest":"sha256:a78f4b260d6399a358b3995e5226bd687cc9ca13f7f85ee7d8c56af8742c551d","observation_id":"1b1f0d2a-e239-4670-9539-70369bca1236","resolution":{"observed_at":"2026-08-06T18:51:31.112415Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2309.06256","last_updated":"2024-10-13T19:27:20Z","snapshot_observed_at":"2026-07-06T16:17:26.924098Z","submitted_at":"2023-09-12T14:16:54Z","title":"Mitigating the Alignment Tax of RLHF","version":4},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2309.06256","snapshot_observed_at":"2026-08-06T18:24:50.745763Z","title":"arXiv preprint arXiv:2309.06256 (2023)","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2507.08357","last_updated":"2025-07-11T07:18:01Z","snapshot_observed_at":"2026-08-06T18:18:06.344554Z","submitted_at":"2025-07-11T07:18:01Z","title":"Cycle Context Verification for In-Context Medical Image Segmentation","version":1},"reference_index":11,"source":"pdf_text","source_observed_at":"2026-08-06T18:24:50.745763Z"},"links":{"cited_paper":"/paper/2309.06256","citing_paper":"/paper/2507.08357"},"observation_digest":"sha256:c6d225888e2929e7d2504d5f1ce59dcbcfc6dd8ae47fbc4572c60d9b75061654","observation_id":"cf35a2c0-128f-4efe-90fa-6dd0c16c0c5a","resolution":{"observed_at":"2026-08-06T18:24:50.745763Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2309.06256","last_updated":"2024-10-13T19:27:20Z","snapshot_observed_at":"2026-07-06T16:17:26.924098Z","submitted_at":"2023-09-12T14:16:54Z","title":"Mitigating the Alignment Tax of RLHF","version":4},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2309.06256","snapshot_observed_at":"2026-08-06T05:29:14.879078Z","title":"Mitigating the alignment tax of rlhf","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2508.01781","last_updated":"2025-08-03T14:37:16Z","snapshot_observed_at":"2026-08-07T21:40:51.636920Z","submitted_at":"2025-08-03T14:37:16Z","title":"A comprehensive taxonomy of hallucinations in Large Language Models","version":1},"reference_index":55,"source":"pdf_text","source_observed_at":"2026-08-06T05:29:14.879078Z"},"links":{"cited_paper":"/paper/2309.06256","citing_paper":"/paper/2508.01781"},"observation_digest":"sha256:31013d97fdc06ab7afe0b74e0f5621879e8512959e0ae9010d5c9a451aeb58a1","observation_id":"46881b43-642c-4772-b5c1-270725d1f111","resolution":{"observed_at":"2026-08-06T05:29:14.879078Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2309.06256","last_updated":"2024-10-13T19:27:20Z","snapshot_observed_at":"2026-07-06T16:17:26.924098Z","submitted_at":"2023-09-12T14:16:54Z","title":"Mitigating the Alignment Tax of RLHF","version":4},"cited_work":{"arxiv_id":"2309.06256","doi":"10.48550/arxiv.2309.06256","metadata_source":"arxiv_reference","pith_arxiv_id":"2309.06256","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"13 CapTrack: Multifaceted Evaluation of Forgetting in LLM Post-Training Liu, D","venue":"arXiv (Cornell University)","work_id":"f1108279-8a48-4a03-82fb-254472b97842","year":2024},"citing_paper":{"arxiv_id":"2509.03403","last_updated":"2026-05-15T21:29:46Z","snapshot_observed_at":"2026-07-06T22:23:00.681940Z","submitted_at":"2025-09-03T15:28:51Z","title":"Beyond Correctness: Harmonizing Process and Outcome Rewards through RL Training","version":2},"reference_index":18,"source":"arxiv_source","source_observed_at":"2026-05-21T22:38:57.833414Z"},"links":{"cited_paper":"/paper/2309.06256","citing_paper":"/paper/2509.03403"},"observation_digest":"sha256:f9b6f91cdbc1a34e9f0af42b7b7adbb5c020e963b015477221c075d0cbb1e56c","observation_id":"652a0f17-da08-49ed-a08c-15566b8bdf03","resolution":{"observed_at":"2026-05-21T22:40:43.282501Z","resolver_source":"arxiv_id","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2309.06256","last_updated":"2024-10-13T19:27:20Z","snapshot_observed_at":"2026-07-06T16:17:26.924098Z","submitted_at":"2023-09-12T14:16:54Z","title":"Mitigating the Alignment Tax of RLHF","version":4},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2309.06256","snapshot_observed_at":"2026-08-04T21:01:50.125128Z","title":null,"venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2509.08255","last_updated":"2025-09-10T03:20:56Z","snapshot_observed_at":"2026-08-07T23:55:40.448001Z","submitted_at":"2025-09-10T03:20:56Z","title":"Mitigating Catastrophic Forgetting in Large Language Models with Forgetting-aware Pruning","version":1},"reference_index":19,"source":"arxiv_source","source_observed_at":"2026-08-04T21:01:50.125128Z"},"links":{"cited_paper":"/paper/2309.06256","citing_paper":"/paper/2509.08255"},"observation_digest":"sha256:8a18f0df7eee7e5954c841990d99a318fce687660783cf3e915bbd8a2cc58e02","observation_id":"1f34472f-2b41-4fc4-af15-a05165cda945","resolution":{"observed_at":"2026-08-04T21:01:50.125128Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2309.06256","last_updated":"2024-10-13T19:27:20Z","snapshot_observed_at":"2026-07-06T16:17:26.924098Z","submitted_at":"2023-09-12T14:16:54Z","title":"Mitigating the Alignment Tax of RLHF","version":4},"cited_work":{"arxiv_id":"2309.06256","doi":"10.48550/arxiv.2309.06256","metadata_source":"arxiv_reference","pith_arxiv_id":"2309.06256","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"13 CapTrack: Multifaceted Evaluation of Forgetting in LLM Post-Training Liu, D","venue":"arXiv (Cornell University)","work_id":"f1108279-8a48-4a03-82fb-254472b97842","year":2024},"citing_paper":{"arxiv_id":"2603.06610","last_updated":"2026-05-22T08:27:37Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2026-02-19T09:46:24Z","title":"CapTrack: Multifaceted Evaluation of Forgetting in LLM Post-Training","version":2},"reference_index":28,"source":"pdf_text","source_observed_at":"2026-05-25T06:40:51.046965Z"},"links":{"cited_paper":"/paper/2309.06256","citing_paper":"/paper/2603.06610"},"observation_digest":"sha256:e221e96388aba1d81d6e87a1e5432d3d8aeef954475cfe43bc9074d9e03e62e6","observation_id":"9c6b2640-fb55-435b-b9a5-497f3548e30e","resolution":{"observed_at":"2026-05-25T06:45:26.492450Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2309.06256","last_updated":"2024-10-13T19:27:20Z","snapshot_observed_at":"2026-07-06T16:17:26.924098Z","submitted_at":"2023-09-12T14:16:54Z","title":"Mitigating the Alignment Tax of RLHF","version":4},"cited_work":{"arxiv_id":"2309.06256","doi":"10.48550/arxiv.2309.06256","metadata_source":"arxiv_reference","pith_arxiv_id":"2309.06256","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"13 CapTrack: Multifaceted Evaluation of Forgetting in LLM Post-Training Liu, D","venue":"arXiv (Cornell University)","work_id":"f1108279-8a48-4a03-82fb-254472b97842","year":2024},"citing_paper":{"arxiv_id":"2604.17497","last_updated":"2026-04-19T15:32:48Z","snapshot_observed_at":"2026-08-04T20:11:37.254234Z","submitted_at":"2026-04-19T15:32:48Z","title":"Generative AI Technologies, Techniques & Tensions: A Primer","version":1},"reference_index":11,"source":"pdf_text","source_observed_at":"2026-05-10T05:39:12.226558Z"},"links":{"cited_paper":"/paper/2309.06256","citing_paper":"/paper/2604.17497"},"observation_digest":"sha256:1fc839c0bfab97cfcbd468935f0a438d4a9c76ddcab793f42b885f7cd2619d78","observation_id":"c9986d8d-5bad-41a8-8adf-dfdb044c48a7","resolution":{"observed_at":"2026-05-10T05:41:01.385099Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2309.06256","last_updated":"2024-10-13T19:27:20Z","snapshot_observed_at":"2026-07-06T16:17:26.924098Z","submitted_at":"2023-09-12T14:16:54Z","title":"Mitigating the Alignment Tax of RLHF","version":4},"cited_work":{"arxiv_id":"2309.06256","doi":"10.48550/arxiv.2309.06256","metadata_source":"arxiv_reference","pith_arxiv_id":"2309.06256","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"13 CapTrack: Multifaceted Evaluation of Forgetting in LLM Post-Training Liu, D","venue":"arXiv (Cornell University)","work_id":"f1108279-8a48-4a03-82fb-254472b97842","year":2024},"citing_paper":{"arxiv_id":"2604.19087","last_updated":"2026-04-21T04:59:37Z","snapshot_observed_at":"2026-07-06T23:05:48.885141Z","submitted_at":"2026-04-21T04:59:37Z","title":"OLLM: Options-based Large Language Models","version":1},"reference_index":10,"source":"pdf_text","source_observed_at":"2026-05-10T02:04:12.174834Z"},"links":{"cited_paper":"/paper/2309.06256","citing_paper":"/paper/2604.19087"},"observation_digest":"sha256:08506fe6b129b10670226332df777651a6cba2a48650545df26e338961fcd727","observation_id":"787e0af1-5a2b-480c-a4b2-6d089dd5ee3c","resolution":{"observed_at":"2026-05-11T13:16:08.545128Z","resolver_source":"arxiv_id","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2309.06256","last_updated":"2024-10-13T19:27:20Z","snapshot_observed_at":"2026-07-06T16:17:26.924098Z","submitted_at":"2023-09-12T14:16:54Z","title":"Mitigating the Alignment Tax of RLHF","version":4},"cited_work":{"arxiv_id":"2309.06256","doi":"10.48550/arxiv.2309.06256","metadata_source":"arxiv_reference","pith_arxiv_id":"2309.06256","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"13 CapTrack: Multifaceted Evaluation of Forgetting in LLM Post-Training Liu, D","venue":"arXiv (Cornell University)","work_id":"f1108279-8a48-4a03-82fb-254472b97842","year":2024},"citing_paper":{"arxiv_id":"2605.11679","last_updated":"2026-05-13T09:28:34Z","snapshot_observed_at":"2026-07-06T23:23:30.499023Z","submitted_at":"2026-05-12T07:38:59Z","title":"Explaining and Breaking the Safety-Helpfulness Ceiling via Preference Dimensional Expansion","version":1},"reference_index":28,"source":"pdf_text","source_observed_at":"2026-05-13T01:03:10.263663Z"},"links":{"cited_paper":"/paper/2309.06256","citing_paper":"/paper/2605.11679"},"observation_digest":"sha256:46a048ebb167c675b6e84573117d40ff8710f4e47d48145514aaa7afc6d7b211","observation_id":"24c97580-deb9-4128-80cc-55e119410d26","resolution":{"observed_at":"2026-05-13T01:57:06.339401Z","resolver_source":"arxiv_id","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2309.06256","last_updated":"2024-10-13T19:27:20Z","snapshot_observed_at":"2026-07-06T16:17:26.924098Z","submitted_at":"2023-09-12T14:16:54Z","title":"Mitigating the Alignment Tax of RLHF","version":4},"cited_work":{"arxiv_id":"2309.06256","doi":"10.48550/arxiv.2309.06256","metadata_source":"arxiv_reference","pith_arxiv_id":"2309.06256","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"13 CapTrack: Multifaceted Evaluation of Forgetting in LLM Post-Training Liu, D","venue":"arXiv (Cornell University)","work_id":"f1108279-8a48-4a03-82fb-254472b97842","year":2024},"citing_paper":{"arxiv_id":"2605.11679","last_updated":"2026-05-13T09:28:34Z","snapshot_observed_at":"2026-07-06T23:23:30.499023Z","submitted_at":"2026-05-12T07:38:59Z","title":"Explaining and Breaking the Safety-Helpfulness Ceiling via Preference Dimensional Expansion","version":2},"reference_index":28,"source":"pdf_text","source_observed_at":"2026-05-14T21:12:06.989077Z"},"links":{"cited_paper":"/paper/2309.06256","citing_paper":"/paper/2605.11679"},"observation_digest":"sha256:9ad016b7df190a2e09be7f2b35be7ca648fe188e5012dcf0a77f67eb15f8b08c","observation_id":"f868b637-308a-4d11-9a44-e1df3252cccb","resolution":{"observed_at":"2026-05-14T21:12:58.908959Z","resolver_source":"arxiv_id","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2309.06256","last_updated":"2024-10-13T19:27:20Z","snapshot_observed_at":"2026-07-06T16:17:26.924098Z","submitted_at":"2023-09-12T14:16:54Z","title":"Mitigating the Alignment Tax of RLHF","version":4},"cited_work":{"arxiv_id":"2309.06256","doi":"10.48550/arxiv.2309.06256","metadata_source":"arxiv_reference","pith_arxiv_id":"2309.06256","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"13 CapTrack: Multifaceted Evaluation of Forgetting in LLM Post-Training Liu, D","venue":"arXiv (Cornell University)","work_id":"f1108279-8a48-4a03-82fb-254472b97842","year":2024},"citing_paper":{"arxiv_id":"2605.12484","last_updated":"2026-05-14T17:49:32Z","snapshot_observed_at":"2026-08-02T06:23:20.104772Z","submitted_at":"2026-05-12T17:58:20Z","title":"Learning, Fast and Slow: Towards LLMs That Adapt Continually","version":1},"reference_index":33,"source":"pdf_text","source_observed_at":"2026-05-13T05:00:31.452781Z"},"links":{"cited_paper":"/paper/2309.06256","citing_paper":"/paper/2605.12484"},"observation_digest":"sha256:e850ee75fb41b955e4cd32093e75ef3e7cc294682a20714b4c59977715b1c9c6","observation_id":"8718851e-af55-470f-9c06-cccf72d2c979","resolution":{"observed_at":"2026-05-13T05:07:18.462997Z","resolver_source":"arxiv_id","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2309.06256","last_updated":"2024-10-13T19:27:20Z","snapshot_observed_at":"2026-07-06T16:17:26.924098Z","submitted_at":"2023-09-12T14:16:54Z","title":"Mitigating the Alignment Tax of RLHF","version":4},"cited_work":{"arxiv_id":"2309.06256","doi":"10.48550/arxiv.2309.06256","metadata_source":"arxiv_reference","pith_arxiv_id":"2309.06256","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"13 CapTrack: Multifaceted Evaluation of Forgetting in LLM Post-Training Liu, D","venue":"arXiv (Cornell University)","work_id":"f1108279-8a48-4a03-82fb-254472b97842","year":2024},"citing_paper":{"arxiv_id":"2605.12484","last_updated":"2026-05-14T17:49:32Z","snapshot_observed_at":"2026-08-02T06:23:20.104772Z","submitted_at":"2026-05-12T17:58:20Z","title":"Learning, Fast and Slow: Towards LLMs That Adapt Continually","version":2},"reference_index":33,"source":"pdf_text","source_observed_at":"2026-05-15T05:19:05.368681Z"},"links":{"cited_paper":"/paper/2309.06256","citing_paper":"/paper/2605.12484"},"observation_digest":"sha256:166fee30d5d1a6aecc1cc836ad09c8b570b6655b94e4be4224707dd120e7429e","observation_id":"4ac69b6d-a040-4ae5-bfad-47abb874b41b","resolution":{"observed_at":"2026-05-15T05:19:45.701852Z","resolver_source":"arxiv_id","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2309.06256","last_updated":"2024-10-13T19:27:20Z","snapshot_observed_at":"2026-07-06T16:17:26.924098Z","submitted_at":"2023-09-12T14:16:54Z","title":"Mitigating the Alignment Tax of RLHF","version":4},"cited_work":{"arxiv_id":"2309.06256","doi":"10.48550/arxiv.2309.06256","metadata_source":"arxiv_reference","pith_arxiv_id":"2309.06256","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"13 CapTrack: Multifaceted Evaluation of Forgetting in LLM Post-Training Liu, D","venue":"arXiv (Cornell University)","work_id":"f1108279-8a48-4a03-82fb-254472b97842","year":2024},"citing_paper":{"arxiv_id":"2606.26102","last_updated":"2026-07-12T17:21:56Z","snapshot_observed_at":"2026-07-16T23:18:45.986624Z","submitted_at":"2026-04-30T17:55:22Z","title":"Helpfulness Hurts: Domain-Dependent Degradation of Mid-Trained Compassion Values Under Post-Training","version":2},"reference_index":13,"source":"pdf_text","source_observed_at":"2026-07-01T08:21:41.008505Z"},"links":{"cited_paper":"/paper/2309.06256","citing_paper":"/paper/2606.26102"},"observation_digest":"sha256:e3ea1abad35263266ceb2a59b37b28d12d8105be9f7021d1fc3e81e9f58d0004","observation_id":"c8d86648-e791-47ff-a418-5c284e9d5afa","resolution":{"observed_at":"2026-07-01T08:25:33.131993Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2309.06256","last_updated":"2024-10-13T19:27:20Z","snapshot_observed_at":"2026-07-06T16:17:26.924098Z","submitted_at":"2023-09-12T14:16:54Z","title":"Mitigating the Alignment Tax of RLHF","version":4},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2309.06256","snapshot_observed_at":"2026-07-14T19:20:54.974570Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2606.26102","last_updated":"2026-07-12T17:21:56Z","snapshot_observed_at":"2026-07-16T23:18:45.986624Z","submitted_at":"2026-04-30T17:55:22Z","title":"Helpfulness Hurts: Domain-Dependent Degradation of Mid-Trained Compassion Values Under Post-Training","version":3},"reference_index":13,"source":"pdf_text","source_observed_at":"2026-07-14T19:20:54.974570Z"},"links":{"cited_paper":"/paper/2309.06256","citing_paper":"/paper/2606.26102"},"observation_digest":"sha256:454b6b7635ac28db8b0b7714ac1f25c0f2c2a924c1fd18767bfa79e6dd7f2afc","observation_id":"8a5c7bdf-874e-4fb0-929b-75d6f025ea84","resolution":{"observed_at":"2026-07-14T19:20:54.974570Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2309.06256","last_updated":"2024-10-13T19:27:20Z","snapshot_observed_at":"2026-07-06T16:17:26.924098Z","submitted_at":"2023-09-12T14:16:54Z","title":"Mitigating the Alignment Tax of RLHF","version":4},"cited_work":{"arxiv_id":"2309.06256","doi":"10.48550/arxiv.2309.06256","metadata_source":"arxiv_reference","pith_arxiv_id":"2309.06256","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"13 CapTrack: Multifaceted Evaluation of Forgetting in LLM Post-Training Liu, D","venue":"arXiv (Cornell University)","work_id":"f1108279-8a48-4a03-82fb-254472b97842","year":2024},"citing_paper":{"arxiv_id":"2606.29706","last_updated":"2026-06-29T02:18:51Z","snapshot_observed_at":"2026-08-08T15:19:57.076234Z","submitted_at":"2026-06-29T02:18:51Z","title":"ARMOR: Adaptive Retriever Optimization for Low-Resource Telecom Question Answering","version":1},"reference_index":23,"source":"pdf_text","source_observed_at":"2026-06-30T04:53:59.203935Z"},"links":{"cited_paper":"/paper/2309.06256","citing_paper":"/paper/2606.29706"},"observation_digest":"sha256:4b9c7b1dceeb9d01b5e1514d9efbd4e3975226671214010e741830c1fc956636","observation_id":"140ebab7-adb2-4c5f-857d-4c936101c097","resolution":{"observed_at":"2026-06-30T04:54:16.249019Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2309.06256","last_updated":"2024-10-13T19:27:20Z","snapshot_observed_at":"2026-07-06T16:17:26.924098Z","submitted_at":"2023-09-12T14:16:54Z","title":"Mitigating the Alignment Tax of RLHF","version":4},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2309.06256","snapshot_observed_at":"2026-07-11T13:53:36.775836Z","title":"arXiv preprint arXiv:2309.06256 , year=","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2607.04763","last_updated":"2026-07-26T14:17:18Z","snapshot_observed_at":"2026-08-02T10:24:43.977557Z","submitted_at":"2026-07-06T07:56:53Z","title":"Multi-Turn On-Policy Distillation with Prefix Replay","version":1},"reference_index":249,"source":"arxiv_source","source_observed_at":"2026-07-11T13:53:36.775836Z"},"links":{"cited_paper":"/paper/2309.06256","citing_paper":"/paper/2607.04763"},"observation_digest":"sha256:95462973407fb9e507cb8609e326c7effcaee71170deb868bd3d1c3d7cd8b028","observation_id":"71b40abc-9420-4341-99ca-719257dc5299","resolution":{"observed_at":"2026-07-11T13:53:36.775836Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2309.06256","last_updated":"2024-10-13T19:27:20Z","snapshot_observed_at":"2026-07-06T16:17:26.924098Z","submitted_at":"2023-09-12T14:16:54Z","title":"Mitigating the Alignment Tax of RLHF","version":4},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2309.06256","snapshot_observed_at":"2026-08-02T08:41:01.426952Z","title":"arXiv preprint arXiv:2309.06256 , year=","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2607.04763","last_updated":"2026-07-26T14:17:18Z","snapshot_observed_at":"2026-08-02T10:24:43.977557Z","submitted_at":"2026-07-06T07:56:53Z","title":"Multi-Turn On-Policy Distillation with Prefix Replay","version":3},"reference_index":250,"source":"arxiv_source","source_observed_at":"2026-08-02T08:41:01.426952Z"},"links":{"cited_paper":"/paper/2309.06256","citing_paper":"/paper/2607.04763"},"observation_digest":"sha256:ecbeec62fe5acbec90267196f7af7f8055f7222b9ccd759823120f39f1328288","observation_id":"737093c4-4e6b-43e1-8156-8c8d0465eb86","resolution":{"observed_at":"2026-08-02T08:41:01.426952Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2309.06256","last_updated":"2024-10-13T19:27:20Z","snapshot_observed_at":"2026-07-06T16:17:26.924098Z","submitted_at":"2023-09-12T14:16:54Z","title":"Mitigating the Alignment Tax of RLHF","version":4},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2309.06256","snapshot_observed_at":"2026-08-02T09:51:03.639953Z","title":"arXiv:2309.06256 , year=","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2607.16252","last_updated":"2026-06-27T03:25:36Z","snapshot_observed_at":"2026-08-08T02:31:05.095460Z","submitted_at":"2026-06-27T03:25:36Z","title":"SOS-LoRA: Static Orthogonal-Subspace Low-Rank Adaptation with Fixed Multi-Scale Scaling","version":1},"reference_index":157,"source":"arxiv_source","source_observed_at":"2026-08-02T09:51:03.639953Z"},"links":{"cited_paper":"/paper/2309.06256","citing_paper":"/paper/2607.16252"},"observation_digest":"sha256:2dd1572b073462b602eabad947ca715c9c734b848f26d6ad56b85a714ef772ab","observation_id":"66d3b4ac-d4b1-49c5-8d8c-4c78377c6b24","resolution":{"observed_at":"2026-08-02T09:51:03.639953Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2309.06256","last_updated":"2024-10-13T19:27:20Z","snapshot_observed_at":"2026-07-06T16:17:26.924098Z","submitted_at":"2023-09-12T14:16:54Z","title":"Mitigating the Alignment Tax of RLHF","version":4},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2309.06256","snapshot_observed_at":"2026-08-04T04:53:01.794216Z","title":"Mitigating the alignment tax of rlhf,","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2608.02553","last_updated":"2026-08-03T17:37:38Z","snapshot_observed_at":"2026-08-06T23:39:10.317481Z","submitted_at":"2026-08-03T17:37:38Z","title":"A Taxonomy of Cognitive Capability Gaps in Generative and Agentic AI","version":1},"reference_index":89,"source":"pdf_text","source_observed_at":"2026-08-04T04:53:01.794216Z"},"links":{"cited_paper":"/paper/2309.06256","citing_paper":"/paper/2608.02553"},"observation_digest":"sha256:fdb33514d57853fb46f24858cb7c4370682c4df29efbdf31cc8821594948c349","observation_id":"2ce3cc3a-5cad-4056-8208-3bf5fef146ab","resolution":{"observed_at":"2026-08-04T04:53:01.794216Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}}],"links":{"evidence":"/evidence","html":"/paper/2309.06256/citation-record","integrity":"/paper/2309.06256/integrity","json":"/paper/2309.06256/citation-record.json","paper":"/paper/2309.06256"},"outbound":[],"paper":{"arxiv_id":"2309.06256","last_updated":"2024-10-13T19:27:20Z","latest_version":4,"primary_category":"cs.LG","snapshot_observed_at":"2026-07-06T16:17:26.924098Z","submitted_at":"2023-09-12T14:16:54Z","title":"Mitigating the Alignment Tax of RLHF"},"reference_resolution":{"displayed":0,"state_counts":{"malformed_identifier":0,"metadata_mismatch":0,"parse_uncertain":0,"unresolved":0,"verified_exact":0,"verified_fuzzy":0},"total_outbound_references":0},"refusal":"A citation records a reference. It does not transfer a finding from one paper to another.","schema":"pith.paper-citation-record.v1","standing_sources":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"thesis":"As of 8 August 2026, this Paper Citation Record lists 0 of 0 outbound references and 22 inbound Pith citation observations for arXiv:2309.06256."}