{"as_of":"2026-08-14T09:55:00Z","caps":{"database_statements":6,"inbound":100,"outbound":100},"context_digest":"sha256:63409d47a692f4e1e4b7d26572613daf6b19aa644bbf7a11b0344be992030662","coverage":[{"denominator":50,"lane":"reference_resolution","note":"Typed states for the displayed outbound observations.","records_observed":50,"source":"paper_references, paper_reference_links","source_observed_at":"2026-08-11T23:27:32.794315Z","state":"measured"},{"denominator":50,"lane":"standing_notices","note":"One-hop event checks from named stored sources.","records_observed":50,"source":"scholarly_work_events, retraction_status_cache","source_observed_at":"2026-08-14T06:32:32.682623+00:00","state":"measured"},{"denominator":0,"lane":"inbound_itemization","note":"Pith citing papers itemized under the disclosed page cap.","records_observed":0,"source":"paper_references, paper_reference_links","source_observed_at":null,"state":"measured"},{"denominator":1,"lane":"external_citation_measurements","note":"A source-named dated measurement, never combined with another source.","records_observed":0,"source":"cited_works","source_observed_at":null,"state":"measured"}],"external_citation_measurements":[],"inbound":[],"links":{"evidence":"/evidence","html":"/paper/2608.09109/citation-record","integrity":"/paper/2608.09109/integrity","json":"/paper/2608.09109/citation-record.json","paper":"/paper/2608.09109"},"outbound":[{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":"10.18653/v1/2022.findings-emnlp.296","metadata_source":"doi_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T23:27:33.044972Z","title":null,"venue":null,"work_id":"9170cb79-abbf-4a13-bbec-735e293da05a","year":2022},"citing_paper":{"arxiv_id":"2608.09109","last_updated":"2026-08-10T04:27:58Z","snapshot_observed_at":"2026-08-14T07:10:14.697062Z","submitted_at":"2026-08-10T04:27:58Z","title":"Different Feedback, Different Updates: Selective Self-Learning from User Interactions for Large Language Models","version":1},"reference_index":1,"source":"pdf_text","source_observed_at":"2026-08-11T23:27:32.605729Z"},"links":{"citing_paper":"/paper/2608.09109"},"observation_digest":"sha256:1bc58295469faeb7bec89f460e9c4650c09f57ff8a1b7a8a67c403dc80e7b4bb","observation_id":"3760904d-7100-4036-b05c-1f494382414b","resolution":{"observed_at":"2026-08-11T23:27:33.048510Z","resolver_source":"doi","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-14T06:32:32.682623+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-14T06:32:32.682623+00:00","source":"crossref"},{"observed_at":"2026-08-14T06:32:18.44784+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2510.17281","last_updated":"2026-06-03T04:15:09Z","snapshot_observed_at":"2026-08-11T13:20:24.725773Z","submitted_at":"2025-10-20T08:16:12Z","title":"MemoryBench: A Benchmark for Memory and Continual Learning in LLM Systems","version":7},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2510.17281","snapshot_observed_at":"2026-08-11T23:27:32.610040Z","title":null,"venue":null,"work_id":null,"year":2026},"citing_paper":{"arxiv_id":"2608.09109","last_updated":"2026-08-10T04:27:58Z","snapshot_observed_at":"2026-08-14T07:10:14.697062Z","submitted_at":"2026-08-10T04:27:58Z","title":"Different Feedback, Different Updates: Selective Self-Learning from User Interactions for Large Language Models","version":1},"reference_index":2,"source":"pdf_text","source_observed_at":"2026-08-11T23:27:32.610040Z"},"links":{"cited_paper":"/paper/2510.17281","citing_paper":"/paper/2608.09109"},"observation_digest":"sha256:e439d47e67a35cedc9c15baf8d4b3678c30422b3440de9a02b953919a0e83e9b","observation_id":"4fcf339f-deca-450b-9596-1f0e2439ffc2","resolution":{"observed_at":"2026-08-11T23:27:32.610040Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T23:27:32.614542Z","title":null,"venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2608.09109","last_updated":"2026-08-10T04:27:58Z","snapshot_observed_at":"2026-08-14T07:10:14.697062Z","submitted_at":"2026-08-10T04:27:58Z","title":"Different Feedback, Different Updates: Selective Self-Learning from User Interactions for Large Language Models","version":1},"reference_index":3,"source":"pdf_text","source_observed_at":"2026-08-11T23:27:32.614542Z"},"links":{"citing_paper":"/paper/2608.09109"},"observation_digest":"sha256:8afcfbd38a6e14fab299e2ae1057d452d58fa8235057807ff7111b07a73e6d94","observation_id":"3d69960b-10b2-406f-a378-dcc050bc3ea2","resolution":{"observed_at":"2026-08-11T23:27:32.614542Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":"10.18653/v1/2026.acl-long.441","metadata_source":"doi_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T23:27:33.027940Z","title":null,"venue":null,"work_id":"e22ef603-1225-4dd0-87ff-1a451f8ec33a","year":2026},"citing_paper":{"arxiv_id":"2608.09109","last_updated":"2026-08-10T04:27:58Z","snapshot_observed_at":"2026-08-14T07:10:14.697062Z","submitted_at":"2026-08-10T04:27:58Z","title":"Different Feedback, Different Updates: Selective Self-Learning from User Interactions for Large Language Models","version":1},"reference_index":4,"source":"pdf_text","source_observed_at":"2026-08-11T23:27:32.617845Z"},"links":{"citing_paper":"/paper/2608.09109"},"observation_digest":"sha256:0f80da14da0077be872129050c670a5a39a2a0d7195a295adb005b4247db9322","observation_id":"9e44dc09-ade2-467d-88f0-60caa9a84b51","resolution":{"observed_at":"2026-08-11T23:27:33.033154Z","resolver_source":"doi","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-14T06:32:32.682623+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-14T06:32:32.682623+00:00","source":"crossref"},{"observed_at":"2026-08-14T06:32:18.44784+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":"10.18653/v1/2025.acl-long.1200","metadata_source":"doi_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T23:27:33.008650Z","title":null,"venue":null,"work_id":"609d43ff-eb88-41b6-ad07-27cbfb2a66d4","year":2025},"citing_paper":{"arxiv_id":"2608.09109","last_updated":"2026-08-10T04:27:58Z","snapshot_observed_at":"2026-08-14T07:10:14.697062Z","submitted_at":"2026-08-10T04:27:58Z","title":"Different Feedback, Different Updates: Selective Self-Learning from User Interactions for Large Language Models","version":1},"reference_index":5,"source":"pdf_text","source_observed_at":"2026-08-11T23:27:32.621106Z"},"links":{"citing_paper":"/paper/2608.09109"},"observation_digest":"sha256:df630ee96ac00557e91db0b015547a7238c61df5d51f160c8377485d4e2414f8","observation_id":"c8b5931f-6879-4638-a8ef-5745600d941b","resolution":{"observed_at":"2026-08-11T23:27:33.012175Z","resolver_source":"doi","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-14T06:32:32.682623+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-14T06:32:32.682623+00:00","source":"crossref"},{"observed_at":"2026-08-14T06:32:18.44784+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2404.04475","last_updated":"2025-03-10T09:27:03Z","snapshot_observed_at":"2026-07-06T17:56:23.317089Z","submitted_at":"2024-04-06T02:29:02Z","title":"Length-Controlled AlpacaEval: A Simple Way to Debias Automatic Evaluators","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2404.04475","snapshot_observed_at":"2026-08-11T23:27:32.625100Z","title":"Hashimoto","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2608.09109","last_updated":"2026-08-10T04:27:58Z","snapshot_observed_at":"2026-08-14T07:10:14.697062Z","submitted_at":"2026-08-10T04:27:58Z","title":"Different Feedback, Different Updates: Selective Self-Learning from User Interactions for Large Language Models","version":1},"reference_index":6,"source":"pdf_text","source_observed_at":"2026-08-11T23:27:32.625100Z"},"links":{"cited_paper":"/paper/2404.04475","citing_paper":"/paper/2608.09109"},"observation_digest":"sha256:35d0eb60faea1562657ff8246eebc42635473bab9602f90c244d3a6c69af128c","observation_id":"862f1f3d-120c-4225-a81f-ffb6d140653d","resolution":{"observed_at":"2026-08-11T23:27:32.625100Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2106.09685","last_updated":"2021-10-16T18:40:34Z","snapshot_observed_at":"2026-08-11T08:20:29.798517Z","submitted_at":"2021-06-17T17:37:18Z","title":"LoRA: Low-Rank Adaptation of Large Language Models","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2106.09685","snapshot_observed_at":"2026-08-11T23:27:32.628831Z","title":"Hu, Yelong Shen, Phillip Wallis, Zeyuan Allen-Zhu, Yuanzhi Li, Shean Wang, Lu Wang, and Weizhu Chen","venue":null,"work_id":null,"year":2021},"citing_paper":{"arxiv_id":"2608.09109","last_updated":"2026-08-10T04:27:58Z","snapshot_observed_at":"2026-08-14T07:10:14.697062Z","submitted_at":"2026-08-10T04:27:58Z","title":"Different Feedback, Different Updates: Selective Self-Learning from User Interactions for Large Language Models","version":1},"reference_index":7,"source":"pdf_text","source_observed_at":"2026-08-11T23:27:32.628831Z"},"links":{"cited_paper":"/paper/2106.09685","citing_paper":"/paper/2608.09109"},"observation_digest":"sha256:65d7472de42e5355e5c53e6b676091a040218cc20ac2eb47b06d5c373973999c","observation_id":"e68d4db8-a732-4991-9c3e-ac967e7fe43c","resolution":{"observed_at":"2026-08-11T23:27:32.628831Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2601.20802","last_updated":"2026-02-16T14:49:34Z","snapshot_observed_at":"2026-08-13T15:28:29.238641Z","submitted_at":"2026-01-28T17:45:12Z","title":"Reinforcement Learning via Self-Distillation","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2601.20802","snapshot_observed_at":"2026-08-11T23:27:32.632709Z","title":null,"venue":null,"work_id":null,"year":2026},"citing_paper":{"arxiv_id":"2608.09109","last_updated":"2026-08-10T04:27:58Z","snapshot_observed_at":"2026-08-14T07:10:14.697062Z","submitted_at":"2026-08-10T04:27:58Z","title":"Different Feedback, Different Updates: Selective Self-Learning from User Interactions for Large Language Models","version":1},"reference_index":8,"source":"pdf_text","source_observed_at":"2026-08-11T23:27:32.632709Z"},"links":{"cited_paper":"/paper/2601.20802","citing_paper":"/paper/2608.09109"},"observation_digest":"sha256:cbdb623f5509f7532a131f7c20f76edb6509b7643e8a1de636bca9b10c37dc6c","observation_id":"3262af93-f0ae-4b20-8a6f-7c34cc069583","resolution":{"observed_at":"2026-08-11T23:27:32.632709Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T23:27:32.636081Z","title":null,"venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2608.09109","last_updated":"2026-08-10T04:27:58Z","snapshot_observed_at":"2026-08-14T07:10:14.697062Z","submitted_at":"2026-08-10T04:27:58Z","title":"Different Feedback, Different Updates: Selective Self-Learning from User Interactions for Large Language Models","version":1},"reference_index":9,"source":"pdf_text","source_observed_at":"2026-08-11T23:27:32.636081Z"},"links":{"citing_paper":"/paper/2608.09109"},"observation_digest":"sha256:94e64a18ec5006f491343c34e25bb1bd86c05bfb961f109e0374cafd465a9b83","observation_id":"f17ab736-15e5-40b3-b4d8-d1f9de2db81f","resolution":{"observed_at":"2026-08-11T23:27:32.636081Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T23:27:32.639753Z","title":null,"venue":null,"work_id":null,"year":2026},"citing_paper":{"arxiv_id":"2608.09109","last_updated":"2026-08-10T04:27:58Z","snapshot_observed_at":"2026-08-14T07:10:14.697062Z","submitted_at":"2026-08-10T04:27:58Z","title":"Different Feedback, Different Updates: Selective Self-Learning from User Interactions for Large Language Models","version":1},"reference_index":10,"source":"pdf_text","source_observed_at":"2026-08-11T23:27:32.639753Z"},"links":{"citing_paper":"/paper/2608.09109"},"observation_digest":"sha256:5fa870b33326156d94ba6107771a1fd7d98b9354bf9570ddba8379fdb9c9ebd2","observation_id":"cb9034fa-85a7-449b-acee-a4d13f34b82d","resolution":{"observed_at":"2026-08-11T23:27:32.639753Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":"10.18653/v1/2026.acl-long.1831","metadata_source":"doi_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T23:27:32.976005Z","title":null,"venue":null,"work_id":"5b540574-94a9-4fcd-9467-d4efba3c0069","year":2026},"citing_paper":{"arxiv_id":"2608.09109","last_updated":"2026-08-10T04:27:58Z","snapshot_observed_at":"2026-08-14T07:10:14.697062Z","submitted_at":"2026-08-10T04:27:58Z","title":"Different Feedback, Different Updates: Selective Self-Learning from User Interactions for Large Language Models","version":1},"reference_index":11,"source":"pdf_text","source_observed_at":"2026-08-11T23:27:32.642822Z"},"links":{"citing_paper":"/paper/2608.09109"},"observation_digest":"sha256:4c2db7759aebc34739e54be8376f7ac8112c2b6cdbf7fc3f36445fc95ab87d65","observation_id":"07b71941-e55f-40cb-b877-9fc5233a429d","resolution":{"observed_at":"2026-08-11T23:27:32.979903Z","resolver_source":"doi","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-14T06:32:32.682623+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-14T06:32:32.682623+00:00","source":"crossref"},{"observed_at":"2026-08-14T06:32:18.44784+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2601.22900","last_updated":"2026-05-30T11:40:38Z","snapshot_observed_at":"2026-08-09T01:46:22.354879Z","submitted_at":"2026-01-30T12:19:54Z","title":"MulFeRL: Enhancing Reinforcement Learning with Verbal Feedback in a Multi-turn Loop","version":2},"cited_work":{"arxiv_id":"2601.22900","doi":null,"metadata_source":"pith","pith_arxiv_id":"2601.22900","snapshot_observed_at":"2026-08-11T23:27:33.239830Z","title":"MulFeRL: Enhancing Reinforcement Learning with Verbal Feedback in a Multi-turn Loop","venue":"cs.AI","work_id":"31a684ff-715a-4f20-a491-249bc316e15e","year":2026},"citing_paper":{"arxiv_id":"2608.09109","last_updated":"2026-08-10T04:27:58Z","snapshot_observed_at":"2026-08-14T07:10:14.697062Z","submitted_at":"2026-08-10T04:27:58Z","title":"Different Feedback, Different Updates: Selective Self-Learning from User Interactions for Large Language Models","version":1},"reference_index":12,"source":"pdf_text","source_observed_at":"2026-08-11T23:27:32.645757Z"},"links":{"cited_paper":"/paper/2601.22900","citing_paper":"/paper/2608.09109"},"observation_digest":"sha256:dd03c268bbc72ac9bfd96e2ad1f0c21c551dcdf52f9b34b4729ca8ec7e8bdf9e","observation_id":"06f62b64-f665-46e8-a1f2-1cfca9c3763d","resolution":{"observed_at":"2026-08-11T23:27:33.244435Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-14T06:32:32.682623+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-14T06:32:32.682623+00:00","source":"crossref"},{"observed_at":"2026-08-14T06:32:18.44784+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T23:27:33.571950Z","title":null,"venue":null,"work_id":"3d7f71b8-52c0-4d98-a9dc-cc730e7f8d05","year":null},"citing_paper":{"arxiv_id":"2608.09109","last_updated":"2026-08-10T04:27:58Z","snapshot_observed_at":"2026-08-14T07:10:14.697062Z","submitted_at":"2026-08-10T04:27:58Z","title":"Different Feedback, Different Updates: Selective Self-Learning from User Interactions for Large Language Models","version":1},"reference_index":13,"source":"pdf_text","source_observed_at":"2026-08-11T23:27:32.648978Z"},"links":{"citing_paper":"/paper/2608.09109"},"observation_digest":"sha256:840455903dd60bd69108f2c71a1fe9fd661087c536192c080f5596015aa9c212","observation_id":"08e3b22d-4321-4978-8ba6-b884493639da","resolution":{"observed_at":"2026-08-11T23:27:33.575241Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-14T06:32:32.682623+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-14T06:32:32.682623+00:00","source":"crossref"},{"observed_at":"2026-08-14T06:32:18.44784+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2601.08584","last_updated":"2026-01-13T14:06:03Z","snapshot_observed_at":"2026-08-12T22:11:58.919874Z","submitted_at":"2026-01-13T14:06:03Z","title":"Ministral 3","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2601.08584","snapshot_observed_at":"2026-08-11T23:27:32.655467Z","title":"Liu, Kartik Khandelwal, Sandeep Subramanian, Victor Jouault, Abhinav Rastogi, et al","venue":null,"work_id":null,"year":2026},"citing_paper":{"arxiv_id":"2608.09109","last_updated":"2026-08-10T04:27:58Z","snapshot_observed_at":"2026-08-14T07:10:14.697062Z","submitted_at":"2026-08-10T04:27:58Z","title":"Different Feedback, Different Updates: Selective Self-Learning from User Interactions for Large Language Models","version":1},"reference_index":14,"source":"pdf_text","source_observed_at":"2026-08-11T23:27:32.655467Z"},"links":{"cited_paper":"/paper/2601.08584","citing_paper":"/paper/2608.09109"},"observation_digest":"sha256:2b8d1ad8bd87d02acdbb3567e5d55023626cb6970e75cf39f40e0bfc98cd643b","observation_id":"d275f545-eae2-41de-a217-de6461c842e9","resolution":{"observed_at":"2026-08-11T23:27:32.655467Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2302.02676","last_updated":"2023-10-18T07:11:12Z","snapshot_observed_at":"2026-08-13T12:50:20.733997Z","submitted_at":"2023-02-06T10:28:16Z","title":"Chain of Hindsight Aligns Language Models with Feedback","version":8},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2302.02676","snapshot_observed_at":"2026-08-11T23:27:32.658523Z","title":null,"venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2608.09109","last_updated":"2026-08-10T04:27:58Z","snapshot_observed_at":"2026-08-14T07:10:14.697062Z","submitted_at":"2026-08-10T04:27:58Z","title":"Different Feedback, Different Updates: Selective Self-Learning from User Interactions for Large Language Models","version":1},"reference_index":15,"source":"pdf_text","source_observed_at":"2026-08-11T23:27:32.658523Z"},"links":{"cited_paper":"/paper/2302.02676","citing_paper":"/paper/2608.09109"},"observation_digest":"sha256:99f8699ef50df9232c92ede384ddf955fafa76f68a39314c398672e5ee907903","observation_id":"20999843-4ea5-49a2-a9bb-a9e225dff3ff","resolution":{"observed_at":"2026-08-11T23:27:32.658523Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T23:27:32.661712Z","title":null,"venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2608.09109","last_updated":"2026-08-10T04:27:58Z","snapshot_observed_at":"2026-08-14T07:10:14.697062Z","submitted_at":"2026-08-10T04:27:58Z","title":"Different Feedback, Different Updates: Selective Self-Learning from User Interactions for Large Language Models","version":1},"reference_index":16,"source":"pdf_text","source_observed_at":"2026-08-11T23:27:32.661712Z"},"links":{"citing_paper":"/paper/2608.09109"},"observation_digest":"sha256:db119c0ebe4b95e9af25862f529bfa4b1b3a834ba691742aeb7ee043d251b96b","observation_id":"beb62dd7-0c5f-49a6-99f8-9d28f8a150d9","resolution":{"observed_at":"2026-08-11T23:27:32.661712Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T23:27:32.664633Z","title":null,"venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2608.09109","last_updated":"2026-08-10T04:27:58Z","snapshot_observed_at":"2026-08-14T07:10:14.697062Z","submitted_at":"2026-08-10T04:27:58Z","title":"Different Feedback, Different Updates: Selective Self-Learning from User Interactions for Large Language Models","version":1},"reference_index":17,"source":"pdf_text","source_observed_at":"2026-08-11T23:27:32.664633Z"},"links":{"citing_paper":"/paper/2608.09109"},"observation_digest":"sha256:9e85687c10b05ab312a187bd62e8d521e53b469bd61f9faf2d1ba0b0a53ffa55","observation_id":"cf0b829f-8c19-4c74-a1ac-610978094f78","resolution":{"observed_at":"2026-08-11T23:27:32.664633Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T23:27:32.667584Z","title":"McClelland, Bruce L","venue":null,"work_id":null,"year":1995},"citing_paper":{"arxiv_id":"2608.09109","last_updated":"2026-08-10T04:27:58Z","snapshot_observed_at":"2026-08-14T07:10:14.697062Z","submitted_at":"2026-08-10T04:27:58Z","title":"Different Feedback, Different Updates: Selective Self-Learning from User Interactions for Large Language Models","version":1},"reference_index":18,"source":"pdf_text","source_observed_at":"2026-08-11T23:27:32.667584Z"},"links":{"citing_paper":"/paper/2608.09109"},"observation_digest":"sha256:246d2163030b97b1437cfde8c88f7df008885d9c723572cb99db22b525335c7e","observation_id":"d82ab8b9-1390-4d32-90d9-6fc0156394d1","resolution":{"observed_at":"2026-08-11T23:27:32.667584Z","resolver_source":null,"status":"malformed_identifier"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T23:27:33.562662Z","title":null,"venue":null,"work_id":"5585ab52-ed0d-4cbd-8b82-a267375c0fee","year":2022},"citing_paper":{"arxiv_id":"2608.09109","last_updated":"2026-08-10T04:27:58Z","snapshot_observed_at":"2026-08-14T07:10:14.697062Z","submitted_at":"2026-08-10T04:27:58Z","title":"Different Feedback, Different Updates: Selective Self-Learning from User Interactions for Large Language Models","version":1},"reference_index":19,"source":"pdf_text","source_observed_at":"2026-08-11T23:27:32.670588Z"},"links":{"citing_paper":"/paper/2608.09109"},"observation_digest":"sha256:27812bcb30b7e251b41e354b9baf0d6d4a4a7b55f379883d8d86cd248f2c84cb","observation_id":"792e1bdd-27c4-4b2b-bad8-0f50417b1dbb","resolution":{"observed_at":"2026-08-11T23:27:33.565778Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-14T06:32:32.682623+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-14T06:32:32.682623+00:00","source":"crossref"},{"observed_at":"2026-08-14T06:32:18.44784+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T23:27:32.673400Z","title":null,"venue":null,"work_id":null,"year":2026},"citing_paper":{"arxiv_id":"2608.09109","last_updated":"2026-08-10T04:27:58Z","snapshot_observed_at":"2026-08-14T07:10:14.697062Z","submitted_at":"2026-08-10T04:27:58Z","title":"Different Feedback, Different Updates: Selective Self-Learning from User Interactions for Large Language Models","version":1},"reference_index":20,"source":"pdf_text","source_observed_at":"2026-08-11T23:27:32.673400Z"},"links":{"citing_paper":"/paper/2608.09109"},"observation_digest":"sha256:9ef733f64dd9f5d8d76383717a040187248f537752c1b9a539cbd22296f65945","observation_id":"9483b4e2-e596-4b83-9c38-57ddb06fed9f","resolution":{"observed_at":"2026-08-11T23:27:32.673400Z","resolver_source":null,"status":"malformed_identifier"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T23:27:32.676269Z","title":null,"venue":null,"work_id":null,"year":2026},"citing_paper":{"arxiv_id":"2608.09109","last_updated":"2026-08-10T04:27:58Z","snapshot_observed_at":"2026-08-14T07:10:14.697062Z","submitted_at":"2026-08-10T04:27:58Z","title":"Different Feedback, Different Updates: Selective Self-Learning from User Interactions for Large Language Models","version":1},"reference_index":21,"source":"pdf_text","source_observed_at":"2026-08-11T23:27:32.676269Z"},"links":{"citing_paper":"/paper/2608.09109"},"observation_digest":"sha256:37090993c55e3231d5c1fa946400e3bab743be9ecbcd0930d388f5c3e255dac0","observation_id":"fbdf20d6-8112-4a7b-91a0-fbf670c61672","resolution":{"observed_at":"2026-08-11T23:27:32.676269Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T23:27:33.547343Z","title":null,"venue":null,"work_id":"90165e89-76e5-417b-8ed2-0a3fbceb8e43","year":2026},"citing_paper":{"arxiv_id":"2608.09109","last_updated":"2026-08-10T04:27:58Z","snapshot_observed_at":"2026-08-14T07:10:14.697062Z","submitted_at":"2026-08-10T04:27:58Z","title":"Different Feedback, Different Updates: Selective Self-Learning from User Interactions for Large Language Models","version":1},"reference_index":22,"source":"pdf_text","source_observed_at":"2026-08-11T23:27:32.679064Z"},"links":{"citing_paper":"/paper/2608.09109"},"observation_digest":"sha256:a0b20ecc7587e4889d338daee0706b02188851618372d8195efce47582f846d0","observation_id":"561e00e5-2a54-4e99-b2fc-accaecab29e2","resolution":{"observed_at":"2026-08-11T23:27:33.550933Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-14T06:32:32.682623+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-14T06:32:32.682623+00:00","source":"crossref"},{"observed_at":"2026-08-14T06:32:18.44784+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2305.18290","last_updated":"2024-07-29T22:26:36Z","snapshot_observed_at":"2026-08-01T16:34:38.795326Z","submitted_at":"2023-05-29T17:57:46Z","title":"Direct Preference Optimization: Your Language Model is Secretly a Reward Model","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2305.18290","snapshot_observed_at":"2026-08-11T23:27:32.681796Z","title":"Manning, and Chelsea Finn","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2608.09109","last_updated":"2026-08-10T04:27:58Z","snapshot_observed_at":"2026-08-14T07:10:14.697062Z","submitted_at":"2026-08-10T04:27:58Z","title":"Different Feedback, Different Updates: Selective Self-Learning from User Interactions for Large Language Models","version":1},"reference_index":23,"source":"pdf_text","source_observed_at":"2026-08-11T23:27:32.681796Z"},"links":{"cited_paper":"/paper/2305.18290","citing_paper":"/paper/2608.09109"},"observation_digest":"sha256:9e1846041cbe3e4e2a4a4779aa9910541aaaa2e05b5a3cadbb9c9b74db0368bc","observation_id":"376331fb-3097-4a80-add6-5f2936b4c265","resolution":{"observed_at":"2026-08-11T23:27:32.681796Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T23:27:32.684861Z","title":"Robertson and Hugo Zaragoza","venue":null,"work_id":null,"year":2009},"citing_paper":{"arxiv_id":"2608.09109","last_updated":"2026-08-10T04:27:58Z","snapshot_observed_at":"2026-08-14T07:10:14.697062Z","submitted_at":"2026-08-10T04:27:58Z","title":"Different Feedback, Different Updates: Selective Self-Learning from User Interactions for Large Language Models","version":1},"reference_index":24,"source":"pdf_text","source_observed_at":"2026-08-11T23:27:32.684861Z"},"links":{"citing_paper":"/paper/2608.09109"},"observation_digest":"sha256:14832fea18802eb88e6ab854dba51e8d5504d7b0c6b3172cdf3a05c73872b566","observation_id":"6526d28b-bb1b-4420-80b3-f2b57548fbaf","resolution":{"observed_at":"2026-08-11T23:27:32.684861Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":"10.18653/v1/2026.acl-long.1701","metadata_source":"doi_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T23:27:33.538081Z","title":null,"venue":null,"work_id":"ebe08e58-4173-4b30-b204-c4f95916cff5","year":2026},"citing_paper":{"arxiv_id":"2608.09109","last_updated":"2026-08-10T04:27:58Z","snapshot_observed_at":"2026-08-14T07:10:14.697062Z","submitted_at":"2026-08-10T04:27:58Z","title":"Different Feedback, Different Updates: Selective Self-Learning from User Interactions for Large Language Models","version":1},"reference_index":25,"source":"pdf_text","source_observed_at":"2026-08-11T23:27:32.688088Z"},"links":{"citing_paper":"/paper/2608.09109"},"observation_digest":"sha256:232721dad5c7b7546152f9285e13f9bf724211364afb5d88cf5700830a049ac2","observation_id":"2a503d5c-ec84-4285-bae6-e2445bf70d43","resolution":{"observed_at":"2026-08-11T23:27:33.541137Z","resolver_source":"doi","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-14T06:32:32.682623+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-14T06:32:32.682623+00:00","source":"crossref"},{"observed_at":"2026-08-14T06:32:18.44784+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T23:27:33.528863Z","title":null,"venue":null,"work_id":"a43e2b2b-1664-4cb8-9388-3186c08f7afd","year":2024},"citing_paper":{"arxiv_id":"2608.09109","last_updated":"2026-08-10T04:27:58Z","snapshot_observed_at":"2026-08-14T07:10:14.697062Z","submitted_at":"2026-08-10T04:27:58Z","title":"Different Feedback, Different Updates: Selective Self-Learning from User Interactions for Large Language Models","version":1},"reference_index":26,"source":"pdf_text","source_observed_at":"2026-08-11T23:27:32.691135Z"},"links":{"citing_paper":"/paper/2608.09109"},"observation_digest":"sha256:38b7db6485e3739bf171531d12fa618a3e12cdc370f14227e46894c901f5ccd7","observation_id":"77eea6f1-0ce0-4e04-950d-194db880b5e5","resolution":{"observed_at":"2026-08-11T23:27:33.531952Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-14T06:32:32.682623+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-14T06:32:32.682623+00:00","source":"crossref"},{"observed_at":"2026-08-14T06:32:18.44784+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2501.05727","last_updated":"2025-08-04T02:14:13Z","snapshot_observed_at":"2026-08-10T21:05:12.257951Z","submitted_at":"2025-01-10T05:51:52Z","title":"Self-Evolving Critique Abilities in Large Language Models","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2501.05727","snapshot_observed_at":"2026-08-11T23:27:32.694101Z","title":null,"venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2608.09109","last_updated":"2026-08-10T04:27:58Z","snapshot_observed_at":"2026-08-14T07:10:14.697062Z","submitted_at":"2026-08-10T04:27:58Z","title":"Different Feedback, Different Updates: Selective Self-Learning from User Interactions for Large Language Models","version":1},"reference_index":27,"source":"pdf_text","source_observed_at":"2026-08-11T23:27:32.694101Z"},"links":{"cited_paper":"/paper/2501.05727","citing_paper":"/paper/2608.09109"},"observation_digest":"sha256:cea053e4d3b91059fdb438518da60854d0efc729134682e3cdc3c42c191dbb07","observation_id":"7822ba08-53ed-4e57-9333-65c5b676de04","resolution":{"observed_at":"2026-08-11T23:27:32.694101Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T23:27:33.519417Z","title":null,"venue":null,"work_id":"e9dea157-d3f1-4abb-9667-b9a0935c8cad","year":2024},"citing_paper":{"arxiv_id":"2608.09109","last_updated":"2026-08-10T04:27:58Z","snapshot_observed_at":"2026-08-14T07:10:14.697062Z","submitted_at":"2026-08-10T04:27:58Z","title":"Different Feedback, Different Updates: Selective Self-Learning from User Interactions for Large Language Models","version":1},"reference_index":28,"source":"pdf_text","source_observed_at":"2026-08-11T23:27:32.697254Z"},"links":{"citing_paper":"/paper/2608.09109"},"observation_digest":"sha256:c1cf10646bb4a397ab5f8ff253fd7326f321035b03a5940bca26a9ed77441a2b","observation_id":"cf9efc98-6fed-4169-9f6c-1a8855bb1b6f","resolution":{"observed_at":"2026-08-11T23:27:33.522787Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-14T06:32:32.682623+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-14T06:32:32.682623+00:00","source":"crossref"},{"observed_at":"2026-08-14T06:32:18.44784+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T23:27:33.510590Z","title":null,"venue":null,"work_id":"ece07ab8-b873-4f59-be58-18c45d4b6cb0","year":2024},"citing_paper":{"arxiv_id":"2608.09109","last_updated":"2026-08-10T04:27:58Z","snapshot_observed_at":"2026-08-14T07:10:14.697062Z","submitted_at":"2026-08-10T04:27:58Z","title":"Different Feedback, Different Updates: Selective Self-Learning from User Interactions for Large Language Models","version":1},"reference_index":29,"source":"pdf_text","source_observed_at":"2026-08-11T23:27:32.705215Z"},"links":{"citing_paper":"/paper/2608.09109"},"observation_digest":"sha256:fc012b19f43f750ea213a2ca25871e233f469590e7f41200df99f215063e02ac","observation_id":"d90a5be7-9302-4bda-9959-a2b3fcc63802","resolution":{"observed_at":"2026-08-11T23:27:33.513581Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-14T06:32:32.682623+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-14T06:32:32.682623+00:00","source":"crossref"},{"observed_at":"2026-08-14T06:32:18.44784+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2511.20857","last_updated":"2026-05-18T16:18:02Z","snapshot_observed_at":"2026-08-11T09:44:01.811561Z","submitted_at":"2025-11-25T21:08:07Z","title":"Evo-Memory: Benchmarking LLM Agent Test-time Learning with Self-Evolving Memory","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2511.20857","snapshot_observed_at":"2026-08-11T23:27:32.708285Z","title":"Chi, Chi Wang, Shuo Chen, Fernando Pereira, Wang-Cheng Kang, and Derek Zhiyuan Cheng","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2608.09109","last_updated":"2026-08-10T04:27:58Z","snapshot_observed_at":"2026-08-14T07:10:14.697062Z","submitted_at":"2026-08-10T04:27:58Z","title":"Different Feedback, Different Updates: Selective Self-Learning from User Interactions for Large Language Models","version":1},"reference_index":30,"source":"pdf_text","source_observed_at":"2026-08-11T23:27:32.708285Z"},"links":{"cited_paper":"/paper/2511.20857","citing_paper":"/paper/2608.09109"},"observation_digest":"sha256:81778260dcf885c4af1bb241981ec76d372523a20a20d42f0f568afe8d10cb7d","observation_id":"415fb5b4-650a-4e1b-bfd7-4e921de5505c","resolution":{"observed_at":"2026-08-11T23:27:32.708285Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T23:27:32.711965Z","title":null,"venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2608.09109","last_updated":"2026-08-10T04:27:58Z","snapshot_observed_at":"2026-08-14T07:10:14.697062Z","submitted_at":"2026-08-10T04:27:58Z","title":"Different Feedback, Different Updates: Selective Self-Learning from User Interactions for Large Language Models","version":1},"reference_index":31,"source":"pdf_text","source_observed_at":"2026-08-11T23:27:32.711965Z"},"links":{"citing_paper":"/paper/2608.09109"},"observation_digest":"sha256:5a0f79a0591272b619ef041bac6dfe7ced9fd9c79b204c0a8cc3cd760b0dc946","observation_id":"a0a85de5-5111-4429-bd86-64ebbaf5692a","resolution":{"observed_at":"2026-08-11T23:27:32.711965Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":"10.18653/v1/2024.findings-acl.818","metadata_source":"doi_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T23:27:32.880135Z","title":null,"venue":null,"work_id":"4edd77da-3c37-499c-a9b6-b4ca2d80408f","year":2024},"citing_paper":{"arxiv_id":"2608.09109","last_updated":"2026-08-10T04:27:58Z","snapshot_observed_at":"2026-08-14T07:10:14.697062Z","submitted_at":"2026-08-10T04:27:58Z","title":"Different Feedback, Different Updates: Selective Self-Learning from User Interactions for Large Language Models","version":1},"reference_index":32,"source":"pdf_text","source_observed_at":"2026-08-11T23:27:32.718549Z"},"links":{"citing_paper":"/paper/2608.09109"},"observation_digest":"sha256:9afd9cc989c9b89b601c9186f8c65765f2617a92d22b8cfddc167da145f0fef3","observation_id":"f4c62592-924e-4c26-b660-556558500d25","resolution":{"observed_at":"2026-08-11T23:27:32.884316Z","resolver_source":"doi","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-14T06:32:32.682623+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-14T06:32:32.682623+00:00","source":"crossref"},{"observed_at":"2026-08-14T06:32:18.44784+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2502.12110","last_updated":"2025-10-08T01:46:37Z","snapshot_observed_at":"2026-08-03T02:27:06.991396Z","submitted_at":"2025-02-17T18:36:14Z","title":"A-MEM: Agentic Memory for LLM Agents","version":11},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2502.12110","snapshot_observed_at":"2026-08-11T23:27:32.715114Z","title":"arXiv:2502.12110 [cs.CL] https://arxiv.org/abs/2502.12110","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2608.09109","last_updated":"2026-08-10T04:27:58Z","snapshot_observed_at":"2026-08-14T07:10:14.697062Z","submitted_at":"2026-08-10T04:27:58Z","title":"Different Feedback, Different Updates: Selective Self-Learning from User Interactions for Large Language Models","version":1},"reference_index":33,"source":"pdf_text","source_observed_at":"2026-08-11T23:27:32.715114Z"},"links":{"cited_paper":"/paper/2502.12110","citing_paper":"/paper/2608.09109"},"observation_digest":"sha256:4881526409ff8bb28d9906fd9a09507a0f3ac8b7790480f7dd3f3605b4ae9822","observation_id":"fc7a41c0-480e-4f1d-87d6-03607d2196b1","resolution":{"observed_at":"2026-08-11T23:27:32.715114Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":"10.18653/v1/2026.findings-acl.22","metadata_source":"doi_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T23:27:32.849535Z","title":null,"venue":null,"work_id":"a0e42ba4-5f01-4346-8db9-ebedbf30b6b4","year":2026},"citing_paper":{"arxiv_id":"2608.09109","last_updated":"2026-08-10T04:27:58Z","snapshot_observed_at":"2026-08-14T07:10:14.697062Z","submitted_at":"2026-08-10T04:27:58Z","title":"Different Feedback, Different Updates: Selective Self-Learning from User Interactions for Large Language Models","version":1},"reference_index":34,"source":"pdf_text","source_observed_at":"2026-08-11T23:27:32.725029Z"},"links":{"citing_paper":"/paper/2608.09109"},"observation_digest":"sha256:15a391b53d6a2d53b89a87944f19a981acdb03df97639588ca11b38f2a161c5d","observation_id":"a388f0ef-b094-4aee-acf6-25a7a29ca21e","resolution":{"observed_at":"2026-08-11T23:27:32.855611Z","resolver_source":"doi","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-14T06:32:32.682623+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-14T06:32:32.682623+00:00","source":"crossref"},{"observed_at":"2026-08-14T06:32:18.44784+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2505.09388","last_updated":"2025-05-14T13:41:34Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2025-05-14T13:41:34Z","title":"Qwen3 Technical Report","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2505.09388","snapshot_observed_at":"2026-08-11T23:27:32.721644Z","title":null,"venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2608.09109","last_updated":"2026-08-10T04:27:58Z","snapshot_observed_at":"2026-08-14T07:10:14.697062Z","submitted_at":"2026-08-10T04:27:58Z","title":"Different Feedback, Different Updates: Selective Self-Learning from User Interactions for Large Language Models","version":1},"reference_index":35,"source":"pdf_text","source_observed_at":"2026-08-11T23:27:32.721644Z"},"links":{"cited_paper":"/paper/2505.09388","citing_paper":"/paper/2608.09109"},"observation_digest":"sha256:da59447aba2f43992372ec3b96a80e8f3ddf8104fd8a23c59cb07a7c299b819f","observation_id":"4ebb84fc-8be7-4135-8d28-d3f571206e98","resolution":{"observed_at":"2026-08-11T23:27:32.721644Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2405.01470","last_updated":"2024-05-02T17:00:02Z","snapshot_observed_at":"2026-08-13T13:43:58.816757Z","submitted_at":"2024-05-02T17:00:02Z","title":"WildChat: 1M ChatGPT Interaction Logs in the Wild","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2405.01470","snapshot_observed_at":"2026-08-11T23:27:32.740590Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2608.09109","last_updated":"2026-08-10T04:27:58Z","snapshot_observed_at":"2026-08-14T07:10:14.697062Z","submitted_at":"2026-08-10T04:27:58Z","title":"Different Feedback, Different Updates: Selective Self-Learning from User Interactions for Large Language Models","version":1},"reference_index":36,"source":"pdf_text","source_observed_at":"2026-08-11T23:27:32.740590Z"},"links":{"cited_paper":"/paper/2405.01470","citing_paper":"/paper/2608.09109"},"observation_digest":"sha256:f11bf98eca440f3340fcd9df5769a33382e2f8c02dc3f36424ec130ba63ade9d","observation_id":"23d260cc-f4c9-4466-acdf-142244cb54d3","resolution":{"observed_at":"2026-08-11T23:27:32.740590Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T23:27:32.728105Z","title":null,"venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2608.09109","last_updated":"2026-08-10T04:27:58Z","snapshot_observed_at":"2026-08-14T07:10:14.697062Z","submitted_at":"2026-08-10T04:27:58Z","title":"Different Feedback, Different Updates: Selective Self-Learning from User Interactions for Large Language Models","version":1},"reference_index":37,"source":"pdf_text","source_observed_at":"2026-08-11T23:27:32.728105Z"},"links":{"citing_paper":"/paper/2608.09109"},"observation_digest":"sha256:aeaf63df07e1c2efc49ee9df01a2c72105657a7be72cf565326cc0f437f592c1","observation_id":"a2726180-dc30-412d-9334-cdfd54036779","resolution":{"observed_at":"2026-08-11T23:27:32.728105Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2506.05176","last_updated":"2025-06-11T02:54:49Z","snapshot_observed_at":"2026-08-13T07:28:32.447439Z","submitted_at":"2025-06-05T15:49:48Z","title":"Qwen3 Embedding: Advancing Text Embedding and Reranking Through Foundation Models","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2506.05176","snapshot_observed_at":"2026-08-11T23:27:32.737128Z","title":"doi:10.48550/arXiv","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2608.09109","last_updated":"2026-08-10T04:27:58Z","snapshot_observed_at":"2026-08-14T07:10:14.697062Z","submitted_at":"2026-08-10T04:27:58Z","title":"Different Feedback, Different Updates: Selective Self-Learning from User Interactions for Large Language Models","version":1},"reference_index":38,"source":"pdf_text","source_observed_at":"2026-08-11T23:27:32.737128Z"},"links":{"cited_paper":"/paper/2506.05176","citing_paper":"/paper/2608.09109"},"observation_digest":"sha256:41208d7cf68334cea66132c463bece1d8b54211cbfdcd8e1d7f7295759c19e3d","observation_id":"47d7215b-6d51-4bfb-b525-ee48dc605ca6","resolution":{"observed_at":"2026-08-11T23:27:32.737128Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2309.11998","last_updated":"2024-03-10T19:34:57Z","snapshot_observed_at":"2026-08-13T10:08:28.290453Z","submitted_at":"2023-09-21T12:13:55Z","title":"LMSYS-Chat-1M: A Large-Scale Real-World LLM Conversation Dataset","version":4},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2309.11998","snapshot_observed_at":"2026-08-11T23:27:32.744082Z","title":"Xing, Joseph E","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2608.09109","last_updated":"2026-08-10T04:27:58Z","snapshot_observed_at":"2026-08-14T07:10:14.697062Z","submitted_at":"2026-08-10T04:27:58Z","title":"Different Feedback, Different Updates: Selective Self-Learning from User Interactions for Large Language Models","version":1},"reference_index":40,"source":"pdf_text","source_observed_at":"2026-08-11T23:27:32.744082Z"},"links":{"cited_paper":"/paper/2309.11998","citing_paper":"/paper/2608.09109"},"observation_digest":"sha256:8ceca593c8d3cf4c8ceee10026f171fad801e62cffd41ad2d72f0c404a285a77","observation_id":"dc7d04a6-0cd6-473b-bfce-59818bc5591e","resolution":{"observed_at":"2026-08-11T23:27:32.744082Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2311.07911","last_updated":"2023-11-14T05:13:55Z","snapshot_observed_at":"2026-07-06T16:47:08.877195Z","submitted_at":"2023-11-14T05:13:55Z","title":"Instruction-Following Evaluation for Large Language Models","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2311.07911","snapshot_observed_at":"2026-08-11T23:27:32.754578Z","title":"components","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2608.09109","last_updated":"2026-08-10T04:27:58Z","snapshot_observed_at":"2026-08-14T07:10:14.697062Z","submitted_at":"2026-08-10T04:27:58Z","title":"Different Feedback, Different Updates: Selective Self-Learning from User Interactions for Large Language Models","version":1},"reference_index":41,"source":"pdf_text","source_observed_at":"2026-08-11T23:27:32.754578Z"},"links":{"cited_paper":"/paper/2311.07911","citing_paper":"/paper/2608.09109"},"observation_digest":"sha256:c85192cc70bc3174ad3db1348b0f3dcca8d9907306e545fcb6dfe99d7b2c66ed","observation_id":"c47baf38-2a0e-4547-ac64-d94e32b777bd","resolution":{"observed_at":"2026-08-11T23:27:32.754578Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T23:27:33.485802Z","title":"This includes ambiguity, irrelevance, conflict with TASK, pure reaction without a requested change, and a separate deliverable","venue":null,"work_id":"d3201be2-6ed9-4c36-aa55-b067c8b076a7","year":null},"citing_paper":{"arxiv_id":"2608.09109","last_updated":"2026-08-10T04:27:58Z","snapshot_observed_at":"2026-08-14T07:10:14.697062Z","submitted_at":"2026-08-10T04:27:58Z","title":"Different Feedback, Different Updates: Selective Self-Learning from User Interactions for Large Language Models","version":1},"reference_index":42,"source":"pdf_text","source_observed_at":"2026-08-11T23:27:32.766082Z"},"links":{"citing_paper":"/paper/2608.09109"},"observation_digest":"sha256:7d6d941077b11590d69aff7d7e9ca9370baea0183aab8a5e518d58aa01b4a9b4","observation_id":"e89f3539-2c63-46e0-8a9b-57d2bd83bb41","resolution":{"observed_at":"2026-08-11T23:27:33.491012Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-14T06:32:32.682623+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-14T06:32:32.682623+00:00","source":"crossref"},{"observed_at":"2026-08-14T06:32:18.44784+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T23:27:33.474755Z","title":null,"venue":null,"work_id":"8fbb450d-fd9f-4f0c-8cec-117c7270652f","year":null},"citing_paper":{"arxiv_id":"2608.09109","last_updated":"2026-08-10T04:27:58Z","snapshot_observed_at":"2026-08-14T07:10:14.697062Z","submitted_at":"2026-08-10T04:27:58Z","title":"Different Feedback, Different Updates: Selective Self-Learning from User Interactions for Large Language Models","version":1},"reference_index":43,"source":"pdf_text","source_observed_at":"2026-08-11T23:27:32.772120Z"},"links":{"citing_paper":"/paper/2608.09109"},"observation_digest":"sha256:93556f38da039dbeb4dd4d48e19e50cd13ea9f968c0e3964b9ee082907394705","observation_id":"78226881-846d-4006-a881-93d01b07ee8f","resolution":{"observed_at":"2026-08-11T23:27:33.479635Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-14T06:32:32.682623+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-14T06:32:32.682623+00:00","source":"crossref"},{"observed_at":"2026-08-14T06:32:18.44784+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T23:27:33.458480Z","title":"role\": \"FIX","venue":null,"work_id":"4c0027d8-4a77-44c6-9c46-07ad7ea18690","year":null},"citing_paper":{"arxiv_id":"2608.09109","last_updated":"2026-08-10T04:27:58Z","snapshot_observed_at":"2026-08-14T07:10:14.697062Z","submitted_at":"2026-08-10T04:27:58Z","title":"Different Feedback, Different Updates: Selective Self-Learning from User Interactions for Large Language Models","version":1},"reference_index":44,"source":"pdf_text","source_observed_at":"2026-08-11T23:27:32.775505Z"},"links":{"citing_paper":"/paper/2608.09109"},"observation_digest":"sha256:b7443bca8b77758360067bf4d0868914d7a20aa87b3a2d6b8ea44afffdd377cf","observation_id":"93bc17d9-e6b9-4cab-a2ba-55ca31263524","resolution":{"observed_at":"2026-08-11T23:27:33.462154Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-14T06:32:32.682623+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-14T06:32:32.682623+00:00","source":"crossref"},{"observed_at":"2026-08-14T06:32:18.44784+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T23:27:33.448203Z","title":"action\":","venue":null,"work_id":"4d6c446c-0200-4a8b-8c50-6cc137dff782","year":2018},"citing_paper":{"arxiv_id":"2608.09109","last_updated":"2026-08-10T04:27:58Z","snapshot_observed_at":"2026-08-14T07:10:14.697062Z","submitted_at":"2026-08-10T04:27:58Z","title":"Different Feedback, Different Updates: Selective Self-Learning from User Interactions for Large Language Models","version":1},"reference_index":45,"source":"pdf_text","source_observed_at":"2026-08-11T23:27:32.778468Z"},"links":{"citing_paper":"/paper/2608.09109"},"observation_digest":"sha256:697a5ed655c5889d8629503f216e10c889a641542bca824859feb6c0e540537e","observation_id":"d5a520ba-7b9a-46e5-aa57-86bd69e418b3","resolution":{"observed_at":"2026-08-11T23:27:33.452113Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-14T06:32:32.682623+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-14T06:32:32.682623+00:00","source":"crossref"},{"observed_at":"2026-08-14T06:32:18.44784+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T23:27:33.437663Z","title":null,"venue":null,"work_id":"924c6014-d010-4574-b66a-72b21a30a6b3","year":null},"citing_paper":{"arxiv_id":"2608.09109","last_updated":"2026-08-10T04:27:58Z","snapshot_observed_at":"2026-08-14T07:10:14.697062Z","submitted_at":"2026-08-10T04:27:58Z","title":"Different Feedback, Different Updates: Selective Self-Learning from User Interactions for Large Language Models","version":1},"reference_index":46,"source":"pdf_text","source_observed_at":"2026-08-11T23:27:32.782294Z"},"links":{"citing_paper":"/paper/2608.09109"},"observation_digest":"sha256:d14542e19d53589eba52936612811c7607695211fa43a4d8602dbd4d226f70d4","observation_id":"f8d15c49-df33-4187-92c7-0206922ec122","resolution":{"observed_at":"2026-08-11T23:27:33.441027Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-14T06:32:32.682623+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-14T06:32:32.682623+00:00","source":"crossref"},{"observed_at":"2026-08-14T06:32:18.44784+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T23:27:33.427764Z","title":"I’m following up on the dashboard I sent for your use and would appreciate your feedback on its clarity, usefulness, and any areas that could be improved","venue":null,"work_id":"81404e79-895d-4cc4-a642-7a7e7857f716","year":null},"citing_paper":{"arxiv_id":"2608.09109","last_updated":"2026-08-10T04:27:58Z","snapshot_observed_at":"2026-08-14T07:10:14.697062Z","submitted_at":"2026-08-10T04:27:58Z","title":"Different Feedback, Different Updates: Selective Self-Learning from User Interactions for Large Language Models","version":1},"reference_index":47,"source":"pdf_text","source_observed_at":"2026-08-11T23:27:32.785187Z"},"links":{"citing_paper":"/paper/2608.09109"},"observation_digest":"sha256:b44f854080ae2e225004779fbb64f173818ea7f3b00b5155d7aba7936a096919","observation_id":"f5705954-ade0-4a35-9309-f2d607e95d4f","resolution":{"observed_at":"2026-08-11T23:27:33.431401Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-14T06:32:32.682623+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-14T06:32:32.682623+00:00","source":"crossref"},{"observed_at":"2026-08-14T06:32:18.44784+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T23:27:33.398137Z","title":"action\":","venue":null,"work_id":"a128af09-93f5-4eb8-b631-007ebec0f95c","year":null},"citing_paper":{"arxiv_id":"2608.09109","last_updated":"2026-08-10T04:27:58Z","snapshot_observed_at":"2026-08-14T07:10:14.697062Z","submitted_at":"2026-08-10T04:27:58Z","title":"Different Feedback, Different Updates: Selective Self-Learning from User Interactions for Large Language Models","version":1},"reference_index":48,"source":"pdf_text","source_observed_at":"2026-08-11T23:27:32.788305Z"},"links":{"citing_paper":"/paper/2608.09109"},"observation_digest":"sha256:f5b7440ae3ac166abf148577af51e81faade0d300d80ef9853819ab6670b945d","observation_id":"1a3b1f5d-66b4-4109-be39-1f2bca14106f","resolution":{"observed_at":"2026-08-11T23:27:33.417328Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-14T06:32:32.682623+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-14T06:32:32.682623+00:00","source":"crossref"},{"observed_at":"2026-08-14T06:32:18.44784+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T23:27:33.385229Z","title":"Thank you for your guidance","venue":null,"work_id":"ea091ef5-71d7-491c-9a63-e2f31a5bc0a9","year":null},"citing_paper":{"arxiv_id":"2608.09109","last_updated":"2026-08-10T04:27:58Z","snapshot_observed_at":"2026-08-14T07:10:14.697062Z","submitted_at":"2026-08-10T04:27:58Z","title":"Different Feedback, Different Updates: Selective Self-Learning from User Interactions for Large Language Models","version":1},"reference_index":49,"source":"pdf_text","source_observed_at":"2026-08-11T23:27:32.791422Z"},"links":{"citing_paper":"/paper/2608.09109"},"observation_digest":"sha256:781131eefd8472a5b95fe60fed76484f4bf9ff8f418b4777ccfde960355ffd4e","observation_id":"6aa74e32-9a8e-4e0f-bb3e-f4d4b87e70de","resolution":{"observed_at":"2026-08-11T23:27:33.388717Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-14T06:32:32.682623+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-14T06:32:32.682623+00:00","source":"crossref"},{"observed_at":"2026-08-14T06:32:18.44784+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T23:27:33.357582Z","title":"Target- consistent","venue":null,"work_id":"1058b4bf-b6d5-481b-8994-1feadcc92adf","year":2018},"citing_paper":{"arxiv_id":"2608.09109","last_updated":"2026-08-10T04:27:58Z","snapshot_observed_at":"2026-08-14T07:10:14.697062Z","submitted_at":"2026-08-10T04:27:58Z","title":"Different Feedback, Different Updates: Selective Self-Learning from User Interactions for Large Language Models","version":1},"reference_index":50,"source":"pdf_text","source_observed_at":"2026-08-11T23:27:32.794315Z"},"links":{"citing_paper":"/paper/2608.09109"},"observation_digest":"sha256:0bcebb12d8e1aef85bb6dd54324fb31a8bc05532f7f490016bd9676b80d98427","observation_id":"c4cb33b6-c72e-483b-8b7b-fe5e7589b7e4","resolution":{"observed_at":"2026-08-11T23:27:33.360911Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-14T06:32:32.682623+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-14T06:32:32.682623+00:00","source":"crossref"},{"observed_at":"2026-08-14T06:32:18.44784+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2507.03724","last_updated":"2025-12-03T03:19:27Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2025-07-04T17:21:46Z","title":"MemOS: A Memory OS for AI System","version":4},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2507.03724","snapshot_observed_at":"2026-08-11T23:27:32.652089Z","title":"doi:10.48550/arXiv.2507.03724","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2608.09109","last_updated":"2026-08-10T04:27:58Z","snapshot_observed_at":"2026-08-14T07:10:14.697062Z","submitted_at":"2026-08-10T04:27:58Z","title":"Different Feedback, Different Updates: Selective Self-Learning from User Interactions for Large Language Models","version":1},"reference_index":2025,"source":"pdf_text","source_observed_at":"2026-08-11T23:27:32.652089Z"},"links":{"cited_paper":"/paper/2507.03724","citing_paper":"/paper/2608.09109"},"observation_digest":"sha256:9e45211c066339e58f66a7a650071f7af37b2b10ae56f5e40185d07010144f7e","observation_id":"5b328f84-959e-4e16-af4c-1bd5bcf88fb3","resolution":{"observed_at":"2026-08-11T23:27:32.652089Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}}],"paper":{"arxiv_id":"2608.09109","last_updated":"2026-08-10T04:27:58Z","latest_version":1,"primary_category":"cs.AI","snapshot_observed_at":"2026-08-14T07:10:14.697062Z","submitted_at":"2026-08-10T04:27:58Z","title":"Different Feedback, Different Updates: Selective Self-Learning from User Interactions for Large Language Models"},"reference_resolution":{"displayed":50,"state_counts":{"malformed_identifier":2,"metadata_mismatch":0,"parse_uncertain":0,"unresolved":33,"verified_exact":8,"verified_fuzzy":7},"total_outbound_references":50},"refusal":"A citation records a reference. It does not transfer a finding from one paper to another.","schema":"pith.paper-citation-record.v1","standing_sources":[{"observed_at":"2026-08-14T06:32:32.682623+00:00","source":"crossref"},{"observed_at":"2026-08-14T06:32:18.44784+00:00","source":"retraction_watch"}],"thesis":"As of 14 August 2026, this Paper Citation Record lists 50 of 50 outbound references and 0 inbound Pith citation observations for arXiv:2608.09109."}