{"as_of":"2026-08-13T12:48:00Z","caps":{"database_statements":6,"inbound":100,"outbound":100},"context_digest":"sha256:14dc547505f31c9a7d864ec56d4aaddd4ec995337c2c3b27e01fdc4066ace1a6","coverage":[{"denominator":31,"lane":"reference_resolution","note":"Typed states for the displayed outbound observations.","records_observed":31,"source":"paper_references, paper_reference_links","source_observed_at":"2026-08-12T13:21:29.811392Z","state":"measured"},{"denominator":31,"lane":"standing_notices","note":"One-hop event checks from named stored sources.","records_observed":31,"source":"scholarly_work_events, retraction_status_cache","source_observed_at":"2026-08-13T06:32:02.005865+00:00","state":"measured"},{"denominator":0,"lane":"inbound_itemization","note":"Pith citing papers itemized under the disclosed page cap.","records_observed":0,"source":"paper_references, paper_reference_links","source_observed_at":null,"state":"measured"},{"denominator":1,"lane":"external_citation_measurements","note":"A source-named dated measurement, never combined with another source.","records_observed":0,"source":"cited_works","source_observed_at":null,"state":"measured"}],"external_citation_measurements":[],"inbound":[],"links":{"evidence":"/evidence","html":"/paper/2411.16305/citation-record","integrity":"/paper/2411.16305/integrity","json":"/paper/2411.16305/citation-record.json","paper":"/paper/2411.16305"},"outbound":[{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-12T13:21:29.713062Z","title":"online\" 'onlinestring :=","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2411.16305","last_updated":"2024-11-25T11:47:31Z","snapshot_observed_at":"2026-08-12T13:13:26.422111Z","submitted_at":"2024-11-25T11:47:31Z","title":"Learning from Relevant Subgoals in Successful Dialogs using Iterative Training for Task-oriented Dialog Systems","version":1},"reference_index":1,"source":"arxiv_source","source_observed_at":"2026-08-12T13:21:29.713062Z"},"links":{"citing_paper":"/paper/2411.16305"},"observation_digest":"sha256:6d8b8c78c1b8e8c317ae5b7c90e9fec515bcfa2d23371d7f12c30731616d6d60","observation_id":"9a1ab78b-4450-499a-bbcd-b08c5a09dfb0","resolution":{"observed_at":"2026-08-12T13:21:29.713062Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-12T13:21:29.716933Z","title":"write newline","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2411.16305","last_updated":"2024-11-25T11:47:31Z","snapshot_observed_at":"2026-08-12T13:13:26.422111Z","submitted_at":"2024-11-25T11:47:31Z","title":"Learning from Relevant Subgoals in Successful Dialogs using Iterative Training for Task-oriented Dialog Systems","version":1},"reference_index":2,"source":"arxiv_source","source_observed_at":"2026-08-12T13:21:29.716933Z"},"links":{"citing_paper":"/paper/2411.16305"},"observation_digest":"sha256:e331c4296eafbd1023d9d4b6f01b940cefc7ceafacf880b7286b9dc05a42f6bf","observation_id":"ee639f2c-5ea7-46a5-8173-3a5cea167f1c","resolution":{"observed_at":"2026-08-12T13:21:29.716933Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-12T13:21:29.720707Z","title":null,"venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2411.16305","last_updated":"2024-11-25T11:47:31Z","snapshot_observed_at":"2026-08-12T13:13:26.422111Z","submitted_at":"2024-11-25T11:47:31Z","title":"Learning from Relevant Subgoals in Successful Dialogs using Iterative Training for Task-oriented Dialog Systems","version":1},"reference_index":3,"source":"arxiv_source","source_observed_at":"2026-08-12T13:21:29.720707Z"},"links":{"citing_paper":"/paper/2411.16305"},"observation_digest":"sha256:4665f2b0e77ba045135064bc564b7b16f73d07617c9d9fb89978fe6e92b9d841","observation_id":"edd2f15f-eeac-4f55-bd68-54ae0e9ce912","resolution":{"observed_at":"2026-08-12T13:21:29.720707Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-12T13:21:30.135031Z","title":null,"venue":null,"work_id":"306f70ca-30c3-4659-b2dc-6b5becfccebc","year":2023},"citing_paper":{"arxiv_id":"2411.16305","last_updated":"2024-11-25T11:47:31Z","snapshot_observed_at":"2026-08-12T13:13:26.422111Z","submitted_at":"2024-11-25T11:47:31Z","title":"Learning from Relevant Subgoals in Successful Dialogs using Iterative Training for Task-oriented Dialog Systems","version":1},"reference_index":4,"source":"arxiv_source","source_observed_at":"2026-08-12T13:21:29.723836Z"},"links":{"citing_paper":"/paper/2411.16305"},"observation_digest":"sha256:f934cf08541c7ae3a704a17edf2e03417d491aa9b7197e18dc28a3feadd58e90","observation_id":"3a0f2923-9c29-473c-860e-1a99341eafe8","resolution":{"observed_at":"2026-08-12T13:21:30.138011Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-13T06:32:02.005865+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-13T06:32:02.005865+00:00","source":"crossref"},{"observed_at":"2026-08-13T06:31:53.387327+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-12T13:21:29.727730Z","title":null,"venue":null,"work_id":null,"year":2017},"citing_paper":{"arxiv_id":"2411.16305","last_updated":"2024-11-25T11:47:31Z","snapshot_observed_at":"2026-08-12T13:13:26.422111Z","submitted_at":"2024-11-25T11:47:31Z","title":"Learning from Relevant Subgoals in Successful Dialogs using Iterative Training for Task-oriented Dialog Systems","version":1},"reference_index":5,"source":"arxiv_source","source_observed_at":"2026-08-12T13:21:29.727730Z"},"links":{"citing_paper":"/paper/2411.16305"},"observation_digest":"sha256:4e264b211f0e413bcf6b939bd1e20d39646f417b85373138e7150d37bd7669cc","observation_id":"370f2bb2-8153-44bb-915d-369bc7023e33","resolution":{"observed_at":"2026-08-12T13:21:29.727730Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2302.10342","last_updated":"2023-02-20T22:10:04Z","snapshot_observed_at":"2026-08-13T12:40:34.701897Z","submitted_at":"2023-02-20T22:10:04Z","title":"Fantastic Rewards and How to Tame Them: A Case Study on Reward Learning for Task-oriented Dialogue Systems","version":1},"cited_work":{"arxiv_id":"2302.10342","doi":null,"metadata_source":"pith","pith_arxiv_id":"2302.10342","snapshot_observed_at":"2026-08-12T13:21:30.038285Z","title":"Fantastic Rewards and How to Tame Them: A Case Study on Reward Learning for Task-oriented Dialogue Systems","venue":"cs.CL","work_id":"e68a42a4-aeb1-43a2-ada7-d465bdfdaaeb","year":2023},"citing_paper":{"arxiv_id":"2411.16305","last_updated":"2024-11-25T11:47:31Z","snapshot_observed_at":"2026-08-12T13:13:26.422111Z","submitted_at":"2024-11-25T11:47:31Z","title":"Learning from Relevant Subgoals in Successful Dialogs using Iterative Training for Task-oriented Dialog Systems","version":1},"reference_index":6,"source":"arxiv_source","source_observed_at":"2026-08-12T13:21:29.731432Z"},"links":{"cited_paper":"/paper/2302.10342","citing_paper":"/paper/2411.16305"},"observation_digest":"sha256:d3df11f502c3be47141031bea08e4a9516512b635190b3a0cfe4cb37182ceceb","observation_id":"85fb7f01-5dfa-4baa-b86f-2a76c5804ca6","resolution":{"observed_at":"2026-08-12T13:21:30.042810Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-13T06:32:02.005865+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-13T06:32:02.005865+00:00","source":"crossref"},{"observed_at":"2026-08-13T06:31:53.387327+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-12T13:21:29.734744Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2411.16305","last_updated":"2024-11-25T11:47:31Z","snapshot_observed_at":"2026-08-12T13:13:26.422111Z","submitted_at":"2024-11-25T11:47:31Z","title":"Learning from Relevant Subgoals in Successful Dialogs using Iterative Training for Task-oriented Dialog Systems","version":1},"reference_index":7,"source":"arxiv_source","source_observed_at":"2026-08-12T13:21:29.734744Z"},"links":{"citing_paper":"/paper/2411.16305"},"observation_digest":"sha256:1f62972b023c92ae96af1949e4d3340e3d649bdbd9e5ba19a22d47eb3b969dfd","observation_id":"d474e5a5-656e-4339-8f74-0770b2699a28","resolution":{"observed_at":"2026-08-12T13:21:29.734744Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2402.04792","last_updated":"2024-02-29T20:59:17Z","snapshot_observed_at":"2026-08-13T04:24:21.093265Z","submitted_at":"2024-02-07T12:31:13Z","title":"Direct Language Model Alignment from Online AI Feedback","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2402.04792","snapshot_observed_at":"2026-08-12T13:21:29.737705Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2411.16305","last_updated":"2024-11-25T11:47:31Z","snapshot_observed_at":"2026-08-12T13:13:26.422111Z","submitted_at":"2024-11-25T11:47:31Z","title":"Learning from Relevant Subgoals in Successful Dialogs using Iterative Training for Task-oriented Dialog Systems","version":1},"reference_index":8,"source":"arxiv_source","source_observed_at":"2026-08-12T13:21:29.737705Z"},"links":{"cited_paper":"/paper/2402.04792","citing_paper":"/paper/2411.16305"},"observation_digest":"sha256:2ba27d8268371821c5d6d3f5cb00e612bcf9b8587831fb74a101f34d065d9909","observation_id":"70d0da9a-847c-4e6d-b030-2a0f79bf87fe","resolution":{"observed_at":"2026-08-12T13:21:29.737705Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-12T13:21:30.116427Z","title":null,"venue":null,"work_id":"3db4a836-b196-40df-8718-94f8ca923f77","year":2022},"citing_paper":{"arxiv_id":"2411.16305","last_updated":"2024-11-25T11:47:31Z","snapshot_observed_at":"2026-08-12T13:13:26.422111Z","submitted_at":"2024-11-25T11:47:31Z","title":"Learning from Relevant Subgoals in Successful Dialogs using Iterative Training for Task-oriented Dialog Systems","version":1},"reference_index":9,"source":"arxiv_source","source_observed_at":"2026-08-12T13:21:29.741474Z"},"links":{"citing_paper":"/paper/2411.16305"},"observation_digest":"sha256:d824a6879cea339d498e75fe51892adf8d49829019613ade21ef105aba3c953e","observation_id":"e25c4f12-dee5-40eb-b64e-c7a50b3108a4","resolution":{"observed_at":"2026-08-12T13:21:30.119508Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-13T06:32:02.005865+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-13T06:32:02.005865+00:00","source":"crossref"},{"observed_at":"2026-08-13T06:31:53.387327+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-12T13:21:30.107701Z","title":null,"venue":null,"work_id":"4ca54420-5e61-48fc-99e3-e1fe91fa7cd4","year":2020},"citing_paper":{"arxiv_id":"2411.16305","last_updated":"2024-11-25T11:47:31Z","snapshot_observed_at":"2026-08-12T13:13:26.422111Z","submitted_at":"2024-11-25T11:47:31Z","title":"Learning from Relevant Subgoals in Successful Dialogs using Iterative Training for Task-oriented Dialog Systems","version":1},"reference_index":10,"source":"arxiv_source","source_observed_at":"2026-08-12T13:21:29.744386Z"},"links":{"citing_paper":"/paper/2411.16305"},"observation_digest":"sha256:686bb6bcd174a4ecfc2bceb67dc24adeedfcdf9d80c5c0890e51dd8426481a2f","observation_id":"b49ed11c-28d4-47cb-8783-72fad357df85","resolution":{"observed_at":"2026-08-12T13:21:30.110775Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-13T06:32:02.005865+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-13T06:32:02.005865+00:00","source":"crossref"},{"observed_at":"2026-08-13T06:31:53.387327+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-12T13:21:29.747563Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2411.16305","last_updated":"2024-11-25T11:47:31Z","snapshot_observed_at":"2026-08-12T13:13:26.422111Z","submitted_at":"2024-11-25T11:47:31Z","title":"Learning from Relevant Subgoals in Successful Dialogs using Iterative Training for Task-oriented Dialog Systems","version":1},"reference_index":11,"source":"arxiv_source","source_observed_at":"2026-08-12T13:21:29.747563Z"},"links":{"citing_paper":"/paper/2411.16305"},"observation_digest":"sha256:2f82704507e8245e9899fbdd90f9bdd7cc4380cfc75811b4f2845b5007ee6800","observation_id":"f0cadb74-aeed-48e5-90f7-b8e7a72f75e5","resolution":{"observed_at":"2026-08-12T13:21:29.747563Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-12T13:21:29.750542Z","title":null,"venue":null,"work_id":null,"year":2021},"citing_paper":{"arxiv_id":"2411.16305","last_updated":"2024-11-25T11:47:31Z","snapshot_observed_at":"2026-08-12T13:13:26.422111Z","submitted_at":"2024-11-25T11:47:31Z","title":"Learning from Relevant Subgoals in Successful Dialogs using Iterative Training for Task-oriented Dialog Systems","version":1},"reference_index":12,"source":"arxiv_source","source_observed_at":"2026-08-12T13:21:29.750542Z"},"links":{"citing_paper":"/paper/2411.16305"},"observation_digest":"sha256:79bd420ecb3a69ddb3813dbb709d463d4d90554507e26259c5e02447e1fd8a75","observation_id":"6a7b6c4a-9986-4f80-8d94-cc473a568bd2","resolution":{"observed_at":"2026-08-12T13:21:29.750542Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":"10.18653/v1/2021.findings-emnlp.112","metadata_source":"doi_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-12T13:21:29.905499Z","title":null,"venue":null,"work_id":"26b9a9e8-9709-4fa1-913c-13defc12baca","year":2021},"citing_paper":{"arxiv_id":"2411.16305","last_updated":"2024-11-25T11:47:31Z","snapshot_observed_at":"2026-08-12T13:13:26.422111Z","submitted_at":"2024-11-25T11:47:31Z","title":"Learning from Relevant Subgoals in Successful Dialogs using Iterative Training for Task-oriented Dialog Systems","version":1},"reference_index":13,"source":"arxiv_source","source_observed_at":"2026-08-12T13:21:29.753658Z"},"links":{"citing_paper":"/paper/2411.16305"},"observation_digest":"sha256:a7d823470a23d398cee48b19e2ba5a4b2e157aede4bd3ed19871f9a4df065e97","observation_id":"cc4ccdb4-716a-41bc-b090-938c3513ae55","resolution":{"observed_at":"2026-08-12T13:21:29.908535Z","resolver_source":"doi","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-13T06:32:02.005865+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-13T06:32:02.005865+00:00","source":"crossref"},{"observed_at":"2026-08-13T06:31:53.387327+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-12T13:21:29.756846Z","title":null,"venue":null,"work_id":null,"year":2020},"citing_paper":{"arxiv_id":"2411.16305","last_updated":"2024-11-25T11:47:31Z","snapshot_observed_at":"2026-08-12T13:13:26.422111Z","submitted_at":"2024-11-25T11:47:31Z","title":"Learning from Relevant Subgoals in Successful Dialogs using Iterative Training for Task-oriented Dialog Systems","version":1},"reference_index":14,"source":"arxiv_source","source_observed_at":"2026-08-12T13:21:29.756846Z"},"links":{"citing_paper":"/paper/2411.16305"},"observation_digest":"sha256:f1fd2437511cd93e9c0c36f75464e33b1f0738e671376851f516c9d5ac74cd71","observation_id":"0ff90f6f-dbd6-4c11-b62f-72f28cfb0af3","resolution":{"observed_at":"2026-08-12T13:21:29.756846Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":"10.18653/v1/2021.gem-1.4","metadata_source":"doi_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-12T13:21:29.890566Z","title":null,"venue":null,"work_id":"0eeba437-367c-4001-9432-7a51ca6efeb8","year":2021},"citing_paper":{"arxiv_id":"2411.16305","last_updated":"2024-11-25T11:47:31Z","snapshot_observed_at":"2026-08-12T13:13:26.422111Z","submitted_at":"2024-11-25T11:47:31Z","title":"Learning from Relevant Subgoals in Successful Dialogs using Iterative Training for Task-oriented Dialog Systems","version":1},"reference_index":15,"source":"arxiv_source","source_observed_at":"2026-08-12T13:21:29.760761Z"},"links":{"citing_paper":"/paper/2411.16305"},"observation_digest":"sha256:54e79f833c93adf4313db38b35b8254fae78ae9a64c2086b254deb051675639f","observation_id":"4787ab63-517d-42cd-a98b-94170976ca8f","resolution":{"observed_at":"2026-08-12T13:21:29.893734Z","resolver_source":"doi","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-13T06:32:02.005865+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-13T06:32:02.005865+00:00","source":"crossref"},{"observed_at":"2026-08-13T06:31:53.387327+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-12T13:21:29.763702Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2411.16305","last_updated":"2024-11-25T11:47:31Z","snapshot_observed_at":"2026-08-12T13:13:26.422111Z","submitted_at":"2024-11-25T11:47:31Z","title":"Learning from Relevant Subgoals in Successful Dialogs using Iterative Training for Task-oriented Dialog Systems","version":1},"reference_index":16,"source":"arxiv_source","source_observed_at":"2026-08-12T13:21:29.763702Z"},"links":{"citing_paper":"/paper/2411.16305"},"observation_digest":"sha256:088bed98651d6fdc11c6cf6a5e214c86d3f0fd273326b77ca0c5011ce89416ab","observation_id":"f8bb8ee9-5596-4b00-a7ca-02a87ce3efc4","resolution":{"observed_at":"2026-08-12T13:21:29.763702Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-12T13:21:30.093073Z","title":null,"venue":null,"work_id":"85e09ff6-f892-4d50-b9ca-0deac3a0f95e","year":2022},"citing_paper":{"arxiv_id":"2411.16305","last_updated":"2024-11-25T11:47:31Z","snapshot_observed_at":"2026-08-12T13:13:26.422111Z","submitted_at":"2024-11-25T11:47:31Z","title":"Learning from Relevant Subgoals in Successful Dialogs using Iterative Training for Task-oriented Dialog Systems","version":1},"reference_index":17,"source":"arxiv_source","source_observed_at":"2026-08-12T13:21:29.767018Z"},"links":{"citing_paper":"/paper/2411.16305"},"observation_digest":"sha256:97f39fe2f30f70e5b394d46d0a20a67413c7858fb9931bb00acbeb6a45eb394c","observation_id":"2b4c96ad-641f-41e5-8ee7-329d9940e8ee","resolution":{"observed_at":"2026-08-12T13:21:30.096285Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-13T06:32:02.005865+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-13T06:32:02.005865+00:00","source":"crossref"},{"observed_at":"2026-08-13T06:31:53.387327+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-12T13:21:30.083960Z","title":null,"venue":null,"work_id":"689d7085-de68-4c23-8cfe-4d80867ab02c","year":2022},"citing_paper":{"arxiv_id":"2411.16305","last_updated":"2024-11-25T11:47:31Z","snapshot_observed_at":"2026-08-12T13:13:26.422111Z","submitted_at":"2024-11-25T11:47:31Z","title":"Learning from Relevant Subgoals in Successful Dialogs using Iterative Training for Task-oriented Dialog Systems","version":1},"reference_index":18,"source":"arxiv_source","source_observed_at":"2026-08-12T13:21:29.770322Z"},"links":{"citing_paper":"/paper/2411.16305"},"observation_digest":"sha256:e1fc3c91671a9a37b032e3dbfa8c07bda4aaaa849b0c69d4c40c949bd1458acb","observation_id":"3c5fc5d2-1b52-4598-ab79-2501be555e4c","resolution":{"observed_at":"2026-08-12T13:21:30.087226Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-13T06:32:02.005865+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-13T06:32:02.005865+00:00","source":"crossref"},{"observed_at":"2026-08-13T06:31:53.387327+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-12T13:21:29.773351Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2411.16305","last_updated":"2024-11-25T11:47:31Z","snapshot_observed_at":"2026-08-12T13:13:26.422111Z","submitted_at":"2024-11-25T11:47:31Z","title":"Learning from Relevant Subgoals in Successful Dialogs using Iterative Training for Task-oriented Dialog Systems","version":1},"reference_index":19,"source":"arxiv_source","source_observed_at":"2026-08-12T13:21:29.773351Z"},"links":{"citing_paper":"/paper/2411.16305"},"observation_digest":"sha256:3fe28e7fb20245a0339b67e966bf98da60b420edfb9696b890e3029b1d7aa9b6","observation_id":"14f000d8-c9f2-418c-9441-e58a0610b521","resolution":{"observed_at":"2026-08-12T13:21:29.773351Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-12T13:21:30.075457Z","title":null,"venue":null,"work_id":"892d4b08-2a87-4135-af2d-c91dea8c7853","year":2024},"citing_paper":{"arxiv_id":"2411.16305","last_updated":"2024-11-25T11:47:31Z","snapshot_observed_at":"2026-08-12T13:13:26.422111Z","submitted_at":"2024-11-25T11:47:31Z","title":"Learning from Relevant Subgoals in Successful Dialogs using Iterative Training for Task-oriented Dialog Systems","version":1},"reference_index":20,"source":"arxiv_source","source_observed_at":"2026-08-12T13:21:29.777098Z"},"links":{"citing_paper":"/paper/2411.16305"},"observation_digest":"sha256:34da41788626436a641304228adf6b785dd2795c0b76993aa51030bf9632ae37","observation_id":"18689bd7-5d58-47bb-a92e-76ee5f5ad2b9","resolution":{"observed_at":"2026-08-12T13:21:30.078402Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-13T06:32:02.005865+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-13T06:32:02.005865+00:00","source":"crossref"},{"observed_at":"2026-08-13T06:31:53.387327+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-12T13:21:30.065963Z","title":"Ziegler, Ryan Lowe, Chelsea Voss, Alec Radford, Dario Amodei, and Paul Christiano","venue":null,"work_id":"a84a28c0-a757-4ef4-92cb-6c6a1341a8f9","year":2020},"citing_paper":{"arxiv_id":"2411.16305","last_updated":"2024-11-25T11:47:31Z","snapshot_observed_at":"2026-08-12T13:13:26.422111Z","submitted_at":"2024-11-25T11:47:31Z","title":"Learning from Relevant Subgoals in Successful Dialogs using Iterative Training for Task-oriented Dialog Systems","version":1},"reference_index":21,"source":"arxiv_source","source_observed_at":"2026-08-12T13:21:29.780448Z"},"links":{"citing_paper":"/paper/2411.16305"},"observation_digest":"sha256:6e59abdf90eb9383c84d03870c47ee606db463f64ada5e99d3f715468842bbe6","observation_id":"147f44a1-524b-4f2e-a26b-1151440ccada","resolution":{"observed_at":"2026-08-12T13:21:30.070122Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-13T06:32:02.005865+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-13T06:32:02.005865+00:00","source":"crossref"},{"observed_at":"2026-08-13T06:31:53.387327+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-12T13:21:29.783528Z","title":null,"venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2411.16305","last_updated":"2024-11-25T11:47:31Z","snapshot_observed_at":"2026-08-12T13:13:26.422111Z","submitted_at":"2024-11-25T11:47:31Z","title":"Learning from Relevant Subgoals in Successful Dialogs using Iterative Training for Task-oriented Dialog Systems","version":1},"reference_index":22,"source":"arxiv_source","source_observed_at":"2026-08-12T13:21:29.783528Z"},"links":{"citing_paper":"/paper/2411.16305"},"observation_digest":"sha256:2637992025d0c4498aef44d74023ca14daf6ccf4e2644abb3e67e777c9918093","observation_id":"b64e833c-860b-46d4-acff-a6ac3725ec08","resolution":{"observed_at":"2026-08-12T13:21:29.783528Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-12T13:21:29.786820Z","title":null,"venue":null,"work_id":null,"year":2021},"citing_paper":{"arxiv_id":"2411.16305","last_updated":"2024-11-25T11:47:31Z","snapshot_observed_at":"2026-08-12T13:13:26.422111Z","submitted_at":"2024-11-25T11:47:31Z","title":"Learning from Relevant Subgoals in Successful Dialogs using Iterative Training for Task-oriented Dialog Systems","version":1},"reference_index":23,"source":"arxiv_source","source_observed_at":"2026-08-12T13:21:29.786820Z"},"links":{"citing_paper":"/paper/2411.16305"},"observation_digest":"sha256:c116edb98acd88b7f4236f067cec3fa288bcf261d0657cdf561de4c566d7b491","observation_id":"cb389fd4-7381-4313-aac2-5ae22a9c550b","resolution":{"observed_at":"2026-08-12T13:21:29.786820Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":"10.1609/aaai.v37i11.26602","metadata_source":"doi_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-12T13:21:29.864164Z","title":null,"venue":null,"work_id":"c79febdd-ba64-460f-bf20-72d826d7281e","year":2023},"citing_paper":{"arxiv_id":"2411.16305","last_updated":"2024-11-25T11:47:31Z","snapshot_observed_at":"2026-08-12T13:13:26.422111Z","submitted_at":"2024-11-25T11:47:31Z","title":"Learning from Relevant Subgoals in Successful Dialogs using Iterative Training for Task-oriented Dialog Systems","version":1},"reference_index":24,"source":"arxiv_source","source_observed_at":"2026-08-12T13:21:29.790007Z"},"links":{"citing_paper":"/paper/2411.16305"},"observation_digest":"sha256:c459ee47ae57b5dabaf5b36df4f27a9d54ce4c38db11df513268b7522c482cdd","observation_id":"58e00e15-7f05-4765-ac24-b6510ddb3bb5","resolution":{"observed_at":"2026-08-12T13:21:29.867782Z","resolver_source":"doi","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-13T06:32:02.005865+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-13T06:32:02.005865+00:00","source":"crossref"},{"observed_at":"2026-08-13T06:31:53.387327+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":"10.18653/v1/2023.sigdial-1.24","metadata_source":"doi_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-12T13:21:29.854124Z","title":null,"venue":null,"work_id":"49691209-d473-4320-ab63-8f36e17d0962","year":2023},"citing_paper":{"arxiv_id":"2411.16305","last_updated":"2024-11-25T11:47:31Z","snapshot_observed_at":"2026-08-12T13:13:26.422111Z","submitted_at":"2024-11-25T11:47:31Z","title":"Learning from Relevant Subgoals in Successful Dialogs using Iterative Training for Task-oriented Dialog Systems","version":1},"reference_index":25,"source":"arxiv_source","source_observed_at":"2026-08-12T13:21:29.792931Z"},"links":{"citing_paper":"/paper/2411.16305"},"observation_digest":"sha256:a70c288b3b20420fed4bbcee7e930f31d5e80e83b0d047d0e2f9251e720ff77f","observation_id":"581f2955-01b1-4b5a-9aef-201958dbdffe","resolution":{"observed_at":"2026-08-12T13:21:29.857513Z","resolver_source":"doi","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-13T06:32:02.005865+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-13T06:32:02.005865+00:00","source":"crossref"},{"observed_at":"2026-08-13T06:31:53.387327+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2312.16682","last_updated":"2024-04-22T22:51:32Z","snapshot_observed_at":"2026-08-13T04:53:15.268224Z","submitted_at":"2023-12-27T18:53:09Z","title":"Some things are more CRINGE than others: Iterative Preference Optimization with the Pairwise Cringe Loss","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2312.16682","snapshot_observed_at":"2026-08-12T13:21:29.795760Z","title":null,"venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2411.16305","last_updated":"2024-11-25T11:47:31Z","snapshot_observed_at":"2026-08-12T13:13:26.422111Z","submitted_at":"2024-11-25T11:47:31Z","title":"Learning from Relevant Subgoals in Successful Dialogs using Iterative Training for Task-oriented Dialog Systems","version":1},"reference_index":26,"source":"arxiv_source","source_observed_at":"2026-08-12T13:21:29.795760Z"},"links":{"cited_paper":"/paper/2312.16682","citing_paper":"/paper/2411.16305"},"observation_digest":"sha256:dddc07578a16c851f161ee86931eae1d6e0dd72638e6eff415c3783e58334df3","observation_id":"a1d66b31-f1f8-4beb-8e96-8bc4f7c167a0","resolution":{"observed_at":"2026-08-12T13:21:29.795760Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":"10.18653/v1/2023.emnlp-main.759","metadata_source":"doi_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-12T13:21:29.843210Z","title":null,"venue":null,"work_id":"448fec26-4a71-4298-9fa9-6e7578474ced","year":2023},"citing_paper":{"arxiv_id":"2411.16305","last_updated":"2024-11-25T11:47:31Z","snapshot_observed_at":"2026-08-12T13:13:26.422111Z","submitted_at":"2024-11-25T11:47:31Z","title":"Learning from Relevant Subgoals in Successful Dialogs using Iterative Training for Task-oriented Dialog Systems","version":1},"reference_index":27,"source":"arxiv_source","source_observed_at":"2026-08-12T13:21:29.799004Z"},"links":{"citing_paper":"/paper/2411.16305"},"observation_digest":"sha256:43836b6e1d8e062a1f4b27bf18e00dc781e1558516a3dd9389e4dda5d1454c25","observation_id":"d3e6ef04-a348-436d-927e-e1944adc3002","resolution":{"observed_at":"2026-08-12T13:21:29.847443Z","resolver_source":"doi","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-13T06:32:02.005865+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-13T06:32:02.005865+00:00","source":"crossref"},{"observed_at":"2026-08-13T06:31:53.387327+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-12T13:21:30.057407Z","title":null,"venue":null,"work_id":"6c2838fa-e1bd-4ea7-aad1-5b2a40f0ec12","year":2023},"citing_paper":{"arxiv_id":"2411.16305","last_updated":"2024-11-25T11:47:31Z","snapshot_observed_at":"2026-08-12T13:13:26.422111Z","submitted_at":"2024-11-25T11:47:31Z","title":"Learning from Relevant Subgoals in Successful Dialogs using Iterative Training for Task-oriented Dialog Systems","version":1},"reference_index":28,"source":"arxiv_source","source_observed_at":"2026-08-12T13:21:29.801980Z"},"links":{"citing_paper":"/paper/2411.16305"},"observation_digest":"sha256:1e47095e6b8477b2f9896d5687b7440c7bb45d1a6d4087a99bb8abc79919b5cd","observation_id":"6a00af7b-fc8f-449f-8ac2-b7bd86da7101","resolution":{"observed_at":"2026-08-12T13:21:30.060366Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-13T06:32:02.005865+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-13T06:32:02.005865+00:00","source":"crossref"},{"observed_at":"2026-08-13T06:31:53.387327+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-12T13:21:30.048601Z","title":null,"venue":null,"work_id":"3fd13a10-a0d3-4053-a44e-e361d59e6936","year":2020},"citing_paper":{"arxiv_id":"2411.16305","last_updated":"2024-11-25T11:47:31Z","snapshot_observed_at":"2026-08-12T13:13:26.422111Z","submitted_at":"2024-11-25T11:47:31Z","title":"Learning from Relevant Subgoals in Successful Dialogs using Iterative Training for Task-oriented Dialog Systems","version":1},"reference_index":29,"source":"arxiv_source","source_observed_at":"2026-08-12T13:21:29.805328Z"},"links":{"citing_paper":"/paper/2411.16305"},"observation_digest":"sha256:d2f3b3c7b6538cd046e551a9676bb07b1607786ee2370ea0838b51cf8d15edf4","observation_id":"73bae19c-4601-4a25-835b-ebb84839b14b","resolution":{"observed_at":"2026-08-12T13:21:30.051751Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-13T06:32:02.005865+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-13T06:32:02.005865+00:00","source":"crossref"},{"observed_at":"2026-08-13T06:31:53.387327+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2201.08904","last_updated":"2022-01-21T22:07:41Z","snapshot_observed_at":"2026-08-04T03:28:33.939819Z","submitted_at":"2022-01-21T22:07:41Z","title":"Description-Driven Task-Oriented Dialog Modeling","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2201.08904","snapshot_observed_at":"2026-08-12T13:21:29.808227Z","title":null,"venue":null,"work_id":null,"year":2022},"citing_paper":{"arxiv_id":"2411.16305","last_updated":"2024-11-25T11:47:31Z","snapshot_observed_at":"2026-08-12T13:13:26.422111Z","submitted_at":"2024-11-25T11:47:31Z","title":"Learning from Relevant Subgoals in Successful Dialogs using Iterative Training for Task-oriented Dialog Systems","version":1},"reference_index":30,"source":"arxiv_source","source_observed_at":"2026-08-12T13:21:29.808227Z"},"links":{"cited_paper":"/paper/2201.08904","citing_paper":"/paper/2411.16305"},"observation_digest":"sha256:3c04b652e6157d732c5483d3c0665beaaa30b1c925c8adb55194bbb016c86f47","observation_id":"88ed6aa7-b314-4e67-a241-a07720b594bf","resolution":{"observed_at":"2026-08-12T13:21:29.808227Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-12T13:21:29.811392Z","title":null,"venue":null,"work_id":null,"year":2019},"citing_paper":{"arxiv_id":"2411.16305","last_updated":"2024-11-25T11:47:31Z","snapshot_observed_at":"2026-08-12T13:13:26.422111Z","submitted_at":"2024-11-25T11:47:31Z","title":"Learning from Relevant Subgoals in Successful Dialogs using Iterative Training for Task-oriented Dialog Systems","version":1},"reference_index":31,"source":"arxiv_source","source_observed_at":"2026-08-12T13:21:29.811392Z"},"links":{"citing_paper":"/paper/2411.16305"},"observation_digest":"sha256:a9cac81b84bebddb21844cf85b06f2383aa902b02d90ee6d0fbdadc5e4a18771","observation_id":"5ba1ee13-5092-4b67-9b48-d6da995e0258","resolution":{"observed_at":"2026-08-12T13:21:29.811392Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}}],"paper":{"arxiv_id":"2411.16305","last_updated":"2024-11-25T11:47:31Z","latest_version":1,"primary_category":"cs.CL","snapshot_observed_at":"2026-08-12T13:13:26.422111Z","submitted_at":"2024-11-25T11:47:31Z","title":"Learning from Relevant Subgoals in Successful Dialogs using Iterative Training for Task-oriented Dialog Systems"},"reference_resolution":{"displayed":31,"state_counts":{"malformed_identifier":0,"metadata_mismatch":0,"parse_uncertain":0,"unresolved":24,"verified_exact":6,"verified_fuzzy":1},"total_outbound_references":31},"refusal":"A citation records a reference. It does not transfer a finding from one paper to another.","schema":"pith.paper-citation-record.v1","standing_sources":[{"observed_at":"2026-08-13T06:32:02.005865+00:00","source":"crossref"},{"observed_at":"2026-08-13T06:31:53.387327+00:00","source":"retraction_watch"}],"thesis":"As of 13 August 2026, this Paper Citation Record lists 31 of 31 outbound references and 0 inbound Pith citation observations for arXiv:2411.16305."}