{"as_of":"2026-08-08T01:02:00Z","caps":{"database_statements":6,"inbound":100,"outbound":100},"context_digest":"sha256:b0243a5597b75a98be552deb6a056433ca09b5b05387e9e9d8261c60de53c978","coverage":[{"denominator":71,"lane":"reference_resolution","note":"Typed states for the displayed outbound observations.","records_observed":71,"source":"paper_references, paper_reference_links","source_observed_at":"2026-08-02T07:20:39.150679Z","state":"measured"},{"denominator":71,"lane":"standing_notices","note":"One-hop event checks from named stored sources.","records_observed":71,"source":"scholarly_work_events, retraction_status_cache","source_observed_at":"2026-08-07T06:34:17.273281+00:00","state":"measured"},{"denominator":0,"lane":"inbound_itemization","note":"Pith citing papers itemized under the disclosed page cap.","records_observed":0,"source":"paper_references, paper_reference_links","source_observed_at":null,"state":"measured"},{"denominator":1,"lane":"external_citation_measurements","note":"A source-named dated measurement, never combined with another source.","records_observed":0,"source":"cited_works","source_observed_at":null,"state":"measured"}],"external_citation_measurements":[],"inbound":[],"links":{"evidence":"/evidence","html":"/paper/2607.10428/citation-record","integrity":"/paper/2607.10428/integrity","json":"/paper/2607.10428/citation-record.json","paper":"/paper/2607.10428"},"outbound":[{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-02T07:20:31.171652Z","title":"arXiv preprint arXiv:2510.13747 , year=","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2607.10428","last_updated":"2026-07-24T14:32:58Z","snapshot_observed_at":"2026-08-06T10:45:46.606756Z","submitted_at":"2026-07-11T18:22:49Z","title":"Enjoy Your Talk: A Human-Centered Benchmark for Multi-Turn Dialogue with Decoupled User Simulation, Target Modeling, and Judging","version":2},"reference_index":1,"source":"arxiv_source","source_observed_at":"2026-08-02T07:20:31.171652Z"},"links":{"citing_paper":"/paper/2607.10428"},"observation_digest":"sha256:3b8bc088a492f9c0371564fe412d5bb1b585c546e8a4ff32cfae546a0b8c00df","observation_id":"b29e8cce-661f-462a-af3e-f8f492285ade","resolution":{"observed_at":"2026-08-02T07:20:31.171652Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2606.25325","last_updated":"2026-07-20T13:20:39Z","snapshot_observed_at":"2026-08-02T10:17:41.360200Z","submitted_at":"2026-06-24T02:43:26Z","title":"Omni-Perception Policy Optimization for Multimodal Emotion Reasoning","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2606.25325","snapshot_observed_at":"2026-08-02T07:20:31.224451Z","title":"arXiv preprint arXiv:2606.25325 , year=","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2607.10428","last_updated":"2026-07-24T14:32:58Z","snapshot_observed_at":"2026-08-06T10:45:46.606756Z","submitted_at":"2026-07-11T18:22:49Z","title":"Enjoy Your Talk: A Human-Centered Benchmark for Multi-Turn Dialogue with Decoupled User Simulation, Target Modeling, and Judging","version":2},"reference_index":2,"source":"arxiv_source","source_observed_at":"2026-08-02T07:20:31.224451Z"},"links":{"cited_paper":"/paper/2606.25325","citing_paper":"/paper/2607.10428"},"observation_digest":"sha256:704b3c5405ba734d41829c1df8d47e242741a01c565b6d400922c4a7a67b93cc","observation_id":"55c21b7c-3287-4720-8d4f-bb96357ec4d5","resolution":{"observed_at":"2026-08-02T07:20:31.224451Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2606.27652","last_updated":"2026-06-26T02:07:53Z","snapshot_observed_at":"2026-08-06T10:45:51.342747Z","submitted_at":"2026-06-26T02:07:53Z","title":"MER-R1: Multimodal Emotion Reasoning via Slow-Fast Thinking Synergy","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2606.27652","snapshot_observed_at":"2026-08-02T07:20:31.305011Z","title":"arXiv preprint arXiv:2606.27652 , year=","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2607.10428","last_updated":"2026-07-24T14:32:58Z","snapshot_observed_at":"2026-08-06T10:45:46.606756Z","submitted_at":"2026-07-11T18:22:49Z","title":"Enjoy Your Talk: A Human-Centered Benchmark for Multi-Turn Dialogue with Decoupled User Simulation, Target Modeling, and Judging","version":2},"reference_index":3,"source":"arxiv_source","source_observed_at":"2026-08-02T07:20:31.305011Z"},"links":{"cited_paper":"/paper/2606.27652","citing_paper":"/paper/2607.10428"},"observation_digest":"sha256:616c9a09ea2e2caf558e8d5397edca3ea9560764632c4e843c0a431b39941fe8","observation_id":"2bbb0d5f-756b-4635-9656-5a9a272bd2bb","resolution":{"observed_at":"2026-08-02T07:20:31.305011Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2402.14762","last_updated":"2024-11-05T16:40:21Z","snapshot_observed_at":"2026-07-06T17:34:07.737296Z","submitted_at":"2024-02-22T18:21:59Z","title":"MT-Bench-101: A Fine-Grained Benchmark for Evaluating Large Language Models in Multi-Turn Dialogues","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2402.14762","snapshot_observed_at":"2026-08-02T07:20:31.426325Z","title":"arXiv preprint arXiv:2402.14762 , year=","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2607.10428","last_updated":"2026-07-24T14:32:58Z","snapshot_observed_at":"2026-08-06T10:45:46.606756Z","submitted_at":"2026-07-11T18:22:49Z","title":"Enjoy Your Talk: A Human-Centered Benchmark for Multi-Turn Dialogue with Decoupled User Simulation, Target Modeling, and Judging","version":2},"reference_index":4,"source":"arxiv_source","source_observed_at":"2026-08-02T07:20:31.426325Z"},"links":{"cited_paper":"/paper/2402.14762","citing_paper":"/paper/2607.10428"},"observation_digest":"sha256:1fe26fa01b8806503149b4b44a325a0030c0519761e44e0f8f68bc792b82b16b","observation_id":"857cc624-b313-4e75-831d-b45d1c04ed67","resolution":{"observed_at":"2026-08-02T07:20:31.426325Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-02T07:20:31.549056Z","title":"arXiv preprint arXiv:2511.00850 , year=","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2607.10428","last_updated":"2026-07-24T14:32:58Z","snapshot_observed_at":"2026-08-06T10:45:46.606756Z","submitted_at":"2026-07-11T18:22:49Z","title":"Enjoy Your Talk: A Human-Centered Benchmark for Multi-Turn Dialogue with Decoupled User Simulation, Target Modeling, and Judging","version":2},"reference_index":5,"source":"arxiv_source","source_observed_at":"2026-08-02T07:20:31.549056Z"},"links":{"citing_paper":"/paper/2607.10428"},"observation_digest":"sha256:8d49cc3b0c874406ea2d3e7951a324db11af401ee068746d952d8870054ea508","observation_id":"c8af4768-d292-42a5-b009-89c6070e8b7c","resolution":{"observed_at":"2026-08-02T07:20:31.549056Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-02T07:20:31.681572Z","title":"arXiv preprint arXiv:2603.00552 , year=","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2607.10428","last_updated":"2026-07-24T14:32:58Z","snapshot_observed_at":"2026-08-06T10:45:46.606756Z","submitted_at":"2026-07-11T18:22:49Z","title":"Enjoy Your Talk: A Human-Centered Benchmark for Multi-Turn Dialogue with Decoupled User Simulation, Target Modeling, and Judging","version":2},"reference_index":6,"source":"arxiv_source","source_observed_at":"2026-08-02T07:20:31.681572Z"},"links":{"citing_paper":"/paper/2607.10428"},"observation_digest":"sha256:7a0d97892a37234097eef5ee17315e68ed0dcf62b7f74eb0c4f103baadd949b6","observation_id":"92730466-ebf1-4529-af61-7a7da7b8e32f","resolution":{"observed_at":"2026-08-02T07:20:31.681572Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-02T07:20:31.829375Z","title":"Findings of the Association for Computational Linguistics: ACL 2025 , pages=","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2607.10428","last_updated":"2026-07-24T14:32:58Z","snapshot_observed_at":"2026-08-06T10:45:46.606756Z","submitted_at":"2026-07-11T18:22:49Z","title":"Enjoy Your Talk: A Human-Centered Benchmark for Multi-Turn Dialogue with Decoupled User Simulation, Target Modeling, and Judging","version":2},"reference_index":7,"source":"arxiv_source","source_observed_at":"2026-08-02T07:20:31.829375Z"},"links":{"citing_paper":"/paper/2607.10428"},"observation_digest":"sha256:886a621cc807b6f924a7e8dff59cde4c3c1f8ddae9895b025c3e4cf313d3dacc","observation_id":"c6aa9472-74c9-45b1-8eb2-53dd7d828ff8","resolution":{"observed_at":"2026-08-02T07:20:31.829375Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-02T07:20:31.994581Z","title":"Proceedings of the AAAI Conference on Artificial Intelligence , volume=","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2607.10428","last_updated":"2026-07-24T14:32:58Z","snapshot_observed_at":"2026-08-06T10:45:46.606756Z","submitted_at":"2026-07-11T18:22:49Z","title":"Enjoy Your Talk: A Human-Centered Benchmark for Multi-Turn Dialogue with Decoupled User Simulation, Target Modeling, and Judging","version":2},"reference_index":8,"source":"arxiv_source","source_observed_at":"2026-08-02T07:20:31.994581Z"},"links":{"citing_paper":"/paper/2607.10428"},"observation_digest":"sha256:e7358ca2f562ba705e2697ff5372e6d3cc53b83a1334d6c34d06fc6c3f0a5742","observation_id":"c6bfee14-8883-4fd6-a50f-c92e2aa18549","resolution":{"observed_at":"2026-08-02T07:20:31.994581Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-02T07:20:32.113588Z","title":"Proceedings of the 62nd Annual Meeting of the Association for Computational Linguistics (Volume 1: Long Papers) , pages=","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2607.10428","last_updated":"2026-07-24T14:32:58Z","snapshot_observed_at":"2026-08-06T10:45:46.606756Z","submitted_at":"2026-07-11T18:22:49Z","title":"Enjoy Your Talk: A Human-Centered Benchmark for Multi-Turn Dialogue with Decoupled User Simulation, Target Modeling, and Judging","version":2},"reference_index":9,"source":"arxiv_source","source_observed_at":"2026-08-02T07:20:32.113588Z"},"links":{"citing_paper":"/paper/2607.10428"},"observation_digest":"sha256:fc4727ddce67dadff81ef4c62f8da4fe5bef18f2d27e0feea8ea04dff1a26d35","observation_id":"fb6be6b9-a78a-4ee3-acdb-e8d3a16e959f","resolution":{"observed_at":"2026-08-02T07:20:32.113588Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-02T07:20:32.313483Z","title":"arXiv preprint arXiv:2510.23182 , year=","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2607.10428","last_updated":"2026-07-24T14:32:58Z","snapshot_observed_at":"2026-08-06T10:45:46.606756Z","submitted_at":"2026-07-11T18:22:49Z","title":"Enjoy Your Talk: A Human-Centered Benchmark for Multi-Turn Dialogue with Decoupled User Simulation, Target Modeling, and Judging","version":2},"reference_index":10,"source":"arxiv_source","source_observed_at":"2026-08-02T07:20:32.313483Z"},"links":{"citing_paper":"/paper/2607.10428"},"observation_digest":"sha256:4cbb17cbb6786691b3d842e6bb28b4f9cf0eb877fa389772f1485062ffd18200","observation_id":"b69f21ba-19bf-4603-9aaf-a436b6410d76","resolution":{"observed_at":"2026-08-02T07:20:32.313483Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-02T07:20:32.450541Z","title":"arXiv preprint arXiv:2509.21856 , year=","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2607.10428","last_updated":"2026-07-24T14:32:58Z","snapshot_observed_at":"2026-08-06T10:45:46.606756Z","submitted_at":"2026-07-11T18:22:49Z","title":"Enjoy Your Talk: A Human-Centered Benchmark for Multi-Turn Dialogue with Decoupled User Simulation, Target Modeling, and Judging","version":2},"reference_index":11,"source":"arxiv_source","source_observed_at":"2026-08-02T07:20:32.450541Z"},"links":{"citing_paper":"/paper/2607.10428"},"observation_digest":"sha256:dbcce44e154c7f5436fadd7b0c55f339bc6bf95d49d3f11545d2ac366fad5de6","observation_id":"8ccd89c7-a12e-4f11-beff-195d2292fdbc","resolution":{"observed_at":"2026-08-02T07:20:32.450541Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2406.12045","last_updated":"2024-06-17T19:33:08Z","snapshot_observed_at":"2026-08-02T22:19:29.043854Z","submitted_at":"2024-06-17T19:33:08Z","title":"$\\tau$-bench: A Benchmark for Tool-Agent-User Interaction in Real-World Domains","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2406.12045","snapshot_observed_at":"2026-08-02T07:20:32.591228Z","title":"arXiv preprint arXiv:2406.12045 , year=","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2607.10428","last_updated":"2026-07-24T14:32:58Z","snapshot_observed_at":"2026-08-06T10:45:46.606756Z","submitted_at":"2026-07-11T18:22:49Z","title":"Enjoy Your Talk: A Human-Centered Benchmark for Multi-Turn Dialogue with Decoupled User Simulation, Target Modeling, and Judging","version":2},"reference_index":12,"source":"arxiv_source","source_observed_at":"2026-08-02T07:20:32.591228Z"},"links":{"cited_paper":"/paper/2406.12045","citing_paper":"/paper/2607.10428"},"observation_digest":"sha256:908151a3d980bef581caf8d550dca954a09707f1c612624f16daa5554b1429e2","observation_id":"193e0292-12ca-437f-92ca-931e570d47fa","resolution":{"observed_at":"2026-08-02T07:20:32.591228Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-02T07:20:32.762133Z","title":"Advances in Neural Information Processing Systems (NeurIPS) , volume=","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2607.10428","last_updated":"2026-07-24T14:32:58Z","snapshot_observed_at":"2026-08-06T10:45:46.606756Z","submitted_at":"2026-07-11T18:22:49Z","title":"Enjoy Your Talk: A Human-Centered Benchmark for Multi-Turn Dialogue with Decoupled User Simulation, Target Modeling, and Judging","version":2},"reference_index":13,"source":"arxiv_source","source_observed_at":"2026-08-02T07:20:32.762133Z"},"links":{"citing_paper":"/paper/2607.10428"},"observation_digest":"sha256:15bb3f8a2015c905e532c6a2eec484ec8a0c5793867e2a9192845c6fea3d0a3e","observation_id":"2b247a99-842e-4672-afeb-dfb59f1c50a5","resolution":{"observed_at":"2026-08-02T07:20:32.762133Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-02T07:20:32.856823Z","title":"Advances in Neural Information Processing Systems (NeurIPS) , year=","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2607.10428","last_updated":"2026-07-24T14:32:58Z","snapshot_observed_at":"2026-08-06T10:45:46.606756Z","submitted_at":"2026-07-11T18:22:49Z","title":"Enjoy Your Talk: A Human-Centered Benchmark for Multi-Turn Dialogue with Decoupled User Simulation, Target Modeling, and Judging","version":2},"reference_index":14,"source":"arxiv_source","source_observed_at":"2026-08-02T07:20:32.856823Z"},"links":{"citing_paper":"/paper/2607.10428"},"observation_digest":"sha256:a37d0f2e8fac6cbadd6d206a5b3a1137effd4b60e8d9600a80571cc0010493dd","observation_id":"3a7eea97-6938-4f68-8ac2-c5ad4ea49af8","resolution":{"observed_at":"2026-08-02T07:20:32.856823Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-02T07:20:32.991307Z","title":"arXiv preprint arXiv:2509.21117 , year=","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2607.10428","last_updated":"2026-07-24T14:32:58Z","snapshot_observed_at":"2026-08-06T10:45:46.606756Z","submitted_at":"2026-07-11T18:22:49Z","title":"Enjoy Your Talk: A Human-Centered Benchmark for Multi-Turn Dialogue with Decoupled User Simulation, Target Modeling, and Judging","version":2},"reference_index":15,"source":"arxiv_source","source_observed_at":"2026-08-02T07:20:32.991307Z"},"links":{"citing_paper":"/paper/2607.10428"},"observation_digest":"sha256:dae00566f328830a45352f077d9dd3e854337a0248869f0b1127721c519a6aaa","observation_id":"e3648069-e6b2-47f4-b737-9d0379010f98","resolution":{"observed_at":"2026-08-02T07:20:32.991307Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2305.17926","last_updated":"2023-08-30T13:22:35Z","snapshot_observed_at":"2026-08-03T19:35:11.838629Z","submitted_at":"2023-05-29T07:41:03Z","title":"Large Language Models are not Fair Evaluators","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2305.17926","snapshot_observed_at":"2026-08-02T07:20:33.074919Z","title":"arXiv preprint arXiv:2305.17926 , year=","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2607.10428","last_updated":"2026-07-24T14:32:58Z","snapshot_observed_at":"2026-08-06T10:45:46.606756Z","submitted_at":"2026-07-11T18:22:49Z","title":"Enjoy Your Talk: A Human-Centered Benchmark for Multi-Turn Dialogue with Decoupled User Simulation, Target Modeling, and Judging","version":2},"reference_index":16,"source":"arxiv_source","source_observed_at":"2026-08-02T07:20:33.074919Z"},"links":{"cited_paper":"/paper/2305.17926","citing_paper":"/paper/2607.10428"},"observation_digest":"sha256:efe505bf6b8860da5a492d62380f04956d2c881a9bb4608d6b6e6e256ead9d1e","observation_id":"6ca12b44-40ab-48e5-8247-3dcd0e92baec","resolution":{"observed_at":"2026-08-02T07:20:33.074919Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-02T07:20:33.212966Z","title":"Proceedings of the 2024 Conference of the North American Chapter of the Association for Computational Linguistics: Human Language Technologies (Volume 1: Long Papers) , pages=","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2607.10428","last_updated":"2026-07-24T14:32:58Z","snapshot_observed_at":"2026-08-06T10:45:46.606756Z","submitted_at":"2026-07-11T18:22:49Z","title":"Enjoy Your Talk: A Human-Centered Benchmark for Multi-Turn Dialogue with Decoupled User Simulation, Target Modeling, and Judging","version":2},"reference_index":17,"source":"arxiv_source","source_observed_at":"2026-08-02T07:20:33.212966Z"},"links":{"citing_paper":"/paper/2607.10428"},"observation_digest":"sha256:0a30f4ed268f0ba947bcdeee987fec6eddf0592493822997d78cdac9468da8f1","observation_id":"bb18b7c2-b37c-4176-b33c-789872b5274b","resolution":{"observed_at":"2026-08-02T07:20:33.212966Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2411.15594","last_updated":"2025-10-19T10:32:43Z","snapshot_observed_at":"2026-08-02T10:23:50.881300Z","submitted_at":"2024-11-23T16:03:35Z","title":"A Survey on LLM-as-a-Judge","version":6},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2411.15594","snapshot_observed_at":"2026-08-02T07:20:33.355853Z","title":"arXiv preprint arXiv:2411.15594 , year=","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2607.10428","last_updated":"2026-07-24T14:32:58Z","snapshot_observed_at":"2026-08-06T10:45:46.606756Z","submitted_at":"2026-07-11T18:22:49Z","title":"Enjoy Your Talk: A Human-Centered Benchmark for Multi-Turn Dialogue with Decoupled User Simulation, Target Modeling, and Judging","version":2},"reference_index":18,"source":"arxiv_source","source_observed_at":"2026-08-02T07:20:33.355853Z"},"links":{"cited_paper":"/paper/2411.15594","citing_paper":"/paper/2607.10428"},"observation_digest":"sha256:eb8b28fe288ff4b10cebb0281611d07ffee25af1652c53983d0f8ab2674829c1","observation_id":"48eda901-299d-4113-8600-2077c593d038","resolution":{"observed_at":"2026-08-02T07:20:33.355853Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-02T07:20:33.487879Z","title":"Proceedings of the AAAI Conference on Artificial Intelligence , volume=","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2607.10428","last_updated":"2026-07-24T14:32:58Z","snapshot_observed_at":"2026-08-06T10:45:46.606756Z","submitted_at":"2026-07-11T18:22:49Z","title":"Enjoy Your Talk: A Human-Centered Benchmark for Multi-Turn Dialogue with Decoupled User Simulation, Target Modeling, and Judging","version":2},"reference_index":19,"source":"arxiv_source","source_observed_at":"2026-08-02T07:20:33.487879Z"},"links":{"citing_paper":"/paper/2607.10428"},"observation_digest":"sha256:419b5603d49ee2eafb3455e7d82ac6fbcddb6c7820098962ca70537b8cc1e940","observation_id":"78bde588-dcc0-4e05-bf39-67906b362d85","resolution":{"observed_at":"2026-08-02T07:20:33.487879Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2506.00777","last_updated":"2025-06-01T02:01:52Z","snapshot_observed_at":"2026-08-07T11:55:52.371442Z","submitted_at":"2025-06-01T02:01:52Z","title":"Improving Automatic Evaluation of Large Language Models (LLMs) in Biomedical Relation Extraction via LLMs-as-the-Judge","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2506.00777","snapshot_observed_at":"2026-08-02T07:20:33.617934Z","title":"arXiv preprint arXiv:2506.00777 , year=","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2607.10428","last_updated":"2026-07-24T14:32:58Z","snapshot_observed_at":"2026-08-06T10:45:46.606756Z","submitted_at":"2026-07-11T18:22:49Z","title":"Enjoy Your Talk: A Human-Centered Benchmark for Multi-Turn Dialogue with Decoupled User Simulation, Target Modeling, and Judging","version":2},"reference_index":20,"source":"arxiv_source","source_observed_at":"2026-08-02T07:20:33.617934Z"},"links":{"cited_paper":"/paper/2506.00777","citing_paper":"/paper/2607.10428"},"observation_digest":"sha256:e91792461f8ff8b4b33b7cff5d5a28e73fddaf8f3d69e4cfc1bbd76274aef789","observation_id":"3d01d9e4-23ce-496d-8e1e-f52f7e92bc81","resolution":{"observed_at":"2026-08-02T07:20:33.617934Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-02T07:20:33.796526Z","title":"Proceedings of the 63rd Annual Meeting of the Association for Computational Linguistics (Volume 1: Long Papers) , pages=","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2607.10428","last_updated":"2026-07-24T14:32:58Z","snapshot_observed_at":"2026-08-06T10:45:46.606756Z","submitted_at":"2026-07-11T18:22:49Z","title":"Enjoy Your Talk: A Human-Centered Benchmark for Multi-Turn Dialogue with Decoupled User Simulation, Target Modeling, and Judging","version":2},"reference_index":21,"source":"arxiv_source","source_observed_at":"2026-08-02T07:20:33.796526Z"},"links":{"citing_paper":"/paper/2607.10428"},"observation_digest":"sha256:d01fcee54125eb64422b6ce68e527dd5489fe0b4e15c35903884d0070e13bf83","observation_id":"97f87e4d-fa84-486d-a966-4a3639feace5","resolution":{"observed_at":"2026-08-02T07:20:33.796526Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-02T07:20:33.944041Z","title":"arXiv preprint arXiv:2510.12462 , year=","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2607.10428","last_updated":"2026-07-24T14:32:58Z","snapshot_observed_at":"2026-08-06T10:45:46.606756Z","submitted_at":"2026-07-11T18:22:49Z","title":"Enjoy Your Talk: A Human-Centered Benchmark for Multi-Turn Dialogue with Decoupled User Simulation, Target Modeling, and Judging","version":2},"reference_index":22,"source":"arxiv_source","source_observed_at":"2026-08-02T07:20:33.944041Z"},"links":{"citing_paper":"/paper/2607.10428"},"observation_digest":"sha256:a77bc7a50a590561589223bf1956b607882cb40d2461612c09b2e4120eaedd5e","observation_id":"96f8aa3b-3c6d-4355-a0f8-455a9d7d07b9","resolution":{"observed_at":"2026-08-02T07:20:33.944041Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-02T07:20:34.123775Z","title":"ACM Transactions on Information Systems (TOIS) , volume=","venue":null,"work_id":null,"year":2020},"citing_paper":{"arxiv_id":"2607.10428","last_updated":"2026-07-24T14:32:58Z","snapshot_observed_at":"2026-08-06T10:45:46.606756Z","submitted_at":"2026-07-11T18:22:49Z","title":"Enjoy Your Talk: A Human-Centered Benchmark for Multi-Turn Dialogue with Decoupled User Simulation, Target Modeling, and Judging","version":2},"reference_index":23,"source":"arxiv_source","source_observed_at":"2026-08-02T07:20:34.123775Z"},"links":{"citing_paper":"/paper/2607.10428"},"observation_digest":"sha256:81a37e025d7649f4b9fa8af66aab09f32b4708b7aa97a41817070a8408474445","observation_id":"5f9a78fb-2292-4bff-99e3-1e94a5542168","resolution":{"observed_at":"2026-08-02T07:20:34.123775Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2107.11904","last_updated":"2021-07-25T22:59:09Z","snapshot_observed_at":"2026-08-06T04:48:24.227713Z","submitted_at":"2021-07-25T22:59:09Z","title":"Transferable Dialogue Systems and User Simulators","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2107.11904","snapshot_observed_at":"2026-08-02T07:20:34.257535Z","title":"arXiv preprint arXiv:2107.11904 , year=","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2607.10428","last_updated":"2026-07-24T14:32:58Z","snapshot_observed_at":"2026-08-06T10:45:46.606756Z","submitted_at":"2026-07-11T18:22:49Z","title":"Enjoy Your Talk: A Human-Centered Benchmark for Multi-Turn Dialogue with Decoupled User Simulation, Target Modeling, and Judging","version":2},"reference_index":24,"source":"arxiv_source","source_observed_at":"2026-08-02T07:20:34.257535Z"},"links":{"cited_paper":"/paper/2107.11904","citing_paper":"/paper/2607.10428"},"observation_digest":"sha256:28d487653645c211288bdebebccc7dd0606327bcc0133a0e81a7563587812a27","observation_id":"773d4df3-b1f0-4e91-90e0-a6a590fd7ab2","resolution":{"observed_at":"2026-08-02T07:20:34.257535Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2505.06120","last_updated":"2025-05-09T15:21:44Z","snapshot_observed_at":"2026-08-02T13:06:07.304813Z","submitted_at":"2025-05-09T15:21:44Z","title":"LLMs Get Lost In Multi-Turn Conversation","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2505.06120","snapshot_observed_at":"2026-08-02T07:20:34.350919Z","title":"arXiv preprint arXiv:2505.06120 , year=","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2607.10428","last_updated":"2026-07-24T14:32:58Z","snapshot_observed_at":"2026-08-06T10:45:46.606756Z","submitted_at":"2026-07-11T18:22:49Z","title":"Enjoy Your Talk: A Human-Centered Benchmark for Multi-Turn Dialogue with Decoupled User Simulation, Target Modeling, and Judging","version":2},"reference_index":25,"source":"arxiv_source","source_observed_at":"2026-08-02T07:20:34.350919Z"},"links":{"cited_paper":"/paper/2505.06120","citing_paper":"/paper/2607.10428"},"observation_digest":"sha256:273eef6cdbcaab876feb36a823bd48a591d1ef7bc33ea70a01c5912b3436be2a","observation_id":"9d39619e-0dc6-40dc-93e8-6bdcdb8bdc43","resolution":{"observed_at":"2026-08-02T07:20:34.350919Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-02T07:20:34.433695Z","title":"Proceedings of the 25th ACM Conference on Economics and Computation (EC '24) , pages=","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2607.10428","last_updated":"2026-07-24T14:32:58Z","snapshot_observed_at":"2026-08-06T10:45:46.606756Z","submitted_at":"2026-07-11T18:22:49Z","title":"Enjoy Your Talk: A Human-Centered Benchmark for Multi-Turn Dialogue with Decoupled User Simulation, Target Modeling, and Judging","version":2},"reference_index":26,"source":"arxiv_source","source_observed_at":"2026-08-02T07:20:34.433695Z"},"links":{"citing_paper":"/paper/2607.10428"},"observation_digest":"sha256:292b540e47d4638595c708ee7e3722686783cfe0796fb9a6e573715b46d2b233","observation_id":"e336d869-81ab-4989-a4ed-10da04f7647c","resolution":{"observed_at":"2026-08-02T07:20:34.433695Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2502.00640","last_updated":"2025-07-29T22:56:43Z","snapshot_observed_at":"2026-07-06T20:29:40.670401Z","submitted_at":"2025-02-02T03:05:52Z","title":"CollabLLM: From Passive Responders to Active Collaborators","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2502.00640","snapshot_observed_at":"2026-08-02T07:20:34.516344Z","title":"arXiv preprint arXiv:2502.00640 , year=","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2607.10428","last_updated":"2026-07-24T14:32:58Z","snapshot_observed_at":"2026-08-06T10:45:46.606756Z","submitted_at":"2026-07-11T18:22:49Z","title":"Enjoy Your Talk: A Human-Centered Benchmark for Multi-Turn Dialogue with Decoupled User Simulation, Target Modeling, and Judging","version":2},"reference_index":27,"source":"arxiv_source","source_observed_at":"2026-08-02T07:20:34.516344Z"},"links":{"cited_paper":"/paper/2502.00640","citing_paper":"/paper/2607.10428"},"observation_digest":"sha256:2d1e2332be62974e1d2dc15ac7f63d225726288b3727d810d382bab075e7a4a1","observation_id":"46bb3c12-ea9f-4189-b737-06dafe98f2c4","resolution":{"observed_at":"2026-08-02T07:20:34.516344Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2504.07114","last_updated":"2025-08-12T01:53:40Z","snapshot_observed_at":"2026-08-07T16:44:44.128963Z","submitted_at":"2025-03-22T01:21:40Z","title":"ChatBench: From Static Benchmarks to Human-AI Evaluation","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2504.07114","snapshot_observed_at":"2026-08-02T07:20:34.591235Z","title":"arXiv preprint arXiv:2504.07114 , year=","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2607.10428","last_updated":"2026-07-24T14:32:58Z","snapshot_observed_at":"2026-08-06T10:45:46.606756Z","submitted_at":"2026-07-11T18:22:49Z","title":"Enjoy Your Talk: A Human-Centered Benchmark for Multi-Turn Dialogue with Decoupled User Simulation, Target Modeling, and Judging","version":2},"reference_index":28,"source":"arxiv_source","source_observed_at":"2026-08-02T07:20:34.591235Z"},"links":{"cited_paper":"/paper/2504.07114","citing_paper":"/paper/2607.10428"},"observation_digest":"sha256:c9d10d9cd2edbb8008db21a1ca24bf980f18261461350f3e2a9048306d444a99","observation_id":"52ea65f1-f84f-45ff-b0fc-ca81773892be","resolution":{"observed_at":"2026-08-02T07:20:34.591235Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2502.16761","last_updated":"2026-04-16T08:22:39Z","snapshot_observed_at":"2026-07-06T20:41:23.980642Z","submitted_at":"2025-02-24T00:31:33Z","title":"Language Model Fine-Tuning on Scaled Survey Data for Predicting Distributions of Public Opinions","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2502.16761","snapshot_observed_at":"2026-08-02T07:20:34.652851Z","title":"arXiv preprint arXiv:2502.16761 , year=","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2607.10428","last_updated":"2026-07-24T14:32:58Z","snapshot_observed_at":"2026-08-06T10:45:46.606756Z","submitted_at":"2026-07-11T18:22:49Z","title":"Enjoy Your Talk: A Human-Centered Benchmark for Multi-Turn Dialogue with Decoupled User Simulation, Target Modeling, and Judging","version":2},"reference_index":29,"source":"arxiv_source","source_observed_at":"2026-08-02T07:20:34.652851Z"},"links":{"cited_paper":"/paper/2502.16761","citing_paper":"/paper/2607.10428"},"observation_digest":"sha256:ef84dc1f36d9299d35f68b7b5d7963adbd2431a878bbb8ff7e296dd25b1eae97","observation_id":"33c16312-11b5-470a-b899-fe6815b41459","resolution":{"observed_at":"2026-08-02T07:20:34.652851Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2505.21362","last_updated":"2025-05-27T15:52:39Z","snapshot_observed_at":"2026-08-07T16:29:56.189099Z","submitted_at":"2025-05-27T15:52:39Z","title":"Evaluating LLM Adaptation to Sociodemographic Factors: User Profile vs. Dialogue History","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2505.21362","snapshot_observed_at":"2026-08-02T07:20:34.719196Z","title":"Dialogue History , author=","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2607.10428","last_updated":"2026-07-24T14:32:58Z","snapshot_observed_at":"2026-08-06T10:45:46.606756Z","submitted_at":"2026-07-11T18:22:49Z","title":"Enjoy Your Talk: A Human-Centered Benchmark for Multi-Turn Dialogue with Decoupled User Simulation, Target Modeling, and Judging","version":2},"reference_index":30,"source":"arxiv_source","source_observed_at":"2026-08-02T07:20:34.719196Z"},"links":{"cited_paper":"/paper/2505.21362","citing_paper":"/paper/2607.10428"},"observation_digest":"sha256:863162715682b42e5f4d5e51ede3e3033133c95c545bd87deec86e63391038cd","observation_id":"9ea9c0ed-ac30-499e-a580-606eabe772ff","resolution":{"observed_at":"2026-08-02T07:20:34.719196Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-02T07:20:34.795329Z","title":"Nature Machine Intelligence , pages=","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2607.10428","last_updated":"2026-07-24T14:32:58Z","snapshot_observed_at":"2026-08-06T10:45:46.606756Z","submitted_at":"2026-07-11T18:22:49Z","title":"Enjoy Your Talk: A Human-Centered Benchmark for Multi-Turn Dialogue with Decoupled User Simulation, Target Modeling, and Judging","version":2},"reference_index":31,"source":"arxiv_source","source_observed_at":"2026-08-02T07:20:34.795329Z"},"links":{"citing_paper":"/paper/2607.10428"},"observation_digest":"sha256:4c03cf43f264b8d043be2baae2ddcd990e2bcd7f361fb90208933d1e574a3265","observation_id":"ef3dc94d-cd8f-4394-b489-3f820e249926","resolution":{"observed_at":"2026-08-02T07:20:34.795329Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-02T07:20:34.853612Z","title":"Human Language Technologies 2007: The Conference of the North American Chapter of the Association for Computational Linguistics; Companion Volume, Short Papers , pages=","venue":null,"work_id":null,"year":2007},"citing_paper":{"arxiv_id":"2607.10428","last_updated":"2026-07-24T14:32:58Z","snapshot_observed_at":"2026-08-06T10:45:46.606756Z","submitted_at":"2026-07-11T18:22:49Z","title":"Enjoy Your Talk: A Human-Centered Benchmark for Multi-Turn Dialogue with Decoupled User Simulation, Target Modeling, and Judging","version":2},"reference_index":32,"source":"arxiv_source","source_observed_at":"2026-08-02T07:20:34.853612Z"},"links":{"citing_paper":"/paper/2607.10428"},"observation_digest":"sha256:63a54154276f3bdada179443a2dbf7c3b6a8972baa5740a15dc20fee43976cf2","observation_id":"7ddeefaa-e93a-4549-bd84-d689eeed2a2e","resolution":{"observed_at":"2026-08-02T07:20:34.853612Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"1612.05688","last_updated":"2017-11-13T05:52:42Z","snapshot_observed_at":"2026-08-02T11:32:13.136863Z","submitted_at":"2016-12-17T01:03:55Z","title":"A User Simulator for Task-Completion Dialogues","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"1612.05688","snapshot_observed_at":"2026-08-02T07:20:34.914259Z","title":"arXiv preprint arXiv:1612.05688 , year=","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2607.10428","last_updated":"2026-07-24T14:32:58Z","snapshot_observed_at":"2026-08-06T10:45:46.606756Z","submitted_at":"2026-07-11T18:22:49Z","title":"Enjoy Your Talk: A Human-Centered Benchmark for Multi-Turn Dialogue with Decoupled User Simulation, Target Modeling, and Judging","version":2},"reference_index":33,"source":"arxiv_source","source_observed_at":"2026-08-02T07:20:34.914259Z"},"links":{"cited_paper":"/paper/1612.05688","citing_paper":"/paper/2607.10428"},"observation_digest":"sha256:d61ce72b56c658716d82f72bdadf32f3c2dea8558327c5af98f8098647d4324c","observation_id":"a8955479-e8c9-4d7a-b1f8-8b397a3bd38d","resolution":{"observed_at":"2026-08-02T07:20:34.914259Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-02T07:20:34.968555Z","title":"Proceedings of the 36th Annual ACM Symposium on User Interface Software and Technology (UIST) , pages=","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2607.10428","last_updated":"2026-07-24T14:32:58Z","snapshot_observed_at":"2026-08-06T10:45:46.606756Z","submitted_at":"2026-07-11T18:22:49Z","title":"Enjoy Your Talk: A Human-Centered Benchmark for Multi-Turn Dialogue with Decoupled User Simulation, Target Modeling, and Judging","version":2},"reference_index":34,"source":"arxiv_source","source_observed_at":"2026-08-02T07:20:34.968555Z"},"links":{"citing_paper":"/paper/2607.10428"},"observation_digest":"sha256:dde3d2bc5683d8c693b7a921caf0297321c97dcc8628e979363a85cbcfb41250","observation_id":"cb1fbcfc-c1e2-419d-ad61-34c7afb2e2a3","resolution":{"observed_at":"2026-08-02T07:20:34.968555Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-02T07:20:35.042805Z","title":"Political Analysis , volume=","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2607.10428","last_updated":"2026-07-24T14:32:58Z","snapshot_observed_at":"2026-08-06T10:45:46.606756Z","submitted_at":"2026-07-11T18:22:49Z","title":"Enjoy Your Talk: A Human-Centered Benchmark for Multi-Turn Dialogue with Decoupled User Simulation, Target Modeling, and Judging","version":2},"reference_index":35,"source":"arxiv_source","source_observed_at":"2026-08-02T07:20:35.042805Z"},"links":{"citing_paper":"/paper/2607.10428"},"observation_digest":"sha256:331532970583ab4ec9a4cc16e0db4c363501adb1cd14f594c855e7d6b4447aaa","observation_id":"c95b2e05-94f6-4a06-bc40-c167e4546bac","resolution":{"observed_at":"2026-08-02T07:20:35.042805Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2406.01633","last_updated":"2024-06-01T15:54:45Z","snapshot_observed_at":"2026-07-06T18:24:38.958815Z","submitted_at":"2024-06-01T15:54:45Z","title":"On Overcoming Miscalibrated Conversational Priors in LLM-based Chatbots","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2406.01633","snapshot_observed_at":"2026-08-02T07:20:35.092841Z","title":"arXiv preprint arXiv:2406.01633 , year=","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2607.10428","last_updated":"2026-07-24T14:32:58Z","snapshot_observed_at":"2026-08-06T10:45:46.606756Z","submitted_at":"2026-07-11T18:22:49Z","title":"Enjoy Your Talk: A Human-Centered Benchmark for Multi-Turn Dialogue with Decoupled User Simulation, Target Modeling, and Judging","version":2},"reference_index":36,"source":"arxiv_source","source_observed_at":"2026-08-02T07:20:35.092841Z"},"links":{"cited_paper":"/paper/2406.01633","citing_paper":"/paper/2607.10428"},"observation_digest":"sha256:87899f5e54348e976a2405334944c8a336bd0aee281bedf0c9eece7d8156a2b8","observation_id":"14ff5429-222b-46bd-a1fa-314cfdb7a200","resolution":{"observed_at":"2026-08-02T07:20:35.092841Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-02T07:20:35.143800Z","title":"ACL , year=","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2607.10428","last_updated":"2026-07-24T14:32:58Z","snapshot_observed_at":"2026-08-06T10:45:46.606756Z","submitted_at":"2026-07-11T18:22:49Z","title":"Enjoy Your Talk: A Human-Centered Benchmark for Multi-Turn Dialogue with Decoupled User Simulation, Target Modeling, and Judging","version":2},"reference_index":37,"source":"arxiv_source","source_observed_at":"2026-08-02T07:20:35.143800Z"},"links":{"citing_paper":"/paper/2607.10428"},"observation_digest":"sha256:df7b7a92e2100c633faf6b4a203dc1ef4bebeb8b3d3abe4e05955f15c2b64164","observation_id":"c438bbe1-b1d3-46d1-8ab6-95241964f935","resolution":{"observed_at":"2026-08-02T07:20:35.143800Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2406.20094","last_updated":"2025-05-08T00:24:02Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2024-06-28T17:59:01Z","title":"Scaling Synthetic Data Creation with 1,000,000,000 Personas","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2406.20094","snapshot_observed_at":"2026-08-02T07:20:35.191658Z","title":"arXiv preprint arXiv:2406.20094 , year=","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2607.10428","last_updated":"2026-07-24T14:32:58Z","snapshot_observed_at":"2026-08-06T10:45:46.606756Z","submitted_at":"2026-07-11T18:22:49Z","title":"Enjoy Your Talk: A Human-Centered Benchmark for Multi-Turn Dialogue with Decoupled User Simulation, Target Modeling, and Judging","version":2},"reference_index":38,"source":"arxiv_source","source_observed_at":"2026-08-02T07:20:35.191658Z"},"links":{"cited_paper":"/paper/2406.20094","citing_paper":"/paper/2607.10428"},"observation_digest":"sha256:b9db42db9c4a22f9ca6e0c93603e17b3b37656069fb9f4956ecdd984c3d61ca3","observation_id":"83084c80-54b0-4d26-897d-9c05356fcc34","resolution":{"observed_at":"2026-08-02T07:20:35.191658Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-02T07:20:35.253158Z","title":"2025 , howpublished=","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2607.10428","last_updated":"2026-07-24T14:32:58Z","snapshot_observed_at":"2026-08-06T10:45:46.606756Z","submitted_at":"2026-07-11T18:22:49Z","title":"Enjoy Your Talk: A Human-Centered Benchmark for Multi-Turn Dialogue with Decoupled User Simulation, Target Modeling, and Judging","version":2},"reference_index":39,"source":"arxiv_source","source_observed_at":"2026-08-02T07:20:35.253158Z"},"links":{"citing_paper":"/paper/2607.10428"},"observation_digest":"sha256:fc54e66ef82caba843752aea96fa9335357ea6c77873f99175e4411e9aed3206","observation_id":"daae4709-4808-47a1-a79c-99c5526ec313","resolution":{"observed_at":"2026-08-02T07:20:35.253158Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-02T07:20:35.324311Z","title":"ACL , year=","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2607.10428","last_updated":"2026-07-24T14:32:58Z","snapshot_observed_at":"2026-08-06T10:45:46.606756Z","submitted_at":"2026-07-11T18:22:49Z","title":"Enjoy Your Talk: A Human-Centered Benchmark for Multi-Turn Dialogue with Decoupled User Simulation, Target Modeling, and Judging","version":2},"reference_index":40,"source":"arxiv_source","source_observed_at":"2026-08-02T07:20:35.324311Z"},"links":{"citing_paper":"/paper/2607.10428"},"observation_digest":"sha256:51482305b34465412e243e2fa0508f6ae24e76679d3e9e8147dc94005ff65bd6","observation_id":"15a7b99b-d08b-435a-af57-3a99c78d10d2","resolution":{"observed_at":"2026-08-02T07:20:35.324311Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-02T07:20:35.390815Z","title":"arXiv preprint arXiv:2512.06688 , year=","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2607.10428","last_updated":"2026-07-24T14:32:58Z","snapshot_observed_at":"2026-08-06T10:45:46.606756Z","submitted_at":"2026-07-11T18:22:49Z","title":"Enjoy Your Talk: A Human-Centered Benchmark for Multi-Turn Dialogue with Decoupled User Simulation, Target Modeling, and Judging","version":2},"reference_index":41,"source":"arxiv_source","source_observed_at":"2026-08-02T07:20:35.390815Z"},"links":{"citing_paper":"/paper/2607.10428"},"observation_digest":"sha256:3fb0eecfaadff9f23d0aa989bd727721a1e975f3effba18ecb3a3ccf7b837a85","observation_id":"24799a9b-9f34-4bbd-83a8-fc2aeeea9829","resolution":{"observed_at":"2026-08-02T07:20:35.390815Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-02T07:20:35.474300Z","title":"Proceedings of the 62nd Annual Meeting of the Association for Computational Linguistics (Volume 1: Long Papers) , pages=","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2607.10428","last_updated":"2026-07-24T14:32:58Z","snapshot_observed_at":"2026-08-06T10:45:46.606756Z","submitted_at":"2026-07-11T18:22:49Z","title":"Enjoy Your Talk: A Human-Centered Benchmark for Multi-Turn Dialogue with Decoupled User Simulation, Target Modeling, and Judging","version":2},"reference_index":42,"source":"arxiv_source","source_observed_at":"2026-08-02T07:20:35.474300Z"},"links":{"citing_paper":"/paper/2607.10428"},"observation_digest":"sha256:8552e8623bfb3f8f83072687328b1f235653d2e5e15448b81eec8a07dbbeece2","observation_id":"7615334d-6bf0-4a0c-b5d3-216caff3a5f6","resolution":{"observed_at":"2026-08-02T07:20:35.474300Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-02T07:20:35.541057Z","title":"ACL , year=","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2607.10428","last_updated":"2026-07-24T14:32:58Z","snapshot_observed_at":"2026-08-06T10:45:46.606756Z","submitted_at":"2026-07-11T18:22:49Z","title":"Enjoy Your Talk: A Human-Centered Benchmark for Multi-Turn Dialogue with Decoupled User Simulation, Target Modeling, and Judging","version":2},"reference_index":43,"source":"arxiv_source","source_observed_at":"2026-08-02T07:20:35.541057Z"},"links":{"citing_paper":"/paper/2607.10428"},"observation_digest":"sha256:e74c1f19bae1e19c68090b7d731738f46bd239e55f5c606a0dbacdf124dd8a67","observation_id":"1d72884e-bf7e-47ed-a4e0-eb2d925b9b95","resolution":{"observed_at":"2026-08-02T07:20:35.541057Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-02T07:20:35.590927Z","title":"ACL , year=","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2607.10428","last_updated":"2026-07-24T14:32:58Z","snapshot_observed_at":"2026-08-06T10:45:46.606756Z","submitted_at":"2026-07-11T18:22:49Z","title":"Enjoy Your Talk: A Human-Centered Benchmark for Multi-Turn Dialogue with Decoupled User Simulation, Target Modeling, and Judging","version":2},"reference_index":44,"source":"arxiv_source","source_observed_at":"2026-08-02T07:20:35.590927Z"},"links":{"citing_paper":"/paper/2607.10428"},"observation_digest":"sha256:7ddd97d20a75494ed00dbf884bf3def883f8bc4ec64c6f0ba7959b3e58142f43","observation_id":"6ac7ba65-6db5-4ec9-8e8a-b58b2bc6b790","resolution":{"observed_at":"2026-08-02T07:20:35.590927Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-02T07:20:35.625127Z","title":"Imagination, Cognition and Personality , year=","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2607.10428","last_updated":"2026-07-24T14:32:58Z","snapshot_observed_at":"2026-08-06T10:45:46.606756Z","submitted_at":"2026-07-11T18:22:49Z","title":"Enjoy Your Talk: A Human-Centered Benchmark for Multi-Turn Dialogue with Decoupled User Simulation, Target Modeling, and Judging","version":2},"reference_index":45,"source":"arxiv_source","source_observed_at":"2026-08-02T07:20:35.625127Z"},"links":{"citing_paper":"/paper/2607.10428"},"observation_digest":"sha256:fc98affe77fd2a0787407dffd25cf44560bfa34f47bb68c69652a0a39353c986","observation_id":"18d7fbcb-828d-4857-b709-925260550192","resolution":{"observed_at":"2026-08-02T07:20:35.625127Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-02T07:20:35.639867Z","title":null,"venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2607.10428","last_updated":"2026-07-24T14:32:58Z","snapshot_observed_at":"2026-08-06T10:45:46.606756Z","submitted_at":"2026-07-11T18:22:49Z","title":"Enjoy Your Talk: A Human-Centered Benchmark for Multi-Turn Dialogue with Decoupled User Simulation, Target Modeling, and Judging","version":2},"reference_index":46,"source":"arxiv_source","source_observed_at":"2026-08-02T07:20:35.639867Z"},"links":{"citing_paper":"/paper/2607.10428"},"observation_digest":"sha256:ce07bf4ad2a019131d18682720a5651de6341443b312d5ea455851791280a8ce","observation_id":"f3eb4455-0440-4bd4-948d-14eb211a41de","resolution":{"observed_at":"2026-08-02T07:20:35.639867Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-02T07:20:35.660515Z","title":"arXiv preprint arXiv:2511.08394 , year=","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2607.10428","last_updated":"2026-07-24T14:32:58Z","snapshot_observed_at":"2026-08-06T10:45:46.606756Z","submitted_at":"2026-07-11T18:22:49Z","title":"Enjoy Your Talk: A Human-Centered Benchmark for Multi-Turn Dialogue with Decoupled User Simulation, Target Modeling, and Judging","version":2},"reference_index":47,"source":"arxiv_source","source_observed_at":"2026-08-02T07:20:35.660515Z"},"links":{"citing_paper":"/paper/2607.10428"},"observation_digest":"sha256:f4e1e453f0a9aca997c587363cfecd20bc7b1dda4545146e123554b83e01da64","observation_id":"46c1c6e5-6d21-4423-bb69-fafe7ca10528","resolution":{"observed_at":"2026-08-02T07:20:35.660515Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-02T07:20:35.814712Z","title":"arXiv preprint , year=","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2607.10428","last_updated":"2026-07-24T14:32:58Z","snapshot_observed_at":"2026-08-06T10:45:46.606756Z","submitted_at":"2026-07-11T18:22:49Z","title":"Enjoy Your Talk: A Human-Centered Benchmark for Multi-Turn Dialogue with Decoupled User Simulation, Target Modeling, and Judging","version":2},"reference_index":48,"source":"arxiv_source","source_observed_at":"2026-08-02T07:20:35.814712Z"},"links":{"citing_paper":"/paper/2607.10428"},"observation_digest":"sha256:4f77e85046223c833934699f9baa5b2a17462262b7b0d8c8ec43f99afe343439","observation_id":"e3f250ee-002f-4f76-9ecc-878570ea017c","resolution":{"observed_at":"2026-08-02T07:20:35.814712Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2504.16939","last_updated":"2025-04-07T21:01:25Z","snapshot_observed_at":"2026-08-07T16:07:47.395679Z","submitted_at":"2025-04-07T21:01:25Z","title":"A Desideratum for Conversational Agents: Capabilities, Challenges, and Future Directions","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2504.16939","snapshot_observed_at":"2026-08-02T07:20:35.969578Z","title":"arXiv preprint arXiv:2504.16939 , year=","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2607.10428","last_updated":"2026-07-24T14:32:58Z","snapshot_observed_at":"2026-08-06T10:45:46.606756Z","submitted_at":"2026-07-11T18:22:49Z","title":"Enjoy Your Talk: A Human-Centered Benchmark for Multi-Turn Dialogue with Decoupled User Simulation, Target Modeling, and Judging","version":2},"reference_index":49,"source":"arxiv_source","source_observed_at":"2026-08-02T07:20:35.969578Z"},"links":{"cited_paper":"/paper/2504.16939","citing_paper":"/paper/2607.10428"},"observation_digest":"sha256:5179b87644899686b23c34b2dbbeee82171101fdca5b839862d5c2fdc9fd554f","observation_id":"d7f8d8ce-bb52-410f-bb69-5d2610228297","resolution":{"observed_at":"2026-08-02T07:20:35.969578Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2403.15115","last_updated":"2024-06-22T12:17:38Z","snapshot_observed_at":"2026-07-06T17:48:59.440247Z","submitted_at":"2024-03-22T11:16:43Z","title":"Language Models in Dialogue: Conversational Maxims for Human-AI Interactions","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2403.15115","snapshot_observed_at":"2026-08-02T07:20:36.120148Z","title":"arXiv preprint arXiv:2403.15115 , year=","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2607.10428","last_updated":"2026-07-24T14:32:58Z","snapshot_observed_at":"2026-08-06T10:45:46.606756Z","submitted_at":"2026-07-11T18:22:49Z","title":"Enjoy Your Talk: A Human-Centered Benchmark for Multi-Turn Dialogue with Decoupled User Simulation, Target Modeling, and Judging","version":2},"reference_index":50,"source":"arxiv_source","source_observed_at":"2026-08-02T07:20:36.120148Z"},"links":{"cited_paper":"/paper/2403.15115","citing_paper":"/paper/2607.10428"},"observation_digest":"sha256:c73f2b03d9bb7ac6fe9ec0931eb28a0b1e817e09ae7ab09790e8a2be618b6446","observation_id":"7ee82542-17ab-4a79-a39c-f5a011d9bbc6","resolution":{"observed_at":"2026-08-02T07:20:36.120148Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2504.04717","last_updated":"2026-04-20T02:55:07Z","snapshot_observed_at":"2026-07-06T21:05:11.303685Z","submitted_at":"2025-04-07T04:00:08Z","title":"Beyond Single-Turn: A Survey on Multi-Turn Interactions with Large Language Models","version":6},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2504.04717","snapshot_observed_at":"2026-08-02T07:20:36.288732Z","title":"arXiv preprint arXiv:2504.04717 , year=","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2607.10428","last_updated":"2026-07-24T14:32:58Z","snapshot_observed_at":"2026-08-06T10:45:46.606756Z","submitted_at":"2026-07-11T18:22:49Z","title":"Enjoy Your Talk: A Human-Centered Benchmark for Multi-Turn Dialogue with Decoupled User Simulation, Target Modeling, and Judging","version":2},"reference_index":51,"source":"arxiv_source","source_observed_at":"2026-08-02T07:20:36.288732Z"},"links":{"cited_paper":"/paper/2504.04717","citing_paper":"/paper/2607.10428"},"observation_digest":"sha256:63cc0bc88cb2c2319b71fa098407152249400d463065649aa620f47e98157737","observation_id":"5e7946b8-bc0d-45ef-9888-68f4b57a0b92","resolution":{"observed_at":"2026-08-02T07:20:36.288732Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-02T07:20:36.454494Z","title":"arXiv preprint arXiv:2512.10493 , year=","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2607.10428","last_updated":"2026-07-24T14:32:58Z","snapshot_observed_at":"2026-08-06T10:45:46.606756Z","submitted_at":"2026-07-11T18:22:49Z","title":"Enjoy Your Talk: A Human-Centered Benchmark for Multi-Turn Dialogue with Decoupled User Simulation, Target Modeling, and Judging","version":2},"reference_index":52,"source":"arxiv_source","source_observed_at":"2026-08-02T07:20:36.454494Z"},"links":{"citing_paper":"/paper/2607.10428"},"observation_digest":"sha256:b1fddf1a4dc67c814a15a033a978e690f465e5e9a59e393a1178812894cb0a0d","observation_id":"3e0a2ab0-f4ce-4d10-b19e-3fd95d254e7d","resolution":{"observed_at":"2026-08-02T07:20:36.454494Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-02T07:20:36.628647Z","title":"Proceedings of the 40th Annual Meeting of the Association for Computational Linguistics , pages=","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2607.10428","last_updated":"2026-07-24T14:32:58Z","snapshot_observed_at":"2026-08-06T10:45:46.606756Z","submitted_at":"2026-07-11T18:22:49Z","title":"Enjoy Your Talk: A Human-Centered Benchmark for Multi-Turn Dialogue with Decoupled User Simulation, Target Modeling, and Judging","version":2},"reference_index":53,"source":"arxiv_source","source_observed_at":"2026-08-02T07:20:36.628647Z"},"links":{"citing_paper":"/paper/2607.10428"},"observation_digest":"sha256:6fbbeb93523bb2f3f761bdf4efb0266dd2aa766e2401e49835c58e98e7201702","observation_id":"94922d0c-4bcb-447b-b79c-888e1811ff9a","resolution":{"observed_at":"2026-08-02T07:20:36.628647Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-02T07:20:36.791237Z","title":"Text Summarization Branches Out , pages=","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2607.10428","last_updated":"2026-07-24T14:32:58Z","snapshot_observed_at":"2026-08-06T10:45:46.606756Z","submitted_at":"2026-07-11T18:22:49Z","title":"Enjoy Your Talk: A Human-Centered Benchmark for Multi-Turn Dialogue with Decoupled User Simulation, Target Modeling, and Judging","version":2},"reference_index":54,"source":"arxiv_source","source_observed_at":"2026-08-02T07:20:36.791237Z"},"links":{"citing_paper":"/paper/2607.10428"},"observation_digest":"sha256:cdf17038dcc82148e77849dddaace1f481a452db71bff4528afcf3426ab3e642","observation_id":"f2bf47b1-24ff-493f-80a7-494da32f4abb","resolution":{"observed_at":"2026-08-02T07:20:36.791237Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-02T07:20:36.957153Z","title":"ACM Transactions on Information Systems , volume=","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2607.10428","last_updated":"2026-07-24T14:32:58Z","snapshot_observed_at":"2026-08-06T10:45:46.606756Z","submitted_at":"2026-07-11T18:22:49Z","title":"Enjoy Your Talk: A Human-Centered Benchmark for Multi-Turn Dialogue with Decoupled User Simulation, Target Modeling, and Judging","version":2},"reference_index":55,"source":"arxiv_source","source_observed_at":"2026-08-02T07:20:36.957153Z"},"links":{"citing_paper":"/paper/2607.10428"},"observation_digest":"sha256:d0ce01712ece41bedc53078b28eb6d93d488d71c0740e320dde1e300197bfa95","observation_id":"92a2c7b7-bfed-46a2-a849-b4b41a12895f","resolution":{"observed_at":"2026-08-02T07:20:36.957153Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-02T07:20:37.125077Z","title":"arXiv preprint arXiv:2511.00222 , year=","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2607.10428","last_updated":"2026-07-24T14:32:58Z","snapshot_observed_at":"2026-08-06T10:45:46.606756Z","submitted_at":"2026-07-11T18:22:49Z","title":"Enjoy Your Talk: A Human-Centered Benchmark for Multi-Turn Dialogue with Decoupled User Simulation, Target Modeling, and Judging","version":2},"reference_index":56,"source":"arxiv_source","source_observed_at":"2026-08-02T07:20:37.125077Z"},"links":{"citing_paper":"/paper/2607.10428"},"observation_digest":"sha256:276c550f5c88e1b1bb5e6d51f82e6f7991796d1410807867ef37946f39fcde42","observation_id":"b7d37a22-8a8f-43fe-b7c2-588206ee8e9f","resolution":{"observed_at":"2026-08-02T07:20:37.125077Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-02T07:20:37.237479Z","title":"arXiv preprint arXiv:2511.03508 , year=","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2607.10428","last_updated":"2026-07-24T14:32:58Z","snapshot_observed_at":"2026-08-06T10:45:46.606756Z","submitted_at":"2026-07-11T18:22:49Z","title":"Enjoy Your Talk: A Human-Centered Benchmark for Multi-Turn Dialogue with Decoupled User Simulation, Target Modeling, and Judging","version":2},"reference_index":57,"source":"arxiv_source","source_observed_at":"2026-08-02T07:20:37.237479Z"},"links":{"citing_paper":"/paper/2607.10428"},"observation_digest":"sha256:fbdec2547bf9c43ada0f8e6db5dc1e0432bc69056b9d4f00679fb75852e61d2f","observation_id":"1b3f6841-c0dd-47cb-bcb8-e6fcbf5726cb","resolution":{"observed_at":"2026-08-02T07:20:37.237479Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-02T07:20:37.445078Z","title":"arXiv preprint arXiv:2507.20152 , year=","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2607.10428","last_updated":"2026-07-24T14:32:58Z","snapshot_observed_at":"2026-08-06T10:45:46.606756Z","submitted_at":"2026-07-11T18:22:49Z","title":"Enjoy Your Talk: A Human-Centered Benchmark for Multi-Turn Dialogue with Decoupled User Simulation, Target Modeling, and Judging","version":2},"reference_index":58,"source":"arxiv_source","source_observed_at":"2026-08-02T07:20:37.445078Z"},"links":{"citing_paper":"/paper/2607.10428"},"observation_digest":"sha256:05d9ddb00226723219b9ff22d2ce447296beca3f5c5f36eee2839e8a1d4c63b9","observation_id":"09d78d45-3650-49a4-ad54-9747f74c41eb","resolution":{"observed_at":"2026-08-02T07:20:37.445078Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-02T07:20:37.609359Z","title":"The Twelfth International Conference on Learning Representations (ICLR) , year=","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2607.10428","last_updated":"2026-07-24T14:32:58Z","snapshot_observed_at":"2026-08-06T10:45:46.606756Z","submitted_at":"2026-07-11T18:22:49Z","title":"Enjoy Your Talk: A Human-Centered Benchmark for Multi-Turn Dialogue with Decoupled User Simulation, Target Modeling, and Judging","version":2},"reference_index":59,"source":"arxiv_source","source_observed_at":"2026-08-02T07:20:37.609359Z"},"links":{"citing_paper":"/paper/2607.10428"},"observation_digest":"sha256:3008ffab0a9038f86f10d6818e69cb04bb84ec98b20c46330966d3472a9dfeba","observation_id":"03370a05-40c7-4988-a3ea-fe504cca4821","resolution":{"observed_at":"2026-08-02T07:20:37.609359Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-02T07:20:37.737009Z","title":"Proceedings of the 29th Symposium on Operating Systems Principles (SOSP) , pages=","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2607.10428","last_updated":"2026-07-24T14:32:58Z","snapshot_observed_at":"2026-08-06T10:45:46.606756Z","submitted_at":"2026-07-11T18:22:49Z","title":"Enjoy Your Talk: A Human-Centered Benchmark for Multi-Turn Dialogue with Decoupled User Simulation, Target Modeling, and Judging","version":2},"reference_index":60,"source":"arxiv_source","source_observed_at":"2026-08-02T07:20:37.737009Z"},"links":{"citing_paper":"/paper/2607.10428"},"observation_digest":"sha256:ce5b34cb50d695c4fc17bc1158c904f3149b09a57cc0e03934d5e101524c2383","observation_id":"51e8b65a-385b-4f80-8447-fb889fe720cb","resolution":{"observed_at":"2026-08-02T07:20:37.737009Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-02T07:20:37.879054Z","title":"Educational and Psychological Measurement , year=","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2607.10428","last_updated":"2026-07-24T14:32:58Z","snapshot_observed_at":"2026-08-06T10:45:46.606756Z","submitted_at":"2026-07-11T18:22:49Z","title":"Enjoy Your Talk: A Human-Centered Benchmark for Multi-Turn Dialogue with Decoupled User Simulation, Target Modeling, and Judging","version":2},"reference_index":61,"source":"arxiv_source","source_observed_at":"2026-08-02T07:20:37.879054Z"},"links":{"citing_paper":"/paper/2607.10428"},"observation_digest":"sha256:1a8d89ccf397b2207c36aa8bc9b46870ee7e782d89b03e9b8c53da97cc99a5dd","observation_id":"e9903963-a031-4626-afc2-9f4c9ab5652d","resolution":{"observed_at":"2026-08-02T07:20:37.879054Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-02T07:20:38.084127Z","title":null,"venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2607.10428","last_updated":"2026-07-24T14:32:58Z","snapshot_observed_at":"2026-08-06T10:45:46.606756Z","submitted_at":"2026-07-11T18:22:49Z","title":"Enjoy Your Talk: A Human-Centered Benchmark for Multi-Turn Dialogue with Decoupled User Simulation, Target Modeling, and Judging","version":2},"reference_index":62,"source":"arxiv_source","source_observed_at":"2026-08-02T07:20:38.084127Z"},"links":{"citing_paper":"/paper/2607.10428"},"observation_digest":"sha256:eb94b5f527a1499d34d284b4404068e658bb811fe93c9a8b3062ea4fdca8014f","observation_id":"4a77c5cf-b515-4eb5-9836-7b3bcf87ff3a","resolution":{"observed_at":"2026-08-02T07:20:38.084127Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-02T07:20:38.255709Z","title":"Proceedings of the Royal Society of London , year=","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2607.10428","last_updated":"2026-07-24T14:32:58Z","snapshot_observed_at":"2026-08-06T10:45:46.606756Z","submitted_at":"2026-07-11T18:22:49Z","title":"Enjoy Your Talk: A Human-Centered Benchmark for Multi-Turn Dialogue with Decoupled User Simulation, Target Modeling, and Judging","version":2},"reference_index":63,"source":"arxiv_source","source_observed_at":"2026-08-02T07:20:38.255709Z"},"links":{"citing_paper":"/paper/2607.10428"},"observation_digest":"sha256:bda514f1006adb5faeef42628c87b3cd9b55cca73accef377774ed432434c142","observation_id":"8e545315-cc2a-4844-9662-ccc7ddb08b15","resolution":{"observed_at":"2026-08-02T07:20:38.255709Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-02T07:20:38.424617Z","title":"Proceedings of the 2024 CHI Conference on Human Factors in Computing Systems , numpages=","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2607.10428","last_updated":"2026-07-24T14:32:58Z","snapshot_observed_at":"2026-08-06T10:45:46.606756Z","submitted_at":"2026-07-11T18:22:49Z","title":"Enjoy Your Talk: A Human-Centered Benchmark for Multi-Turn Dialogue with Decoupled User Simulation, Target Modeling, and Judging","version":2},"reference_index":64,"source":"arxiv_source","source_observed_at":"2026-08-02T07:20:38.424617Z"},"links":{"citing_paper":"/paper/2607.10428"},"observation_digest":"sha256:d5efc12ef69960731805f085941d76d6ced3a8c9e2518cc8e06525613baedd06","observation_id":"d0bad826-6dd5-45f8-8297-100e7214ee0b","resolution":{"observed_at":"2026-08-02T07:20:38.424617Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2508.10142","last_updated":"2025-08-24T20:24:03Z","snapshot_observed_at":"2026-08-07T14:02:28.974363Z","submitted_at":"2025-08-13T19:14:45Z","title":"Multi-Turn Puzzles: Evaluating Interactive Reasoning and Strategic Dialogue in LLMs","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2508.10142","snapshot_observed_at":"2026-08-02T07:20:38.536898Z","title":"arXiv preprint arXiv:2508.10142 , year=","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2607.10428","last_updated":"2026-07-24T14:32:58Z","snapshot_observed_at":"2026-08-06T10:45:46.606756Z","submitted_at":"2026-07-11T18:22:49Z","title":"Enjoy Your Talk: A Human-Centered Benchmark for Multi-Turn Dialogue with Decoupled User Simulation, Target Modeling, and Judging","version":2},"reference_index":65,"source":"arxiv_source","source_observed_at":"2026-08-02T07:20:38.536898Z"},"links":{"cited_paper":"/paper/2508.10142","citing_paper":"/paper/2607.10428"},"observation_digest":"sha256:3f6e48d6cbe7471a1affdbbb0d5f05f86f208b604320f30bd656ee71def58b7a","observation_id":"01fd8100-fde8-4243-8187-2de2f14d5396","resolution":{"observed_at":"2026-08-02T07:20:38.536898Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2505.02018","last_updated":"2025-05-04T07:48:36Z","snapshot_observed_at":"2026-08-07T15:56:40.603452Z","submitted_at":"2025-05-04T07:48:36Z","title":"R-Bench: Graduate-level Multi-disciplinary Benchmarks for LLM & MLLM Complex Reasoning Evaluation","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2505.02018","snapshot_observed_at":"2026-08-02T07:20:38.710219Z","title":"arXiv preprint arXiv:2505.02018 , year=","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2607.10428","last_updated":"2026-07-24T14:32:58Z","snapshot_observed_at":"2026-08-06T10:45:46.606756Z","submitted_at":"2026-07-11T18:22:49Z","title":"Enjoy Your Talk: A Human-Centered Benchmark for Multi-Turn Dialogue with Decoupled User Simulation, Target Modeling, and Judging","version":2},"reference_index":66,"source":"arxiv_source","source_observed_at":"2026-08-02T07:20:38.710219Z"},"links":{"cited_paper":"/paper/2505.02018","citing_paper":"/paper/2607.10428"},"observation_digest":"sha256:b5396532170f7b4490238faa46851da4f85c3a2ed3893282ec8ecfa2e0ea37ab","observation_id":"9f756069-7f36-4809-bd2d-a468123d613e","resolution":{"observed_at":"2026-08-02T07:20:38.710219Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2311.07911","last_updated":"2023-11-14T05:13:55Z","snapshot_observed_at":"2026-07-06T16:47:08.877195Z","submitted_at":"2023-11-14T05:13:55Z","title":"Instruction-Following Evaluation for Large Language Models","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2311.07911","snapshot_observed_at":"2026-08-02T07:20:38.736690Z","title":"arXiv preprint arXiv:2311.07911 , year=","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2607.10428","last_updated":"2026-07-24T14:32:58Z","snapshot_observed_at":"2026-08-06T10:45:46.606756Z","submitted_at":"2026-07-11T18:22:49Z","title":"Enjoy Your Talk: A Human-Centered Benchmark for Multi-Turn Dialogue with Decoupled User Simulation, Target Modeling, and Judging","version":2},"reference_index":67,"source":"arxiv_source","source_observed_at":"2026-08-02T07:20:38.736690Z"},"links":{"cited_paper":"/paper/2311.07911","citing_paper":"/paper/2607.10428"},"observation_digest":"sha256:636b2358a4fc8155268d13cdd4d44b7f985e995d3d0b4d74127a8e2cae1649e6","observation_id":"a34bba85-9614-40c4-b482-2a785ec08613","resolution":{"observed_at":"2026-08-02T07:20:38.736690Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-02T07:20:38.785314Z","title":"Advances in Neural Information Processing Systems , volume=","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2607.10428","last_updated":"2026-07-24T14:32:58Z","snapshot_observed_at":"2026-08-06T10:45:46.606756Z","submitted_at":"2026-07-11T18:22:49Z","title":"Enjoy Your Talk: A Human-Centered Benchmark for Multi-Turn Dialogue with Decoupled User Simulation, Target Modeling, and Judging","version":2},"reference_index":68,"source":"arxiv_source","source_observed_at":"2026-08-02T07:20:38.785314Z"},"links":{"citing_paper":"/paper/2607.10428"},"observation_digest":"sha256:cc82c7dceb3b9cb3c2d3751769338e920a40830ca0af9947bdcb108d4205c3a5","observation_id":"c6bb1e5c-48bc-48fc-82a3-9e23f4d58850","resolution":{"observed_at":"2026-08-02T07:20:38.785314Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-02T07:20:38.901702Z","title":"arXiv preprint arXiv:2508.06196 , year=","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2607.10428","last_updated":"2026-07-24T14:32:58Z","snapshot_observed_at":"2026-08-06T10:45:46.606756Z","submitted_at":"2026-07-11T18:22:49Z","title":"Enjoy Your Talk: A Human-Centered Benchmark for Multi-Turn Dialogue with Decoupled User Simulation, Target Modeling, and Judging","version":2},"reference_index":69,"source":"arxiv_source","source_observed_at":"2026-08-02T07:20:38.901702Z"},"links":{"citing_paper":"/paper/2607.10428"},"observation_digest":"sha256:5a43dbb0fb2bbe0af482e9e2b4e9f9f5b75ade51e7bd7749e0b68d9113a5ddb2","observation_id":"3e70b0f7-86e4-43b5-8eae-d3315b547d8e","resolution":{"observed_at":"2026-08-02T07:20:38.901702Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-02T07:20:39.009425Z","title":"arXiv preprint arXiv:2505.23810 , year=","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2607.10428","last_updated":"2026-07-24T14:32:58Z","snapshot_observed_at":"2026-08-06T10:45:46.606756Z","submitted_at":"2026-07-11T18:22:49Z","title":"Enjoy Your Talk: A Human-Centered Benchmark for Multi-Turn Dialogue with Decoupled User Simulation, Target Modeling, and Judging","version":2},"reference_index":70,"source":"arxiv_source","source_observed_at":"2026-08-02T07:20:39.009425Z"},"links":{"citing_paper":"/paper/2607.10428"},"observation_digest":"sha256:9f71fd528b61e91f74407b30a9d8d02da4fd737da50013ef94564dd496042cfa","observation_id":"2b6073aa-7313-4b74-a449-82a3b70b7577","resolution":{"observed_at":"2026-08-02T07:20:39.009425Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-02T07:20:39.150679Z","title":"Proceedings of the ACM on Management of Data , volume=","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2607.10428","last_updated":"2026-07-24T14:32:58Z","snapshot_observed_at":"2026-08-06T10:45:46.606756Z","submitted_at":"2026-07-11T18:22:49Z","title":"Enjoy Your Talk: A Human-Centered Benchmark for Multi-Turn Dialogue with Decoupled User Simulation, Target Modeling, and Judging","version":2},"reference_index":71,"source":"arxiv_source","source_observed_at":"2026-08-02T07:20:39.150679Z"},"links":{"citing_paper":"/paper/2607.10428"},"observation_digest":"sha256:3cb0e01b13b71d9fa2613f286a5d0542ff67209cb00dd44d7c2ef0c4d73f1359","observation_id":"3b1e17c1-ed7f-44d1-b266-fe1c15f483f3","resolution":{"observed_at":"2026-08-02T07:20:39.150679Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}}],"paper":{"arxiv_id":"2607.10428","last_updated":"2026-07-24T14:32:58Z","latest_version":2,"primary_category":"cs.CL","snapshot_observed_at":"2026-08-06T10:45:46.606756Z","submitted_at":"2026-07-11T18:22:49Z","title":"Enjoy Your Talk: A Human-Centered Benchmark for Multi-Turn Dialogue with Decoupled User Simulation, Target Modeling, and Judging"},"reference_resolution":{"displayed":71,"state_counts":{"malformed_identifier":0,"metadata_mismatch":0,"parse_uncertain":0,"unresolved":71,"verified_exact":0,"verified_fuzzy":0},"total_outbound_references":71},"refusal":"A citation records a reference. It does not transfer a finding from one paper to another.","schema":"pith.paper-citation-record.v1","standing_sources":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"thesis":"As of 8 August 2026, this Paper Citation Record lists 71 of 71 outbound references and 0 inbound Pith citation observations for arXiv:2607.10428."}