{"as_of":"2026-08-05T07:56:00Z","caps":{"database_statements":6,"inbound":100,"outbound":100},"context_digest":"sha256:043721627f65eec5c0782908fd055fbe90a31c04363c39783ca908044e4094fd","coverage":[{"denominator":29,"lane":"reference_resolution","note":"Typed states for the displayed outbound observations.","records_observed":29,"source":"paper_references, paper_reference_links","source_observed_at":"2026-05-16T00:03:55.599967Z","state":"measured"},{"denominator":103,"lane":"standing_notices","note":"One-hop event checks from named stored sources.","records_observed":103,"source":"scholarly_work_events, retraction_status_cache","source_observed_at":"2026-08-05T06:32:48.257954+00:00","state":"measured"},{"denominator":74,"lane":"inbound_itemization","note":"Pith citing papers itemized under the disclosed page cap.","records_observed":74,"source":"paper_references, paper_reference_links","source_observed_at":"2026-08-04T22:24:42.425486Z","state":"measured"},{"denominator":1,"lane":"external_citation_measurements","note":"A source-named dated measurement, never combined with another source.","records_observed":1,"source":"pith","source_observed_at":"2026-08-05T02:28:24.338817Z","state":"measured"}],"external_citation_measurements":[{"count":8,"observed_at":"2026-08-05T02:28:24.338817Z","source":"pith"}],"inbound":[{"citation":{"cited_paper":{"arxiv_id":"2406.20094","last_updated":"2025-05-08T00:24:02Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2024-06-28T17:59:01Z","title":"Scaling Synthetic Data Creation with 1,000,000,000 Personas","version":3},"cited_work":{"arxiv_id":"2406.20094","doi":"10.48550/arxiv.2406.20094","metadata_source":"pith","pith_arxiv_id":"2406.20094","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Scaling Synthetic Data Creation with 1,000,000,000 Personas","venue":"cs.CL","work_id":"4acdd1d7-4a74-460c-a32c-01c386a5c58a","year":2024},"citing_paper":{"arxiv_id":"2409.17146","last_updated":"2024-12-05T14:28:40Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2024-09-25T17:59:51Z","title":"Molmo and PixMo: Open Weights and Open Data for State-of-the-Art Vision-Language Models","version":2},"reference_index":35,"source":"pdf_text","source_observed_at":"2026-05-15T01:55:12.501409Z"},"links":{"cited_paper":"/paper/2406.20094","citing_paper":"/paper/2409.17146"},"observation_digest":"sha256:d7bff9a9a23ccc26e62a676e6e48f40b14fe43c4504befe5a28f6beba96b4586","observation_id":"6681f19c-668e-4764-9aab-1cf21dff7002","resolution":{"observed_at":"2026-05-16T00:03:55.809567Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-07-13T23:49:58.855418+00:00","source":"crossref_status_cache"},{"observed_at":"2026-07-13T23:49:58.855418+00:00","source":"openalex_status_cache"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2406.20094","last_updated":"2025-05-08T00:24:02Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2024-06-28T17:59:01Z","title":"Scaling Synthetic Data Creation with 1,000,000,000 Personas","version":3},"cited_work":{"arxiv_id":"2406.20094","doi":"10.48550/arxiv.2406.20094","metadata_source":"pith","pith_arxiv_id":"2406.20094","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Scaling Synthetic Data Creation with 1,000,000,000 Personas","venue":"cs.CL","work_id":"4acdd1d7-4a74-460c-a32c-01c386a5c58a","year":2024},"citing_paper":{"arxiv_id":"2502.02737","last_updated":"2025-02-04T21:43:16Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2025-02-04T21:43:16Z","title":"SmolLM2: When Smol Goes Big -- Data-Centric Training of a Small Language Model","version":1},"reference_index":173,"source":"arxiv_source","source_observed_at":"2026-05-13T17:30:02.803757Z"},"links":{"cited_paper":"/paper/2406.20094","citing_paper":"/paper/2502.02737"},"observation_digest":"sha256:fcb844f27d7a05fe29509237fe14c11014700759fd8e23e6922baa8645f78468","observation_id":"d6785153-bac2-49d4-bb66-7ce27142f9e0","resolution":{"observed_at":"2026-05-16T00:03:55.809567Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-07-13T23:49:58.855418+00:00","source":"crossref_status_cache"},{"observed_at":"2026-07-13T23:49:58.855418+00:00","source":"openalex_status_cache"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2406.20094","last_updated":"2025-05-08T00:24:02Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2024-06-28T17:59:01Z","title":"Scaling Synthetic Data Creation with 1,000,000,000 Personas","version":3},"cited_work":{"arxiv_id":"2406.20094","doi":"10.48550/arxiv.2406.20094","metadata_source":"pith","pith_arxiv_id":"2406.20094","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Scaling Synthetic Data Creation with 1,000,000,000 Personas","venue":"cs.CL","work_id":"4acdd1d7-4a74-460c-a32c-01c386a5c58a","year":2024},"citing_paper":{"arxiv_id":"2504.20605","last_updated":"2026-05-02T05:59:29Z","snapshot_observed_at":"2026-07-06T21:16:23.106510Z","submitted_at":"2025-04-29T10:15:28Z","title":"TF1-EN-3M: Three Million Synthetic Moral Fables for Training Small, Open Language Models","version":2},"reference_index":31,"source":"pdf_text","source_observed_at":"2026-05-22T19:01:42.307514Z"},"links":{"cited_paper":"/paper/2406.20094","citing_paper":"/paper/2504.20605"},"observation_digest":"sha256:77d0cdbc88fe0ccd4e1e288acb0dc7621fd1c1c449a315a73971fcb6c885a336","observation_id":"4c78b7e8-09e4-4b1b-ab35-0a80bf7224b2","resolution":{"observed_at":"2026-05-22T19:01:57.809507Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-07-13T23:49:58.855418+00:00","source":"crossref_status_cache"},{"observed_at":"2026-07-13T23:49:58.855418+00:00","source":"openalex_status_cache"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2406.20094","last_updated":"2025-05-08T00:24:02Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2024-06-28T17:59:01Z","title":"Scaling Synthetic Data Creation with 1,000,000,000 Personas","version":3},"cited_work":{"arxiv_id":"2406.20094","doi":"10.48550/arxiv.2406.20094","metadata_source":"pith","pith_arxiv_id":"2406.20094","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Scaling Synthetic Data Creation with 1,000,000,000 Personas","venue":"cs.CL","work_id":"4acdd1d7-4a74-460c-a32c-01c386a5c58a","year":2024},"citing_paper":{"arxiv_id":"2506.05176","last_updated":"2025-06-11T02:54:49Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2025-06-05T15:49:48Z","title":"Qwen3 Embedding: Advancing Text Embedding and Reranking Through Foundation Models","version":3},"reference_index":3,"source":"pdf_text","source_observed_at":"2026-05-10T13:46:09.851496Z"},"links":{"cited_paper":"/paper/2406.20094","citing_paper":"/paper/2506.05176"},"observation_digest":"sha256:11499609e13bf16a4a2df7188b7de2730ead9c07a95fd357fd14a2adeb2140a1","observation_id":"6ab87079-812a-4fd4-96c1-6ba54f8b0da8","resolution":{"observed_at":"2026-05-16T00:03:55.809567Z","resolver_source":"arxiv_id","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-07-13T23:49:58.855418+00:00","source":"crossref_status_cache"},{"observed_at":"2026-07-13T23:49:58.855418+00:00","source":"openalex_status_cache"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2406.20094","last_updated":"2025-05-08T00:24:02Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2024-06-28T17:59:01Z","title":"Scaling Synthetic Data Creation with 1,000,000,000 Personas","version":3},"cited_work":{"arxiv_id":"2406.20094","doi":"10.48550/arxiv.2406.20094","metadata_source":"pith","pith_arxiv_id":"2406.20094","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Scaling Synthetic Data Creation with 1,000,000,000 Personas","venue":"cs.CL","work_id":"4acdd1d7-4a74-460c-a32c-01c386a5c58a","year":2024},"citing_paper":{"arxiv_id":"2508.00414","last_updated":"2026-04-22T08:14:17Z","snapshot_observed_at":"2026-08-02T16:42:25.825514Z","submitted_at":"2025-08-01T08:11:31Z","title":"Cognitive Kernel-Pro: A Framework for Deep Research Agents and Agent Foundation Models Training","version":3},"reference_index":4,"source":"pdf_text","source_observed_at":"2026-05-19T01:50:30.626249Z"},"links":{"cited_paper":"/paper/2406.20094","citing_paper":"/paper/2508.00414"},"observation_digest":"sha256:47e9cd28e7c0f953432a341d786f54b423205c59d57470edf10d9117269fc91b","observation_id":"5b18aa0c-f28e-4536-8598-07fa6a1e197f","resolution":{"observed_at":"2026-05-19T01:51:57.872811Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-07-13T23:49:58.855418+00:00","source":"crossref_status_cache"},{"observed_at":"2026-07-13T23:49:58.855418+00:00","source":"openalex_status_cache"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2406.20094","last_updated":"2025-05-08T00:24:02Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2024-06-28T17:59:01Z","title":"Scaling Synthetic Data Creation with 1,000,000,000 Personas","version":3},"cited_work":{"arxiv_id":"2406.20094","doi":"10.48550/arxiv.2406.20094","metadata_source":"pith","pith_arxiv_id":"2406.20094","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Scaling Synthetic Data Creation with 1,000,000,000 Personas","venue":"cs.CL","work_id":"4acdd1d7-4a74-460c-a32c-01c386a5c58a","year":2024},"citing_paper":{"arxiv_id":"2508.18167","last_updated":"2026-05-14T18:20:00Z","snapshot_observed_at":"2026-08-04T20:09:49.064679Z","submitted_at":"2025-08-25T16:16:42Z","title":"DiscussLLM: Teaching Large Language Models When to Speak","version":2},"reference_index":35,"source":"pdf_text","source_observed_at":"2026-05-21T22:09:03.740109Z"},"links":{"cited_paper":"/paper/2406.20094","citing_paper":"/paper/2508.18167"},"observation_digest":"sha256:4a3dd48f8378f98776b6069e18a973a06be759623b7854ee774d5c3826b280c7","observation_id":"4201e59c-c7e9-43b3-b835-7556fe4a9dac","resolution":{"observed_at":"2026-05-21T22:10:42.309144Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-07-13T23:49:58.855418+00:00","source":"crossref_status_cache"},{"observed_at":"2026-07-13T23:49:58.855418+00:00","source":"openalex_status_cache"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2406.20094","last_updated":"2025-05-08T00:24:02Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2024-06-28T17:59:01Z","title":"Scaling Synthetic Data Creation with 1,000,000,000 Personas","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2406.20094","snapshot_observed_at":"2026-08-04T22:24:42.425486Z","title":"arXiv preprint arXiv:2406.20094","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2509.07370","last_updated":"2025-09-11T09:42:02Z","snapshot_observed_at":"2026-08-04T22:24:33.245386Z","submitted_at":"2025-09-09T03:39:28Z","title":"PersonaFuse: A Personality Activation-Driven Framework for Enhancing Human-LLM Interactions","version":2},"reference_index":28,"source":"arxiv_source","source_observed_at":"2026-08-04T22:24:42.425486Z"},"links":{"cited_paper":"/paper/2406.20094","citing_paper":"/paper/2509.07370"},"observation_digest":"sha256:a439704d4afb08d295cab650c85d0b3398aebdf5026ad28c00a27bb9daa33d1f","observation_id":"e135bbac-9535-4d6b-ad29-f91497864662","resolution":{"observed_at":"2026-08-04T22:24:42.425486Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2406.20094","last_updated":"2025-05-08T00:24:02Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2024-06-28T17:59:01Z","title":"Scaling Synthetic Data Creation with 1,000,000,000 Personas","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2406.20094","snapshot_observed_at":"2026-08-04T14:43:49.646148Z","title":"Scaling synthetic data creation with 1,000,000,000 personas","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2509.23694","last_updated":"2026-06-03T04:13:57Z","snapshot_observed_at":"2026-08-04T14:43:33.365212Z","submitted_at":"2025-09-28T07:05:17Z","title":"SafeSearch: Automated Red-Teaming of LLM-Based Search Agents","version":6},"reference_index":12,"source":"arxiv_source","source_observed_at":"2026-08-04T14:43:49.646148Z"},"links":{"cited_paper":"/paper/2406.20094","citing_paper":"/paper/2509.23694"},"observation_digest":"sha256:de8fc113cde6ab1271d55ae30c7ea21cca5ad8b52157b35ed6d9d1c66bf91957","observation_id":"29c26475-3d34-4bcc-8a68-025dfb8c14ac","resolution":{"observed_at":"2026-08-04T14:43:49.646148Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2406.20094","last_updated":"2025-05-08T00:24:02Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2024-06-28T17:59:01Z","title":"Scaling Synthetic Data Creation with 1,000,000,000 Personas","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2406.20094","snapshot_observed_at":"2026-08-04T10:26:10.742274Z","title":null,"venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2510.09905","last_updated":"2026-06-16T21:38:14Z","snapshot_observed_at":"2026-08-04T10:26:08.424231Z","submitted_at":"2025-10-10T22:39:37Z","title":"The Personalization Trap: How User Memory Alters Emotional Reasoning in LLMs","version":2},"reference_index":5,"source":"arxiv_source","source_observed_at":"2026-08-04T10:26:10.742274Z"},"links":{"cited_paper":"/paper/2406.20094","citing_paper":"/paper/2510.09905"},"observation_digest":"sha256:12e13f6e858b2c124e2ba9e27fa9e73a0b373d8c7cdb3f9a28a9dab8b0d610d0","observation_id":"e6203da6-5f4c-4682-87e2-4dd881572d55","resolution":{"observed_at":"2026-08-04T10:26:10.742274Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2406.20094","last_updated":"2025-05-08T00:24:02Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2024-06-28T17:59:01Z","title":"Scaling Synthetic Data Creation with 1,000,000,000 Personas","version":3},"cited_work":{"arxiv_id":"2406.20094","doi":"10.48550/arxiv.2406.20094","metadata_source":"pith","pith_arxiv_id":"2406.20094","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Scaling Synthetic Data Creation with 1,000,000,000 Personas","venue":"cs.CL","work_id":"4acdd1d7-4a74-460c-a32c-01c386a5c58a","year":2024},"citing_paper":{"arxiv_id":"2511.08565","last_updated":"2026-05-14T13:52:44Z","snapshot_observed_at":"2026-07-06T22:35:31.750339Z","submitted_at":"2025-11-11T18:47:44Z","title":"Moral Susceptibility and Robustness under Persona Role-Play in Large Language Models","version":3},"reference_index":13,"source":"pdf_text","source_observed_at":"2026-05-17T23:16:56.957905Z"},"links":{"cited_paper":"/paper/2406.20094","citing_paper":"/paper/2511.08565"},"observation_digest":"sha256:1210c388ad0d6c673762c235ab7a722ce5f9979349f211a9a18e10477f0b32b4","observation_id":"8fe0745c-5eb8-4809-bd3c-4ebdc3cfbe48","resolution":{"observed_at":"2026-05-17T23:20:27.782409Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-07-13T23:49:58.855418+00:00","source":"crossref_status_cache"},{"observed_at":"2026-07-13T23:49:58.855418+00:00","source":"openalex_status_cache"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2406.20094","last_updated":"2025-05-08T00:24:02Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2024-06-28T17:59:01Z","title":"Scaling Synthetic Data Creation with 1,000,000,000 Personas","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2406.20094","snapshot_observed_at":"2026-08-03T21:09:19.299709Z","title":null,"venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2511.16478","last_updated":"2026-07-20T11:57:11Z","snapshot_observed_at":"2026-08-03T21:09:04.724152Z","submitted_at":"2025-11-20T15:46:27Z","title":"Music Recommendation with Large Language Models: Challenges, Opportunities, and Evaluation","version":2},"reference_index":59,"source":"pdf_text","source_observed_at":"2026-08-03T21:09:19.299709Z"},"links":{"cited_paper":"/paper/2406.20094","citing_paper":"/paper/2511.16478"},"observation_digest":"sha256:bc79cdbedb8c588e0d5c938427e4393b3a0b8d521cda642c67ceceb6ce7ce21e","observation_id":"e3ea5326-418d-4b40-a7ca-f94597f04329","resolution":{"observed_at":"2026-08-03T21:09:19.299709Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2406.20094","last_updated":"2025-05-08T00:24:02Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2024-06-28T17:59:01Z","title":"Scaling Synthetic Data Creation with 1,000,000,000 Personas","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2406.20094","snapshot_observed_at":"2026-08-03T13:38:44.519215Z","title":"Scaling synthetic data creation with 1,000,000,000 personas.CoRR, abs/2406.20094, 2024","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2512.23601","last_updated":"2026-06-20T14:05:44Z","snapshot_observed_at":"2026-08-04T05:31:46.653286Z","submitted_at":"2025-12-29T16:53:48Z","title":"Enhancing Diversity of LLM-Generated Educational Tasks","version":2},"reference_index":9,"source":"pdf_text","source_observed_at":"2026-08-03T13:38:44.519215Z"},"links":{"cited_paper":"/paper/2406.20094","citing_paper":"/paper/2512.23601"},"observation_digest":"sha256:79d14d90d279153eafbbff57f35c07c527b0224a60e4602ea4b06738078549b2","observation_id":"bb044c02-ebe6-4fe3-a145-ad040bc2a053","resolution":{"observed_at":"2026-08-03T13:38:44.519215Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2406.20094","last_updated":"2025-05-08T00:24:02Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2024-06-28T17:59:01Z","title":"Scaling Synthetic Data Creation with 1,000,000,000 Personas","version":3},"cited_work":{"arxiv_id":"2406.20094","doi":"10.48550/arxiv.2406.20094","metadata_source":"pith","pith_arxiv_id":"2406.20094","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Scaling Synthetic Data Creation with 1,000,000,000 Personas","venue":"cs.CL","work_id":"4acdd1d7-4a74-460c-a32c-01c386a5c58a","year":2024},"citing_paper":{"arxiv_id":"2602.08167","last_updated":"2026-05-16T05:42:57Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2026-02-09T00:10:17Z","title":"Self-Supervised Bootstrapping of Action-Predictive Embodied Reasoning","version":2},"reference_index":62,"source":"pdf_text","source_observed_at":"2026-05-21T14:03:48.795572Z"},"links":{"cited_paper":"/paper/2406.20094","citing_paper":"/paper/2602.08167"},"observation_digest":"sha256:85c0c41180c5533cc5632bd447e572ac4729ee9f36c3373983636d1287a19754","observation_id":"a8094486-1029-441b-a8c2-1866168bbebe","resolution":{"observed_at":"2026-05-21T14:04:11.894069Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-07-13T23:49:58.855418+00:00","source":"crossref_status_cache"},{"observed_at":"2026-07-13T23:49:58.855418+00:00","source":"openalex_status_cache"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2406.20094","last_updated":"2025-05-08T00:24:02Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2024-06-28T17:59:01Z","title":"Scaling Synthetic Data Creation with 1,000,000,000 Personas","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2406.20094","snapshot_observed_at":"2026-08-03T01:17:12.004920Z","title":"Scaling synthetic data creation with 1,000,000,000 personas.arXiv preprint arXiv:2406.20094, 2024","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2602.10388","last_updated":"2026-05-29T05:05:29Z","snapshot_observed_at":"2026-08-03T01:17:09.494206Z","submitted_at":"2026-02-11T00:23:13Z","title":"Less is Enough: Synthesizing Diverse Data in LLM Feature Space with Sparse Autoencoders","version":4},"reference_index":59,"source":"pdf_text","source_observed_at":"2026-08-03T01:17:12.004920Z"},"links":{"cited_paper":"/paper/2406.20094","citing_paper":"/paper/2602.10388"},"observation_digest":"sha256:d74a89b533fadfbe0f69911d81079e66ebdbc7cdbabeba3c7d19c14ce916675d","observation_id":"9e73e183-cabb-4450-ac9f-cc1bc4f1de11","resolution":{"observed_at":"2026-08-03T01:17:12.004920Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2406.20094","last_updated":"2025-05-08T00:24:02Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2024-06-28T17:59:01Z","title":"Scaling Synthetic Data Creation with 1,000,000,000 Personas","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2406.20094","snapshot_observed_at":"2026-08-02T23:53:01.051590Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2602.12394","last_updated":"2026-06-28T15:54:56Z","snapshot_observed_at":"2026-08-02T23:52:54.482047Z","submitted_at":"2026-02-12T20:41:22Z","title":"Synthetic Interaction Data for Scalable Personalization in Large Language Models","version":2},"reference_index":7,"source":"pdf_text","source_observed_at":"2026-08-02T23:53:01.051590Z"},"links":{"cited_paper":"/paper/2406.20094","citing_paper":"/paper/2602.12394"},"observation_digest":"sha256:2f93706c1faad86e5874f4965ce0842178b0123f4252f6dc2d8d93b15cfcab48","observation_id":"f3cf075b-86b6-49f4-8464-430c4b889e5e","resolution":{"observed_at":"2026-08-02T23:53:01.051590Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2406.20094","last_updated":"2025-05-08T00:24:02Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2024-06-28T17:59:01Z","title":"Scaling Synthetic Data Creation with 1,000,000,000 Personas","version":3},"cited_work":{"arxiv_id":"2406.20094","doi":"10.48550/arxiv.2406.20094","metadata_source":"pith","pith_arxiv_id":"2406.20094","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Scaling Synthetic Data Creation with 1,000,000,000 Personas","venue":"cs.CL","work_id":"4acdd1d7-4a74-460c-a32c-01c386a5c58a","year":2024},"citing_paper":{"arxiv_id":"2603.24326","last_updated":"2026-04-03T06:49:01Z","snapshot_observed_at":"2026-07-06T22:50:28.640006Z","submitted_at":"2026-03-25T14:08:56Z","title":"Boosting Document Parsing Efficiency and Performance with Coarse-to-Fine Visual Processing","version":2},"reference_index":14,"source":"pdf_text","source_observed_at":"2026-05-15T00:25:19.782732Z"},"links":{"cited_paper":"/paper/2406.20094","citing_paper":"/paper/2603.24326"},"observation_digest":"sha256:a03a852131a4f5053e4f3c9b2cf9fd0298d835f1a78bd00a971cb94ca0356091","observation_id":"c8f5f7af-6802-4b69-9537-ef53c120b5bd","resolution":{"observed_at":"2026-05-16T00:03:55.809567Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-07-13T23:49:58.855418+00:00","source":"crossref_status_cache"},{"observed_at":"2026-07-13T23:49:58.855418+00:00","source":"openalex_status_cache"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2406.20094","last_updated":"2025-05-08T00:24:02Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2024-06-28T17:59:01Z","title":"Scaling Synthetic Data Creation with 1,000,000,000 Personas","version":3},"cited_work":{"arxiv_id":"2406.20094","doi":"10.48550/arxiv.2406.20094","metadata_source":"pith","pith_arxiv_id":"2406.20094","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Scaling Synthetic Data Creation with 1,000,000,000 Personas","venue":"cs.CL","work_id":"4acdd1d7-4a74-460c-a32c-01c386a5c58a","year":2024},"citing_paper":{"arxiv_id":"2604.02522","last_updated":"2026-04-02T21:23:00Z","snapshot_observed_at":"2026-08-03T01:43:37.179078Z","submitted_at":"2026-04-02T21:23:00Z","title":"Opal: Private Memory for Personal AI","version":1},"reference_index":81,"source":"pdf_text","source_observed_at":"2026-05-13T21:09:06.320543Z"},"links":{"cited_paper":"/paper/2406.20094","citing_paper":"/paper/2604.02522"},"observation_digest":"sha256:667cc87a74cc1dbb4bb6c71c28d6d9ed1794254f42a1b0711195acacc9f08154","observation_id":"46cb934c-e54d-4e2b-8312-b7f52442b861","resolution":{"observed_at":"2026-05-16T00:03:55.809567Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-07-13T23:49:58.855418+00:00","source":"crossref_status_cache"},{"observed_at":"2026-07-13T23:49:58.855418+00:00","source":"openalex_status_cache"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2406.20094","last_updated":"2025-05-08T00:24:02Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2024-06-28T17:59:01Z","title":"Scaling Synthetic Data Creation with 1,000,000,000 Personas","version":3},"cited_work":{"arxiv_id":"2406.20094","doi":"10.48550/arxiv.2406.20094","metadata_source":"pith","pith_arxiv_id":"2406.20094","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Scaling Synthetic Data Creation with 1,000,000,000 Personas","venue":"cs.CL","work_id":"4acdd1d7-4a74-460c-a32c-01c386a5c58a","year":2024},"citing_paper":{"arxiv_id":"2604.02794","last_updated":"2026-04-03T07:02:13Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2026-04-03T07:02:13Z","title":"CharTool: Tool-Integrated Visual Reasoning for Chart Understanding","version":1},"reference_index":8,"source":"arxiv_source","source_observed_at":"2026-05-13T20:07:23.153064Z"},"links":{"cited_paper":"/paper/2406.20094","citing_paper":"/paper/2604.02794"},"observation_digest":"sha256:b2c09bbc0f01177f06fa119eac7b27b0132e587cfa35cb84365034350e338d28","observation_id":"e4cd8f29-0e3a-40ec-a43f-f9392fd8bb24","resolution":{"observed_at":"2026-05-16T00:03:55.809567Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-07-13T23:49:58.855418+00:00","source":"crossref_status_cache"},{"observed_at":"2026-07-13T23:49:58.855418+00:00","source":"openalex_status_cache"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2406.20094","last_updated":"2025-05-08T00:24:02Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2024-06-28T17:59:01Z","title":"Scaling Synthetic Data Creation with 1,000,000,000 Personas","version":3},"cited_work":{"arxiv_id":"2406.20094","doi":"10.48550/arxiv.2406.20094","metadata_source":"pith","pith_arxiv_id":"2406.20094","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Scaling Synthetic Data Creation with 1,000,000,000 Personas","venue":"cs.CL","work_id":"4acdd1d7-4a74-460c-a32c-01c386a5c58a","year":2024},"citing_paper":{"arxiv_id":"2604.06204","last_updated":"2026-03-15T09:21:21Z","snapshot_observed_at":"2026-08-02T08:11:49.656958Z","submitted_at":"2026-03-15T09:21:21Z","title":"SensorPersona: An LLM-Empowered System for Continual Persona Extraction from Longitudinal Mobile Sensor Streams","version":1},"reference_index":16,"source":"pdf_text","source_observed_at":"2026-05-15T11:45:33.612135Z"},"links":{"cited_paper":"/paper/2406.20094","citing_paper":"/paper/2604.06204"},"observation_digest":"sha256:f8ecafe83aaa43f137f755d2fce7771ee03891dae45baaf535beddaea943e167","observation_id":"81fd992b-0e34-427a-89f0-c2be7ad79168","resolution":{"observed_at":"2026-05-16T00:03:55.809567Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-07-13T23:49:58.855418+00:00","source":"crossref_status_cache"},{"observed_at":"2026-07-13T23:49:58.855418+00:00","source":"openalex_status_cache"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2406.20094","last_updated":"2025-05-08T00:24:02Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2024-06-28T17:59:01Z","title":"Scaling Synthetic Data Creation with 1,000,000,000 Personas","version":3},"cited_work":{"arxiv_id":"2406.20094","doi":"10.48550/arxiv.2406.20094","metadata_source":"pith","pith_arxiv_id":"2406.20094","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Scaling Synthetic Data Creation with 1,000,000,000 Personas","venue":"cs.CL","work_id":"4acdd1d7-4a74-460c-a32c-01c386a5c58a","year":2024},"citing_paper":{"arxiv_id":"2604.08516","last_updated":"2026-04-09T17:54:02Z","snapshot_observed_at":"2026-07-06T22:57:35.435713Z","submitted_at":"2026-04-09T17:54:02Z","title":"MolmoWeb: Open Visual Web Agent and Open Data for the Open Web","version":1},"reference_index":73,"source":"pdf_text","source_observed_at":"2026-05-10T18:00:34.401698Z"},"links":{"cited_paper":"/paper/2406.20094","citing_paper":"/paper/2604.08516"},"observation_digest":"sha256:a7851879549f75dec2142a0b8ecde8800666b3341675096ff183a13970a69579","observation_id":"1419206a-e062-43e1-ba23-154244cb423b","resolution":{"observed_at":"2026-05-16T00:03:55.809567Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-07-13T23:49:58.855418+00:00","source":"crossref_status_cache"},{"observed_at":"2026-07-13T23:49:58.855418+00:00","source":"openalex_status_cache"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2406.20094","last_updated":"2025-05-08T00:24:02Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2024-06-28T17:59:01Z","title":"Scaling Synthetic Data Creation with 1,000,000,000 Personas","version":3},"cited_work":{"arxiv_id":"2406.20094","doi":"10.48550/arxiv.2406.20094","metadata_source":"pith","pith_arxiv_id":"2406.20094","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Scaling Synthetic Data Creation with 1,000,000,000 Personas","venue":"cs.CL","work_id":"4acdd1d7-4a74-460c-a32c-01c386a5c58a","year":2024},"citing_paper":{"arxiv_id":"2604.10733","last_updated":"2026-04-12T17:12:55Z","snapshot_observed_at":"2026-07-06T22:59:18.756089Z","submitted_at":"2026-04-12T17:12:55Z","title":"Too Nice to Tell the Truth: Quantifying Agreeableness-Driven Sycophancy in Role-Playing Language Models","version":1},"reference_index":18,"source":"arxiv_source","source_observed_at":"2026-05-10T16:05:09.033412Z"},"links":{"cited_paper":"/paper/2406.20094","citing_paper":"/paper/2604.10733"},"observation_digest":"sha256:2d0ebfc54aafbc69b6467dd8596e89e4a805d41fab0405f82abea439295401a8","observation_id":"4e8e1731-03b8-4779-8fa5-82ee9ab98b8a","resolution":{"observed_at":"2026-05-16T00:03:55.809567Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-07-13T23:49:58.855418+00:00","source":"crossref_status_cache"},{"observed_at":"2026-07-13T23:49:58.855418+00:00","source":"openalex_status_cache"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2406.20094","last_updated":"2025-05-08T00:24:02Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2024-06-28T17:59:01Z","title":"Scaling Synthetic Data Creation with 1,000,000,000 Personas","version":3},"cited_work":{"arxiv_id":"2406.20094","doi":"10.48550/arxiv.2406.20094","metadata_source":"pith","pith_arxiv_id":"2406.20094","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Scaling Synthetic Data Creation with 1,000,000,000 Personas","venue":"cs.CL","work_id":"4acdd1d7-4a74-460c-a32c-01c386a5c58a","year":2024},"citing_paper":{"arxiv_id":"2604.13074","last_updated":"2026-03-20T17:59:57Z","snapshot_observed_at":"2026-07-06T23:01:10.333101Z","submitted_at":"2026-03-20T17:59:57Z","title":"PersonaVLM: Long-Term Personalized Multimodal LLMs","version":1},"reference_index":9,"source":"pdf_text","source_observed_at":"2026-05-15T08:02:10.327523Z"},"links":{"cited_paper":"/paper/2406.20094","citing_paper":"/paper/2604.13074"},"observation_digest":"sha256:71b3d0fa31c53a782c1d4bf9c9ff1962c4f515abcc2cd83428f4096fe75c79dc","observation_id":"15e9ee39-a79d-49df-8528-0a3ca7819958","resolution":{"observed_at":"2026-05-16T00:03:55.809567Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-07-13T23:49:58.855418+00:00","source":"crossref_status_cache"},{"observed_at":"2026-07-13T23:49:58.855418+00:00","source":"openalex_status_cache"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2406.20094","last_updated":"2025-05-08T00:24:02Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2024-06-28T17:59:01Z","title":"Scaling Synthetic Data Creation with 1,000,000,000 Personas","version":3},"cited_work":{"arxiv_id":"2406.20094","doi":"10.48550/arxiv.2406.20094","metadata_source":"pith","pith_arxiv_id":"2406.20094","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Scaling Synthetic Data Creation with 1,000,000,000 Personas","venue":"cs.CL","work_id":"4acdd1d7-4a74-460c-a32c-01c386a5c58a","year":2024},"citing_paper":{"arxiv_id":"2604.15675","last_updated":"2026-04-17T03:54:12Z","snapshot_observed_at":"2026-07-06T23:03:12.000381Z","submitted_at":"2026-04-17T03:54:12Z","title":"C-Mining: Unsupervised Discovery of Seeds for Cultural Data Synthesis via Geometric Misalignment","version":1},"reference_index":10,"source":"pdf_text","source_observed_at":"2026-05-10T09:49:20.724437Z"},"links":{"cited_paper":"/paper/2406.20094","citing_paper":"/paper/2604.15675"},"observation_digest":"sha256:fda4f301dfd447625ec27ce2f2e6d995b268a3894ff64fdeace08fb793f70efc","observation_id":"5b2dde8d-8a8d-49da-be26-f936913a66c1","resolution":{"observed_at":"2026-05-16T00:03:55.809567Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-07-13T23:49:58.855418+00:00","source":"crossref_status_cache"},{"observed_at":"2026-07-13T23:49:58.855418+00:00","source":"openalex_status_cache"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2406.20094","last_updated":"2025-05-08T00:24:02Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2024-06-28T17:59:01Z","title":"Scaling Synthetic Data Creation with 1,000,000,000 Personas","version":3},"cited_work":{"arxiv_id":"2406.20094","doi":"10.48550/arxiv.2406.20094","metadata_source":"pith","pith_arxiv_id":"2406.20094","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Scaling Synthetic Data Creation with 1,000,000,000 Personas","venue":"cs.CL","work_id":"4acdd1d7-4a74-460c-a32c-01c386a5c58a","year":2024},"citing_paper":{"arxiv_id":"2605.03547","last_updated":"2026-05-05T09:18:22Z","snapshot_observed_at":"2026-07-06T23:16:25.301792Z","submitted_at":"2026-05-05T09:18:22Z","title":"Erase Persona, Forget Lore: Benchmarking Multimodal Copyright Unlearning in Large Vision Language Models","version":1},"reference_index":19,"source":"arxiv_source","source_observed_at":"2026-05-07T17:58:17.879345Z"},"links":{"cited_paper":"/paper/2406.20094","citing_paper":"/paper/2605.03547"},"observation_digest":"sha256:4565c656756f564b8310e3cfff7efbba45b2d54be254ff99bac45128997f1b03","observation_id":"5f7d8027-2bcd-45ec-bc22-27608bd2cbce","resolution":{"observed_at":"2026-05-16T00:03:55.809567Z","resolver_source":"arxiv_id","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-07-13T23:49:58.855418+00:00","source":"crossref_status_cache"},{"observed_at":"2026-07-13T23:49:58.855418+00:00","source":"openalex_status_cache"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2406.20094","last_updated":"2025-05-08T00:24:02Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2024-06-28T17:59:01Z","title":"Scaling Synthetic Data Creation with 1,000,000,000 Personas","version":3},"cited_work":{"arxiv_id":"2406.20094","doi":"10.48550/arxiv.2406.20094","metadata_source":"pith","pith_arxiv_id":"2406.20094","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Scaling Synthetic Data Creation with 1,000,000,000 Personas","venue":"cs.CL","work_id":"4acdd1d7-4a74-460c-a32c-01c386a5c58a","year":2024},"citing_paper":{"arxiv_id":"2605.04018","last_updated":"2026-05-05T17:42:50Z","snapshot_observed_at":"2026-07-06T23:16:49.553583Z","submitted_at":"2026-05-05T17:42:50Z","title":"Rethinking Reasoning-Intensive Retrieval: Evaluating and Advancing Retrievers in Agentic Search Systems","version":1},"reference_index":21,"source":"pdf_text","source_observed_at":"2026-05-07T16:06:51.876295Z"},"links":{"cited_paper":"/paper/2406.20094","citing_paper":"/paper/2605.04018"},"observation_digest":"sha256:648e380246c6c92bc31d3381180175959abee2842b835a2c82db9c6dbe564e09","observation_id":"30487319-1ea2-4bae-89bb-36a4d78cead6","resolution":{"observed_at":"2026-05-16T00:03:55.809567Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-07-13T23:49:58.855418+00:00","source":"crossref_status_cache"},{"observed_at":"2026-07-13T23:49:58.855418+00:00","source":"openalex_status_cache"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2406.20094","last_updated":"2025-05-08T00:24:02Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2024-06-28T17:59:01Z","title":"Scaling Synthetic Data Creation with 1,000,000,000 Personas","version":3},"cited_work":{"arxiv_id":"2406.20094","doi":"10.48550/arxiv.2406.20094","metadata_source":"pith","pith_arxiv_id":"2406.20094","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Scaling Synthetic Data Creation with 1,000,000,000 Personas","venue":"cs.CL","work_id":"4acdd1d7-4a74-460c-a32c-01c386a5c58a","year":2024},"citing_paper":{"arxiv_id":"2605.07699","last_updated":"2026-07-31T14:41:42Z","snapshot_observed_at":"2026-08-05T07:17:48.989888Z","submitted_at":"2026-05-08T13:10:49Z","title":"DRIP-R: A Benchmark for Decision-Making and Reasoning Under Real-World Policy Ambiguity in the Retail Domain","version":1},"reference_index":3,"source":"arxiv_source","source_observed_at":"2026-05-11T02:40:27.234973Z"},"links":{"cited_paper":"/paper/2406.20094","citing_paper":"/paper/2605.07699"},"observation_digest":"sha256:377ed6832c3fe82a90d943a8b86465607dfbc58362230a6f0c487eddf4e26370","observation_id":"df269eda-67cb-4254-ab5e-ed4a05930d0a","resolution":{"observed_at":"2026-05-16T00:03:55.809567Z","resolver_source":"arxiv_id","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-07-13T23:49:58.855418+00:00","source":"crossref_status_cache"},{"observed_at":"2026-07-13T23:49:58.855418+00:00","source":"openalex_status_cache"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2406.20094","last_updated":"2025-05-08T00:24:02Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2024-06-28T17:59:01Z","title":"Scaling Synthetic Data Creation with 1,000,000,000 Personas","version":3},"cited_work":{"arxiv_id":"2406.20094","doi":"10.48550/arxiv.2406.20094","metadata_source":"pith","pith_arxiv_id":"2406.20094","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Scaling Synthetic Data Creation with 1,000,000,000 Personas","venue":"cs.CL","work_id":"4acdd1d7-4a74-460c-a32c-01c386a5c58a","year":2024},"citing_paper":{"arxiv_id":"2605.08766","last_updated":"2026-05-09T07:51:03Z","snapshot_observed_at":"2026-08-02T11:50:07.900845Z","submitted_at":"2026-05-09T07:51:03Z","title":"UserGPT Technical Report","version":1},"reference_index":62,"source":"arxiv_source","source_observed_at":"2026-05-12T03:10:51.555653Z"},"links":{"cited_paper":"/paper/2406.20094","citing_paper":"/paper/2605.08766"},"observation_digest":"sha256:2837dede862db993aec67cb43965bceae48e2a4334a4b0cfe096225bd1081c8b","observation_id":"1a0e6fba-c3c9-4a8c-a369-96ea774e7c20","resolution":{"observed_at":"2026-05-16T00:03:55.809567Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-07-13T23:49:58.855418+00:00","source":"crossref_status_cache"},{"observed_at":"2026-07-13T23:49:58.855418+00:00","source":"openalex_status_cache"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2406.20094","last_updated":"2025-05-08T00:24:02Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2024-06-28T17:59:01Z","title":"Scaling Synthetic Data Creation with 1,000,000,000 Personas","version":3},"cited_work":{"arxiv_id":"2406.20094","doi":"10.48550/arxiv.2406.20094","metadata_source":"pith","pith_arxiv_id":"2406.20094","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Scaling Synthetic Data Creation with 1,000,000,000 Personas","venue":"cs.CL","work_id":"4acdd1d7-4a74-460c-a32c-01c386a5c58a","year":2024},"citing_paper":{"arxiv_id":"2605.09530","last_updated":"2026-05-14T14:55:04Z","snapshot_observed_at":"2026-07-06T23:21:35.327066Z","submitted_at":"2026-05-10T13:31:58Z","title":"MemPrivacy: Privacy-Preserving Personalized Memory Management for Edge-Cloud Agents","version":1},"reference_index":12,"source":"pdf_text","source_observed_at":"2026-05-12T04:07:27.168365Z"},"links":{"cited_paper":"/paper/2406.20094","citing_paper":"/paper/2605.09530"},"observation_digest":"sha256:c87a32920429b55c20b1527fc74fdf550838f42276630908a506a8931ab4a792","observation_id":"6bb0310a-56d3-4bf1-9bcf-ee7f1aac9d6e","resolution":{"observed_at":"2026-05-16T00:03:55.809567Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-07-13T23:49:58.855418+00:00","source":"crossref_status_cache"},{"observed_at":"2026-07-13T23:49:58.855418+00:00","source":"openalex_status_cache"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2406.20094","last_updated":"2025-05-08T00:24:02Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2024-06-28T17:59:01Z","title":"Scaling Synthetic Data Creation with 1,000,000,000 Personas","version":3},"cited_work":{"arxiv_id":"2406.20094","doi":"10.48550/arxiv.2406.20094","metadata_source":"pith","pith_arxiv_id":"2406.20094","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Scaling Synthetic Data Creation with 1,000,000,000 Personas","venue":"cs.CL","work_id":"4acdd1d7-4a74-460c-a32c-01c386a5c58a","year":2024},"citing_paper":{"arxiv_id":"2605.09530","last_updated":"2026-05-14T14:55:04Z","snapshot_observed_at":"2026-07-06T23:21:35.327066Z","submitted_at":"2026-05-10T13:31:58Z","title":"MemPrivacy: Privacy-Preserving Personalized Memory Management for Edge-Cloud Agents","version":2},"reference_index":12,"source":"pdf_text","source_observed_at":"2026-05-13T07:23:35.728424Z"},"links":{"cited_paper":"/paper/2406.20094","citing_paper":"/paper/2605.09530"},"observation_digest":"sha256:91437f231254abe8db33e7b9345f5c4e642a293f248d7e3ca4d444f9e9ed3167","observation_id":"60c44765-f7d1-45cd-acce-451c152a8f74","resolution":{"observed_at":"2026-05-16T00:03:55.809567Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-07-13T23:49:58.855418+00:00","source":"crossref_status_cache"},{"observed_at":"2026-07-13T23:49:58.855418+00:00","source":"openalex_status_cache"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2406.20094","last_updated":"2025-05-08T00:24:02Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2024-06-28T17:59:01Z","title":"Scaling Synthetic Data Creation with 1,000,000,000 Personas","version":3},"cited_work":{"arxiv_id":"2406.20094","doi":"10.48550/arxiv.2406.20094","metadata_source":"pith","pith_arxiv_id":"2406.20094","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Scaling Synthetic Data Creation with 1,000,000,000 Personas","venue":"cs.CL","work_id":"4acdd1d7-4a74-460c-a32c-01c386a5c58a","year":2024},"citing_paper":{"arxiv_id":"2605.09530","last_updated":"2026-05-14T14:55:04Z","snapshot_observed_at":"2026-07-06T23:21:35.327066Z","submitted_at":"2026-05-10T13:31:58Z","title":"MemPrivacy: Privacy-Preserving Personalized Memory Management for Edge-Cloud Agents","version":3},"reference_index":12,"source":"pdf_text","source_observed_at":"2026-05-15T05:44:33.434389Z"},"links":{"cited_paper":"/paper/2406.20094","citing_paper":"/paper/2605.09530"},"observation_digest":"sha256:d788df55dfd0cf7c3c69b566da404b2f7a3b8b6acbd1937996948de33674e54f","observation_id":"028d9d72-a45b-4cff-8f03-4e295234ab8a","resolution":{"observed_at":"2026-05-16T00:03:55.809567Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-07-13T23:49:58.855418+00:00","source":"crossref_status_cache"},{"observed_at":"2026-07-13T23:49:58.855418+00:00","source":"openalex_status_cache"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2406.20094","last_updated":"2025-05-08T00:24:02Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2024-06-28T17:59:01Z","title":"Scaling Synthetic Data Creation with 1,000,000,000 Personas","version":3},"cited_work":{"arxiv_id":"2406.20094","doi":"10.48550/arxiv.2406.20094","metadata_source":"pith","pith_arxiv_id":"2406.20094","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Scaling Synthetic Data Creation with 1,000,000,000 Personas","venue":"cs.CL","work_id":"4acdd1d7-4a74-460c-a32c-01c386a5c58a","year":2024},"citing_paper":{"arxiv_id":"2605.09808","last_updated":"2026-05-10T23:06:24Z","snapshot_observed_at":"2026-08-02T12:27:41.948284Z","submitted_at":"2026-05-10T23:06:24Z","title":"Quantifying the Utility of User Simulators for Building Collaborative LLM Assistants","version":1},"reference_index":21,"source":"pdf_text","source_observed_at":"2026-05-12T02:28:13.317630Z"},"links":{"cited_paper":"/paper/2406.20094","citing_paper":"/paper/2605.09808"},"observation_digest":"sha256:03bd3e62f1ce363b0d70051297798622e9697f6dcfe230ecae335bb865f789b3","observation_id":"2fdc7cf3-434f-4d61-b2df-e2e0e80f69b7","resolution":{"observed_at":"2026-05-16T00:03:55.809567Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-07-13T23:49:58.855418+00:00","source":"crossref_status_cache"},{"observed_at":"2026-07-13T23:49:58.855418+00:00","source":"openalex_status_cache"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2406.20094","last_updated":"2025-05-08T00:24:02Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2024-06-28T17:59:01Z","title":"Scaling Synthetic Data Creation with 1,000,000,000 Personas","version":3},"cited_work":{"arxiv_id":"2406.20094","doi":"10.48550/arxiv.2406.20094","metadata_source":"pith","pith_arxiv_id":"2406.20094","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Scaling Synthetic Data Creation with 1,000,000,000 Personas","venue":"cs.CL","work_id":"4acdd1d7-4a74-460c-a32c-01c386a5c58a","year":2024},"citing_paper":{"arxiv_id":"2605.12894","last_updated":"2026-05-13T02:16:51Z","snapshot_observed_at":"2026-08-02T12:05:13.301743Z","submitted_at":"2026-05-13T02:16:51Z","title":"Beyond Cooperative Simulators: Generating Realistic User Personas for Robust Evaluation of LLM Agents","version":1},"reference_index":4,"source":"pdf_text","source_observed_at":"2026-05-14T20:24:58.741662Z"},"links":{"cited_paper":"/paper/2406.20094","citing_paper":"/paper/2605.12894"},"observation_digest":"sha256:45c76dadaadeab76d4229cae35a7516a04550243db24586dfdb3fcba9167f798","observation_id":"dc8dbb9d-f663-496b-8c10-49a5b51a571c","resolution":{"observed_at":"2026-05-16T00:03:55.809567Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-07-13T23:49:58.855418+00:00","source":"crossref_status_cache"},{"observed_at":"2026-07-13T23:49:58.855418+00:00","source":"openalex_status_cache"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2406.20094","last_updated":"2025-05-08T00:24:02Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2024-06-28T17:59:01Z","title":"Scaling Synthetic Data Creation with 1,000,000,000 Personas","version":3},"cited_work":{"arxiv_id":"2406.20094","doi":"10.48550/arxiv.2406.20094","metadata_source":"pith","pith_arxiv_id":"2406.20094","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Scaling Synthetic Data Creation with 1,000,000,000 Personas","venue":"cs.CL","work_id":"4acdd1d7-4a74-460c-a32c-01c386a5c58a","year":2024},"citing_paper":{"arxiv_id":"2605.20876","last_updated":"2026-05-20T08:14:51Z","snapshot_observed_at":"2026-07-06T23:31:23.407780Z","submitted_at":"2026-05-20T08:14:51Z","title":"Terminal-World: Scaling Terminal-Agent Environments via Agent Skills","version":1},"reference_index":24,"source":"pdf_text","source_observed_at":"2026-05-21T04:54:17.662339Z"},"links":{"cited_paper":"/paper/2406.20094","citing_paper":"/paper/2605.20876"},"observation_digest":"sha256:b6dcf1bf842a53f722cdd905c9fd3c3d2729d3468ae48db8091794db6ba56a96","observation_id":"b444fac8-5ebf-4b38-b6eb-a17f6ac10a7f","resolution":{"observed_at":"2026-05-21T04:54:35.815250Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-07-13T23:49:58.855418+00:00","source":"crossref_status_cache"},{"observed_at":"2026-07-13T23:49:58.855418+00:00","source":"openalex_status_cache"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2406.20094","last_updated":"2025-05-08T00:24:02Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2024-06-28T17:59:01Z","title":"Scaling Synthetic Data Creation with 1,000,000,000 Personas","version":3},"cited_work":{"arxiv_id":"2406.20094","doi":"10.48550/arxiv.2406.20094","metadata_source":"pith","pith_arxiv_id":"2406.20094","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Scaling Synthetic Data Creation with 1,000,000,000 Personas","venue":"cs.CL","work_id":"4acdd1d7-4a74-460c-a32c-01c386a5c58a","year":2024},"citing_paper":{"arxiv_id":"2605.25052","last_updated":"2026-05-24T12:57:01Z","snapshot_observed_at":"2026-07-06T23:35:01.860096Z","submitted_at":"2026-05-24T12:57:01Z","title":"Faithfulness Metrics Don't Measure Faithfulness: A Meta-Evaluation with Ground Truth","version":1},"reference_index":59,"source":"pdf_text","source_observed_at":"2026-06-30T11:56:53.355299Z"},"links":{"cited_paper":"/paper/2406.20094","citing_paper":"/paper/2605.25052"},"observation_digest":"sha256:ee38d5c792bd96c22b8e076279258527f6c2520f01837875a52532f5ecb12ffc","observation_id":"2dbd1c1c-dd4f-400d-9708-2266eb49206e","resolution":{"observed_at":"2026-06-30T12:04:39.022762Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-07-13T23:49:58.855418+00:00","source":"crossref_status_cache"},{"observed_at":"2026-07-13T23:49:58.855418+00:00","source":"openalex_status_cache"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2406.20094","last_updated":"2025-05-08T00:24:02Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2024-06-28T17:59:01Z","title":"Scaling Synthetic Data Creation with 1,000,000,000 Personas","version":3},"cited_work":{"arxiv_id":"2406.20094","doi":"10.48550/arxiv.2406.20094","metadata_source":"pith","pith_arxiv_id":"2406.20094","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Scaling Synthetic Data Creation with 1,000,000,000 Personas","venue":"cs.CL","work_id":"4acdd1d7-4a74-460c-a32c-01c386a5c58a","year":2024},"citing_paper":{"arxiv_id":"2605.25680","last_updated":"2026-05-25T10:39:08Z","snapshot_observed_at":"2026-08-02T02:16:17.525704Z","submitted_at":"2026-05-25T10:39:08Z","title":"Simulating Human Memory with Language Models","version":1},"reference_index":47,"source":"pdf_text","source_observed_at":"2026-06-29T22:00:26.703406Z"},"links":{"cited_paper":"/paper/2406.20094","citing_paper":"/paper/2605.25680"},"observation_digest":"sha256:e84f5f2c2af8bbd0bd0d5c441873b0a17cc221bd4bcf0eeaca93648c84e9a83a","observation_id":"e15ef9f5-fccb-4cf4-b287-52398e463bb4","resolution":{"observed_at":"2026-06-29T22:04:00.318452Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-07-13T23:49:58.855418+00:00","source":"crossref_status_cache"},{"observed_at":"2026-07-13T23:49:58.855418+00:00","source":"openalex_status_cache"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2406.20094","last_updated":"2025-05-08T00:24:02Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2024-06-28T17:59:01Z","title":"Scaling Synthetic Data Creation with 1,000,000,000 Personas","version":3},"cited_work":{"arxiv_id":"2406.20094","doi":"10.48550/arxiv.2406.20094","metadata_source":"pith","pith_arxiv_id":"2406.20094","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Scaling Synthetic Data Creation with 1,000,000,000 Personas","venue":"cs.CL","work_id":"4acdd1d7-4a74-460c-a32c-01c386a5c58a","year":2024},"citing_paper":{"arxiv_id":"2605.26046","last_updated":"2026-06-03T18:01:05Z","snapshot_observed_at":"2026-07-06T23:35:55.988443Z","submitted_at":"2026-05-25T17:08:55Z","title":"When Gradients Collide: Failure Modes of Multi-Objective Prompt Optimization for LLM Judges","version":2},"reference_index":1,"source":"pdf_text","source_observed_at":"2026-06-29T21:13:44.154803Z"},"links":{"cited_paper":"/paper/2406.20094","citing_paper":"/paper/2605.26046"},"observation_digest":"sha256:0a9f4d8643ab0f8bdf4708539db7eaa200d11d84098e87cbb18b3dfd88d710dc","observation_id":"ed2660eb-aa86-4444-a15b-e0f4c351a8a7","resolution":{"observed_at":"2026-06-29T21:13:59.407908Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-07-13T23:49:58.855418+00:00","source":"crossref_status_cache"},{"observed_at":"2026-07-13T23:49:58.855418+00:00","source":"openalex_status_cache"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2406.20094","last_updated":"2025-05-08T00:24:02Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2024-06-28T17:59:01Z","title":"Scaling Synthetic Data Creation with 1,000,000,000 Personas","version":3},"cited_work":{"arxiv_id":"2406.20094","doi":"10.48550/arxiv.2406.20094","metadata_source":"pith","pith_arxiv_id":"2406.20094","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Scaling Synthetic Data Creation with 1,000,000,000 Personas","venue":"cs.CL","work_id":"4acdd1d7-4a74-460c-a32c-01c386a5c58a","year":2024},"citing_paper":{"arxiv_id":"2605.28806","last_updated":"2026-05-27T17:56:11Z","snapshot_observed_at":"2026-07-06T23:38:19.096349Z","submitted_at":"2026-05-27T17:56:11Z","title":"Personal Visual Memory from Explicit and Implicit Evidence","version":1},"reference_index":35,"source":"pdf_text","source_observed_at":"2026-06-29T12:40:34.740804Z"},"links":{"cited_paper":"/paper/2406.20094","citing_paper":"/paper/2605.28806"},"observation_digest":"sha256:76df8f683de7443c1354bf8895cc0d57eba066bbad577839ff2eb6071e7f85ce","observation_id":"10d70c7d-e17a-48db-822b-86b75ab1d859","resolution":{"observed_at":"2026-06-29T12:43:25.427765Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-07-13T23:49:58.855418+00:00","source":"crossref_status_cache"},{"observed_at":"2026-07-13T23:49:58.855418+00:00","source":"openalex_status_cache"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2406.20094","last_updated":"2025-05-08T00:24:02Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2024-06-28T17:59:01Z","title":"Scaling Synthetic Data Creation with 1,000,000,000 Personas","version":3},"cited_work":{"arxiv_id":"2406.20094","doi":"10.48550/arxiv.2406.20094","metadata_source":"pith","pith_arxiv_id":"2406.20094","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Scaling Synthetic Data Creation with 1,000,000,000 Personas","venue":"cs.CL","work_id":"4acdd1d7-4a74-460c-a32c-01c386a5c58a","year":2024},"citing_paper":{"arxiv_id":"2605.30058","last_updated":"2026-05-28T15:08:03Z","snapshot_observed_at":"2026-08-04T07:31:13.672921Z","submitted_at":"2026-05-28T15:08:03Z","title":"HEART-Bench: Do LLM Agents Exhibit Human-like Psychology?","version":1},"reference_index":17,"source":"pdf_text","source_observed_at":"2026-06-29T07:18:52.517792Z"},"links":{"cited_paper":"/paper/2406.20094","citing_paper":"/paper/2605.30058"},"observation_digest":"sha256:1f7dbe8d214656f6ec2f66e709959215e8eac71215ad8bcea8133e06e6ebf054","observation_id":"3dcb5c2f-6afd-4f98-b934-09f553f0b761","resolution":{"observed_at":"2026-06-29T07:23:12.808699Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-07-13T23:49:58.855418+00:00","source":"crossref_status_cache"},{"observed_at":"2026-07-13T23:49:58.855418+00:00","source":"openalex_status_cache"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2406.20094","last_updated":"2025-05-08T00:24:02Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2024-06-28T17:59:01Z","title":"Scaling Synthetic Data Creation with 1,000,000,000 Personas","version":3},"cited_work":{"arxiv_id":"2406.20094","doi":"10.48550/arxiv.2406.20094","metadata_source":"pith","pith_arxiv_id":"2406.20094","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Scaling Synthetic Data Creation with 1,000,000,000 Personas","venue":"cs.CL","work_id":"4acdd1d7-4a74-460c-a32c-01c386a5c58a","year":2024},"citing_paper":{"arxiv_id":"2605.30169","last_updated":"2026-07-02T23:58:01Z","snapshot_observed_at":"2026-08-01T18:48:30.510137Z","submitted_at":"2026-05-28T16:20:19Z","title":"Dissociative Identity: Language Model Agents Lack Grounding for Reputation Mechanisms","version":2},"reference_index":55,"source":"pdf_text","source_observed_at":"2026-06-29T00:26:54.019256Z"},"links":{"cited_paper":"/paper/2406.20094","citing_paper":"/paper/2605.30169"},"observation_digest":"sha256:a067028fe3120360ef1b69830d211367f0f5106ee1b8636603aabf202f9df728","observation_id":"a929e5a6-9e8d-418a-ac40-fa27115d136d","resolution":{"observed_at":"2026-06-29T00:32:53.172356Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-07-13T23:49:58.855418+00:00","source":"crossref_status_cache"},{"observed_at":"2026-07-13T23:49:58.855418+00:00","source":"openalex_status_cache"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2406.20094","last_updated":"2025-05-08T00:24:02Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2024-06-28T17:59:01Z","title":"Scaling Synthetic Data Creation with 1,000,000,000 Personas","version":3},"cited_work":{"arxiv_id":"2406.20094","doi":"10.48550/arxiv.2406.20094","metadata_source":"pith","pith_arxiv_id":"2406.20094","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Scaling Synthetic Data Creation with 1,000,000,000 Personas","venue":"cs.CL","work_id":"4acdd1d7-4a74-460c-a32c-01c386a5c58a","year":2024},"citing_paper":{"arxiv_id":"2605.30207","last_updated":"2026-05-28T16:43:38Z","snapshot_observed_at":"2026-07-30T00:42:25.159554Z","submitted_at":"2026-05-28T16:43:38Z","title":"Persona Conditioning of Brand Recommendations in Retrieval-Augmented Commercial Chat: A Prominence-Stratified Cross-Provider Audit","version":1},"reference_index":8,"source":"pdf_text","source_observed_at":"2026-06-29T07:20:32.538522Z"},"links":{"cited_paper":"/paper/2406.20094","citing_paper":"/paper/2605.30207"},"observation_digest":"sha256:7adc84e30f156bced9bf8785826953f68a76a5f71d871288b5cd0d8107a0d7b1","observation_id":"2f3e26ed-c5f6-45b8-bf7b-42c69ef23b23","resolution":{"observed_at":"2026-06-29T07:23:12.706660Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-07-13T23:49:58.855418+00:00","source":"crossref_status_cache"},{"observed_at":"2026-07-13T23:49:58.855418+00:00","source":"openalex_status_cache"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2406.20094","last_updated":"2025-05-08T00:24:02Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2024-06-28T17:59:01Z","title":"Scaling Synthetic Data Creation with 1,000,000,000 Personas","version":3},"cited_work":{"arxiv_id":"2406.20094","doi":"10.48550/arxiv.2406.20094","metadata_source":"pith","pith_arxiv_id":"2406.20094","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Scaling Synthetic Data Creation with 1,000,000,000 Personas","venue":"cs.CL","work_id":"4acdd1d7-4a74-460c-a32c-01c386a5c58a","year":2024},"citing_paper":{"arxiv_id":"2606.00860","last_updated":"2026-05-30T19:20:26Z","snapshot_observed_at":"2026-07-10T23:18:19.841675Z","submitted_at":"2026-05-30T19:20:26Z","title":"GenPT: Beyond Self-Report for Reliable LLM Psychometrics via Generative Projective Testing","version":1},"reference_index":55,"source":"arxiv_source","source_observed_at":"2026-06-28T17:47:27.191611Z"},"links":{"cited_paper":"/paper/2406.20094","citing_paper":"/paper/2606.00860"},"observation_digest":"sha256:6bf42a054611cde838845fc1e9d7a4f80310e2ffa1f0721fd7ec843c4e8ace33","observation_id":"d9f6851a-e2c2-4b50-8466-e4e9318fd2dc","resolution":{"observed_at":"2026-06-28T17:52:26.535763Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-07-13T23:49:58.855418+00:00","source":"crossref_status_cache"},{"observed_at":"2026-07-13T23:49:58.855418+00:00","source":"openalex_status_cache"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2406.20094","last_updated":"2025-05-08T00:24:02Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2024-06-28T17:59:01Z","title":"Scaling Synthetic Data Creation with 1,000,000,000 Personas","version":3},"cited_work":{"arxiv_id":"2406.20094","doi":"10.48550/arxiv.2406.20094","metadata_source":"pith","pith_arxiv_id":"2406.20094","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Scaling Synthetic Data Creation with 1,000,000,000 Personas","venue":"cs.CL","work_id":"4acdd1d7-4a74-460c-a32c-01c386a5c58a","year":2024},"citing_paper":{"arxiv_id":"2606.00869","last_updated":"2026-05-30T19:53:19Z","snapshot_observed_at":"2026-08-01T18:48:30.860504Z","submitted_at":"2026-05-30T19:53:19Z","title":"Enhancing LLM Metacognition via Cognitive Pairwise Training","version":1},"reference_index":85,"source":"pdf_text","source_observed_at":"2026-06-28T19:01:18.153145Z"},"links":{"cited_paper":"/paper/2406.20094","citing_paper":"/paper/2606.00869"},"observation_digest":"sha256:e2a9e9e2ad48ee3b6a5c6256bd6392fe45699ea1c0db6a90999ec67911c1481b","observation_id":"9bd659e4-fe25-4360-b389-5b52ea14fdef","resolution":{"observed_at":"2026-06-28T19:02:33.900775Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-07-13T23:49:58.855418+00:00","source":"crossref_status_cache"},{"observed_at":"2026-07-13T23:49:58.855418+00:00","source":"openalex_status_cache"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2406.20094","last_updated":"2025-05-08T00:24:02Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2024-06-28T17:59:01Z","title":"Scaling Synthetic Data Creation with 1,000,000,000 Personas","version":3},"cited_work":{"arxiv_id":"2406.20094","doi":"10.48550/arxiv.2406.20094","metadata_source":"pith","pith_arxiv_id":"2406.20094","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Scaling Synthetic Data Creation with 1,000,000,000 Personas","venue":"cs.CL","work_id":"4acdd1d7-4a74-460c-a32c-01c386a5c58a","year":2024},"citing_paper":{"arxiv_id":"2606.01552","last_updated":"2026-06-01T01:56:27Z","snapshot_observed_at":"2026-08-04T14:55:57.156200Z","submitted_at":"2026-06-01T01:56:27Z","title":"RoleCDE:Benchmarking and Mitigating Role-Alignment Trade-offs in Role-Playing Agents","version":1},"reference_index":24,"source":"arxiv_source","source_observed_at":"2026-06-28T15:00:57.851111Z"},"links":{"cited_paper":"/paper/2406.20094","citing_paper":"/paper/2606.01552"},"observation_digest":"sha256:da6349528e761cf3a0a42c1af4f2f20db4e994be69be60cf57aa9193e12eb74d","observation_id":"bce990d5-8bfc-4d51-9028-8017f6fec5e2","resolution":{"observed_at":"2026-07-01T22:46:20.041422Z","resolver_source":"local_arxiv","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-07-13T23:49:58.855418+00:00","source":"crossref_status_cache"},{"observed_at":"2026-07-13T23:49:58.855418+00:00","source":"openalex_status_cache"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2406.20094","last_updated":"2025-05-08T00:24:02Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2024-06-28T17:59:01Z","title":"Scaling Synthetic Data Creation with 1,000,000,000 Personas","version":3},"cited_work":{"arxiv_id":"2406.20094","doi":"10.48550/arxiv.2406.20094","metadata_source":"pith","pith_arxiv_id":"2406.20094","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Scaling Synthetic Data Creation with 1,000,000,000 Personas","venue":"cs.CL","work_id":"4acdd1d7-4a74-460c-a32c-01c386a5c58a","year":2024},"citing_paper":{"arxiv_id":"2606.01976","last_updated":"2026-06-01T09:37:51Z","snapshot_observed_at":"2026-07-31T12:24:36.790524Z","submitted_at":"2026-06-01T09:37:51Z","title":"AutoBG: A Board Game Design Assistant with Interactive Ideation, Iterative Rulebook Generation, and Individualized Feedback","version":1},"reference_index":4,"source":"pdf_text","source_observed_at":"2026-06-28T13:02:12.843333Z"},"links":{"cited_paper":"/paper/2406.20094","citing_paper":"/paper/2606.01976"},"observation_digest":"sha256:a1c25d844d08484834d39d8e89f06ef0c1e4eb6f46cd21c1a309724703f59eb6","observation_id":"d4e6e580-c810-4ece-8e4e-6abb5860b802","resolution":{"observed_at":"2026-07-02T00:56:25.298093Z","resolver_source":"local_arxiv","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-07-13T23:49:58.855418+00:00","source":"crossref_status_cache"},{"observed_at":"2026-07-13T23:49:58.855418+00:00","source":"openalex_status_cache"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2406.20094","last_updated":"2025-05-08T00:24:02Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2024-06-28T17:59:01Z","title":"Scaling Synthetic Data Creation with 1,000,000,000 Personas","version":3},"cited_work":{"arxiv_id":"2406.20094","doi":"10.48550/arxiv.2406.20094","metadata_source":"pith","pith_arxiv_id":"2406.20094","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Scaling Synthetic Data Creation with 1,000,000,000 Personas","venue":"cs.CL","work_id":"4acdd1d7-4a74-460c-a32c-01c386a5c58a","year":2024},"citing_paper":{"arxiv_id":"2606.02001","last_updated":"2026-06-01T09:57:52Z","snapshot_observed_at":"2026-07-06T23:42:30.555538Z","submitted_at":"2026-06-01T09:57:52Z","title":"Scaling Agentic Capabilities via Grounded Interaction Synthesis","version":1},"reference_index":5,"source":"pdf_text","source_observed_at":"2026-06-28T15:04:54.247779Z"},"links":{"cited_paper":"/paper/2406.20094","citing_paper":"/paper/2606.02001"},"observation_digest":"sha256:f08d1317bc5e655e59ccbc48276dc04f76ea3cb8763725ffd864a8408c549368","observation_id":"3021555f-dcda-4e6d-ae33-a9417450654d","resolution":{"observed_at":"2026-06-28T15:12:18.670221Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-07-13T23:49:58.855418+00:00","source":"crossref_status_cache"},{"observed_at":"2026-07-13T23:49:58.855418+00:00","source":"openalex_status_cache"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2406.20094","last_updated":"2025-05-08T00:24:02Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2024-06-28T17:59:01Z","title":"Scaling Synthetic Data Creation with 1,000,000,000 Personas","version":3},"cited_work":{"arxiv_id":"2406.20094","doi":"10.48550/arxiv.2406.20094","metadata_source":"pith","pith_arxiv_id":"2406.20094","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Scaling Synthetic Data Creation with 1,000,000,000 Personas","venue":"cs.CL","work_id":"4acdd1d7-4a74-460c-a32c-01c386a5c58a","year":2024},"citing_paper":{"arxiv_id":"2606.02300","last_updated":"2026-06-01T14:23:17Z","snapshot_observed_at":"2026-07-06T23:42:44.661608Z","submitted_at":"2026-06-01T14:23:17Z","title":"Beyond Isolated Behaviors: Hierarchical User Modeling for LLM Personalization","version":1},"reference_index":93,"source":"arxiv_source","source_observed_at":"2026-06-28T14:51:27.110839Z"},"links":{"cited_paper":"/paper/2406.20094","citing_paper":"/paper/2606.02300"},"observation_digest":"sha256:e4fc74f2f23f458adcbd8da2accfc3cf316bf07273ae9458fc7f497ede541509","observation_id":"ee1875f6-8f01-4832-9451-810c39caf7d9","resolution":{"observed_at":"2026-07-01T22:56:20.455973Z","resolver_source":"local_arxiv","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-07-13T23:49:58.855418+00:00","source":"crossref_status_cache"},{"observed_at":"2026-07-13T23:49:58.855418+00:00","source":"openalex_status_cache"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2406.20094","last_updated":"2025-05-08T00:24:02Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2024-06-28T17:59:01Z","title":"Scaling Synthetic Data Creation with 1,000,000,000 Personas","version":3},"cited_work":{"arxiv_id":"2406.20094","doi":"10.48550/arxiv.2406.20094","metadata_source":"pith","pith_arxiv_id":"2406.20094","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Scaling Synthetic Data Creation with 1,000,000,000 Personas","venue":"cs.CL","work_id":"4acdd1d7-4a74-460c-a32c-01c386a5c58a","year":2024},"citing_paper":{"arxiv_id":"2606.03635","last_updated":"2026-06-02T13:31:57Z","snapshot_observed_at":"2026-08-05T07:22:25.874494Z","submitted_at":"2026-06-02T13:31:57Z","title":"VidMsg: A Benchmark for Implicit Message Inference in Short Videos","version":1},"reference_index":9,"source":"pdf_text","source_observed_at":"2026-06-28T10:25:06.594946Z"},"links":{"cited_paper":"/paper/2406.20094","citing_paper":"/paper/2606.03635"},"observation_digest":"sha256:f0d3e2d0f62bf2f5598c388e0de1d10ec503b59ccb364af37429dd20a1f01b31","observation_id":"12d1daaa-bfde-4b68-bd61-24bcbc61eb15","resolution":{"observed_at":"2026-07-02T03:06:29.020642Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-07-13T23:49:58.855418+00:00","source":"crossref_status_cache"},{"observed_at":"2026-07-13T23:49:58.855418+00:00","source":"openalex_status_cache"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2406.20094","last_updated":"2025-05-08T00:24:02Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2024-06-28T17:59:01Z","title":"Scaling Synthetic Data Creation with 1,000,000,000 Personas","version":3},"cited_work":{"arxiv_id":"2406.20094","doi":"10.48550/arxiv.2406.20094","metadata_source":"pith","pith_arxiv_id":"2406.20094","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Scaling Synthetic Data Creation with 1,000,000,000 Personas","venue":"cs.CL","work_id":"4acdd1d7-4a74-460c-a32c-01c386a5c58a","year":2024},"citing_paper":{"arxiv_id":"2606.04155","last_updated":"2026-06-02T19:20:54Z","snapshot_observed_at":"2026-08-02T18:30:05.252168Z","submitted_at":"2026-06-02T19:20:54Z","title":"SocialCoach: Personalized Social Skill Learning with RL-based Agentic Tutoring and Practice","version":1},"reference_index":13,"source":"pdf_text","source_observed_at":"2026-06-28T08:04:17.224987Z"},"links":{"cited_paper":"/paper/2406.20094","citing_paper":"/paper/2606.04155"},"observation_digest":"sha256:478996cccb771087caf4217ca528bab37a8220be6a6cad26dd42d96f380ddc8c","observation_id":"ae4aba05-b0f3-4fed-83ae-a8220120f40b","resolution":{"observed_at":"2026-07-02T05:46:41.070794Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-07-13T23:49:58.855418+00:00","source":"crossref_status_cache"},{"observed_at":"2026-07-13T23:49:58.855418+00:00","source":"openalex_status_cache"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2406.20094","last_updated":"2025-05-08T00:24:02Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2024-06-28T17:59:01Z","title":"Scaling Synthetic Data Creation with 1,000,000,000 Personas","version":3},"cited_work":{"arxiv_id":"2406.20094","doi":"10.48550/arxiv.2406.20094","metadata_source":"pith","pith_arxiv_id":"2406.20094","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Scaling Synthetic Data Creation with 1,000,000,000 Personas","venue":"cs.CL","work_id":"4acdd1d7-4a74-460c-a32c-01c386a5c58a","year":2024},"citing_paper":{"arxiv_id":"2606.08841","last_updated":"2026-06-07T21:11:47Z","snapshot_observed_at":"2026-08-05T02:02:53.704575Z","submitted_at":"2026-06-07T21:11:47Z","title":"ZIPP:Zero-shot Image Personalization from Personas","version":1},"reference_index":25,"source":"arxiv_source","source_observed_at":"2026-06-27T18:11:22.987938Z"},"links":{"cited_paper":"/paper/2406.20094","citing_paper":"/paper/2606.08841"},"observation_digest":"sha256:469ae2f6d9cd9c6f9bc001905b8326e56556fa53fff9a7fc4df22f09b201bc88","observation_id":"b1bde8e8-ddf8-4017-ba44-365b0c1c0438","resolution":{"observed_at":"2026-07-02T23:27:27.965202Z","resolver_source":"local_arxiv","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-07-13T23:49:58.855418+00:00","source":"crossref_status_cache"},{"observed_at":"2026-07-13T23:49:58.855418+00:00","source":"openalex_status_cache"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2406.20094","last_updated":"2025-05-08T00:24:02Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2024-06-28T17:59:01Z","title":"Scaling Synthetic Data Creation with 1,000,000,000 Personas","version":3},"cited_work":{"arxiv_id":"2406.20094","doi":"10.48550/arxiv.2406.20094","metadata_source":"pith","pith_arxiv_id":"2406.20094","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Scaling Synthetic Data Creation with 1,000,000,000 Personas","venue":"cs.CL","work_id":"4acdd1d7-4a74-460c-a32c-01c386a5c58a","year":2024},"citing_paper":{"arxiv_id":"2606.09038","last_updated":"2026-06-08T05:10:05Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2026-06-08T05:10:05Z","title":"Personalization Meets Safety:Mechanisms,Risks,and Mitigations in Personalized LLMs","version":1},"reference_index":46,"source":"pdf_text","source_observed_at":"2026-06-27T16:49:14.243931Z"},"links":{"cited_paper":"/paper/2406.20094","citing_paper":"/paper/2606.09038"},"observation_digest":"sha256:4902fef22f5849e8b609741f1b5083ce6dbe0a8955b5ec316b0694a5b31201d4","observation_id":"5207e9cd-a3b4-4f71-90df-4174ca437ee5","resolution":{"observed_at":"2026-07-03T01:07:30.287835Z","resolver_source":"local_arxiv","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-07-13T23:49:58.855418+00:00","source":"crossref_status_cache"},{"observed_at":"2026-07-13T23:49:58.855418+00:00","source":"openalex_status_cache"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2406.20094","last_updated":"2025-05-08T00:24:02Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2024-06-28T17:59:01Z","title":"Scaling Synthetic Data Creation with 1,000,000,000 Personas","version":3},"cited_work":{"arxiv_id":"2406.20094","doi":"10.48550/arxiv.2406.20094","metadata_source":"pith","pith_arxiv_id":"2406.20094","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Scaling Synthetic Data Creation with 1,000,000,000 Personas","venue":"cs.CL","work_id":"4acdd1d7-4a74-460c-a32c-01c386a5c58a","year":2024},"citing_paper":{"arxiv_id":"2606.12433","last_updated":"2026-05-15T17:42:39Z","snapshot_observed_at":"2026-07-06T23:51:22.219590Z","submitted_at":"2026-05-15T17:42:39Z","title":"Marginal Alignment Does Not Guarantee Joint-Distribution Fidelity: An Official-Reference Audit of Nemotron-Personas-Korea with Cross-Locale Replication","version":1},"reference_index":50,"source":"arxiv_source","source_observed_at":"2026-06-30T19:04:36.443335Z"},"links":{"cited_paper":"/paper/2406.20094","citing_paper":"/paper/2606.12433"},"observation_digest":"sha256:6acdaf6e1abc825b61ae1a26de2041fdbba7dc4829a7815b64e74ed9676d7c59","observation_id":"062113dc-2e32-4d43-ae90-ee560cb0ee85","resolution":{"observed_at":"2026-06-30T19:05:00.229554Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-07-13T23:49:58.855418+00:00","source":"crossref_status_cache"},{"observed_at":"2026-07-13T23:49:58.855418+00:00","source":"openalex_status_cache"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2406.20094","last_updated":"2025-05-08T00:24:02Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2024-06-28T17:59:01Z","title":"Scaling Synthetic Data Creation with 1,000,000,000 Personas","version":3},"cited_work":{"arxiv_id":"2406.20094","doi":"10.48550/arxiv.2406.20094","metadata_source":"pith","pith_arxiv_id":"2406.20094","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Scaling Synthetic Data Creation with 1,000,000,000 Personas","venue":"cs.CL","work_id":"4acdd1d7-4a74-460c-a32c-01c386a5c58a","year":2024},"citing_paper":{"arxiv_id":"2606.12730","last_updated":"2026-06-10T22:28:53Z","snapshot_observed_at":"2026-07-06T23:51:35.904251Z","submitted_at":"2026-06-10T22:28:53Z","title":"Rethinking Psychometric Evaluation of LLMs: When and Why Self-Reports Predict Behavior","version":1},"reference_index":32,"source":"pdf_text","source_observed_at":"2026-06-27T09:36:36.821058Z"},"links":{"cited_paper":"/paper/2406.20094","citing_paper":"/paper/2606.12730"},"observation_digest":"sha256:4ecd8688edf5a7252f34b12a6f81d37fd1a7333f3878d52c701f2f36069eb72a","observation_id":"6dc52b19-58e3-48b0-a065-2335919f37a0","resolution":{"observed_at":"2026-07-03T11:18:03.739962Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-07-13T23:49:58.855418+00:00","source":"crossref_status_cache"},{"observed_at":"2026-07-13T23:49:58.855418+00:00","source":"openalex_status_cache"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2406.20094","last_updated":"2025-05-08T00:24:02Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2024-06-28T17:59:01Z","title":"Scaling Synthetic Data Creation with 1,000,000,000 Personas","version":3},"cited_work":{"arxiv_id":"2406.20094","doi":"10.48550/arxiv.2406.20094","metadata_source":"pith","pith_arxiv_id":"2406.20094","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Scaling Synthetic Data Creation with 1,000,000,000 Personas","venue":"cs.CL","work_id":"4acdd1d7-4a74-460c-a32c-01c386a5c58a","year":2024},"citing_paper":{"arxiv_id":"2606.13681","last_updated":"2026-06-17T04:10:52Z","snapshot_observed_at":"2026-07-06T23:52:23.981568Z","submitted_at":"2026-06-11T17:59:59Z","title":"EvoArena: Tracking Memory Evolution for Robust LLM Agents in Dynamic Environments","version":2},"reference_index":4,"source":"pdf_text","source_observed_at":"2026-06-27T06:28:04.239773Z"},"links":{"cited_paper":"/paper/2406.20094","citing_paper":"/paper/2606.13681"},"observation_digest":"sha256:bb146a73187ae5d02af2fe28498345fcbfd4621fc0725f078ce13cbfb5ab347d","observation_id":"9e3ac820-679b-4ef8-9a86-f691ef50b849","resolution":{"observed_at":"2026-07-03T15:28:34.406759Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-07-13T23:49:58.855418+00:00","source":"crossref_status_cache"},{"observed_at":"2026-07-13T23:49:58.855418+00:00","source":"openalex_status_cache"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2406.20094","last_updated":"2025-05-08T00:24:02Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2024-06-28T17:59:01Z","title":"Scaling Synthetic Data Creation with 1,000,000,000 Personas","version":3},"cited_work":{"arxiv_id":"2406.20094","doi":"10.48550/arxiv.2406.20094","metadata_source":"pith","pith_arxiv_id":"2406.20094","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Scaling Synthetic Data Creation with 1,000,000,000 Personas","venue":"cs.CL","work_id":"4acdd1d7-4a74-460c-a32c-01c386a5c58a","year":2024},"citing_paper":{"arxiv_id":"2606.18263","last_updated":"2026-05-12T10:33:48Z","snapshot_observed_at":"2026-08-02T05:46:15.140773Z","submitted_at":"2026-05-12T10:33:48Z","title":"How Well Do Large Language Models Capture Human Personality?","version":1},"reference_index":10,"source":"pdf_text","source_observed_at":"2026-06-30T22:20:10.160763Z"},"links":{"cited_paper":"/paper/2406.20094","citing_paper":"/paper/2606.18263"},"observation_digest":"sha256:d791913678ff52e8878e4db468336775c3dee9e36615b32b6e508e14019ceeb1","observation_id":"9bcd4766-166c-4054-a139-6fbb5f9fbcd1","resolution":{"observed_at":"2026-07-01T14:05:46.457757Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-07-13T23:49:58.855418+00:00","source":"crossref_status_cache"},{"observed_at":"2026-07-13T23:49:58.855418+00:00","source":"openalex_status_cache"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2406.20094","last_updated":"2025-05-08T00:24:02Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2024-06-28T17:59:01Z","title":"Scaling Synthetic Data Creation with 1,000,000,000 Personas","version":3},"cited_work":{"arxiv_id":"2406.20094","doi":"10.48550/arxiv.2406.20094","metadata_source":"pith","pith_arxiv_id":"2406.20094","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Scaling Synthetic Data Creation with 1,000,000,000 Personas","venue":"cs.CL","work_id":"4acdd1d7-4a74-460c-a32c-01c386a5c58a","year":2024},"citing_paper":{"arxiv_id":"2606.19640","last_updated":"2026-06-17T22:36:06Z","snapshot_observed_at":"2026-07-06T23:54:56.685241Z","submitted_at":"2026-06-17T22:36:06Z","title":"Creating Multilingual Mental Health Dialogue Datasets: Limits of Persona-Based Localization via Nationality and Language","version":1},"reference_index":32,"source":"arxiv_source","source_observed_at":"2026-06-26T20:27:41.561992Z"},"links":{"cited_paper":"/paper/2406.20094","citing_paper":"/paper/2606.19640"},"observation_digest":"sha256:835da064c0701abcb98ed624b24f5de12af2af67e167c65f7ef68698742af80f","observation_id":"ca285ca4-b445-4e44-8b0f-25d651bd244f","resolution":{"observed_at":"2026-06-26T20:29:57.481501Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-07-13T23:49:58.855418+00:00","source":"crossref_status_cache"},{"observed_at":"2026-07-13T23:49:58.855418+00:00","source":"openalex_status_cache"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2406.20094","last_updated":"2025-05-08T00:24:02Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2024-06-28T17:59:01Z","title":"Scaling Synthetic Data Creation with 1,000,000,000 Personas","version":3},"cited_work":{"arxiv_id":"2406.20094","doi":"10.48550/arxiv.2406.20094","metadata_source":"pith","pith_arxiv_id":"2406.20094","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Scaling Synthetic Data Creation with 1,000,000,000 Personas","venue":"cs.CL","work_id":"4acdd1d7-4a74-460c-a32c-01c386a5c58a","year":2024},"citing_paper":{"arxiv_id":"2606.20400","last_updated":"2026-06-18T15:53:22Z","snapshot_observed_at":"2026-08-03T10:59:22.252735Z","submitted_at":"2026-06-18T15:53:22Z","title":"The Significance of Style Diversity in Annotation-Free Synthetic Data Generation","version":1},"reference_index":1,"source":"pdf_text","source_observed_at":"2026-06-26T18:14:48.021786Z"},"links":{"cited_paper":"/paper/2406.20094","citing_paper":"/paper/2606.20400"},"observation_digest":"sha256:eabc2cd067f6cd38368a1e801763439590ba13f90192a3af67a7260efb1c0885","observation_id":"e7dcb06f-7848-486d-94e1-5aeae753cb1a","resolution":{"observed_at":"2026-07-04T03:19:30.637083Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-07-13T23:49:58.855418+00:00","source":"crossref_status_cache"},{"observed_at":"2026-07-13T23:49:58.855418+00:00","source":"openalex_status_cache"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2406.20094","last_updated":"2025-05-08T00:24:02Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2024-06-28T17:59:01Z","title":"Scaling Synthetic Data Creation with 1,000,000,000 Personas","version":3},"cited_work":{"arxiv_id":"2406.20094","doi":"10.48550/arxiv.2406.20094","metadata_source":"pith","pith_arxiv_id":"2406.20094","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Scaling Synthetic Data Creation with 1,000,000,000 Personas","venue":"cs.CL","work_id":"4acdd1d7-4a74-460c-a32c-01c386a5c58a","year":2024},"citing_paper":{"arxiv_id":"2606.29793","last_updated":"2026-06-30T06:50:50Z","snapshot_observed_at":"2026-07-07T00:03:41.292432Z","submitted_at":"2026-06-29T05:14:43Z","title":"Fund2Persona: A Framework for Building and Refining Financial Advisor Personas from Fund Disclosure Data","version":1},"reference_index":16,"source":"arxiv_source","source_observed_at":"2026-06-30T06:25:43.114521Z"},"links":{"cited_paper":"/paper/2406.20094","citing_paper":"/paper/2606.29793"},"observation_digest":"sha256:7efff249b0206bce198a7092952178c90764b91c2939300c93d217348b5185e3","observation_id":"e9b09111-a562-4ffa-a2b3-7a11caefe5e5","resolution":{"observed_at":"2026-06-30T06:34:18.541736Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-07-13T23:49:58.855418+00:00","source":"crossref_status_cache"},{"observed_at":"2026-07-13T23:49:58.855418+00:00","source":"openalex_status_cache"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2406.20094","last_updated":"2025-05-08T00:24:02Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2024-06-28T17:59:01Z","title":"Scaling Synthetic Data Creation with 1,000,000,000 Personas","version":3},"cited_work":{"arxiv_id":"2406.20094","doi":"10.48550/arxiv.2406.20094","metadata_source":"pith","pith_arxiv_id":"2406.20094","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Scaling Synthetic Data Creation with 1,000,000,000 Personas","venue":"cs.CL","work_id":"4acdd1d7-4a74-460c-a32c-01c386a5c58a","year":2024},"citing_paper":{"arxiv_id":"2606.29793","last_updated":"2026-06-30T06:50:50Z","snapshot_observed_at":"2026-07-07T00:03:41.292432Z","submitted_at":"2026-06-29T05:14:43Z","title":"Fund2Persona: A Framework for Building and Refining Financial Advisor Personas from Fund Disclosure Data","version":2},"reference_index":16,"source":"arxiv_source","source_observed_at":"2026-07-01T07:05:00.917347Z"},"links":{"cited_paper":"/paper/2406.20094","citing_paper":"/paper/2606.29793"},"observation_digest":"sha256:6114d89acb6bcd6efc79f3374b50525eb3e3c0677bc846e499ce9912c8901cf2","observation_id":"2d04e9c7-26e8-4a1f-b10b-7772bbc01239","resolution":{"observed_at":"2026-07-01T07:05:27.441665Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-07-13T23:49:58.855418+00:00","source":"crossref_status_cache"},{"observed_at":"2026-07-13T23:49:58.855418+00:00","source":"openalex_status_cache"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2406.20094","last_updated":"2025-05-08T00:24:02Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2024-06-28T17:59:01Z","title":"Scaling Synthetic Data Creation with 1,000,000,000 Personas","version":3},"cited_work":{"arxiv_id":"2406.20094","doi":"10.48550/arxiv.2406.20094","metadata_source":"pith","pith_arxiv_id":"2406.20094","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Scaling Synthetic Data Creation with 1,000,000,000 Personas","venue":"cs.CL","work_id":"4acdd1d7-4a74-460c-a32c-01c386a5c58a","year":2024},"citing_paper":{"arxiv_id":"2607.01727","last_updated":"2026-07-02T05:31:36Z","snapshot_observed_at":"2026-07-07T00:07:13.555571Z","submitted_at":"2026-07-02T05:31:36Z","title":"When Does Generating More Help? Disentangling Fixed-Source Synthesis from Source Expansion in Synthetic Data Scaling","version":1},"reference_index":3,"source":"arxiv_source","source_observed_at":"2026-07-03T15:20:51.474398Z"},"links":{"cited_paper":"/paper/2406.20094","citing_paper":"/paper/2607.01727"},"observation_digest":"sha256:07d89731841c428fe6ff21aded6542de89ff3c981770554d63c23e50169051bd","observation_id":"d2f2137d-ac20-48ea-8e71-b399d8840290","resolution":{"observed_at":"2026-07-03T15:28:33.815702Z","resolver_source":"local_arxiv","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-07-13T23:49:58.855418+00:00","source":"crossref_status_cache"},{"observed_at":"2026-07-13T23:49:58.855418+00:00","source":"openalex_status_cache"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2406.20094","last_updated":"2025-05-08T00:24:02Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2024-06-28T17:59:01Z","title":"Scaling Synthetic Data Creation with 1,000,000,000 Personas","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2406.20094","snapshot_observed_at":"2026-07-12T05:44:33.099337Z","title":"Scaling Synthetic Data Creation with 1,000,000,000 Personas , publisher =","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2607.02972","last_updated":"2026-07-03T05:28:43Z","snapshot_observed_at":"2026-08-04T09:23:39.890170Z","submitted_at":"2026-07-03T05:28:43Z","title":"A Scalable Approach to Evaluating Moral Sensitivity in LLMs","version":1},"reference_index":98,"source":"arxiv_source","source_observed_at":"2026-07-12T05:44:33.099337Z"},"links":{"cited_paper":"/paper/2406.20094","citing_paper":"/paper/2607.02972"},"observation_digest":"sha256:a8cefcd3fa513e5ac6834129251811233e24324392b5eaf9a6f70ec9527a134f","observation_id":"b3308365-28e4-4d31-8cde-7ce2408f14d0","resolution":{"observed_at":"2026-07-12T05:44:33.099337Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2406.20094","last_updated":"2025-05-08T00:24:02Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2024-06-28T17:59:01Z","title":"Scaling Synthetic Data Creation with 1,000,000,000 Personas","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2406.20094","snapshot_observed_at":"2026-07-11T16:02:00.920066Z","title":"arXiv preprint arXiv:2406.20094 , year=","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2607.04636","last_updated":"2026-07-06T03:39:08Z","snapshot_observed_at":"2026-07-11T16:01:58.025782Z","submitted_at":"2026-07-06T03:39:08Z","title":"Enhancing Large Multimodal Models in Key Information Extraction via Scene-Aware Document Synthesis","version":1},"reference_index":70,"source":"arxiv_source","source_observed_at":"2026-07-11T16:02:00.920066Z"},"links":{"cited_paper":"/paper/2406.20094","citing_paper":"/paper/2607.04636"},"observation_digest":"sha256:e7532886b886430a2b7d91a3c6f179158ec2c7e4d37dae6a047183a85e916103","observation_id":"2a438e36-0e83-4d7c-8070-751240fcdc1f","resolution":{"observed_at":"2026-07-11T16:02:00.920066Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2406.20094","last_updated":"2025-05-08T00:24:02Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2024-06-28T17:59:01Z","title":"Scaling Synthetic Data Creation with 1,000,000,000 Personas","version":3},"cited_work":{"arxiv_id":"2406.20094","doi":"10.48550/arxiv.2406.20094","metadata_source":"pith","pith_arxiv_id":"2406.20094","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Scaling Synthetic Data Creation with 1,000,000,000 Personas","venue":"cs.CL","work_id":"4acdd1d7-4a74-460c-a32c-01c386a5c58a","year":2024},"citing_paper":{"arxiv_id":"2607.05761","last_updated":"2026-07-07T02:38:41Z","snapshot_observed_at":"2026-07-11T02:22:41.502718Z","submitted_at":"2026-07-07T02:38:41Z","title":"Synthetic Consumer Insight Generation with Large Language Models","version":1},"reference_index":2,"source":"pdf_text","source_observed_at":"2026-07-11T02:22:45.862673Z"},"links":{"cited_paper":"/paper/2406.20094","citing_paper":"/paper/2607.05761"},"observation_digest":"sha256:23aa1a130e1427de5b4745ad2126539db472db9e18f5568b1e6fde815dd5c547","observation_id":"c2a36032-4b82-4c48-92a5-445426993370","resolution":{"observed_at":"2026-07-11T02:27:47.641300Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-07-13T23:49:58.855418+00:00","source":"crossref_status_cache"},{"observed_at":"2026-07-13T23:49:58.855418+00:00","source":"openalex_status_cache"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2406.20094","last_updated":"2025-05-08T00:24:02Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2024-06-28T17:59:01Z","title":"Scaling Synthetic Data Creation with 1,000,000,000 Personas","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2406.20094","snapshot_observed_at":"2026-07-14T11:49:31.595873Z","title":"arXiv preprint arXiv:2406.20094 , year=","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2607.10428","last_updated":"2026-07-24T14:32:58Z","snapshot_observed_at":"2026-08-02T07:20:28.030282Z","submitted_at":"2026-07-11T18:22:49Z","title":"Enjoy Your Talk: A Human-Centered Benchmark for Multi-Turn Dialogue with Decoupled User Simulation, Target Modeling, and Judging","version":1},"reference_index":38,"source":"arxiv_source","source_observed_at":"2026-07-14T11:49:31.595873Z"},"links":{"cited_paper":"/paper/2406.20094","citing_paper":"/paper/2607.10428"},"observation_digest":"sha256:294a3dee95290144051360800d5cdad4f3742f7a4fb51b831f963602dad3e65c","observation_id":"192fd6ce-ceb3-412b-99bd-c6547b65e256","resolution":{"observed_at":"2026-07-14T11:49:31.595873Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2406.20094","last_updated":"2025-05-08T00:24:02Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2024-06-28T17:59:01Z","title":"Scaling Synthetic Data Creation with 1,000,000,000 Personas","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2406.20094","snapshot_observed_at":"2026-08-02T07:20:35.191658Z","title":"arXiv preprint arXiv:2406.20094 , year=","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2607.10428","last_updated":"2026-07-24T14:32:58Z","snapshot_observed_at":"2026-08-02T07:20:28.030282Z","submitted_at":"2026-07-11T18:22:49Z","title":"Enjoy Your Talk: A Human-Centered Benchmark for Multi-Turn Dialogue with Decoupled User Simulation, Target Modeling, and Judging","version":2},"reference_index":38,"source":"arxiv_source","source_observed_at":"2026-08-02T07:20:35.191658Z"},"links":{"cited_paper":"/paper/2406.20094","citing_paper":"/paper/2607.10428"},"observation_digest":"sha256:40ce184b2580a2c655bb78b103fedd9a982c8a0609928d97ce3bb5978107e7fe","observation_id":"83084c80-54b0-4d26-897d-9c05356fcc34","resolution":{"observed_at":"2026-08-02T07:20:35.191658Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2406.20094","last_updated":"2025-05-08T00:24:02Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2024-06-28T17:59:01Z","title":"Scaling Synthetic Data Creation with 1,000,000,000 Personas","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2406.20094","snapshot_observed_at":"2026-08-01T23:43:18.659490Z","title":"Scaling synthetic data creation with 1,000,000,000 personas, 2025","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2607.15277","last_updated":"2026-07-16T17:59:31Z","snapshot_observed_at":"2026-08-03T17:08:01.241612Z","submitted_at":"2026-07-16T17:59:31Z","title":"Partition, Prompt, Aggregate: Statistical Self-Consistency in Language Models","version":1},"reference_index":133,"source":"arxiv_source","source_observed_at":"2026-08-01T23:43:18.659490Z"},"links":{"cited_paper":"/paper/2406.20094","citing_paper":"/paper/2607.15277"},"observation_digest":"sha256:9ec3c11a5bb495d20d0463bb1506854f083703aa84f0f687f751e33c3431f7f6","observation_id":"81029224-b9b0-4199-a72a-13566746d629","resolution":{"observed_at":"2026-08-01T23:43:18.659490Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2406.20094","last_updated":"2025-05-08T00:24:02Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2024-06-28T17:59:01Z","title":"Scaling Synthetic Data Creation with 1,000,000,000 Personas","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2406.20094","snapshot_observed_at":"2026-08-02T14:50:11.096699Z","title":"Scaling synthetic data creation with 1,000,000,000 personas.arXiv preprint arXiv:2406.20094, 2024","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2607.16204","last_updated":"2026-05-07T00:40:32Z","snapshot_observed_at":"2026-08-02T14:50:06.224987Z","submitted_at":"2026-05-07T00:40:32Z","title":"Masked Diffusion Language Models are Strong and Steerable Text-Based World Models for Agentic RL","version":1},"reference_index":33,"source":"pdf_text","source_observed_at":"2026-08-02T14:50:11.096699Z"},"links":{"cited_paper":"/paper/2406.20094","citing_paper":"/paper/2607.16204"},"observation_digest":"sha256:dc86f07c9f9d9eb298eb4f5a29a3a56afb127d09a26e8a9c495f76fe816b5b94","observation_id":"12537a41-d23c-4f5a-ba9a-b0ba4af3a73c","resolution":{"observed_at":"2026-08-02T14:50:11.096699Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2406.20094","last_updated":"2025-05-08T00:24:02Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2024-06-28T17:59:01Z","title":"Scaling Synthetic Data Creation with 1,000,000,000 Personas","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2406.20094","snapshot_observed_at":"2026-08-02T14:33:32.207950Z","title":"Scaling synthetic data creation with 1,000,000,000 personas.arXiv preprint arXiv:2406.20094, 2024","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2607.20429","last_updated":"2026-05-10T00:52:25Z","snapshot_observed_at":"2026-08-02T17:38:53.828313Z","submitted_at":"2026-05-10T00:52:25Z","title":"More Is Not More: What Matters for Diversity in LLM Opinions?","version":1},"reference_index":18,"source":"pdf_text","source_observed_at":"2026-08-02T14:33:32.207950Z"},"links":{"cited_paper":"/paper/2406.20094","citing_paper":"/paper/2607.20429"},"observation_digest":"sha256:2d7cb54414b6a2980bf821c5dfed4151c7bf7e19cc213cf7c1136faeb6c19437","observation_id":"eb9838b7-5ef4-4a22-b790-961661f07fec","resolution":{"observed_at":"2026-08-02T14:33:32.207950Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2406.20094","last_updated":"2025-05-08T00:24:02Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2024-06-28T17:59:01Z","title":"Scaling Synthetic Data Creation with 1,000,000,000 Personas","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2406.20094","snapshot_observed_at":"2026-08-02T12:43:06.781115Z","title":"InThe Thirty-ninth Annual Conference on Neural Informa- tion Processing Systems Datasets and Benchmarks Track","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2607.20482","last_updated":"2026-08-04T07:30:15Z","snapshot_observed_at":"2026-08-05T07:47:49.966048Z","submitted_at":"2026-05-30T13:27:55Z","title":"PersonaTrail: Benchmarking Personalized Web Agents through Browsing Trails","version":1},"reference_index":2025,"source":"pdf_text","source_observed_at":"2026-08-02T12:43:06.781115Z"},"links":{"cited_paper":"/paper/2406.20094","citing_paper":"/paper/2607.20482"},"observation_digest":"sha256:19a3a6c8d9cfc74ea7e9d2648aeaa3c3001083733c40493747789b3bb059fff7","observation_id":"c999e900-e83e-4823-9e9e-16598178dec7","resolution":{"observed_at":"2026-08-02T12:43:06.781115Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2406.20094","last_updated":"2025-05-08T00:24:02Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2024-06-28T17:59:01Z","title":"Scaling Synthetic Data Creation with 1,000,000,000 Personas","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2406.20094","snapshot_observed_at":"2026-08-02T08:30:03.086317Z","title":"Scaling synthetic data creation with 1,000,000,000 personas.https://arxiv.org/abs/2406.20094, 2025","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2607.22671","last_updated":"2026-07-06T17:26:46Z","snapshot_observed_at":"2026-08-04T23:59:24.522067Z","submitted_at":"2026-07-06T17:26:46Z","title":"AIR-BENCH Live: An Evolving Safety Benchmark for Foundation Models","version":1},"reference_index":12,"source":"pdf_text","source_observed_at":"2026-08-02T08:30:03.086317Z"},"links":{"cited_paper":"/paper/2406.20094","citing_paper":"/paper/2607.22671"},"observation_digest":"sha256:53b319c41172b6f775513d4c6789fb100a1b7b9f8e83ef01e58363528428bcc2","observation_id":"e6f1fdf6-ac2d-4467-938a-a90226758841","resolution":{"observed_at":"2026-08-02T08:30:03.086317Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2406.20094","last_updated":"2025-05-08T00:24:02Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2024-06-28T17:59:01Z","title":"Scaling Synthetic Data Creation with 1,000,000,000 Personas","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2406.20094","snapshot_observed_at":"2026-08-01T08:29:25.281882Z","title":"arXiv preprint arXiv:2406.20094 , year=","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2607.27379","last_updated":"2026-07-29T18:37:14Z","snapshot_observed_at":"2026-08-03T17:05:06.191495Z","submitted_at":"2026-07-29T18:37:14Z","title":"HSS-Synth: Humanities and Social Sciences Data Synthesis for LLMs","version":1},"reference_index":34,"source":"arxiv_source","source_observed_at":"2026-08-01T08:29:25.281882Z"},"links":{"cited_paper":"/paper/2406.20094","citing_paper":"/paper/2607.27379"},"observation_digest":"sha256:891edf19f5cc8862c0a475eecb69bd5e3ae46304179ade10d445880203aca77c","observation_id":"64541786-994b-4dd2-9e61-723b495f2370","resolution":{"observed_at":"2026-08-01T08:29:25.281882Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2406.20094","last_updated":"2025-05-08T00:24:02Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2024-06-28T17:59:01Z","title":"Scaling Synthetic Data Creation with 1,000,000,000 Personas","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2406.20094","snapshot_observed_at":"2026-08-03T00:55:41.850618Z","title":"arXiv preprint arXiv:2406.20094 , year=","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2607.28636","last_updated":"2026-05-19T13:56:13Z","snapshot_observed_at":"2026-08-05T07:12:04.791657Z","submitted_at":"2026-05-19T13:56:13Z","title":"Chain-of-Models: Cross-Model Auditing for Bias-Robust LLM Judges","version":1},"reference_index":282,"source":"arxiv_source","source_observed_at":"2026-08-03T00:55:41.850618Z"},"links":{"cited_paper":"/paper/2406.20094","citing_paper":"/paper/2607.28636"},"observation_digest":"sha256:414efaa525e4d0ec205c60d0cf328068de6f572eb78ba89466e852ff93d5759a","observation_id":"8919fb3b-845c-4e3a-9c9e-a991121bcb0c","resolution":{"observed_at":"2026-08-03T00:55:41.850618Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2406.20094","last_updated":"2025-05-08T00:24:02Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2024-06-28T17:59:01Z","title":"Scaling Synthetic Data Creation with 1,000,000,000 Personas","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2406.20094","snapshot_observed_at":"2026-08-03T00:22:46.381946Z","title":"Scaling synthetic data creation with 1,000,000,000 personas.arXiv preprint arXiv:2406.20094,","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2607.28818","last_updated":"2026-07-30T20:18:39Z","snapshot_observed_at":"2026-08-05T07:13:06.086069Z","submitted_at":"2026-07-30T20:18:39Z","title":"Best Friends, Not Forever: Evaluating Long-Horizon Persona Collapse and Behavioral Drift in AI Companions","version":1},"reference_index":1999,"source":"pdf_text","source_observed_at":"2026-08-03T00:22:46.381946Z"},"links":{"cited_paper":"/paper/2406.20094","citing_paper":"/paper/2607.28818"},"observation_digest":"sha256:d66a696a1ad8761b2708086e29422325dd4244abaad9ba89520a0f48eaa5760a","observation_id":"c2a468a7-ac07-476b-b897-1796092ef27a","resolution":{"observed_at":"2026-08-03T00:22:46.381946Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2406.20094","last_updated":"2025-05-08T00:24:02Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2024-06-28T17:59:01Z","title":"Scaling Synthetic Data Creation with 1,000,000,000 Personas","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2406.20094","snapshot_observed_at":"2026-08-04T02:49:04.839414Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2608.00007","last_updated":"2026-06-10T10:25:25Z","snapshot_observed_at":"2026-08-05T07:22:40.250584Z","submitted_at":"2026-06-10T10:25:25Z","title":"MemoryForge: Synthesize Lifelong Memory for Human-Like LLM Agents","version":1},"reference_index":10,"source":"arxiv_source","source_observed_at":"2026-08-04T02:49:04.839414Z"},"links":{"cited_paper":"/paper/2406.20094","citing_paper":"/paper/2608.00007"},"observation_digest":"sha256:826d2874f2159bef14f0193fd4500182ac55e8975c129cea05606a6ca6c0c21a","observation_id":"493e73ab-5ded-4324-9b6e-dcc707e9239e","resolution":{"observed_at":"2026-08-04T02:49:04.839414Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2406.20094","last_updated":"2025-05-08T00:24:02Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2024-06-28T17:59:01Z","title":"Scaling Synthetic Data Creation with 1,000,000,000 Personas","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2406.20094","snapshot_observed_at":"2026-08-04T06:17:17.123404Z","title":"Tao Ge, Xin Chan, Xiaoyang Wang, Dian Yu, Haitao Mi, and Dong Yu","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2608.02491","last_updated":"2026-08-03T16:55:50Z","snapshot_observed_at":"2026-08-05T07:43:56.869459Z","submitted_at":"2026-08-03T16:55:50Z","title":"Long-term Measurements: Towards a Longitudinal Understanding of Human-AI Interactions","version":1},"reference_index":2024,"source":"pdf_text","source_observed_at":"2026-08-04T06:17:17.123404Z"},"links":{"cited_paper":"/paper/2406.20094","citing_paper":"/paper/2608.02491"},"observation_digest":"sha256:07aa9969ed8f03c27665d6a3a4345fd79fb346c6e5331f7eb1bbae4ffbc04542","observation_id":"3f90f562-4730-4199-974f-fe5e34d900dd","resolution":{"observed_at":"2026-08-04T06:17:17.123404Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}}],"links":{"evidence":"/evidence","html":"/paper/2406.20094/citation-record","integrity":"/paper/2406.20094/integrity","json":"/paper/2406.20094/citation-record.json","paper":"/paper/2406.20094"},"outbound":[{"citation":{"cited_paper":{"arxiv_id":"2404.14219","last_updated":"2024-08-30T21:17:17Z","snapshot_observed_at":"2026-07-06T18:03:47.096406Z","submitted_at":"2024-04-22T14:32:33Z","title":"Phi-3 Technical Report: A Highly Capable Language Model Locally on Your Phone","version":4},"cited_work":{"arxiv_id":"2404.14219","doi":"10.48550/arxiv.2404.14219","metadata_source":"pith","pith_arxiv_id":"2404.14219","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Phi-3 Technical Report: A Highly Capable Language Model Locally on Your Phone","venue":"cs.CL","work_id":"feef9556-a016-493c-abd2-0c97a23a7ebf","year":2024},"citing_paper":{"arxiv_id":"2406.20094","last_updated":"2025-05-08T00:24:02Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2024-06-28T17:59:01Z","title":"Scaling Synthetic Data Creation with 1,000,000,000 Personas","version":3},"reference_index":1,"source":"pdf_text","source_observed_at":"2026-05-16T00:03:55.599967Z"},"links":{"cited_paper":"/paper/2404.14219","citing_paper":"/paper/2406.20094"},"observation_digest":"sha256:521cc7f0c6530ba7aa0bb8775da0aa2ed8207d5347aa0240ced62f2180e7360e","observation_id":"48bf3066-0a57-44fe-a7b5-133be8d4861d","resolution":{"observed_at":"2026-05-16T00:03:55.666645Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2303.08774","last_updated":"2024-03-04T06:01:33Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2023-03-15T17:15:04Z","title":"GPT-4 Technical Report","version":6},"cited_work":{"arxiv_id":"2303.08774","doi":"10.1002/tea.20265","metadata_source":"pith","pith_arxiv_id":"2303.08774","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"GPT-4 Technical Report","venue":"cs.CL","work_id":"b928e041-6991-4c08-8c81-0359e4097c7b","year":2023},"citing_paper":{"arxiv_id":"2406.20094","last_updated":"2025-05-08T00:24:02Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2024-06-28T17:59:01Z","title":"Scaling Synthetic Data Creation with 1,000,000,000 Personas","version":3},"reference_index":2,"source":"pdf_text","source_observed_at":"2026-05-16T00:03:55.599967Z"},"links":{"cited_paper":"/paper/2303.08774","citing_paper":"/paper/2406.20094"},"observation_digest":"sha256:46203b425c0a4557dfddd2af00ebcfe8a2625303df5cfffa82a6007596d2ef7c","observation_id":"d146f7e4-e526-4a3d-83ae-ed9e88e16965","resolution":{"observed_at":"2026-05-16T00:03:55.633202Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2403.18058","last_updated":"2024-11-02T11:08:49Z","snapshot_observed_at":"2026-07-06T17:51:36.160539Z","submitted_at":"2024-03-26T19:24:18Z","title":"COIG-CQIA: Quality is All You Need for Chinese Instruction Fine-tuning","version":2},"cited_work":{"arxiv_id":"2403.18058","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2403.18058","snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Coig-cqia: Quality is all you need for chinese instruction fine-tuning","venue":null,"work_id":"4d58f189-c89b-4190-b567-6fc2fba5e0b7","year":2024},"citing_paper":{"arxiv_id":"2406.20094","last_updated":"2025-05-08T00:24:02Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2024-06-28T17:59:01Z","title":"Scaling Synthetic Data Creation with 1,000,000,000 Personas","version":3},"reference_index":3,"source":"pdf_text","source_observed_at":"2026-05-16T00:03:55.599967Z"},"links":{"cited_paper":"/paper/2403.18058","citing_paper":"/paper/2406.20094"},"observation_digest":"sha256:2ed11ebcf7e73929a7d570e24a5de369f34fa3632c83d9cf5cbefab509231050","observation_id":"1a5d2d6b-9641-4569-be44-208fb3d3ed81","resolution":{"observed_at":"2026-05-16T00:03:55.640199Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2401.02524","last_updated":"2024-02-01T22:06:51Z","snapshot_observed_at":"2026-07-06T17:11:46.822922Z","submitted_at":"2024-01-04T20:23:51Z","title":"Comprehensive Exploration of Synthetic Data Generation: A Survey","version":2},"cited_work":{"arxiv_id":"2401.02524","doi":"10.48550/arxiv.2401.02524","metadata_source":"arxiv_reference","pith_arxiv_id":"2401.02524","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"arXiv preprint arXiv:2401.02524 (2024)","venue":"arXiv (Cornell University)","work_id":"7c3b201d-cc93-4fcb-bcb6-2583f29edbbf","year":2024},"citing_paper":{"arxiv_id":"2406.20094","last_updated":"2025-05-08T00:24:02Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2024-06-28T17:59:01Z","title":"Scaling Synthetic Data Creation with 1,000,000,000 Personas","version":3},"reference_index":4,"source":"pdf_text","source_observed_at":"2026-05-16T00:03:55.599967Z"},"links":{"cited_paper":"/paper/2401.02524","citing_paper":"/paper/2406.20094"},"observation_digest":"sha256:1d3ebbd35aefe59aa58bcb734f0bc069817f47de9850fa08d5fede61d87e5ab7","observation_id":"cb638a36-de58-4443-b8ae-e9f5f07fd327","resolution":{"observed_at":"2026-05-16T00:03:55.647045Z","resolver_source":"arxiv_id","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2401.02954","last_updated":"2024-01-05T18:59:13Z","snapshot_observed_at":"2026-08-02T13:11:16.882565Z","submitted_at":"2024-01-05T18:59:13Z","title":"DeepSeek LLM: Scaling Open-Source Language Models with Longtermism","version":1},"cited_work":{"arxiv_id":"2401.02954","doi":"10.48550/arxiv.2401.02954","metadata_source":"pith","pith_arxiv_id":"2401.02954","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"DeepSeek LLM: Scaling Open-Source Language Models with Longtermism","venue":"cs.CL","work_id":"01b10587-025b-499d-8ba3-7a538d24c2d6","year":2024},"citing_paper":{"arxiv_id":"2406.20094","last_updated":"2025-05-08T00:24:02Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2024-06-28T17:59:01Z","title":"Scaling Synthetic Data Creation with 1,000,000,000 Personas","version":3},"reference_index":5,"source":"pdf_text","source_observed_at":"2026-05-16T00:03:55.599967Z"},"links":{"cited_paper":"/paper/2401.02954","citing_paper":"/paper/2406.20094"},"observation_digest":"sha256:94da4f09c0c98d6bc2a8334c3134479bf2fac1094e2fd97eecf084184e9c7946","observation_id":"da49297d-d1d0-4f6b-be57-6ce305d822b7","resolution":{"observed_at":"2026-05-16T00:03:55.652918Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"On the resemblance and containment of documents","venue":null,"work_id":"633cbdd4-286a-4e14-bdad-921d7e65ecee","year":1997},"citing_paper":{"arxiv_id":"2406.20094","last_updated":"2025-05-08T00:24:02Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2024-06-28T17:59:01Z","title":"Scaling Synthetic Data Creation with 1,000,000,000 Personas","version":3},"reference_index":6,"source":"pdf_text","source_observed_at":"2026-05-16T00:03:55.599967Z"},"links":{"citing_paper":"/paper/2406.20094"},"observation_digest":"sha256:cc9501e05105137d3170032b728eeb53a509dd32640c68cacad2c33262e07206","observation_id":"33948438-b57a-4ef6-9ce0-80ed70f2d0b4","resolution":{"observed_at":"2026-05-16T00:03:55.804178Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2305.17126","last_updated":"2024-03-11T01:15:09Z","snapshot_observed_at":"2026-08-05T01:49:57.023215Z","submitted_at":"2023-05-26T17:50:11Z","title":"Large Language Models as Tool Makers","version":2},"cited_work":{"arxiv_id":"2305.17126","doi":"10.48550/arxiv.2305.17126","metadata_source":"arxiv_reference","pith_arxiv_id":"2305.17126","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Large language models as tool makers","venue":"arXiv (Cornell University)","work_id":"1447b78e-0a79-4af6-8cd4-93220e680d2b","year":2023},"citing_paper":{"arxiv_id":"2406.20094","last_updated":"2025-05-08T00:24:02Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2024-06-28T17:59:01Z","title":"Scaling Synthetic Data Creation with 1,000,000,000 Personas","version":3},"reference_index":7,"source":"pdf_text","source_observed_at":"2026-05-16T00:03:55.599967Z"},"links":{"cited_paper":"/paper/2305.17126","citing_paper":"/paper/2406.20094"},"observation_digest":"sha256:ac4d0c56069dae582354798f2fff9ff4f34b95700a832f844d4937f9c2ed52b5","observation_id":"2762c771-0c5b-4b9d-a5ec-815ba18a20d1","resolution":{"observed_at":"2026-05-16T00:03:55.674398Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2304.04736","last_updated":"2023-10-02T20:56:35Z","snapshot_observed_at":"2026-07-06T15:14:06.328855Z","submitted_at":"2023-04-10T17:47:39Z","title":"On the Possibilities of AI-Generated Text Detection","version":3},"cited_work":{"arxiv_id":"2304.04736","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2304.04736","snapshot_observed_at":"2026-07-03T09:17:48.211003Z","title":"org/CorpusID:261660497","venue":null,"work_id":"ecd9bcaf-8465-4a71-8229-27cc0d08920d","year":2023},"citing_paper":{"arxiv_id":"2406.20094","last_updated":"2025-05-08T00:24:02Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2024-06-28T17:59:01Z","title":"Scaling Synthetic Data Creation with 1,000,000,000 Personas","version":3},"reference_index":8,"source":"pdf_text","source_observed_at":"2026-05-16T00:03:55.599967Z"},"links":{"cited_paper":"/paper/2304.04736","citing_paper":"/paper/2406.20094"},"observation_digest":"sha256:09881aa2ec21cc979b2b87290294caca0f98f5cc82f34fc9aaf81df017943c65","observation_id":"52d3d290-0195-4c01-87aa-2614dc28fd57","resolution":{"observed_at":"2026-05-16T00:03:55.681990Z","resolver_source":"arxiv_id","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2211.12588","last_updated":"2023-10-23T01:27:38Z","snapshot_observed_at":"2026-08-02T13:06:11.850456Z","submitted_at":"2022-11-22T21:06:00Z","title":"Program of Thoughts Prompting: Disentangling Computation from Reasoning for Numerical Reasoning Tasks","version":4},"cited_work":{"arxiv_id":"2211.12588","doi":"10.48550/arxiv.2211.12588","metadata_source":"pith","pith_arxiv_id":"2211.12588","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Program of Thoughts Prompting: Disentangling Computation from Reasoning for Numerical Reasoning Tasks","venue":"cs.CL","work_id":"618aa44c-a6c6-425c-abce-8aa8aa842921","year":2022},"citing_paper":{"arxiv_id":"2406.20094","last_updated":"2025-05-08T00:24:02Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2024-06-28T17:59:01Z","title":"Scaling Synthetic Data Creation with 1,000,000,000 Personas","version":3},"reference_index":9,"source":"pdf_text","source_observed_at":"2026-05-16T00:03:55.599967Z"},"links":{"cited_paper":"/paper/2211.12588","citing_paper":"/paper/2406.20094"},"observation_digest":"sha256:5f941b850613685559384835dc35642bdcf5f4959f0280d1c8e2bde236a59fe9","observation_id":"30613c27-989e-4aaa-a11c-3d0d8a7a9d53","resolution":{"observed_at":"2026-05-16T00:03:55.688697Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2309.10668","last_updated":"2024-03-18T23:15:47Z","snapshot_observed_at":"2026-07-06T16:20:36.866744Z","submitted_at":"2023-09-19T14:50:38Z","title":"Language Modeling Is Compression","version":2},"cited_work":{"arxiv_id":"2309.10668","doi":"10.48550/arxiv.2309.10668","metadata_source":"pith","pith_arxiv_id":"2309.10668","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Language Modeling Is Compression","venue":"cs.LG","work_id":"e9f96f8e-ac74-44e1-bfd5-ac6ff083003f","year":2023},"citing_paper":{"arxiv_id":"2406.20094","last_updated":"2025-05-08T00:24:02Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2024-06-28T17:59:01Z","title":"Scaling Synthetic Data Creation with 1,000,000,000 Personas","version":3},"reference_index":10,"source":"pdf_text","source_observed_at":"2026-05-16T00:03:55.599967Z"},"links":{"cited_paper":"/paper/2309.10668","citing_paper":"/paper/2406.20094"},"observation_digest":"sha256:6afd02c223f59b4ae1d1f8eea77fb6930a043fbdf69b14a4f6d38b7c0a421f32","observation_id":"017bb499-cdaf-4d86-a560-c9627e89d884","resolution":{"observed_at":"2026-05-17T22:36:13.497932Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2402.07043","last_updated":"2024-05-31T12:27:52Z","snapshot_observed_at":"2026-07-06T17:28:27.220300Z","submitted_at":"2024-02-10T21:06:34Z","title":"A Tale of Tails: Model Collapse as a Change of Scaling Laws","version":2},"cited_work":{"arxiv_id":"2402.07043","doi":"10.48550/arxiv.2402.07043","metadata_source":"arxiv_reference","pith_arxiv_id":"2402.07043","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"A tale of tails: Model collapse as a change of scaling laws","venue":"arXiv (Cornell University)","work_id":"45b166f4-4936-4df8-b88b-16e32d9220ce","year":2024},"citing_paper":{"arxiv_id":"2406.20094","last_updated":"2025-05-08T00:24:02Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2024-06-28T17:59:01Z","title":"Scaling Synthetic Data Creation with 1,000,000,000 Personas","version":3},"reference_index":11,"source":"pdf_text","source_observed_at":"2026-05-16T00:03:55.599967Z"},"links":{"cited_paper":"/paper/2402.07043","citing_paper":"/paper/2406.20094"},"observation_digest":"sha256:a96fc419ea2aab89e0f2152cc11fc2888ce22b0e36d4d0b17b1e5301bbc1c2a1","observation_id":"f9bc22b5-804e-46ec-b7b2-c4e3fa27385e","resolution":{"observed_at":"2026-05-16T00:03:55.702551Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2305.19165","last_updated":"2023-05-30T16:09:19Z","snapshot_observed_at":"2026-07-06T15:35:30.017377Z","submitted_at":"2023-05-30T16:09:19Z","title":"Strategic Reasoning with Language Models","version":1},"cited_work":{"arxiv_id":"2305.19165","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2305.19165","snapshot_observed_at":"2026-07-04T00:39:17.311582Z","title":"Gandhi, D","venue":null,"work_id":"8d0b2824-8733-4c36-9a45-dab1f9a28e4f","year":2023},"citing_paper":{"arxiv_id":"2406.20094","last_updated":"2025-05-08T00:24:02Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2024-06-28T17:59:01Z","title":"Scaling Synthetic Data Creation with 1,000,000,000 Personas","version":3},"reference_index":12,"source":"pdf_text","source_observed_at":"2026-05-16T00:03:55.599967Z"},"links":{"cited_paper":"/paper/2305.19165","citing_paper":"/paper/2406.20094"},"observation_digest":"sha256:a0a987de3a36d7a8d7c72443502811b55ad0d02def7937c00d0c2aeea7b0d395","observation_id":"1a3ddce9-52b0-45d3-bc7c-e7becced4370","resolution":{"observed_at":"2026-05-16T00:03:55.708814Z","resolver_source":"arxiv_id","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2103.03874","last_updated":"2021-11-08T21:30:18Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2021-03-05T18:59:39Z","title":"Measuring Mathematical Problem Solving With the MATH Dataset","version":2},"cited_work":{"arxiv_id":"2103.03874","doi":"10.48550/arxiv.2103.03874","metadata_source":"pith","pith_arxiv_id":"2103.03874","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Measuring Mathematical Problem Solving With the MATH Dataset","venue":"cs.LG","work_id":"50652ac6-fb7c-4675-a2c2-159c241feb17","year":2021},"citing_paper":{"arxiv_id":"2406.20094","last_updated":"2025-05-08T00:24:02Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2024-06-28T17:59:01Z","title":"Scaling Synthetic Data Creation with 1,000,000,000 Personas","version":3},"reference_index":13,"source":"pdf_text","source_observed_at":"2026-05-16T00:03:55.599967Z"},"links":{"cited_paper":"/paper/2103.03874","citing_paper":"/paper/2406.20094"},"observation_digest":"sha256:f4110ca22cb9050aee18976853826fe5f429334f0cc16c39d3ce719eaf973514","observation_id":"12cfabf5-9a20-49c9-bfde-f398fe9c1b63","resolution":{"observed_at":"2026-05-16T00:03:55.715644Z","resolver_source":"local_arxiv","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-07-14T18:20:22.649941+00:00","source":"crossref_status_cache"},{"observed_at":"2026-07-14T18:20:22.649941+00:00","source":"openalex_status_cache"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2403.02333","last_updated":"2024-05-08T01:48:46Z","snapshot_observed_at":"2026-07-06T17:39:24.823145Z","submitted_at":"2024-03-04T18:58:30Z","title":"Key-Point-Driven Data Synthesis with its Enhancement on Mathematical Reasoning","version":3},"cited_work":{"arxiv_id":"2403.02333","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2403.02333","snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"arXiv preprint arXiv:2403.02333 , year=","venue":null,"work_id":"52d56266-7e01-478c-8b8f-1a30b87b2094","year":2024},"citing_paper":{"arxiv_id":"2406.20094","last_updated":"2025-05-08T00:24:02Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2024-06-28T17:59:01Z","title":"Scaling Synthetic Data Creation with 1,000,000,000 Personas","version":3},"reference_index":14,"source":"pdf_text","source_observed_at":"2026-05-16T00:03:55.599967Z"},"links":{"cited_paper":"/paper/2403.02333","citing_paper":"/paper/2406.20094"},"observation_digest":"sha256:e68b0755206508301527bdc01b5d29aafbe652e2c8929c88ac18cef073f41c14","observation_id":"5fabff6e-65fe-42c3-95a3-b3e367d387b0","resolution":{"observed_at":"2026-05-16T00:03:55.722916Z","resolver_source":"arxiv_id","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2312.10007","last_updated":"2023-12-15T18:23:50Z","snapshot_observed_at":"2026-08-03T18:40:17.286021Z","submitted_at":"2023-12-15T18:23:50Z","title":"Faithful Persona-based Conversational Dataset Generation with Large Language Models","version":1},"cited_work":{"arxiv_id":"2312.10007","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2312.10007","snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Faithful persona-based conversational dataset generation with large language models","venue":null,"work_id":"38ab3f5d-70ec-44c1-a4f1-a1a8d1690517","year":2023},"citing_paper":{"arxiv_id":"2406.20094","last_updated":"2025-05-08T00:24:02Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2024-06-28T17:59:01Z","title":"Scaling Synthetic Data Creation with 1,000,000,000 Personas","version":3},"reference_index":15,"source":"pdf_text","source_observed_at":"2026-05-16T00:03:55.599967Z"},"links":{"cited_paper":"/paper/2312.10007","citing_paper":"/paper/2406.20094"},"observation_digest":"sha256:12ec9b8a91305af7422b68fa6ace7837d12248adc12c071ecf1f6ae2b4df4905","observation_id":"8f7a7dda-a131-49c6-9428-66480bb2184d","resolution":{"observed_at":"2026-05-16T00:03:55.731116Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2001.08361","last_updated":"2020-01-23T03:59:20Z","snapshot_observed_at":"2026-07-06T08:52:12.656082Z","submitted_at":"2020-01-23T03:59:20Z","title":"Scaling Laws for Neural Language Models","version":1},"cited_work":{"arxiv_id":"2001.08361","doi":"10.1145/3616855.3635845","metadata_source":"pith","pith_arxiv_id":"2001.08361","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Scaling Laws for Neural Language Models","venue":"cs.LG","work_id":"b7dd8749-9c45-4977-ab9b-64478dce1ae8","year":2020},"citing_paper":{"arxiv_id":"2406.20094","last_updated":"2025-05-08T00:24:02Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2024-06-28T17:59:01Z","title":"Scaling Synthetic Data Creation with 1,000,000,000 Personas","version":3},"reference_index":16,"source":"pdf_text","source_observed_at":"2026-05-16T00:03:55.599967Z"},"links":{"cited_paper":"/paper/2001.08361","citing_paper":"/paper/2406.20094"},"observation_digest":"sha256:e0f14741bac6407f135e037d040ef66117add37fa0ccf2fa9aa1f18a50398a2d","observation_id":"29cea877-8fe4-4e00-91c4-e99cac05bae1","resolution":{"observed_at":"2026-05-16T00:03:55.740286Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2403.04706","last_updated":"2024-03-07T18:00:40Z","snapshot_observed_at":"2026-07-06T17:41:10.677334Z","submitted_at":"2024-03-07T18:00:40Z","title":"Common 7B Language Models Already Possess Strong Math Capabilities","version":1},"cited_work":{"arxiv_id":"2403.04706","doi":"10.48550/arxiv.2403.04706","metadata_source":"arxiv_reference","pith_arxiv_id":"2403.04706","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Common 7b language models already possess strong math capabilities","venue":"arXiv (Cornell University)","work_id":"1fa5eff0-85a8-43ed-a52d-877fbdde5dab","year":2024},"citing_paper":{"arxiv_id":"2406.20094","last_updated":"2025-05-08T00:24:02Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2024-06-28T17:59:01Z","title":"Scaling Synthetic Data Creation with 1,000,000,000 Personas","version":3},"reference_index":17,"source":"pdf_text","source_observed_at":"2026-05-16T00:03:55.599967Z"},"links":{"cited_paper":"/paper/2403.04706","citing_paper":"/paper/2406.20094"},"observation_digest":"sha256:a97f0cfac91c726a26291178a0c91cf96b0762fa15790829bc14101b85c194d4","observation_id":"851ef560-07cf-4904-b041-711b48a2b200","resolution":{"observed_at":"2026-05-16T00:03:55.750774Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2310.02170","last_updated":"2024-11-15T04:30:04Z","snapshot_observed_at":"2026-07-06T16:27:11.151651Z","submitted_at":"2023-10-03T16:05:48Z","title":"A Dynamic LLM-Powered Agent Network for Task-Oriented Agent Collaboration","version":2},"cited_work":{"arxiv_id":"2310.02170","doi":null,"metadata_source":"pith","pith_arxiv_id":"2310.02170","snapshot_observed_at":"2026-07-03T20:48:56.429753Z","title":"A Dynamic LLM-Powered Agent Network for Task-Oriented Agent Collaboration","venue":"cs.CL","work_id":"34a2258f-5686-44be-9205-50522661f454","year":2023},"citing_paper":{"arxiv_id":"2406.20094","last_updated":"2025-05-08T00:24:02Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2024-06-28T17:59:01Z","title":"Scaling Synthetic Data Creation with 1,000,000,000 Personas","version":3},"reference_index":18,"source":"pdf_text","source_observed_at":"2026-05-16T00:03:55.599967Z"},"links":{"cited_paper":"/paper/2310.02170","citing_paper":"/paper/2406.20094"},"observation_digest":"sha256:01af17679bfecff04101ef84c8eb247ef59369c266101b62af068b0c0c17a400","observation_id":"488dde72-99e6-4ddf-be4a-60d3a260de87","resolution":{"observed_at":"2026-05-21T21:05:26.036961Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2401.16380","last_updated":"2024-01-29T18:19:08Z","snapshot_observed_at":"2026-07-06T17:21:58.983048Z","submitted_at":"2024-01-29T18:19:08Z","title":"Rephrasing the Web: A Recipe for Compute and Data-Efficient Language Modeling","version":1},"cited_work":{"arxiv_id":"2401.16380","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2401.16380","snapshot_observed_at":"2026-07-04T18:40:03.338957Z","title":"Rephrasing the web: A recipe for compute and data-efficient language modeling","venue":null,"work_id":"da0d38ca-36c3-4cdb-90c1-2525a4978a48","year":2024},"citing_paper":{"arxiv_id":"2406.20094","last_updated":"2025-05-08T00:24:02Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2024-06-28T17:59:01Z","title":"Scaling Synthetic Data Creation with 1,000,000,000 Personas","version":3},"reference_index":19,"source":"pdf_text","source_observed_at":"2026-05-16T00:03:55.599967Z"},"links":{"cited_paper":"/paper/2401.16380","citing_paper":"/paper/2406.20094"},"observation_digest":"sha256:479693d02f011c736b5b7bb475f1a68ef335a54d1439165a416ab70521bb0ae3","observation_id":"37c7d081-73c9-4f4d-8dc9-55aaf38085f6","resolution":{"observed_at":"2026-05-16T00:03:55.762502Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2305.13661","last_updated":"2023-10-26T20:45:39Z","snapshot_observed_at":"2026-07-06T15:31:05.250637Z","submitted_at":"2023-05-23T04:10:26Z","title":"On the Risk of Misinformation Pollution with Large Language Models","version":2},"cited_work":{"arxiv_id":"2305.13661","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2305.13661","snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"On the risk of misinformation pollution with large language models","venue":null,"work_id":"ef5c25a1-204c-4260-86c7-d49bd9372d08","year":2023},"citing_paper":{"arxiv_id":"2406.20094","last_updated":"2025-05-08T00:24:02Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2024-06-28T17:59:01Z","title":"Scaling Synthetic Data Creation with 1,000,000,000 Personas","version":3},"reference_index":20,"source":"pdf_text","source_observed_at":"2026-05-16T00:03:55.599967Z"},"links":{"cited_paper":"/paper/2305.13661","citing_paper":"/paper/2406.20094"},"observation_digest":"sha256:6cc128702b596c4f06908f991af44e176e8def87334194b84348e0be4c1101b3","observation_id":"369c9f5d-19ac-4f14-9f20-aba51671f76a","resolution":{"observed_at":"2026-05-16T00:03:55.768468Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2305.17493","last_updated":"2024-04-14T05:20:10Z","snapshot_observed_at":"2026-07-06T15:34:21.966171Z","submitted_at":"2023-05-27T15:10:41Z","title":"The Curse of Recursion: Training on Generated Data Makes Models Forget","version":3},"cited_work":{"arxiv_id":"2305.17493","doi":"10.48550/arxiv.2305.17493","metadata_source":"pith","pith_arxiv_id":"2305.17493","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"The Curse of Recursion: Training on Generated Data Makes Models Forget","venue":"cs.LG","work_id":"14bc9b3b-d8dd-4c25-a0be-6c178d31be0e","year":2023},"citing_paper":{"arxiv_id":"2406.20094","last_updated":"2025-05-08T00:24:02Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2024-06-28T17:59:01Z","title":"Scaling Synthetic Data Creation with 1,000,000,000 Personas","version":3},"reference_index":21,"source":"pdf_text","source_observed_at":"2026-05-16T00:03:55.599967Z"},"links":{"cited_paper":"/paper/2305.17493","citing_paper":"/paper/2406.20094"},"observation_digest":"sha256:f5c981e982e25450891f1e6108653fec489cd809efe73e2086dc33df6dcabc22","observation_id":"46a5c8b2-a400-42f2-b81c-7819b46e4343","resolution":{"observed_at":"2026-05-20T14:04:58.330652Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-05-20T14:22:09.312587+00:00","source":"crossref_status_cache"},{"observed_at":"2026-05-20T14:22:09.312587+00:00","source":"openalex_status_cache"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2310.03731","last_updated":"2023-10-05T17:52:09Z","snapshot_observed_at":"2026-07-06T16:28:22.350574Z","submitted_at":"2023-10-05T17:52:09Z","title":"MathCoder: Seamless Code Integration in LLMs for Enhanced Mathematical Reasoning","version":1},"cited_work":{"arxiv_id":"2310.03731","doi":"10.48550/arxiv.2310.03731","metadata_source":"arxiv_reference","pith_arxiv_id":"2310.03731","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Mathcoder: Seamless code integration in llms for enhanced mathematical reasoning","venue":"arXiv (Cornell University)","work_id":"36308722-f027-40ce-b4e0-8d2c26bb0f9e","year":2023},"citing_paper":{"arxiv_id":"2406.20094","last_updated":"2025-05-08T00:24:02Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2024-06-28T17:59:01Z","title":"Scaling Synthetic Data Creation with 1,000,000,000 Personas","version":3},"reference_index":22,"source":"pdf_text","source_observed_at":"2026-05-16T00:03:55.599967Z"},"links":{"cited_paper":"/paper/2310.03731","citing_paper":"/paper/2406.20094"},"observation_digest":"sha256:2ad7363fd16cc7b54a04dcf6be5b9b222f7dc39a426ada3276f44d0ed1338da9","observation_id":"6ea04555-ccdf-4642-b5d7-875164cde317","resolution":{"observed_at":"2026-05-16T00:03:55.779793Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2212.10560","last_updated":"2023-05-25T23:50:07Z","snapshot_observed_at":"2026-07-06T14:33:11.945106Z","submitted_at":"2022-12-20T18:59:19Z","title":"Self-Instruct: Aligning Language Models with Self-Generated Instructions","version":2},"cited_work":{"arxiv_id":"2212.10560","doi":"10.1145/3209978.3210080","metadata_source":"pith","pith_arxiv_id":"2212.10560","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Self-Instruct: Aligning Language Models with Self-Generated Instructions","venue":"cs.CL","work_id":"d0018767-775d-406e-861d-539ed681ff73","year":2022},"citing_paper":{"arxiv_id":"2406.20094","last_updated":"2025-05-08T00:24:02Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2024-06-28T17:59:01Z","title":"Scaling Synthetic Data Creation with 1,000,000,000 Personas","version":3},"reference_index":23,"source":"pdf_text","source_observed_at":"2026-05-16T00:03:55.599967Z"},"links":{"cited_paper":"/paper/2212.10560","citing_paper":"/paper/2406.20094"},"observation_digest":"sha256:9e07c5aaca67ba09d9475b4c4347336e4f37331404619ec6589ff8e4a247e75f","observation_id":"95a360fd-5c36-4ba0-816f-4bbb0314a033","resolution":{"observed_at":"2026-05-16T00:03:55.784962Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-05-25T23:23:20.678+00:00","source":"crossref_status_cache"},{"observed_at":"2026-05-25T23:23:20.678+00:00","source":"openalex_status_cache"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Unleashing the emergent cognitive synergy in large language models: A task-solving agent through multi- persona self-collaboration","venue":null,"work_id":"a67bebb1-e7bd-4526-ba12-996d21fc43ea","year":2024},"citing_paper":{"arxiv_id":"2406.20094","last_updated":"2025-05-08T00:24:02Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2024-06-28T17:59:01Z","title":"Scaling Synthetic Data Creation with 1,000,000,000 Personas","version":3},"reference_index":24,"source":"pdf_text","source_observed_at":"2026-05-16T00:03:55.599967Z"},"links":{"citing_paper":"/paper/2406.20094"},"observation_digest":"sha256:58368fca493f031dfe59d3b8206c452c97f51c881aee3adf834a7c40442fefb5","observation_id":"7a3f031f-02a2-4cb7-b9ab-8d4fb4841384","resolution":{"observed_at":"2026-05-16T00:03:55.808297Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2401.11817","last_updated":"2025-02-13T08:11:25Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2024-01-22T10:26:14Z","title":"Hallucination is Inevitable: An Innate Limitation of Large Language Models","version":2},"cited_work":{"arxiv_id":"2401.11817","doi":"10.48550/arxiv.2401.11817","metadata_source":"pith","pith_arxiv_id":"2401.11817","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Hallucination is Inevitable: An Innate Limitation of Large Language Models","venue":"cs.CL","work_id":"9297385f-b00c-4c3e-b65b-204d2fdbbfe3","year":2024},"citing_paper":{"arxiv_id":"2406.20094","last_updated":"2025-05-08T00:24:02Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2024-06-28T17:59:01Z","title":"Scaling Synthetic Data Creation with 1,000,000,000 Personas","version":3},"reference_index":25,"source":"pdf_text","source_observed_at":"2026-05-16T00:03:55.599967Z"},"links":{"cited_paper":"/paper/2401.11817","citing_paper":"/paper/2406.20094"},"observation_digest":"sha256:6a01d94d7680a1590338adb989cc37e59f9ec2e9e2f79276884f4136a601b1ba","observation_id":"63994169-d189-4dfd-a856-2b4ee05b0678","resolution":{"observed_at":"2026-05-16T00:03:55.790579Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2403.04652","last_updated":"2025-01-21T10:12:05Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2024-03-07T16:52:49Z","title":"Yi: Open Foundation Models by 01.AI","version":3},"cited_work":{"arxiv_id":"2403.04652","doi":"10.48550/arxiv.2403.04652","metadata_source":"pith","pith_arxiv_id":"2403.04652","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Yi: Open Foundation Models by 01.AI","venue":"cs.CL","work_id":"8efee8a1-5e3c-4851-9c65-18e3d1d9e769","year":2024},"citing_paper":{"arxiv_id":"2406.20094","last_updated":"2025-05-08T00:24:02Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2024-06-28T17:59:01Z","title":"Scaling Synthetic Data Creation with 1,000,000,000 Personas","version":3},"reference_index":26,"source":"pdf_text","source_observed_at":"2026-05-16T00:03:55.599967Z"},"links":{"cited_paper":"/paper/2403.04652","citing_paper":"/paper/2406.20094"},"observation_digest":"sha256:3eb19db43cd044f61518ddf09bddea7f894911ebb201774962592bdc084c274d","observation_id":"82e1358a-134e-4c47-93b5-12daca19527d","resolution":{"observed_at":"2026-05-16T00:03:55.795468Z","resolver_source":"local_arxiv","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2309.12284","last_updated":"2024-05-03T17:36:07Z","snapshot_observed_at":"2026-08-02T15:00:50.388422Z","submitted_at":"2023-09-21T17:45:42Z","title":"MetaMath: Bootstrap Your Own Mathematical Questions for Large Language Models","version":4},"cited_work":{"arxiv_id":"2309.12284","doi":"10.48550/arxiv.2309.12284","metadata_source":"pith","pith_arxiv_id":"2309.12284","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"MetaMath: Bootstrap Your Own Mathematical Questions for Large Language Models","venue":"cs.CL","work_id":"e791a2f0-3b75-4c5b-84e1-d297d96e42f0","year":2023},"citing_paper":{"arxiv_id":"2406.20094","last_updated":"2025-05-08T00:24:02Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2024-06-28T17:59:01Z","title":"Scaling Synthetic Data Creation with 1,000,000,000 Personas","version":3},"reference_index":27,"source":"pdf_text","source_observed_at":"2026-05-16T00:03:55.599967Z"},"links":{"cited_paper":"/paper/2309.12284","citing_paper":"/paper/2406.20094"},"observation_digest":"sha256:0a42371cdde4196071078aa1606576d3b3c297eec5e622b7ac42bd18fb309cfa","observation_id":"d65a5ea4-362f-4056-8f6e-888cf512b8f6","resolution":{"observed_at":"2026-05-16T00:03:55.800476Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2404.01230","last_updated":"2024-04-01T16:50:54Z","snapshot_observed_at":"2026-07-06T17:54:03.923842Z","submitted_at":"2024-04-01T16:50:54Z","title":"LLM as a Mastermind: A Survey of Strategic Reasoning with Large Language Models","version":1},"cited_work":{"arxiv_id":"2404.01230","doi":"10.48550/arxiv.2404.01230","metadata_source":"arxiv_reference","pith_arxiv_id":"2404.01230","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Llm as a mastermind: A survey of strategic reasoning with large language models","venue":"arXiv (Cornell University)","work_id":"1691976b-ba8d-472d-adfd-f9e5805c56b9","year":2024},"citing_paper":{"arxiv_id":"2406.20094","last_updated":"2025-05-08T00:24:02Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2024-06-28T17:59:01Z","title":"Scaling Synthetic Data Creation with 1,000,000,000 Personas","version":3},"reference_index":28,"source":"pdf_text","source_observed_at":"2026-05-16T00:03:55.599967Z"},"links":{"cited_paper":"/paper/2404.01230","citing_paper":"/paper/2406.20094"},"observation_digest":"sha256:1892e32ae9055829f50480c848a662edc7d7f44cb99aaa1509e50902e88d6268","observation_id":"b9d270ea-5da7-4e29-9721-df54b58721cd","resolution":{"observed_at":"2026-05-16T00:03:55.660291Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2406.11931","last_updated":"2024-06-17T13:51:35Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2024-06-17T13:51:35Z","title":"DeepSeek-Coder-V2: Breaking the Barrier of Closed-Source Models in Code Intelligence","version":1},"cited_work":{"arxiv_id":"2406.11931","doi":"10.48550/arxiv.2406.11931","metadata_source":"pith","pith_arxiv_id":"2406.11931","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"DeepSeek-Coder-V2: Breaking the Barrier of Closed-Source Models in Code Intelligence","venue":"cs.SE","work_id":"713a39be-8c4e-4ee3-8671-11cf9379262f","year":2024},"citing_paper":{"arxiv_id":"2406.20094","last_updated":"2025-05-08T00:24:02Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2024-06-28T17:59:01Z","title":"Scaling Synthetic Data Creation with 1,000,000,000 Personas","version":3},"reference_index":29,"source":"pdf_text","source_observed_at":"2026-05-16T00:03:55.599967Z"},"links":{"cited_paper":"/paper/2406.11931","citing_paper":"/paper/2406.20094"},"observation_digest":"sha256:b60e1e699e9a4b97573ab70449f1a0a371f70dbdeefc45bb5697ffec04218bb1","observation_id":"742740b9-16c7-43ad-8ffd-cb5b1eb57471","resolution":{"observed_at":"2026-05-16T01:06:07.964742Z","resolver_source":"arxiv_id","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}}],"paper":{"arxiv_id":"2406.20094","last_updated":"2025-05-08T00:24:02Z","latest_version":3,"primary_category":"cs.CL","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2024-06-28T17:59:01Z","title":"Scaling Synthetic Data Creation with 1,000,000,000 Personas"},"reference_resolution":{"displayed":29,"state_counts":{"malformed_identifier":0,"metadata_mismatch":7,"parse_uncertain":0,"unresolved":0,"verified_exact":20,"verified_fuzzy":2},"total_outbound_references":29},"refusal":"A citation records a reference. It does not transfer a finding from one paper to another.","schema":"pith.paper-citation-record.v1","standing_sources":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"thesis":"As of 5 August 2026, this Paper Citation Record lists 29 of 29 outbound references and 74 inbound Pith citation observations for arXiv:2406.20094."}