{"as_of":"2026-08-07T21:04:00Z","caps":{"database_statements":6,"inbound":100,"outbound":100},"context_digest":"sha256:5a52e0a1d6860f450eaf79e2bb4248452aa84a4d74600463e3f6a41fb0c383ce","coverage":[{"denominator":43,"lane":"reference_resolution","note":"Typed states for the displayed outbound observations.","records_observed":43,"source":"paper_references, paper_reference_links","source_observed_at":"2026-08-07T14:00:33.402147Z","state":"measured"},{"denominator":50,"lane":"standing_notices","note":"One-hop event checks from named stored sources.","records_observed":50,"source":"scholarly_work_events, retraction_status_cache","source_observed_at":"2026-08-07T06:34:17.273281+00:00","state":"measured"},{"denominator":7,"lane":"inbound_itemization","note":"Pith citing papers itemized under the disclosed page cap.","records_observed":7,"source":"paper_references, paper_reference_links","source_observed_at":"2026-08-06T20:00:56.113141Z","state":"measured"},{"denominator":1,"lane":"external_citation_measurements","note":"A source-named dated measurement, never combined with another source.","records_observed":1,"source":"arxiv_reference","source_observed_at":"2026-08-05T02:28:24.338817Z","state":"measured"}],"external_citation_measurements":[{"count":0,"observed_at":"2026-08-05T02:28:24.338817Z","source":"arxiv_reference"}],"inbound":[{"citation":{"cited_paper":{"arxiv_id":"2505.20416","last_updated":"2025-05-26T18:06:50Z","snapshot_observed_at":"2026-08-07T13:52:46.937698Z","submitted_at":"2025-05-26T18:06:50Z","title":"GraphGen: Enhancing Supervised Fine-Tuning for LLMs with Knowledge-Driven Synthetic Data Generation","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2505.20416","snapshot_observed_at":"2026-08-06T20:00:56.113141Z","title":"arXiv preprint arXiv:2505.20416","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2507.04009","last_updated":"2025-07-05T11:38:59Z","snapshot_observed_at":"2026-08-06T19:55:08.952460Z","submitted_at":"2025-07-05T11:38:59Z","title":"Easy Dataset: A Unified and Extensible Framework for Synthesizing LLM Fine-Tuning Data from Unstructured Documents","version":1},"reference_index":2025,"source":"pdf_text","source_observed_at":"2026-08-06T20:00:56.113141Z"},"links":{"cited_paper":"/paper/2505.20416","citing_paper":"/paper/2507.04009"},"observation_digest":"sha256:6818bfb15342afacac2ad115f48b34e99657155ce726f1820cf28bfafd35a08d","observation_id":"dde05f2d-f7b4-4d0d-9660-c04418cba721","resolution":{"observed_at":"2026-08-06T20:00:56.113141Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2505.20416","last_updated":"2025-05-26T18:06:50Z","snapshot_observed_at":"2026-08-07T13:52:46.937698Z","submitted_at":"2025-05-26T18:06:50Z","title":"GraphGen: Enhancing Supervised Fine-Tuning for LLMs with Knowledge-Driven Synthetic Data Generation","version":1},"cited_work":{"arxiv_id":"2505.20416","doi":"10.48550/arxiv.2505.20416","metadata_source":"arxiv_reference","pith_arxiv_id":"2505.20416","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Graphgen: Enhancing supervised fine-tuning for llms with knowledge-driven synthetic data generation","venue":"ArXiv.org","work_id":"e01e8e3e-f11f-449e-b079-e093026cec97","year":2025},"citing_paper":{"arxiv_id":"2602.12705","last_updated":"2026-04-07T11:35:36Z","snapshot_observed_at":"2026-07-06T22:45:44.135338Z","submitted_at":"2026-02-13T08:19:38Z","title":"MedXIAOHE: A Comprehensive Recipe for Building Medical MLLMs","version":4},"reference_index":10,"source":"pdf_text","source_observed_at":"2026-05-15T22:52:30.992054Z"},"links":{"cited_paper":"/paper/2505.20416","citing_paper":"/paper/2602.12705"},"observation_digest":"sha256:d7e511d858ca1e3acecda82626a68462e463a2f41d38634a919665aa502b2b35","observation_id":"541e15e3-fab2-4269-9ca0-a4bc0a88c7e0","resolution":{"observed_at":"2026-05-15T22:56:50.406292Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-05-20T14:22:04.374145+00:00","source":"crossref_status_cache"},{"observed_at":"2026-05-20T14:22:04.374145+00:00","source":"openalex_status_cache"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2505.20416","last_updated":"2025-05-26T18:06:50Z","snapshot_observed_at":"2026-08-07T13:52:46.937698Z","submitted_at":"2025-05-26T18:06:50Z","title":"GraphGen: Enhancing Supervised Fine-Tuning for LLMs with Knowledge-Driven Synthetic Data Generation","version":1},"cited_work":{"arxiv_id":"2505.20416","doi":"10.48550/arxiv.2505.20416","metadata_source":"arxiv_reference","pith_arxiv_id":"2505.20416","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Graphgen: Enhancing supervised fine-tuning for llms with knowledge-driven synthetic data generation","venue":"ArXiv.org","work_id":"e01e8e3e-f11f-449e-b079-e093026cec97","year":2025},"citing_paper":{"arxiv_id":"2605.08709","last_updated":"2026-05-09T05:44:29Z","snapshot_observed_at":"2026-07-06T23:20:52.521341Z","submitted_at":"2026-05-09T05:44:29Z","title":"UniShield: Unified Face Attack Detection via KG-Informed Multimodal Reasoning","version":1},"reference_index":4,"source":"pdf_text","source_observed_at":"2026-05-12T01:10:21.661074Z"},"links":{"cited_paper":"/paper/2505.20416","citing_paper":"/paper/2605.08709"},"observation_digest":"sha256:5cf8aa40905a9ca7116a1723085f233a51f46777b0fde45a0962758130b392f7","observation_id":"a015fdf5-7c8f-43ca-92e5-4f94a3fb8400","resolution":{"observed_at":"2026-05-12T08:26:24.434798Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-05-20T14:22:04.374145+00:00","source":"crossref_status_cache"},{"observed_at":"2026-05-20T14:22:04.374145+00:00","source":"openalex_status_cache"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2505.20416","last_updated":"2025-05-26T18:06:50Z","snapshot_observed_at":"2026-08-07T13:52:46.937698Z","submitted_at":"2025-05-26T18:06:50Z","title":"GraphGen: Enhancing Supervised Fine-Tuning for LLMs with Knowledge-Driven Synthetic Data Generation","version":1},"cited_work":{"arxiv_id":"2505.20416","doi":"10.48550/arxiv.2505.20416","metadata_source":"arxiv_reference","pith_arxiv_id":"2505.20416","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Graphgen: Enhancing supervised fine-tuning for llms with knowledge-driven synthetic data generation","venue":"ArXiv.org","work_id":"e01e8e3e-f11f-449e-b079-e093026cec97","year":2025},"citing_paper":{"arxiv_id":"2605.18025","last_updated":"2026-05-18T08:14:49Z","snapshot_observed_at":"2026-07-06T23:28:55.079333Z","submitted_at":"2026-05-18T08:14:49Z","title":"TeleCom-Bench: How Far Are Large Language Models from Industrial Telecommunication Applications?","version":1},"reference_index":6,"source":"pdf_text","source_observed_at":"2026-05-20T11:23:34.256279Z"},"links":{"cited_paper":"/paper/2505.20416","citing_paper":"/paper/2605.18025"},"observation_digest":"sha256:20425d0d8967ce36572ae97501ea355cc1fa0ca4f06c077be3bd32e10c3ffc74","observation_id":"33337c6f-9e3a-452c-a134-5004acaa7301","resolution":{"observed_at":"2026-05-20T11:28:14.626268Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-05-20T14:22:04.374145+00:00","source":"crossref_status_cache"},{"observed_at":"2026-05-20T14:22:04.374145+00:00","source":"openalex_status_cache"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2505.20416","last_updated":"2025-05-26T18:06:50Z","snapshot_observed_at":"2026-08-07T13:52:46.937698Z","submitted_at":"2025-05-26T18:06:50Z","title":"GraphGen: Enhancing Supervised Fine-Tuning for LLMs with Knowledge-Driven Synthetic Data Generation","version":1},"cited_work":{"arxiv_id":"2505.20416","doi":"10.48550/arxiv.2505.20416","metadata_source":"arxiv_reference","pith_arxiv_id":"2505.20416","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Graphgen: Enhancing supervised fine-tuning for llms with knowledge-driven synthetic data generation","venue":"ArXiv.org","work_id":"e01e8e3e-f11f-449e-b079-e093026cec97","year":2025},"citing_paper":{"arxiv_id":"2605.19394","last_updated":"2026-05-19T05:40:12Z","snapshot_observed_at":"2026-08-02T23:18:14.648216Z","submitted_at":"2026-05-19T05:40:12Z","title":"EmbGen: Teaching with Reassembled Corpora","version":1},"reference_index":4,"source":"pdf_text","source_observed_at":"2026-05-20T06:34:39.739666Z"},"links":{"cited_paper":"/paper/2505.20416","citing_paper":"/paper/2605.19394"},"observation_digest":"sha256:0996acba22c0dd1cd3a327cb3e5c9250f0a3d5afebcadaa25129166085d055b2","observation_id":"9337b38c-61da-4dc5-be05-4cee59c5f96c","resolution":{"observed_at":"2026-05-20T06:38:05.530130Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-05-20T14:22:04.374145+00:00","source":"crossref_status_cache"},{"observed_at":"2026-05-20T14:22:04.374145+00:00","source":"openalex_status_cache"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2505.20416","last_updated":"2025-05-26T18:06:50Z","snapshot_observed_at":"2026-08-07T13:52:46.937698Z","submitted_at":"2025-05-26T18:06:50Z","title":"GraphGen: Enhancing Supervised Fine-Tuning for LLMs with Knowledge-Driven Synthetic Data Generation","version":1},"cited_work":{"arxiv_id":"2505.20416","doi":"10.48550/arxiv.2505.20416","metadata_source":"arxiv_reference","pith_arxiv_id":"2505.20416","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Graphgen: Enhancing supervised fine-tuning for llms with knowledge-driven synthetic data generation","venue":"ArXiv.org","work_id":"e01e8e3e-f11f-449e-b079-e093026cec97","year":2025},"citing_paper":{"arxiv_id":"2606.12087","last_updated":"2026-06-10T13:49:11Z","snapshot_observed_at":"2026-08-02T15:03:50.335992Z","submitted_at":"2026-06-10T13:49:11Z","title":"FORT-Searcher: Synthesizing Shortcut-Resistant Search Tasks for Training Deep Search Agents","version":1},"reference_index":46,"source":"arxiv_source","source_observed_at":"2026-06-27T10:01:45.332920Z"},"links":{"cited_paper":"/paper/2505.20416","citing_paper":"/paper/2606.12087"},"observation_digest":"sha256:3d38c244ea157c50fb8ffdea7ecf4a3cb43a5606e0ff59c079f46011a4cb9b5d","observation_id":"88e06038-75e0-43e4-a20f-4a9faf677921","resolution":{"observed_at":"2026-07-03T10:27:56.494259Z","resolver_source":"arxiv_id","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-05-20T14:22:04.374145+00:00","source":"crossref_status_cache"},{"observed_at":"2026-05-20T14:22:04.374145+00:00","source":"openalex_status_cache"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2505.20416","last_updated":"2025-05-26T18:06:50Z","snapshot_observed_at":"2026-08-07T13:52:46.937698Z","submitted_at":"2025-05-26T18:06:50Z","title":"GraphGen: Enhancing Supervised Fine-Tuning for LLMs with Knowledge-Driven Synthetic Data Generation","version":1},"cited_work":{"arxiv_id":"2505.20416","doi":"10.48550/arxiv.2505.20416","metadata_source":"arxiv_reference","pith_arxiv_id":"2505.20416","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Graphgen: Enhancing supervised fine-tuning for llms with knowledge-driven synthetic data generation","venue":"ArXiv.org","work_id":"e01e8e3e-f11f-449e-b079-e093026cec97","year":2025},"citing_paper":{"arxiv_id":"2606.12837","last_updated":"2026-06-17T03:34:47Z","snapshot_observed_at":"2026-07-06T23:51:40.528967Z","submitted_at":"2026-06-11T03:04:32Z","title":"LoHoSearch: Benchmarking Long-Horizon Search Agents Beyond the Human Difficulty Ceiling","version":2},"reference_index":11,"source":"arxiv_source","source_observed_at":"2026-06-27T07:02:40.293652Z"},"links":{"cited_paper":"/paper/2505.20416","citing_paper":"/paper/2606.12837"},"observation_digest":"sha256:134cf274bdc7783fabcd30ab2f5d04c64ee8be6cb2fd3a16e4ff3d79dd2ed2d1","observation_id":"baccde30-33c1-41f3-b1e9-c690b5ce62f0","resolution":{"observed_at":"2026-07-03T14:28:31.706678Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-05-20T14:22:04.374145+00:00","source":"crossref_status_cache"},{"observed_at":"2026-05-20T14:22:04.374145+00:00","source":"openalex_status_cache"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}}],"links":{"evidence":"/evidence","html":"/paper/2505.20416/citation-record","integrity":"/paper/2505.20416/integrity","json":"/paper/2505.20416/citation-record.json","paper":"/paper/2505.20416"},"outbound":[{"citation":{"cited_paper":{"arxiv_id":"2407.05015","last_updated":"2024-07-06T09:10:05Z","snapshot_observed_at":"2026-07-06T18:42:18.289966Z","submitted_at":"2024-07-06T09:10:05Z","title":"How do you know that? Teaching Generative Language Models to Reference Answers to Biomedical Questions","version":1},"cited_work":{"arxiv_id":"2407.05015","doi":null,"metadata_source":"pith","pith_arxiv_id":"2407.05015","snapshot_observed_at":"2026-08-07T14:00:34.351519Z","title":"How do you know that? Teaching Generative Language Models to Reference Answers to Biomedical Questions","venue":"cs.CL","work_id":"c41d0524-5d49-4f4f-8e16-5f6a7e45925b","year":2024},"citing_paper":{"arxiv_id":"2505.20416","last_updated":"2025-05-26T18:06:50Z","snapshot_observed_at":"2026-08-07T13:52:46.937698Z","submitted_at":"2025-05-26T18:06:50Z","title":"GraphGen: Enhancing Supervised Fine-Tuning for LLMs with Knowledge-Driven Synthetic Data Generation","version":1},"reference_index":1,"source":"arxiv_source","source_observed_at":"2026-08-07T14:00:29.497217Z"},"links":{"cited_paper":"/paper/2407.05015","citing_paper":"/paper/2505.20416"},"observation_digest":"sha256:7f6d647135377c9ed84108e79c1ae6d94c93dddff7a18ef536eaceb3705f3fcf","observation_id":"5de4fea8-03f8-4057-8ed6-db863ca1c556","resolution":{"observed_at":"2026-08-07T14:00:34.410611Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2307.06290","last_updated":"2024-07-26T18:09:11Z","snapshot_observed_at":"2026-07-06T15:53:19.485466Z","submitted_at":"2023-07-12T16:37:31Z","title":"Instruction Mining: Instruction Data Selection for Tuning Large Language Models","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2307.06290","snapshot_observed_at":"2026-08-07T14:00:29.600015Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2505.20416","last_updated":"2025-05-26T18:06:50Z","snapshot_observed_at":"2026-08-07T13:52:46.937698Z","submitted_at":"2025-05-26T18:06:50Z","title":"GraphGen: Enhancing Supervised Fine-Tuning for LLMs with Knowledge-Driven Synthetic Data Generation","version":1},"reference_index":2,"source":"arxiv_source","source_observed_at":"2026-08-07T14:00:29.600015Z"},"links":{"cited_paper":"/paper/2307.06290","citing_paper":"/paper/2505.20416"},"observation_digest":"sha256:b747aa60b19f45042f6f3fd79b63937538a7b1d092a0e1d8a1410f30a7bf58ee","observation_id":"c2e9f62e-708f-4c21-84bc-ea1c79d19ced","resolution":{"observed_at":"2026-08-07T14:00:29.600015Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T14:00:37.488225Z","title":null,"venue":null,"work_id":"d6686898-1f2d-4b9f-8874-d11324c522c2","year":2023},"citing_paper":{"arxiv_id":"2505.20416","last_updated":"2025-05-26T18:06:50Z","snapshot_observed_at":"2026-08-07T13:52:46.937698Z","submitted_at":"2025-05-26T18:06:50Z","title":"GraphGen: Enhancing Supervised Fine-Tuning for LLMs with Knowledge-Driven Synthetic Data Generation","version":1},"reference_index":3,"source":"arxiv_source","source_observed_at":"2026-08-07T14:00:29.738159Z"},"links":{"citing_paper":"/paper/2505.20416"},"observation_digest":"sha256:bf247fa7e2017ce91efa13dbac86a5725920235fc5c8dab84b7c4009232f9e56","observation_id":"886f756d-3c45-4b6a-90e8-addf5704fe5b","resolution":{"observed_at":"2026-08-07T14:00:37.588231Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T14:00:37.343660Z","title":null,"venue":null,"work_id":"dcc814b6-3083-4e7b-9e41-321d49ecc83e","year":2017},"citing_paper":{"arxiv_id":"2505.20416","last_updated":"2025-05-26T18:06:50Z","snapshot_observed_at":"2026-08-07T13:52:46.937698Z","submitted_at":"2025-05-26T18:06:50Z","title":"GraphGen: Enhancing Supervised Fine-Tuning for LLMs with Knowledge-Driven Synthetic Data Generation","version":1},"reference_index":4,"source":"arxiv_source","source_observed_at":"2026-08-07T14:00:29.914295Z"},"links":{"citing_paper":"/paper/2505.20416"},"observation_digest":"sha256:cc7b6a09c48279f88feac1618ed24931412e0bdf55f89b6e2a2b59a66182a68e","observation_id":"ea6c4a00-bfde-4796-8c33-c4182977a0a2","resolution":{"observed_at":"2026-08-07T14:00:37.424749Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T14:00:37.177911Z","title":null,"venue":null,"work_id":"5841b7c9-c1e4-4450-a5c5-47df84155b38","year":2022},"citing_paper":{"arxiv_id":"2505.20416","last_updated":"2025-05-26T18:06:50Z","snapshot_observed_at":"2026-08-07T13:52:46.937698Z","submitted_at":"2025-05-26T18:06:50Z","title":"GraphGen: Enhancing Supervised Fine-Tuning for LLMs with Knowledge-Driven Synthetic Data Generation","version":1},"reference_index":5,"source":"arxiv_source","source_observed_at":"2026-08-07T14:00:30.089910Z"},"links":{"citing_paper":"/paper/2505.20416"},"observation_digest":"sha256:565c296ea0f026cd7f30c3ca915c69509080e23ef40cea58efb9e0e4d8b328d4","observation_id":"8059f304-dd4d-4f30-8fe7-cd4313097b88","resolution":{"observed_at":"2026-08-07T14:00:37.264612Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T14:00:36.991495Z","title":null,"venue":null,"work_id":"d67978ac-de11-4a4b-b9dc-c2de3af69ac1","year":2024},"citing_paper":{"arxiv_id":"2505.20416","last_updated":"2025-05-26T18:06:50Z","snapshot_observed_at":"2026-08-07T13:52:46.937698Z","submitted_at":"2025-05-26T18:06:50Z","title":"GraphGen: Enhancing Supervised Fine-Tuning for LLMs with Knowledge-Driven Synthetic Data Generation","version":1},"reference_index":6,"source":"arxiv_source","source_observed_at":"2026-08-07T14:00:30.238178Z"},"links":{"citing_paper":"/paper/2505.20416"},"observation_digest":"sha256:f715d01106124754773fa931d109e0669e5f55d03c73ac05de1b7b10efa023e8","observation_id":"43373927-271c-4ea3-a94e-82f052489631","resolution":{"observed_at":"2026-08-07T14:00:37.080345Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T14:00:36.803646Z","title":null,"venue":null,"work_id":"857815fe-e2e8-414f-9e13-adddb0a28f79","year":2017},"citing_paper":{"arxiv_id":"2505.20416","last_updated":"2025-05-26T18:06:50Z","snapshot_observed_at":"2026-08-07T13:52:46.937698Z","submitted_at":"2025-05-26T18:06:50Z","title":"GraphGen: Enhancing Supervised Fine-Tuning for LLMs with Knowledge-Driven Synthetic Data Generation","version":1},"reference_index":7,"source":"arxiv_source","source_observed_at":"2026-08-07T14:00:30.288868Z"},"links":{"citing_paper":"/paper/2505.20416"},"observation_digest":"sha256:d8b2c9c1c8c862add8246c6247af8628b9f201d85d9f2c10757be4cabe7430a5","observation_id":"87f6604b-84cd-46fd-bb3d-71ee7d5a4266","resolution":{"observed_at":"2026-08-07T14:00:36.907393Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T14:00:36.608237Z","title":null,"venue":null,"work_id":"c7821090-84b1-4bfe-9d84-c4f07e792678","year":2024},"citing_paper":{"arxiv_id":"2505.20416","last_updated":"2025-05-26T18:06:50Z","snapshot_observed_at":"2026-08-07T13:52:46.937698Z","submitted_at":"2025-05-26T18:06:50Z","title":"GraphGen: Enhancing Supervised Fine-Tuning for LLMs with Knowledge-Driven Synthetic Data Generation","version":1},"reference_index":8,"source":"arxiv_source","source_observed_at":"2026-08-07T14:00:30.359664Z"},"links":{"citing_paper":"/paper/2505.20416"},"observation_digest":"sha256:612294df07a0c1ad5d3b1a2fb1e790af4caf2c78a472e865e1c25d0684afa146","observation_id":"11dd96d6-3b39-4e9d-aa74-11f3d5cafb96","resolution":{"observed_at":"2026-08-07T14:00:36.703652Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2410.05779","last_updated":"2025-04-28T17:36:27Z","snapshot_observed_at":"2026-07-31T06:48:41.810549Z","submitted_at":"2024-10-08T08:00:12Z","title":"LightRAG: Simple and Fast Retrieval-Augmented Generation","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2410.05779","snapshot_observed_at":"2026-08-07T14:00:30.443560Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2505.20416","last_updated":"2025-05-26T18:06:50Z","snapshot_observed_at":"2026-08-07T13:52:46.937698Z","submitted_at":"2025-05-26T18:06:50Z","title":"GraphGen: Enhancing Supervised Fine-Tuning for LLMs with Knowledge-Driven Synthetic Data Generation","version":1},"reference_index":9,"source":"arxiv_source","source_observed_at":"2026-08-07T14:00:30.443560Z"},"links":{"cited_paper":"/paper/2410.05779","citing_paper":"/paper/2505.20416"},"observation_digest":"sha256:42ca2b9520d65014e40f1dfa00887ea781439a001e12f54fec64baf9f5fc6184","observation_id":"78d230ca-c325-4625-838c-04acee48a7dd","resolution":{"observed_at":"2026-08-07T14:00:30.443560Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T14:00:30.526434Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2505.20416","last_updated":"2025-05-26T18:06:50Z","snapshot_observed_at":"2026-08-07T13:52:46.937698Z","submitted_at":"2025-05-26T18:06:50Z","title":"GraphGen: Enhancing Supervised Fine-Tuning for LLMs with Knowledge-Driven Synthetic Data Generation","version":1},"reference_index":10,"source":"arxiv_source","source_observed_at":"2026-08-07T14:00:30.526434Z"},"links":{"citing_paper":"/paper/2505.20416"},"observation_digest":"sha256:0f9e2f75d4a40c37105e219148cc216ddcc20d552a0d436343e7f9976911391d","observation_id":"8eca1af2-ef67-49d1-9194-12b445dcac64","resolution":{"observed_at":"2026-08-07T14:00:30.526434Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T14:00:36.388682Z","title":null,"venue":null,"work_id":"70d331e0-c766-46cd-ae06-df6f470886d1","year":2017},"citing_paper":{"arxiv_id":"2505.20416","last_updated":"2025-05-26T18:06:50Z","snapshot_observed_at":"2026-08-07T13:52:46.937698Z","submitted_at":"2025-05-26T18:06:50Z","title":"GraphGen: Enhancing Supervised Fine-Tuning for LLMs with Knowledge-Driven Synthetic Data Generation","version":1},"reference_index":11,"source":"arxiv_source","source_observed_at":"2026-08-07T14:00:30.603101Z"},"links":{"citing_paper":"/paper/2505.20416"},"observation_digest":"sha256:f52254f2a6fbfe243c709a92e7e8a9645f3ccc5bf0633b348d31a99b729f9140","observation_id":"7543da6e-9cc2-42fc-a5f8-0df8677c5d6b","resolution":{"observed_at":"2026-08-07T14:00:36.481755Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T14:00:30.678313Z","title":null,"venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2505.20416","last_updated":"2025-05-26T18:06:50Z","snapshot_observed_at":"2026-08-07T13:52:46.937698Z","submitted_at":"2025-05-26T18:06:50Z","title":"GraphGen: Enhancing Supervised Fine-Tuning for LLMs with Knowledge-Driven Synthetic Data Generation","version":1},"reference_index":12,"source":"arxiv_source","source_observed_at":"2026-08-07T14:00:30.678313Z"},"links":{"citing_paper":"/paper/2505.20416"},"observation_digest":"sha256:bb35d2e838a2cc37631e9e81b8e53f3f4de23e49c31b71f13ff6e0077a25d371","observation_id":"8d1afdfb-ef07-4ec2-9b6b-296b9bfead64","resolution":{"observed_at":"2026-08-07T14:00:30.678313Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T14:00:36.137393Z","title":null,"venue":null,"work_id":"4ad4744e-d42c-4b87-a57f-463deda86144","year":2016},"citing_paper":{"arxiv_id":"2505.20416","last_updated":"2025-05-26T18:06:50Z","snapshot_observed_at":"2026-08-07T13:52:46.937698Z","submitted_at":"2025-05-26T18:06:50Z","title":"GraphGen: Enhancing Supervised Fine-Tuning for LLMs with Knowledge-Driven Synthetic Data Generation","version":1},"reference_index":13,"source":"arxiv_source","source_observed_at":"2026-08-07T14:00:30.782135Z"},"links":{"citing_paper":"/paper/2505.20416"},"observation_digest":"sha256:b34526727ac27f0e5f8cad80a2f3180ca814b563382810e4b5a0988acdc1ef44","observation_id":"7e5587e7-c131-4354-af18-75673abbe7f0","resolution":{"observed_at":"2026-08-07T14:00:36.235786Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T14:00:30.845629Z","title":null,"venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2505.20416","last_updated":"2025-05-26T18:06:50Z","snapshot_observed_at":"2026-08-07T13:52:46.937698Z","submitted_at":"2025-05-26T18:06:50Z","title":"GraphGen: Enhancing Supervised Fine-Tuning for LLMs with Knowledge-Driven Synthetic Data Generation","version":1},"reference_index":14,"source":"arxiv_source","source_observed_at":"2026-08-07T14:00:30.845629Z"},"links":{"citing_paper":"/paper/2505.20416"},"observation_digest":"sha256:c368b5a082ab7a02188314a495e70d7bcf45f5fcb987c1c9663665beac6192de","observation_id":"7e841c88-6ea8-4a66-8005-302c34e511b1","resolution":{"observed_at":"2026-08-07T14:00:30.845629Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2001.08361","last_updated":"2020-01-23T03:59:20Z","snapshot_observed_at":"2026-07-06T08:52:12.656082Z","submitted_at":"2020-01-23T03:59:20Z","title":"Scaling Laws for Neural Language Models","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2001.08361","snapshot_observed_at":"2026-08-07T14:00:30.901388Z","title":"Brown, Benjamin Chess, Rewon Child, Scott Gray, Alec Radford, Jeffrey Wu, and Dario Amodei","venue":null,"work_id":null,"year":2020},"citing_paper":{"arxiv_id":"2505.20416","last_updated":"2025-05-26T18:06:50Z","snapshot_observed_at":"2026-08-07T13:52:46.937698Z","submitted_at":"2025-05-26T18:06:50Z","title":"GraphGen: Enhancing Supervised Fine-Tuning for LLMs with Knowledge-Driven Synthetic Data Generation","version":1},"reference_index":15,"source":"arxiv_source","source_observed_at":"2026-08-07T14:00:30.901388Z"},"links":{"cited_paper":"/paper/2001.08361","citing_paper":"/paper/2505.20416"},"observation_digest":"sha256:9bbdf90e6546ae2c31b2d99cdc2c0c56e03239ff5e6b1a287412ffa06af281cb","observation_id":"3b045b9e-ddd9-4621-945e-658300ccf784","resolution":{"observed_at":"2026-08-07T14:00:30.901388Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T14:00:35.919708Z","title":null,"venue":null,"work_id":"c18e4f2f-e0c3-4faa-bee5-8b6c21bd1042","year":2025},"citing_paper":{"arxiv_id":"2505.20416","last_updated":"2025-05-26T18:06:50Z","snapshot_observed_at":"2026-08-07T13:52:46.937698Z","submitted_at":"2025-05-26T18:06:50Z","title":"GraphGen: Enhancing Supervised Fine-Tuning for LLMs with Knowledge-Driven Synthetic Data Generation","version":1},"reference_index":16,"source":"arxiv_source","source_observed_at":"2026-08-07T14:00:30.981018Z"},"links":{"citing_paper":"/paper/2505.20416"},"observation_digest":"sha256:6d67fc5337242603c6141844441499f42d7ded2f7256370c34101433c5b6a837","observation_id":"431cfa69-32a5-4461-a638-58ecba4f8a38","resolution":{"observed_at":"2026-08-07T14:00:36.014540Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2304.08460","last_updated":"2024-10-03T15:46:13Z","snapshot_observed_at":"2026-08-04T19:43:12.423655Z","submitted_at":"2023-04-17T17:36:35Z","title":"LongForm: Effective Instruction Tuning with Reverse Instructions","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2304.08460","snapshot_observed_at":"2026-08-07T14:00:31.089767Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2505.20416","last_updated":"2025-05-26T18:06:50Z","snapshot_observed_at":"2026-08-07T13:52:46.937698Z","submitted_at":"2025-05-26T18:06:50Z","title":"GraphGen: Enhancing Supervised Fine-Tuning for LLMs with Knowledge-Driven Synthetic Data Generation","version":1},"reference_index":17,"source":"arxiv_source","source_observed_at":"2026-08-07T14:00:31.089767Z"},"links":{"cited_paper":"/paper/2304.08460","citing_paper":"/paper/2505.20416"},"observation_digest":"sha256:71ad9a398d95541bde60e3f4447c2c303381714221e7b92c4af92352813c1f5b","observation_id":"31c18690-2cb5-4afd-981b-b22ec3caf84b","resolution":{"observed_at":"2026-08-07T14:00:31.089767Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2406.16367","last_updated":"2024-06-24T07:17:59Z","snapshot_observed_at":"2026-07-06T18:35:49.991098Z","submitted_at":"2024-06-24T07:17:59Z","title":"On the Role of Long-tail Knowledge in Retrieval Augmented Large Language Models","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2406.16367","snapshot_observed_at":"2026-08-07T14:00:31.170015Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2505.20416","last_updated":"2025-05-26T18:06:50Z","snapshot_observed_at":"2026-08-07T13:52:46.937698Z","submitted_at":"2025-05-26T18:06:50Z","title":"GraphGen: Enhancing Supervised Fine-Tuning for LLMs with Knowledge-Driven Synthetic Data Generation","version":1},"reference_index":18,"source":"arxiv_source","source_observed_at":"2026-08-07T14:00:31.170015Z"},"links":{"cited_paper":"/paper/2406.16367","citing_paper":"/paper/2505.20416"},"observation_digest":"sha256:39c58dec59e1bb31101c01c8b8e5fcb01d73ad8a05703c5550b1a21a5ad6ab4e","observation_id":"50ed8937-6ff9-4060-bc55-45804d103389","resolution":{"observed_at":"2026-08-07T14:00:31.170015Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T14:00:35.726363Z","title":null,"venue":null,"work_id":"4db8df0b-3380-4721-8ed9-2bf9e156b5a2","year":2024},"citing_paper":{"arxiv_id":"2505.20416","last_updated":"2025-05-26T18:06:50Z","snapshot_observed_at":"2026-08-07T13:52:46.937698Z","submitted_at":"2025-05-26T18:06:50Z","title":"GraphGen: Enhancing Supervised Fine-Tuning for LLMs with Knowledge-Driven Synthetic Data Generation","version":1},"reference_index":19,"source":"arxiv_source","source_observed_at":"2026-08-07T14:00:31.239806Z"},"links":{"citing_paper":"/paper/2505.20416"},"observation_digest":"sha256:90bd69e00a7fbd5cb159ab7c40733f94931c6ed4becc62221476e6a592dfba95","observation_id":"a1a95673-f78e-48fd-a18f-11386af5f299","resolution":{"observed_at":"2026-08-07T14:00:35.787443Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T14:00:35.530840Z","title":null,"venue":null,"work_id":"f72c6c9b-9619-4082-b59b-3dbfed2f0065","year":2023},"citing_paper":{"arxiv_id":"2505.20416","last_updated":"2025-05-26T18:06:50Z","snapshot_observed_at":"2026-08-07T13:52:46.937698Z","submitted_at":"2025-05-26T18:06:50Z","title":"GraphGen: Enhancing Supervised Fine-Tuning for LLMs with Knowledge-Driven Synthetic Data Generation","version":1},"reference_index":20,"source":"arxiv_source","source_observed_at":"2026-08-07T14:00:31.324373Z"},"links":{"citing_paper":"/paper/2505.20416"},"observation_digest":"sha256:3a9bc78a6e57036666f166507580e1529f5b652242d5cf4ea9deb2caeddee76a","observation_id":"9a39f985-5418-4a9f-bb1a-0dd0e521f7bc","resolution":{"observed_at":"2026-08-07T14:00:35.612378Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2404.07503","last_updated":"2024-08-10T20:46:47Z","snapshot_observed_at":"2026-07-06T17:58:39.485508Z","submitted_at":"2024-04-11T06:34:17Z","title":"Best Practices and Lessons Learned on Synthetic Data","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2404.07503","snapshot_observed_at":"2026-08-07T14:00:31.405117Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2505.20416","last_updated":"2025-05-26T18:06:50Z","snapshot_observed_at":"2026-08-07T13:52:46.937698Z","submitted_at":"2025-05-26T18:06:50Z","title":"GraphGen: Enhancing Supervised Fine-Tuning for LLMs with Knowledge-Driven Synthetic Data Generation","version":1},"reference_index":21,"source":"arxiv_source","source_observed_at":"2026-08-07T14:00:31.405117Z"},"links":{"cited_paper":"/paper/2404.07503","citing_paper":"/paper/2505.20416"},"observation_digest":"sha256:d7620bb242f29074409cdd9ab347f7991d2a9b0c559f9a57be23b13dece29ec1","observation_id":"5f8b9035-31b8-482b-abc4-c5df9f956dd4","resolution":{"observed_at":"2026-08-07T14:00:31.405117Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2406.15126","last_updated":"2024-06-14T07:47:09Z","snapshot_observed_at":"2026-07-06T18:34:51.641949Z","submitted_at":"2024-06-14T07:47:09Z","title":"On LLMs-Driven Synthetic Data Generation, Curation, and Evaluation: A Survey","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2406.15126","snapshot_observed_at":"2026-08-07T14:00:31.486151Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2505.20416","last_updated":"2025-05-26T18:06:50Z","snapshot_observed_at":"2026-08-07T13:52:46.937698Z","submitted_at":"2025-05-26T18:06:50Z","title":"GraphGen: Enhancing Supervised Fine-Tuning for LLMs with Knowledge-Driven Synthetic Data Generation","version":1},"reference_index":22,"source":"arxiv_source","source_observed_at":"2026-08-07T14:00:31.486151Z"},"links":{"cited_paper":"/paper/2406.15126","citing_paper":"/paper/2505.20416"},"observation_digest":"sha256:2f6a5fda80ffa126269423e1bf77315a5783533e021248e235a4092aadfcb550","observation_id":"e988690b-12a9-44ef-9570-b074cb7ce78d","resolution":{"observed_at":"2026-08-07T14:00:31.486151Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2411.11266","last_updated":"2025-05-19T14:51:44Z","snapshot_observed_at":"2026-07-06T19:51:41.646058Z","submitted_at":"2024-11-18T03:45:34Z","title":"VersaTune: An Efficient Data Composition Framework for Training Multi-Capability LLMs","version":5},"cited_work":{"arxiv_id":"2411.11266","doi":null,"metadata_source":"pith","pith_arxiv_id":"2411.11266","snapshot_observed_at":"2026-08-07T14:00:34.048977Z","title":"VersaTune: An Efficient Data Composition Framework for Training Multi-Capability LLMs","venue":"cs.CL","work_id":"bb296ddf-e814-4b00-bc87-8c69f59866c6","year":2024},"citing_paper":{"arxiv_id":"2505.20416","last_updated":"2025-05-26T18:06:50Z","snapshot_observed_at":"2026-08-07T13:52:46.937698Z","submitted_at":"2025-05-26T18:06:50Z","title":"GraphGen: Enhancing Supervised Fine-Tuning for LLMs with Knowledge-Driven Synthetic Data Generation","version":1},"reference_index":23,"source":"arxiv_source","source_observed_at":"2026-08-07T14:00:31.489767Z"},"links":{"cited_paper":"/paper/2411.11266","citing_paper":"/paper/2505.20416"},"observation_digest":"sha256:48a40414f9403987faa7bfae2be7f5b82287d085770c51d160adcec730fc9780","observation_id":"8057caba-38b6-4025-b1b0-9f71da005ca1","resolution":{"observed_at":"2026-08-07T14:00:34.140189Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2401.16380","last_updated":"2024-01-29T18:19:08Z","snapshot_observed_at":"2026-07-06T17:21:58.983048Z","submitted_at":"2024-01-29T18:19:08Z","title":"Rephrasing the Web: A Recipe for Compute and Data-Efficient Language Modeling","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2401.16380","snapshot_observed_at":"2026-08-07T14:00:31.493097Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2505.20416","last_updated":"2025-05-26T18:06:50Z","snapshot_observed_at":"2026-08-07T13:52:46.937698Z","submitted_at":"2025-05-26T18:06:50Z","title":"GraphGen: Enhancing Supervised Fine-Tuning for LLMs with Knowledge-Driven Synthetic Data Generation","version":1},"reference_index":24,"source":"arxiv_source","source_observed_at":"2026-08-07T14:00:31.493097Z"},"links":{"cited_paper":"/paper/2401.16380","citing_paper":"/paper/2505.20416"},"observation_digest":"sha256:a012cd3f6d72389414a16bdb154de2290a8c0766650cd2ca225baadece2f94bc","observation_id":"af7264e8-e402-4814-abb7-82e5a8c83fc3","resolution":{"observed_at":"2026-08-07T14:00:31.493097Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T14:00:31.504032Z","title":null,"venue":null,"work_id":null,"year":2010},"citing_paper":{"arxiv_id":"2505.20416","last_updated":"2025-05-26T18:06:50Z","snapshot_observed_at":"2026-08-07T13:52:46.937698Z","submitted_at":"2025-05-26T18:06:50Z","title":"GraphGen: Enhancing Supervised Fine-Tuning for LLMs with Knowledge-Driven Synthetic Data Generation","version":1},"reference_index":25,"source":"arxiv_source","source_observed_at":"2026-08-07T14:00:31.504032Z"},"links":{"citing_paper":"/paper/2505.20416"},"observation_digest":"sha256:c0162211ce868ab4ddb2ac99b3194ddfdd96e7a2f1a3aa9191f77c989d48f968","observation_id":"fddf5425-92bf-4879-ba64-9875b5c4c91c","resolution":{"observed_at":"2026-08-07T14:00:31.504032Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2404.00213","last_updated":"2024-04-02T20:09:45Z","snapshot_observed_at":"2026-07-06T17:53:19.196215Z","submitted_at":"2024-03-30T01:56:07Z","title":"Injecting New Knowledge into Large Language Models via Supervised Fine-Tuning","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2404.00213","snapshot_observed_at":"2026-08-07T14:00:31.608166Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2505.20416","last_updated":"2025-05-26T18:06:50Z","snapshot_observed_at":"2026-08-07T13:52:46.937698Z","submitted_at":"2025-05-26T18:06:50Z","title":"GraphGen: Enhancing Supervised Fine-Tuning for LLMs with Knowledge-Driven Synthetic Data Generation","version":1},"reference_index":26,"source":"arxiv_source","source_observed_at":"2026-08-07T14:00:31.608166Z"},"links":{"cited_paper":"/paper/2404.00213","citing_paper":"/paper/2505.20416"},"observation_digest":"sha256:67e228fe02eec7dbbf36d107dc0511d53d816406da0e4ef709645a6b0a8b0685","observation_id":"69709be5-7cdb-4e3c-8d42-31ce3339071c","resolution":{"observed_at":"2026-08-07T14:00:31.608166Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2408.13296","last_updated":"2024-10-30T01:04:15Z","snapshot_observed_at":"2026-08-07T02:13:06.243888Z","submitted_at":"2024-08-23T14:48:02Z","title":"The Ultimate Guide to Fine-Tuning LLMs from Basics to Breakthroughs: An Exhaustive Review of Technologies, Research, Best Practices, Applied Research Challenges and Opportunities","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2408.13296","snapshot_observed_at":"2026-08-07T14:00:31.719031Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2505.20416","last_updated":"2025-05-26T18:06:50Z","snapshot_observed_at":"2026-08-07T13:52:46.937698Z","submitted_at":"2025-05-26T18:06:50Z","title":"GraphGen: Enhancing Supervised Fine-Tuning for LLMs with Knowledge-Driven Synthetic Data Generation","version":1},"reference_index":27,"source":"arxiv_source","source_observed_at":"2026-08-07T14:00:31.719031Z"},"links":{"cited_paper":"/paper/2408.13296","citing_paper":"/paper/2505.20416"},"observation_digest":"sha256:8d36ad2d54b9474881e965a28fbaceaddbbce33ddbcede069bf26c4561ebc131","observation_id":"ff3899ac-635a-4ea3-bb34-96360efe4ed9","resolution":{"observed_at":"2026-08-07T14:00:31.719031Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T14:00:35.309263Z","title":null,"venue":null,"work_id":"229e328e-84f1-40fa-9e0e-b289c1dff73a","year":2017},"citing_paper":{"arxiv_id":"2505.20416","last_updated":"2025-05-26T18:06:50Z","snapshot_observed_at":"2026-08-07T13:52:46.937698Z","submitted_at":"2025-05-26T18:06:50Z","title":"GraphGen: Enhancing Supervised Fine-Tuning for LLMs with Knowledge-Driven Synthetic Data Generation","version":1},"reference_index":28,"source":"arxiv_source","source_observed_at":"2026-08-07T14:00:31.788518Z"},"links":{"citing_paper":"/paper/2505.20416"},"observation_digest":"sha256:2e8f1dcd717cff39fa6375387f5d3739bb48983b49208f9ed54be20dc9ec5b9d","observation_id":"b76f68ed-9d71-4b64-bc9f-80cfe12e1aad","resolution":{"observed_at":"2026-08-07T14:00:35.444372Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T14:00:35.156300Z","title":null,"venue":null,"work_id":"53679ece-3037-4538-9600-a38a709523ca","year":2016},"citing_paper":{"arxiv_id":"2505.20416","last_updated":"2025-05-26T18:06:50Z","snapshot_observed_at":"2026-08-07T13:52:46.937698Z","submitted_at":"2025-05-26T18:06:50Z","title":"GraphGen: Enhancing Supervised Fine-Tuning for LLMs with Knowledge-Driven Synthetic Data Generation","version":1},"reference_index":29,"source":"arxiv_source","source_observed_at":"2026-08-07T14:00:31.906363Z"},"links":{"citing_paper":"/paper/2505.20416"},"observation_digest":"sha256:0181381eeb64e882ad23e92e8f35620a1b2cfd1e9380af6d905fd0a0dc93b512","observation_id":"9bacbbd9-88f6-45a6-82b1-0e3fa842bd0b","resolution":{"observed_at":"2026-08-07T14:00:35.213071Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T14:00:32.002622Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2505.20416","last_updated":"2025-05-26T18:06:50Z","snapshot_observed_at":"2026-08-07T13:52:46.937698Z","submitted_at":"2025-05-26T18:06:50Z","title":"GraphGen: Enhancing Supervised Fine-Tuning for LLMs with Knowledge-Driven Synthetic Data Generation","version":1},"reference_index":30,"source":"arxiv_source","source_observed_at":"2026-08-07T14:00:32.002622Z"},"links":{"citing_paper":"/paper/2505.20416"},"observation_digest":"sha256:5ead79d564c7ef5315723c5b220beb52976c28734d3322c568cea95b080eec85","observation_id":"f6af1e35-d7fd-435c-a3cf-c74b1af20c08","resolution":{"observed_at":"2026-08-07T14:00:32.002622Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T14:00:32.094121Z","title":null,"venue":null,"work_id":null,"year":2022},"citing_paper":{"arxiv_id":"2505.20416","last_updated":"2025-05-26T18:06:50Z","snapshot_observed_at":"2026-08-07T13:52:46.937698Z","submitted_at":"2025-05-26T18:06:50Z","title":"GraphGen: Enhancing Supervised Fine-Tuning for LLMs with Knowledge-Driven Synthetic Data Generation","version":1},"reference_index":31,"source":"arxiv_source","source_observed_at":"2026-08-07T14:00:32.094121Z"},"links":{"citing_paper":"/paper/2505.20416"},"observation_digest":"sha256:444cfed8f5c8f0c646005d2b8a3b4e2b2b1bc35deddc709448137123fa74f127","observation_id":"84166ce3-1c08-44f4-a324-b31d1afe893e","resolution":{"observed_at":"2026-08-07T14:00:32.094121Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T14:00:34.936372Z","title":null,"venue":null,"work_id":"0432fea6-39a1-4c28-a83f-38a7da1a4c65","year":2022},"citing_paper":{"arxiv_id":"2505.20416","last_updated":"2025-05-26T18:06:50Z","snapshot_observed_at":"2026-08-07T13:52:46.937698Z","submitted_at":"2025-05-26T18:06:50Z","title":"GraphGen: Enhancing Supervised Fine-Tuning for LLMs with Knowledge-Driven Synthetic Data Generation","version":1},"reference_index":32,"source":"arxiv_source","source_observed_at":"2026-08-07T14:00:32.183983Z"},"links":{"citing_paper":"/paper/2505.20416"},"observation_digest":"sha256:41f4c7e5ac0639d38e006e30ce298aa00df14a0c68f88319039d2895bd27784c","observation_id":"cb82ee31-057f-4611-9cca-968d857a2265","resolution":{"observed_at":"2026-08-07T14:00:35.071765Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2412.15115","last_updated":"2025-01-03T02:18:21Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2024-12-19T17:56:09Z","title":"Qwen2.5 Technical Report","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2412.15115","snapshot_observed_at":"2026-08-07T14:00:32.277975Z","title":null,"venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2505.20416","last_updated":"2025-05-26T18:06:50Z","snapshot_observed_at":"2026-08-07T13:52:46.937698Z","submitted_at":"2025-05-26T18:06:50Z","title":"GraphGen: Enhancing Supervised Fine-Tuning for LLMs with Knowledge-Driven Synthetic Data Generation","version":1},"reference_index":33,"source":"arxiv_source","source_observed_at":"2026-08-07T14:00:32.277975Z"},"links":{"cited_paper":"/paper/2412.15115","citing_paper":"/paper/2505.20416"},"observation_digest":"sha256:b237465bfa0cc8b26967f0281825abe87681371c8e1c9d0ca82eb90fe49c2576","observation_id":"7dfa25dd-518a-434d-a244-2d729dd91989","resolution":{"observed_at":"2026-08-07T14:00:32.277975Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"1809.09600","last_updated":"2018-09-25T17:28:20Z","snapshot_observed_at":"2026-07-31T11:19:06.421362Z","submitted_at":"2018-09-25T17:28:20Z","title":"HotpotQA: A Dataset for Diverse, Explainable Multi-hop Question Answering","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"1809.09600","snapshot_observed_at":"2026-08-07T14:00:32.393107Z","title":"Cohen, Ruslan Salakhutdinov, and Christopher D","venue":null,"work_id":null,"year":2018},"citing_paper":{"arxiv_id":"2505.20416","last_updated":"2025-05-26T18:06:50Z","snapshot_observed_at":"2026-08-07T13:52:46.937698Z","submitted_at":"2025-05-26T18:06:50Z","title":"GraphGen: Enhancing Supervised Fine-Tuning for LLMs with Knowledge-Driven Synthetic Data Generation","version":1},"reference_index":34,"source":"arxiv_source","source_observed_at":"2026-08-07T14:00:32.393107Z"},"links":{"cited_paper":"/paper/1809.09600","citing_paper":"/paper/2505.20416"},"observation_digest":"sha256:7f679f9438c69f052cb88663de5a216a0d1b9bcacd67d7382fd3edb03d04a70d","observation_id":"bf9b34c4-40cd-437b-9424-1955689571af","resolution":{"observed_at":"2026-08-07T14:00:32.393107Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2409.07431","last_updated":"2024-10-03T13:07:25Z","snapshot_observed_at":"2026-07-06T19:13:50.684379Z","submitted_at":"2024-09-11T17:21:59Z","title":"Synthetic continued pretraining","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2409.07431","snapshot_observed_at":"2026-08-07T14:00:32.484816Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2505.20416","last_updated":"2025-05-26T18:06:50Z","snapshot_observed_at":"2026-08-07T13:52:46.937698Z","submitted_at":"2025-05-26T18:06:50Z","title":"GraphGen: Enhancing Supervised Fine-Tuning for LLMs with Knowledge-Driven Synthetic Data Generation","version":1},"reference_index":35,"source":"arxiv_source","source_observed_at":"2026-08-07T14:00:32.484816Z"},"links":{"cited_paper":"/paper/2409.07431","citing_paper":"/paper/2505.20416"},"observation_digest":"sha256:2b52088f4d7756c83c5feba7235c666df056ffa0413c2d108cbde5b5e828b0e0","observation_id":"022b3d1b-0f37-4f47-b71a-2ca76e38d9a1","resolution":{"observed_at":"2026-08-07T14:00:32.484816Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2401.14367","last_updated":"2024-01-25T18:14:57Z","snapshot_observed_at":"2026-07-06T17:20:33.356092Z","submitted_at":"2024-01-25T18:14:57Z","title":"Genie: Achieving Human Parity in Content-Grounded Datasets Generation","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2401.14367","snapshot_observed_at":"2026-08-07T14:00:32.574415Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2505.20416","last_updated":"2025-05-26T18:06:50Z","snapshot_observed_at":"2026-08-07T13:52:46.937698Z","submitted_at":"2025-05-26T18:06:50Z","title":"GraphGen: Enhancing Supervised Fine-Tuning for LLMs with Knowledge-Driven Synthetic Data Generation","version":1},"reference_index":36,"source":"arxiv_source","source_observed_at":"2026-08-07T14:00:32.574415Z"},"links":{"cited_paper":"/paper/2401.14367","citing_paper":"/paper/2505.20416"},"observation_digest":"sha256:4ecab512a99c0f38011f51c1ffea069933e9af3bc6c8f675f89ab549287de39b","observation_id":"b1f5342f-4abe-4484-a55d-015b20cf901a","resolution":{"observed_at":"2026-08-07T14:00:32.574415Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2505.13220","last_updated":"2025-05-19T15:02:59Z","snapshot_observed_at":"2026-08-07T15:43:09.026441Z","submitted_at":"2025-05-19T15:02:59Z","title":"SeedBench: A Multi-task Benchmark for Evaluating Large Language Models in Seed Science","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2505.13220","snapshot_observed_at":"2026-08-07T14:00:32.650100Z","title":null,"venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2505.20416","last_updated":"2025-05-26T18:06:50Z","snapshot_observed_at":"2026-08-07T13:52:46.937698Z","submitted_at":"2025-05-26T18:06:50Z","title":"GraphGen: Enhancing Supervised Fine-Tuning for LLMs with Knowledge-Driven Synthetic Data Generation","version":1},"reference_index":37,"source":"arxiv_source","source_observed_at":"2026-08-07T14:00:32.650100Z"},"links":{"cited_paper":"/paper/2505.13220","citing_paper":"/paper/2505.20416"},"observation_digest":"sha256:9fe0d87714be027901b1aed346ac3327ed1e1baf09b7d52b4e31849c045eb243","observation_id":"236b59b8-3ae9-4f02-940c-250ec8b3a712","resolution":{"observed_at":"2026-08-07T14:00:32.650100Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2305.11952","last_updated":"2023-05-19T18:26:26Z","snapshot_observed_at":"2026-07-06T15:29:54.267559Z","submitted_at":"2023-05-19T18:26:26Z","title":"Self-QA: Unsupervised Knowledge Guided Language Model Alignment","version":1},"cited_work":{"arxiv_id":"2305.11952","doi":null,"metadata_source":"pith","pith_arxiv_id":"2305.11952","snapshot_observed_at":"2026-08-07T14:00:33.651259Z","title":"Self-QA: Unsupervised Knowledge Guided Language Model Alignment","venue":"cs.CL","work_id":"8e728853-820c-43d6-9348-4e7ad55220c2","year":2023},"citing_paper":{"arxiv_id":"2505.20416","last_updated":"2025-05-26T18:06:50Z","snapshot_observed_at":"2026-08-07T13:52:46.937698Z","submitted_at":"2025-05-26T18:06:50Z","title":"GraphGen: Enhancing Supervised Fine-Tuning for LLMs with Knowledge-Driven Synthetic Data Generation","version":1},"reference_index":39,"source":"arxiv_source","source_observed_at":"2026-08-07T14:00:32.810716Z"},"links":{"cited_paper":"/paper/2305.11952","citing_paper":"/paper/2505.20416"},"observation_digest":"sha256:018c15c399db1c4aa4ae0f4530e7c50e04776eef78489394b77b74fff163acba","observation_id":"68bb9f2a-28dc-4380-98f6-bf1af4989348","resolution":{"observed_at":"2026-08-07T14:00:33.801103Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T14:00:34.718841Z","title":null,"venue":null,"work_id":"18812071-8230-4134-83f3-de82b8a46d73","year":2024},"citing_paper":{"arxiv_id":"2505.20416","last_updated":"2025-05-26T18:06:50Z","snapshot_observed_at":"2026-08-07T13:52:46.937698Z","submitted_at":"2025-05-26T18:06:50Z","title":"GraphGen: Enhancing Supervised Fine-Tuning for LLMs with Knowledge-Driven Synthetic Data Generation","version":1},"reference_index":40,"source":"arxiv_source","source_observed_at":"2026-08-07T14:00:32.905991Z"},"links":{"citing_paper":"/paper/2505.20416"},"observation_digest":"sha256:0d5c093347efb7ede774aa726ffeccbb41b4fa01d1f60941a17dc3b5a4ff405c","observation_id":"a7252192-a611-4aea-a82c-f95842ff96ab","resolution":{"observed_at":"2026-08-07T14:00:34.822257Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2411.14405","last_updated":"2024-11-25T17:57:55Z","snapshot_observed_at":"2026-07-06T19:53:56.407834Z","submitted_at":"2024-11-21T18:37:33Z","title":"Marco-o1: Towards Open Reasoning Models for Open-Ended Solutions","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2411.14405","snapshot_observed_at":"2026-08-07T14:00:33.005432Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2505.20416","last_updated":"2025-05-26T18:06:50Z","snapshot_observed_at":"2026-08-07T13:52:46.937698Z","submitted_at":"2025-05-26T18:06:50Z","title":"GraphGen: Enhancing Supervised Fine-Tuning for LLMs with Knowledge-Driven Synthetic Data Generation","version":1},"reference_index":41,"source":"arxiv_source","source_observed_at":"2026-08-07T14:00:33.005432Z"},"links":{"cited_paper":"/paper/2411.14405","citing_paper":"/paper/2505.20416"},"observation_digest":"sha256:826da104f6a12d43024d83586ecbe09e5b5dddf66c0eb6e8caf84cd96e88667a","observation_id":"72fa98e2-b9cb-43a5-8a90-30b7875a2b80","resolution":{"observed_at":"2026-08-07T14:00:33.005432Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T14:00:34.521747Z","title":null,"venue":null,"work_id":"3773a2f2-dbf4-466e-b3b5-ad2f6e6c1a58","year":2022},"citing_paper":{"arxiv_id":"2505.20416","last_updated":"2025-05-26T18:06:50Z","snapshot_observed_at":"2026-08-07T13:52:46.937698Z","submitted_at":"2025-05-26T18:06:50Z","title":"GraphGen: Enhancing Supervised Fine-Tuning for LLMs with Knowledge-Driven Synthetic Data Generation","version":1},"reference_index":42,"source":"arxiv_source","source_observed_at":"2026-08-07T14:00:33.103094Z"},"links":{"citing_paper":"/paper/2505.20416"},"observation_digest":"sha256:996ab8244721485107031b293251f66c237c76cf3ad2de841ceb5461b88c0da2","observation_id":"69f78243-dde9-4186-a8a4-3affc971f670","resolution":{"observed_at":"2026-08-07T14:00:34.598822Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T14:00:33.241436Z","title":"online\" 'onlinestring :=","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2505.20416","last_updated":"2025-05-26T18:06:50Z","snapshot_observed_at":"2026-08-07T13:52:46.937698Z","submitted_at":"2025-05-26T18:06:50Z","title":"GraphGen: Enhancing Supervised Fine-Tuning for LLMs with Knowledge-Driven Synthetic Data Generation","version":1},"reference_index":43,"source":"arxiv_source","source_observed_at":"2026-08-07T14:00:33.241436Z"},"links":{"citing_paper":"/paper/2505.20416"},"observation_digest":"sha256:56c458793e0737279ae5a5bdd168306bf7e22aa86a363203f06f839963b006ce","observation_id":"7b4ab52d-ed8e-4ec1-83cb-30967a0334a3","resolution":{"observed_at":"2026-08-07T14:00:33.241436Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T14:00:33.402147Z","title":"write newline","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2505.20416","last_updated":"2025-05-26T18:06:50Z","snapshot_observed_at":"2026-08-07T13:52:46.937698Z","submitted_at":"2025-05-26T18:06:50Z","title":"GraphGen: Enhancing Supervised Fine-Tuning for LLMs with Knowledge-Driven Synthetic Data Generation","version":1},"reference_index":44,"source":"arxiv_source","source_observed_at":"2026-08-07T14:00:33.402147Z"},"links":{"citing_paper":"/paper/2505.20416"},"observation_digest":"sha256:05c29455018ec8d9d5792d829a26d72ec2fd59556ec3ace05c8e5e15b12d06eb","observation_id":"1d6430d5-9f3f-444a-a2a8-a6536ad0d4f7","resolution":{"observed_at":"2026-08-07T14:00:33.402147Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}}],"paper":{"arxiv_id":"2505.20416","last_updated":"2025-05-26T18:06:50Z","latest_version":1,"primary_category":"cs.CL","snapshot_observed_at":"2026-08-07T13:52:46.937698Z","submitted_at":"2025-05-26T18:06:50Z","title":"GraphGen: Enhancing Supervised Fine-Tuning for LLMs with Knowledge-Driven Synthetic Data Generation"},"reference_resolution":{"displayed":43,"state_counts":{"malformed_identifier":0,"metadata_mismatch":0,"parse_uncertain":0,"unresolved":40,"verified_exact":3,"verified_fuzzy":0},"total_outbound_references":43},"refusal":"A citation records a reference. It does not transfer a finding from one paper to another.","schema":"pith.paper-citation-record.v1","standing_sources":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"thesis":"As of 7 August 2026, this Paper Citation Record lists 43 of 43 outbound references and 7 inbound Pith citation observations for arXiv:2505.20416."}