{"as_of":"2026-08-08T15:47:00Z","caps":{"database_statements":6,"inbound":100,"outbound":100},"context_digest":"sha256:cecc8f4bd56a1579ef71aec7c05fd3c80d32d4a29f73a835219543158bd9e035","coverage":[{"denominator":38,"lane":"reference_resolution","note":"Typed states for the displayed outbound observations.","records_observed":38,"source":"paper_references, paper_reference_links","source_observed_at":"2026-08-07T11:51:10.305405Z","state":"measured"},{"denominator":38,"lane":"standing_notices","note":"One-hop event checks from named stored sources.","records_observed":38,"source":"scholarly_work_events, retraction_status_cache","source_observed_at":"2026-08-08T06:32:00.761636+00:00","state":"measured"},{"denominator":0,"lane":"inbound_itemization","note":"Pith citing papers itemized under the disclosed page cap.","records_observed":0,"source":"paper_references, paper_reference_links","source_observed_at":null,"state":"measured"},{"denominator":1,"lane":"external_citation_measurements","note":"A source-named dated measurement, never combined with another source.","records_observed":0,"source":"cited_works","source_observed_at":null,"state":"measured"}],"external_citation_measurements":[],"inbound":[],"links":{"evidence":"/evidence","html":"/paper/2506.01305/citation-record","integrity":"/paper/2506.01305/integrity","json":"/paper/2506.01305/citation-record.json","paper":"/paper/2506.01305"},"outbound":[{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T11:51:08.461306Z","title":"Hewett, Mojan Javaheripi, Piero Kauffmann, James R","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2506.01305","last_updated":"2025-06-13T12:40:58Z","snapshot_observed_at":"2026-08-07T11:42:31.914563Z","submitted_at":"2025-06-02T04:32:15Z","title":"VM14K: First Vietnamese Medical Benchmark","version":2},"reference_index":1,"source":"pdf_text","source_observed_at":"2026-08-07T11:51:08.461306Z"},"links":{"citing_paper":"/paper/2506.01305"},"observation_digest":"sha256:3c99a2782d359f4f204e57d5eee429cabd25cbf98968ed9b3f25d4c17503df29","observation_id":"380f8195-8561-46c2-bd0c-69cff7223beb","resolution":{"observed_at":"2026-08-07T11:51:08.461306Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T11:51:13.835999Z","title":"The llama 4 herd: The beginning of a new era of natively multimodal ai innovation, 2025","venue":null,"work_id":"6f7c3576-6519-48a5-9db5-30797fb5d7db","year":2025},"citing_paper":{"arxiv_id":"2506.01305","last_updated":"2025-06-13T12:40:58Z","snapshot_observed_at":"2026-08-07T11:42:31.914563Z","submitted_at":"2025-06-02T04:32:15Z","title":"VM14K: First Vietnamese Medical Benchmark","version":2},"reference_index":2,"source":"pdf_text","source_observed_at":"2026-08-07T11:51:08.494052Z"},"links":{"citing_paper":"/paper/2506.01305"},"observation_digest":"sha256:0e2222c3ed931c9bfd7c45cf10707dfe5e880af346b15939f1cf891b9b2a12fe","observation_id":"615eb08d-1d5d-4590-b656-1b227e2cd828","resolution":{"observed_at":"2026-08-07T11:51:13.979615Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T11:51:13.640191Z","title":null,"venue":null,"work_id":"6675eb7e-63bd-49ab-a579-458c780a7646","year":2025},"citing_paper":{"arxiv_id":"2506.01305","last_updated":"2025-06-13T12:40:58Z","snapshot_observed_at":"2026-08-07T11:42:31.914563Z","submitted_at":"2025-06-02T04:32:15Z","title":"VM14K: First Vietnamese Medical Benchmark","version":2},"reference_index":3,"source":"pdf_text","source_observed_at":"2026-08-07T11:51:08.555771Z"},"links":{"citing_paper":"/paper/2506.01305"},"observation_digest":"sha256:adab101c294c59e7ae4ab6f2aa049c5051109b653a20433360c237900c31472c","observation_id":"0cfd8d42-2824-4530-9cc6-35fb5a4f2a8e","resolution":{"observed_at":"2026-08-07T11:51:13.674451Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T11:51:08.624771Z","title":"Claude 3.5 sonnet, 2024","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2506.01305","last_updated":"2025-06-13T12:40:58Z","snapshot_observed_at":"2026-08-07T11:42:31.914563Z","submitted_at":"2025-06-02T04:32:15Z","title":"VM14K: First Vietnamese Medical Benchmark","version":2},"reference_index":4,"source":"pdf_text","source_observed_at":"2026-08-07T11:51:08.624771Z"},"links":{"citing_paper":"/paper/2506.01305"},"observation_digest":"sha256:add68052c64899daf7c821ba911aef63c828167cec68630bf52a6fa03c5e0735","observation_id":"d8a33649-207b-4168-bc49-ad5d515870f6","resolution":{"observed_at":"2026-08-07T11:51:08.624771Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T11:51:13.507464Z","title":"Overview of the medical question answering task at trec 2017 liveqa","venue":null,"work_id":"f0cb7bbd-2a4c-4cbf-a641-eadee783872e","year":2017},"citing_paper":{"arxiv_id":"2506.01305","last_updated":"2025-06-13T12:40:58Z","snapshot_observed_at":"2026-08-07T11:42:31.914563Z","submitted_at":"2025-06-02T04:32:15Z","title":"VM14K: First Vietnamese Medical Benchmark","version":2},"reference_index":5,"source":"pdf_text","source_observed_at":"2026-08-07T11:51:08.687707Z"},"links":{"citing_paper":"/paper/2506.01305"},"observation_digest":"sha256:d2834097bb7891cd74ab034319b5dc55dbb2b2708fc1d88e1f2859c9208809eb","observation_id":"0321aab0-547b-46be-aae4-90240b24c85f","resolution":{"observed_at":"2026-08-07T11:51:13.558649Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T11:51:08.742786Z","title":"Huatuogpt-o1, towards medical complex reasoning with llms, 2024","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2506.01305","last_updated":"2025-06-13T12:40:58Z","snapshot_observed_at":"2026-08-07T11:42:31.914563Z","submitted_at":"2025-06-02T04:32:15Z","title":"VM14K: First Vietnamese Medical Benchmark","version":2},"reference_index":6,"source":"pdf_text","source_observed_at":"2026-08-07T11:51:08.742786Z"},"links":{"citing_paper":"/paper/2506.01305"},"observation_digest":"sha256:3d7911bdc015fd9fa06c408c7efe966c410acfc9756dc584083a6597b3c9caf9","observation_id":"d055fc9a-5816-4095-b2d1-7ef5236213b3","resolution":{"observed_at":"2026-08-07T11:51:08.742786Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T11:51:13.366039Z","title":"Meditron- 70b: Scaling medical pretraining for large language models, 2023","venue":null,"work_id":"ce4e4955-74a6-418c-a485-51e14ab91400","year":2023},"citing_paper":{"arxiv_id":"2506.01305","last_updated":"2025-06-13T12:40:58Z","snapshot_observed_at":"2026-08-07T11:42:31.914563Z","submitted_at":"2025-06-02T04:32:15Z","title":"VM14K: First Vietnamese Medical Benchmark","version":2},"reference_index":7,"source":"pdf_text","source_observed_at":"2026-08-07T11:51:08.801348Z"},"links":{"citing_paper":"/paper/2506.01305"},"observation_digest":"sha256:c5874b6e9bdf85a53ef84e21fe1510a673d6edbdd6b8747199208de4a350ab19","observation_id":"71da623e-d700-41af-a92f-11ec8b62ebab","resolution":{"observed_at":"2026-08-07T11:51:13.419932Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T11:51:13.252074Z","title":"The llama 3 herd of models, 07 2024","venue":null,"work_id":"135ded32-3f2f-4236-b4f4-584572a154b5","year":2024},"citing_paper":{"arxiv_id":"2506.01305","last_updated":"2025-06-13T12:40:58Z","snapshot_observed_at":"2026-08-07T11:42:31.914563Z","submitted_at":"2025-06-02T04:32:15Z","title":"VM14K: First Vietnamese Medical Benchmark","version":2},"reference_index":8,"source":"pdf_text","source_observed_at":"2026-08-07T11:51:08.844928Z"},"links":{"citing_paper":"/paper/2506.01305"},"observation_digest":"sha256:ff8fa47b286f58b1d0dcea9677798be71f1e1af1484579db51866eee336e2a21","observation_id":"77b8d524-7221-40b3-b664-380c3bc9c014","resolution":{"observed_at":"2026-08-07T11:51:13.292735Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T11:51:13.183315Z","title":"Introducing gemini 2.0: our new ai model for the agentic era, 2024","venue":null,"work_id":"ad2de585-ddac-4a81-88ad-c19274f6e86d","year":2024},"citing_paper":{"arxiv_id":"2506.01305","last_updated":"2025-06-13T12:40:58Z","snapshot_observed_at":"2026-08-07T11:42:31.914563Z","submitted_at":"2025-06-02T04:32:15Z","title":"VM14K: First Vietnamese Medical Benchmark","version":2},"reference_index":9,"source":"pdf_text","source_observed_at":"2026-08-07T11:51:08.901297Z"},"links":{"citing_paper":"/paper/2506.01305"},"observation_digest":"sha256:64f5dec6f92a9ef6091618232e71edebc1fb5e17e20437b23f72c2d4ede7ab5f","observation_id":"d3c79ff2-77fd-4f96-8769-39fce978dd8d","resolution":{"observed_at":"2026-08-07T11:51:13.211115Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T11:51:13.094577Z","title":"Impact of translation on biomedical information extraction from real-life clinical notes, 2023","venue":null,"work_id":"d68fa39d-9fad-477b-9435-a726e8c1e36a","year":2023},"citing_paper":{"arxiv_id":"2506.01305","last_updated":"2025-06-13T12:40:58Z","snapshot_observed_at":"2026-08-07T11:42:31.914563Z","submitted_at":"2025-06-02T04:32:15Z","title":"VM14K: First Vietnamese Medical Benchmark","version":2},"reference_index":10,"source":"pdf_text","source_observed_at":"2026-08-07T11:51:08.952261Z"},"links":{"citing_paper":"/paper/2506.01305"},"observation_digest":"sha256:39175c51b6440ae96b54dac3b020d9b499f70c0a60991893ab738803faf32681","observation_id":"27faf19a-380a-43f2-9483-0e009e6c368a","resolution":{"observed_at":"2026-08-07T11:51:13.169582Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T11:51:08.996002Z","title":"Aligning ai with shared human values.Proceedings of the International Conference on Learning Representations (ICLR), 2021","venue":null,"work_id":null,"year":2021},"citing_paper":{"arxiv_id":"2506.01305","last_updated":"2025-06-13T12:40:58Z","snapshot_observed_at":"2026-08-07T11:42:31.914563Z","submitted_at":"2025-06-02T04:32:15Z","title":"VM14K: First Vietnamese Medical Benchmark","version":2},"reference_index":11,"source":"pdf_text","source_observed_at":"2026-08-07T11:51:08.996002Z"},"links":{"citing_paper":"/paper/2506.01305"},"observation_digest":"sha256:d10b1d4eae6233fc401d4646b42e10e9e926afcc888031fb5c5099e4fb137e9c","observation_id":"e3e3dd3a-e5a1-4aad-b577-bae1985755c8","resolution":{"observed_at":"2026-08-07T11:51:08.996002Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T11:51:09.067486Z","title":"Measuring massive multitask language understanding.Proceedings of the International Conference on Learning Representations (ICLR), 2021","venue":null,"work_id":null,"year":2021},"citing_paper":{"arxiv_id":"2506.01305","last_updated":"2025-06-13T12:40:58Z","snapshot_observed_at":"2026-08-07T11:42:31.914563Z","submitted_at":"2025-06-02T04:32:15Z","title":"VM14K: First Vietnamese Medical Benchmark","version":2},"reference_index":12,"source":"pdf_text","source_observed_at":"2026-08-07T11:51:09.067486Z"},"links":{"citing_paper":"/paper/2506.01305"},"observation_digest":"sha256:192188ac9c03cec2091f483d85e1ff72f424e89ef5300c617533b4364a192813","observation_id":"02162ad0-6e2b-4575-aee9-328deaa5f978","resolution":{"observed_at":"2026-08-07T11:51:09.067486Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2009.13081","last_updated":"2020-09-28T05:07:51Z","snapshot_observed_at":"2026-08-06T03:17:52.286711Z","submitted_at":"2020-09-28T05:07:51Z","title":"What Disease does this Patient Have? A Large-scale Open Domain Question Answering Dataset from Medical Exams","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2009.13081","snapshot_observed_at":"2026-08-07T11:51:09.122543Z","title":"What disease does this patient have? a large-scale open domain question answering dataset from medical exams.arXiv preprint arXiv:2009.13081, 2020","venue":null,"work_id":null,"year":2009},"citing_paper":{"arxiv_id":"2506.01305","last_updated":"2025-06-13T12:40:58Z","snapshot_observed_at":"2026-08-07T11:42:31.914563Z","submitted_at":"2025-06-02T04:32:15Z","title":"VM14K: First Vietnamese Medical Benchmark","version":2},"reference_index":13,"source":"pdf_text","source_observed_at":"2026-08-07T11:51:09.122543Z"},"links":{"cited_paper":"/paper/2009.13081","citing_paper":"/paper/2506.01305"},"observation_digest":"sha256:64198aa5e852e978814bc3f0ed25d8316a9cce8aba6a156d8896c4594f432f29","observation_id":"e596fe38-3d8e-4b14-8c9c-415db111cb44","resolution":{"observed_at":"2026-08-07T11:51:09.122543Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T11:51:09.162560Z","title":"Pubmedqa: A dataset for biomedical research question answering","venue":null,"work_id":null,"year":2019},"citing_paper":{"arxiv_id":"2506.01305","last_updated":"2025-06-13T12:40:58Z","snapshot_observed_at":"2026-08-07T11:42:31.914563Z","submitted_at":"2025-06-02T04:32:15Z","title":"VM14K: First Vietnamese Medical Benchmark","version":2},"reference_index":14,"source":"pdf_text","source_observed_at":"2026-08-07T11:51:09.162560Z"},"links":{"citing_paper":"/paper/2506.01305"},"observation_digest":"sha256:e646bcd827a1365e5f7d122a83be44f5371dc0d91de42da70ef4401a3e0ba49f","observation_id":"826edc70-9a6c-4f60-9d43-0eda5b7ce85f","resolution":{"observed_at":"2026-08-07T11:51:09.162560Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T11:51:12.972450Z","title":"Better to ask in english: Cross-lingual evaluation of large language models for healthcare queries, 2023","venue":null,"work_id":"bb26fdde-a10f-4f5e-84f6-2e80c2a0caf4","year":2023},"citing_paper":{"arxiv_id":"2506.01305","last_updated":"2025-06-13T12:40:58Z","snapshot_observed_at":"2026-08-07T11:42:31.914563Z","submitted_at":"2025-06-02T04:32:15Z","title":"VM14K: First Vietnamese Medical Benchmark","version":2},"reference_index":15,"source":"pdf_text","source_observed_at":"2026-08-07T11:51:09.196470Z"},"links":{"citing_paper":"/paper/2506.01305"},"observation_digest":"sha256:ebfc0bb83c18fdfba9b7cd307b80607f9885c2096355738420497490efaf2903","observation_id":"a55a677e-6d24-4874-adc6-ca2e0c518050","resolution":{"observed_at":"2026-08-07T11:51:12.998743Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T11:51:09.241336Z","title":"Evaluating gpt-4 and chatgpt on japanese medical licensing examinations, 2023","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2506.01305","last_updated":"2025-06-13T12:40:58Z","snapshot_observed_at":"2026-08-07T11:42:31.914563Z","submitted_at":"2025-06-02T04:32:15Z","title":"VM14K: First Vietnamese Medical Benchmark","version":2},"reference_index":16,"source":"pdf_text","source_observed_at":"2026-08-07T11:51:09.241336Z"},"links":{"citing_paper":"/paper/2506.01305"},"observation_digest":"sha256:414732ee8e4c4203a52482651800201a0cf1e2a0260b2cc5b7f9ccf85f995c06","observation_id":"700c5839-3223-432d-9a86-f8a763c89419","resolution":{"observed_at":"2026-08-07T11:51:09.241336Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T11:51:12.842248Z","title":"Kormedmcqa: Multi-choice question answering benchmark for korean healthcare professional licensing examinations, 2024","venue":null,"work_id":"f7f7b4bc-c493-4cde-beb7-efbbc0a7de75","year":2024},"citing_paper":{"arxiv_id":"2506.01305","last_updated":"2025-06-13T12:40:58Z","snapshot_observed_at":"2026-08-07T11:42:31.914563Z","submitted_at":"2025-06-02T04:32:15Z","title":"VM14K: First Vietnamese Medical Benchmark","version":2},"reference_index":17,"source":"pdf_text","source_observed_at":"2026-08-07T11:51:09.306762Z"},"links":{"citing_paper":"/paper/2506.01305"},"observation_digest":"sha256:e275540f433b96d4e5f11b8eff581c8a5cc3f7ccbb8162a6914e524a043f7f7c","observation_id":"9a97e72f-fd0d-4cd6-b6fc-0010677ebb33","resolution":{"observed_at":"2026-08-07T11:51:12.912889Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T11:51:12.596797Z","title":null,"venue":null,"work_id":"98596c8e-ba6c-4c28-aa41-2eb4d9378bb1","year":2020},"citing_paper":{"arxiv_id":"2506.01305","last_updated":"2025-06-13T12:40:58Z","snapshot_observed_at":"2026-08-07T11:42:31.914563Z","submitted_at":"2025-06-02T04:32:15Z","title":"VM14K: First Vietnamese Medical Benchmark","version":2},"reference_index":18,"source":"pdf_text","source_observed_at":"2026-08-07T11:51:09.359074Z"},"links":{"citing_paper":"/paper/2506.01305"},"observation_digest":"sha256:0277fac6b9bc85b13017b4ebd53345ab9c888fdd0042b23406bb98449a70cfe9","observation_id":"1aec9a07-0ba1-410e-8c4f-add5ce6c1327","resolution":{"observed_at":"2026-08-07T11:51:12.718056Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T11:51:12.359588Z","title":"Benchmarking large language models on cmexam – a comprehensive chinese medical exam dataset, 2023","venue":null,"work_id":"d89573a7-52c5-43a6-a8d5-eee373a4a3b0","year":2023},"citing_paper":{"arxiv_id":"2506.01305","last_updated":"2025-06-13T12:40:58Z","snapshot_observed_at":"2026-08-07T11:42:31.914563Z","submitted_at":"2025-06-02T04:32:15Z","title":"VM14K: First Vietnamese Medical Benchmark","version":2},"reference_index":19,"source":"pdf_text","source_observed_at":"2026-08-07T11:51:09.418313Z"},"links":{"citing_paper":"/paper/2506.01305"},"observation_digest":"sha256:75363f5d3fa885b020bc07ec710a509ef8fc5c0dac18dd5f78229f71a9cdcbdd","observation_id":"d7e4af28-3ea2-4546-982c-36bd290d7f41","resolution":{"observed_at":"2026-08-07T11:51:12.476366Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T11:51:09.478712Z","title":"A survey on medical large language models: Technology, application, trustworthiness, and future directions, 2024","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2506.01305","last_updated":"2025-06-13T12:40:58Z","snapshot_observed_at":"2026-08-07T11:42:31.914563Z","submitted_at":"2025-06-02T04:32:15Z","title":"VM14K: First Vietnamese Medical Benchmark","version":2},"reference_index":20,"source":"pdf_text","source_observed_at":"2026-08-07T11:51:09.478712Z"},"links":{"citing_paper":"/paper/2506.01305"},"observation_digest":"sha256:ec5289c5dbf39418e7be1adeafd0581475cecc54158dc2fa03114addb5a00410","observation_id":"8d09ae20-f7b3-4d53-a642-3e0edde1bd34","resolution":{"observed_at":"2026-08-07T11:51:09.478712Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T11:51:09.537690Z","title":"Hello GPT-4o, 2024","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2506.01305","last_updated":"2025-06-13T12:40:58Z","snapshot_observed_at":"2026-08-07T11:42:31.914563Z","submitted_at":"2025-06-02T04:32:15Z","title":"VM14K: First Vietnamese Medical Benchmark","version":2},"reference_index":21,"source":"pdf_text","source_observed_at":"2026-08-07T11:51:09.537690Z"},"links":{"citing_paper":"/paper/2506.01305"},"observation_digest":"sha256:6d76a1108e6f9dfde43bd4ba19cd03651d45fe706c8af68d617806b8995bfaab","observation_id":"80bfa282-7ffe-4b8a-97da-daf8c957ca79","resolution":{"observed_at":"2026-08-07T11:51:09.537690Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T11:51:12.121731Z","title":"Openai o3-mini pushing the frontier of cost-effective reasoning, 2025","venue":null,"work_id":"5305c362-a541-4fde-bae5-cb9d41f9a431","year":2025},"citing_paper":{"arxiv_id":"2506.01305","last_updated":"2025-06-13T12:40:58Z","snapshot_observed_at":"2026-08-07T11:42:31.914563Z","submitted_at":"2025-06-02T04:32:15Z","title":"VM14K: First Vietnamese Medical Benchmark","version":2},"reference_index":22,"source":"pdf_text","source_observed_at":"2026-08-07T11:51:09.594266Z"},"links":{"citing_paper":"/paper/2506.01305"},"observation_digest":"sha256:d8bf93f7c6b75590bb6e4b4c6239a7b107d74fa143ed970ef76998dbe511a764","observation_id":"e81a3841-616a-4752-a7f0-af82fa899603","resolution":{"observed_at":"2026-08-07T11:51:12.252428Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T11:51:09.640596Z","title":"Medmcqa: A large-scale multi-subject multi-choice dataset for medical domain question answering","venue":null,"work_id":null,"year":2022},"citing_paper":{"arxiv_id":"2506.01305","last_updated":"2025-06-13T12:40:58Z","snapshot_observed_at":"2026-08-07T11:42:31.914563Z","submitted_at":"2025-06-02T04:32:15Z","title":"VM14K: First Vietnamese Medical Benchmark","version":2},"reference_index":23,"source":"pdf_text","source_observed_at":"2026-08-07T11:51:09.640596Z"},"links":{"citing_paper":"/paper/2506.01305"},"observation_digest":"sha256:26c7de451be32f7656d6d4d1b2eca324c6607ad1d3baa92080a21e9af7e67e07","observation_id":"ad301650-4847-4ebb-a4f0-c0e502ab813f","resolution":{"observed_at":"2026-08-07T11:51:09.640596Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T11:51:11.901887Z","title":"Assessing the research landscape and clinical utility of large language models: a scoping review.BMC Medical Informatics and Decision Making, 24(72), 2024","venue":null,"work_id":"8d928fc2-654b-4726-9471-6a8495c944ef","year":2024},"citing_paper":{"arxiv_id":"2506.01305","last_updated":"2025-06-13T12:40:58Z","snapshot_observed_at":"2026-08-07T11:42:31.914563Z","submitted_at":"2025-06-02T04:32:15Z","title":"VM14K: First Vietnamese Medical Benchmark","version":2},"reference_index":24,"source":"pdf_text","source_observed_at":"2026-08-07T11:51:09.706822Z"},"links":{"citing_paper":"/paper/2506.01305"},"observation_digest":"sha256:e5d30905eb3553bb79be2905f2a41050b34bf1b5c13b5331195ebb2abeaa3ef8","observation_id":"81bdf47d-c6a9-4894-b3e3-1ee049228e4c","resolution":{"observed_at":"2026-08-07T11:51:11.975972Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T11:51:11.824086Z","title":"Opportunities and challenges for large language models in primary health care.J Prim Care Community Health, 16:21501319241312571, Jan-Dec 2025","venue":null,"work_id":"e687416a-d649-4c5b-9331-4ebd179f8c51","year":2025},"citing_paper":{"arxiv_id":"2506.01305","last_updated":"2025-06-13T12:40:58Z","snapshot_observed_at":"2026-08-07T11:42:31.914563Z","submitted_at":"2025-06-02T04:32:15Z","title":"VM14K: First Vietnamese Medical Benchmark","version":2},"reference_index":25,"source":"pdf_text","source_observed_at":"2026-08-07T11:51:09.750822Z"},"links":{"citing_paper":"/paper/2506.01305"},"observation_digest":"sha256:98f0a4aaa4c0526eb99d9328c249ae977a1ae2ee1e4509c418bc985267bb8330","observation_id":"15a0347d-043e-47a1-ac71-d49af9886ecb","resolution":{"observed_at":"2026-08-07T11:51:11.851283Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T11:51:09.816883Z","title":"Towards building multilingual language model for medicine, 2024","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2506.01305","last_updated":"2025-06-13T12:40:58Z","snapshot_observed_at":"2026-08-07T11:42:31.914563Z","submitted_at":"2025-06-02T04:32:15Z","title":"VM14K: First Vietnamese Medical Benchmark","version":2},"reference_index":26,"source":"pdf_text","source_observed_at":"2026-08-07T11:51:09.816883Z"},"links":{"citing_paper":"/paper/2506.01305"},"observation_digest":"sha256:dc31ed2ab6fe29ac15eca92c94d62139b652771eae63f0b930bf1a06449a1bcb","observation_id":"f8622691-00c8-4666-83fb-e970e8846b26","resolution":{"observed_at":"2026-08-07T11:51:09.816883Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T11:51:11.741612Z","title":"Holistic evaluation of large language models for medical applications","venue":null,"work_id":"1d1768c2-bc0a-4e98-b17b-efdb09cf62f8","year":2024},"citing_paper":{"arxiv_id":"2506.01305","last_updated":"2025-06-13T12:40:58Z","snapshot_observed_at":"2026-08-07T11:42:31.914563Z","submitted_at":"2025-06-02T04:32:15Z","title":"VM14K: First Vietnamese Medical Benchmark","version":2},"reference_index":27,"source":"pdf_text","source_observed_at":"2026-08-07T11:51:09.860112Z"},"links":{"citing_paper":"/paper/2506.01305"},"observation_digest":"sha256:894e3aed88b274b49f5e15c717472dffcc7cc78049586a4c10aeeee518d039a4","observation_id":"51539731-2efe-470e-b605-2a2b5f96acb2","resolution":{"observed_at":"2026-08-07T11:51:11.780489Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T11:51:09.899522Z","title":null,"venue":null,"work_id":null,"year":2022},"citing_paper":{"arxiv_id":"2506.01305","last_updated":"2025-06-13T12:40:58Z","snapshot_observed_at":"2026-08-07T11:42:31.914563Z","submitted_at":"2025-06-02T04:32:15Z","title":"VM14K: First Vietnamese Medical Benchmark","version":2},"reference_index":28,"source":"pdf_text","source_observed_at":"2026-08-07T11:51:09.899522Z"},"links":{"citing_paper":"/paper/2506.01305"},"observation_digest":"sha256:821193e4e1b6a8eedda3d7a8a3083b698e05f93c2569f28c3c33eefe724b7d23","observation_id":"d7e3c252-2aa6-496c-ac8c-d22c134127da","resolution":{"observed_at":"2026-08-07T11:51:09.899522Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T11:51:11.589725Z","title":"Gemma 3 technical report, 03 2025","venue":null,"work_id":"2c12fb3a-f7fb-45da-a4f8-84ab4bc9e0de","year":2025},"citing_paper":{"arxiv_id":"2506.01305","last_updated":"2025-06-13T12:40:58Z","snapshot_observed_at":"2026-08-07T11:42:31.914563Z","submitted_at":"2025-06-02T04:32:15Z","title":"VM14K: First Vietnamese Medical Benchmark","version":2},"reference_index":29,"source":"pdf_text","source_observed_at":"2026-08-07T11:51:09.938457Z"},"links":{"citing_paper":"/paper/2506.01305"},"observation_digest":"sha256:be2b355e22f3bfc6637e6e4685bd77257c7c1ee7a22848732e38d279bb63446f","observation_id":"b140afa9-976b-4cb3-8d6b-5791dba4e6c3","resolution":{"observed_at":"2026-08-07T11:51:11.665446Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T11:51:11.445897Z","title":"ChiMed: A Chinese medical corpus for question answering","venue":null,"work_id":"644a506e-a341-4a99-8b6d-57b4b345c0c3","year":2019},"citing_paper":{"arxiv_id":"2506.01305","last_updated":"2025-06-13T12:40:58Z","snapshot_observed_at":"2026-08-07T11:42:31.914563Z","submitted_at":"2025-06-02T04:32:15Z","title":"VM14K: First Vietnamese Medical Benchmark","version":2},"reference_index":30,"source":"pdf_text","source_observed_at":"2026-08-07T11:51:09.973642Z"},"links":{"citing_paper":"/paper/2506.01305"},"observation_digest":"sha256:7a841ab64f33ae7bd7236447b7e395fd3f93c1599ba3dab044d57c4dea059ce6","observation_id":"1913a8a7-2141-4ad8-8320-a5082ee5cdb5","resolution":{"observed_at":"2026-08-07T11:51:11.502164Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T11:51:11.278060Z","title":"Medqa benchmark, 2025","venue":null,"work_id":"6cb4e2ed-0905-47ec-ad67-68bc5f2b3e49","year":2025},"citing_paper":{"arxiv_id":"2506.01305","last_updated":"2025-06-13T12:40:58Z","snapshot_observed_at":"2026-08-07T11:42:31.914563Z","submitted_at":"2025-06-02T04:32:15Z","title":"VM14K: First Vietnamese Medical Benchmark","version":2},"reference_index":31,"source":"pdf_text","source_observed_at":"2026-08-07T11:51:10.016963Z"},"links":{"citing_paper":"/paper/2506.01305"},"observation_digest":"sha256:a3d0f108e7b4d45d448ba4baf57d9b6337672946e7610dc808f8a8a53b01bf61","observation_id":"71ad8ab4-ce75-4061-a3b4-7398a6c2bd5a","resolution":{"observed_at":"2026-08-07T11:51:11.372335Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T11:51:11.107101Z","title":"Bitterman, Ling Pan, Ching-Yu Cheng, James Zou, and Dianbo Liu","venue":null,"work_id":"45d0c33b-21db-4186-99d0-cbb5d83d46ef","year":2025},"citing_paper":{"arxiv_id":"2506.01305","last_updated":"2025-06-13T12:40:58Z","snapshot_observed_at":"2026-08-07T11:42:31.914563Z","submitted_at":"2025-06-02T04:32:15Z","title":"VM14K: First Vietnamese Medical Benchmark","version":2},"reference_index":32,"source":"pdf_text","source_observed_at":"2026-08-07T11:51:10.060621Z"},"links":{"citing_paper":"/paper/2506.01305"},"observation_digest":"sha256:956014347c0daa930d7bc0122af55be33b078447ae735ec52b67bf3349481525","observation_id":"06f4c4fc-2ed0-4808-a49a-1e4aa9b8ccbc","resolution":{"observed_at":"2026-08-07T11:51:11.169546Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T11:51:10.969647Z","title":"Qwen3 technical report, 2025","venue":null,"work_id":"76803076-834c-4f6e-a7fc-9531ba49dcf9","year":2025},"citing_paper":{"arxiv_id":"2506.01305","last_updated":"2025-06-13T12:40:58Z","snapshot_observed_at":"2026-08-07T11:42:31.914563Z","submitted_at":"2025-06-02T04:32:15Z","title":"VM14K: First Vietnamese Medical Benchmark","version":2},"reference_index":33,"source":"pdf_text","source_observed_at":"2026-08-07T11:51:10.100763Z"},"links":{"citing_paper":"/paper/2506.01305"},"observation_digest":"sha256:34ad6fcf9ff65e5efb7d5771a00359a05d90d65a4781274ab4c48c665ea410d4","observation_id":"5230d645-e6f6-45b1-bc41-b0debcbb529d","resolution":{"observed_at":"2026-08-07T11:51:11.039522Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T11:51:10.797608Z","title":"Tcmd: A traditional chinese medicine qa dataset for evaluating large language models, 2024","venue":null,"work_id":"0bacce80-5cf0-48e1-9098-5be937b29a72","year":2024},"citing_paper":{"arxiv_id":"2506.01305","last_updated":"2025-06-13T12:40:58Z","snapshot_observed_at":"2026-08-07T11:42:31.914563Z","submitted_at":"2025-06-02T04:32:15Z","title":"VM14K: First Vietnamese Medical Benchmark","version":2},"reference_index":34,"source":"pdf_text","source_observed_at":"2026-08-07T11:51:10.141336Z"},"links":{"citing_paper":"/paper/2506.01305"},"observation_digest":"sha256:747c5d1f5b85cfcd50f965630b30c48dc385492c48a7d28fe3313d06d3a0d5fb","observation_id":"3346c145-d44b-4fbd-ad49-d1b9ac7bb9ae","resolution":{"observed_at":"2026-08-07T11:51:10.872433Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T11:51:10.664636Z","title":"A survey of datasets in medicine for large language models.Intelligence & Robotics, 4(4), 2024","venue":null,"work_id":"7961a743-489b-489c-b2e4-f01746202493","year":2024},"citing_paper":{"arxiv_id":"2506.01305","last_updated":"2025-06-13T12:40:58Z","snapshot_observed_at":"2026-08-07T11:42:31.914563Z","submitted_at":"2025-06-02T04:32:15Z","title":"VM14K: First Vietnamese Medical Benchmark","version":2},"reference_index":35,"source":"pdf_text","source_observed_at":"2026-08-07T11:51:10.179628Z"},"links":{"citing_paper":"/paper/2506.01305"},"observation_digest":"sha256:1fbce51ff84703e26ae8c243df0c89fc0869f657c5d758961ede665dbb22d5dd","observation_id":"8ede7b1c-6013-4e05-9273-811c9268acff","resolution":{"observed_at":"2026-08-07T11:51:10.709614Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T11:51:10.212996Z","title":"Ultramedical: Building specialized generalists in biomedicine, 2024","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2506.01305","last_updated":"2025-06-13T12:40:58Z","snapshot_observed_at":"2026-08-07T11:42:31.914563Z","submitted_at":"2025-06-02T04:32:15Z","title":"VM14K: First Vietnamese Medical Benchmark","version":2},"reference_index":36,"source":"pdf_text","source_observed_at":"2026-08-07T11:51:10.212996Z"},"links":{"citing_paper":"/paper/2506.01305"},"observation_digest":"sha256:fa7ff534e1fd5638c1e604463f8d849b6ebfe12532a148eebc45ae198aac2231","observation_id":"a7a3c9ed-db01-46fd-9c39-03fe674be967","resolution":{"observed_at":"2026-08-07T11:51:10.212996Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T11:51:10.511934Z","title":"Zhang, X","venue":null,"work_id":"7371bfe2-ea4d-4089-8a3c-31bc55433d21","year":2018},"citing_paper":{"arxiv_id":"2506.01305","last_updated":"2025-06-13T12:40:58Z","snapshot_observed_at":"2026-08-07T11:42:31.914563Z","submitted_at":"2025-06-02T04:32:15Z","title":"VM14K: First Vietnamese Medical Benchmark","version":2},"reference_index":37,"source":"pdf_text","source_observed_at":"2026-08-07T11:51:10.259506Z"},"links":{"citing_paper":"/paper/2506.01305"},"observation_digest":"sha256:e964005d50a290ff0aa8336298ef0b07795bc2afbb6eb3af5888755ea5e0caaa","observation_id":"4370fa76-848f-4540-aaff-c48abbb594e5","resolution":{"observed_at":"2026-08-07T11:51:10.584976Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T11:51:10.404253Z","title":"Critical care studies using large language models based on electronic healthcare records: A technical note.J Intensive Med, 5(2):137–150, 2024","venue":null,"work_id":"1ed3cbc5-1620-48a6-a571-c31aebc9a1c4","year":2024},"citing_paper":{"arxiv_id":"2506.01305","last_updated":"2025-06-13T12:40:58Z","snapshot_observed_at":"2026-08-07T11:42:31.914563Z","submitted_at":"2025-06-02T04:32:15Z","title":"VM14K: First Vietnamese Medical Benchmark","version":2},"reference_index":38,"source":"pdf_text","source_observed_at":"2026-08-07T11:51:10.305405Z"},"links":{"citing_paper":"/paper/2506.01305"},"observation_digest":"sha256:a6ad751db8be153149bc34bd20c794857e51310aa198c1428c9b120bf40961ee","observation_id":"fa28b376-7e8b-4e64-8b84-ab59155af243","resolution":{"observed_at":"2026-08-07T11:51:10.457341Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}}],"paper":{"arxiv_id":"2506.01305","last_updated":"2025-06-13T12:40:58Z","latest_version":2,"primary_category":"cs.CL","snapshot_observed_at":"2026-08-07T11:42:31.914563Z","submitted_at":"2025-06-02T04:32:15Z","title":"VM14K: First Vietnamese Medical Benchmark"},"reference_resolution":{"displayed":38,"state_counts":{"malformed_identifier":0,"metadata_mismatch":0,"parse_uncertain":0,"unresolved":16,"verified_exact":0,"verified_fuzzy":22},"total_outbound_references":38},"refusal":"A citation records a reference. It does not transfer a finding from one paper to another.","schema":"pith.paper-citation-record.v1","standing_sources":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"thesis":"As of 8 August 2026, this Paper Citation Record lists 38 of 38 outbound references and 0 inbound Pith citation observations for arXiv:2506.01305."}