{"as_of":"2026-08-09T14:04:00Z","caps":{"database_statements":6,"inbound":100,"outbound":100},"context_digest":"sha256:2f8b763ebc9f1e8fc991c425dcea18318de2326bf8f1570784fc10888cc25dd0","coverage":[{"denominator":17,"lane":"reference_resolution","note":"Typed states for the displayed outbound observations.","records_observed":17,"source":"paper_references, paper_reference_links","source_observed_at":"2026-08-02T06:50:47.352443Z","state":"measured"},{"denominator":17,"lane":"standing_notices","note":"One-hop event checks from named stored sources.","records_observed":17,"source":"scholarly_work_events, retraction_status_cache","source_observed_at":"2026-08-09T06:31:02.800959+00:00","state":"measured"},{"denominator":0,"lane":"inbound_itemization","note":"Pith citing papers itemized under the disclosed page cap.","records_observed":0,"source":"paper_references, paper_reference_links","source_observed_at":null,"state":"measured"},{"denominator":1,"lane":"external_citation_measurements","note":"A source-named dated measurement, never combined with another source.","records_observed":0,"source":"cited_works","source_observed_at":null,"state":"measured"}],"external_citation_measurements":[],"inbound":[],"links":{"evidence":"/evidence","html":"/paper/2607.11997/citation-record","integrity":"/paper/2607.11997/integrity","json":"/paper/2607.11997/citation-record.json","paper":"/paper/2607.11997"},"outbound":[{"citation":{"cited_paper":{"arxiv_id":"2404.14219","last_updated":"2024-08-30T21:17:17Z","snapshot_observed_at":"2026-07-06T18:03:47.096406Z","submitted_at":"2024-04-22T14:32:33Z","title":"Phi-3 Technical Report: A Highly Capable Language Model Locally on Your Phone","version":4},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2404.14219","snapshot_observed_at":"2026-08-02T06:50:46.016497Z","title":"A., Bach, N., Bahree, A., Bakhtiari, A., Bao, J., Behl, H., et al","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2607.11997","last_updated":"2026-07-15T11:14:47Z","snapshot_observed_at":"2026-08-02T06:50:45.588647Z","submitted_at":"2026-07-13T16:08:30Z","title":"Are we Merging the Right Models? Impact of Expert Training Duration on Model Merging for LLMs","version":2},"reference_index":1,"source":"pdf_text","source_observed_at":"2026-08-02T06:50:46.016497Z"},"links":{"cited_paper":"/paper/2404.14219","citing_paper":"/paper/2607.11997"},"observation_digest":"sha256:25fc2ace13995d7e053fd8d78319faa21060bb02b1cb01808a3ee97fa37df372","observation_id":"b796874c-4845-4b22-9b05-eac8c00850f0","resolution":{"observed_at":"2026-08-02T06:50:46.016497Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-02T06:50:46.370017Z","title":null,"venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2607.11997","last_updated":"2026-07-15T11:14:47Z","snapshot_observed_at":"2026-08-02T06:50:45.588647Z","submitted_at":"2026-07-13T16:08:30Z","title":"Are we Merging the Right Models? Impact of Expert Training Duration on Model Merging for LLMs","version":2},"reference_index":6,"source":"pdf_text","source_observed_at":"2026-08-02T06:50:46.370017Z"},"links":{"citing_paper":"/paper/2607.11997"},"observation_digest":"sha256:e614524562efbc5ffc4801e6fc1ef1706e680a0e294fcc2fbb43cf77ac2174fb","observation_id":"9a6b27ab-4c0b-4105-930e-40b89b306940","resolution":{"observed_at":"2026-08-02T06:50:46.370017Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2312.02120","last_updated":"2024-06-07T02:50:56Z","snapshot_observed_at":"2026-07-06T16:56:46.051958Z","submitted_at":"2023-12-04T18:50:35Z","title":"Magicoder: Empowering Code Generation with OSS-Instruct","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2312.02120","snapshot_observed_at":"2026-08-02T06:50:46.625517Z","title":"Magicoder: empowering code generation with OSS-instruct.arXiv preprint arXiv:2312.02120,","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2607.11997","last_updated":"2026-07-15T11:14:47Z","snapshot_observed_at":"2026-08-02T06:50:45.588647Z","submitted_at":"2026-07-13T16:08:30Z","title":"Are we Merging the Right Models? Impact of Expert Training Duration on Model Merging for LLMs","version":2},"reference_index":9,"source":"pdf_text","source_observed_at":"2026-08-02T06:50:46.625517Z"},"links":{"cited_paper":"/paper/2312.02120","citing_paper":"/paper/2607.11997"},"observation_digest":"sha256:a85d085e4b9ecce54696467b9ad147b88aff83b9540630dcdc74df3a4800ffc7","observation_id":"ff15ddd1-6727-4afc-9872-ad964480fa97","resolution":{"observed_at":"2026-08-02T06:50:46.625517Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2402.04249","last_updated":"2024-02-27T04:43:08Z","snapshot_observed_at":"2026-07-06T17:26:23.067923Z","submitted_at":"2024-02-06T18:59:08Z","title":"HarmBench: A Standardized Evaluation Framework for Automated Red Teaming and Robust Refusal","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2402.04249","snapshot_observed_at":"2026-08-02T06:50:46.683366Z","title":"Harm- Bench: a standardized evaluation framework for auto- mated red teaming and robust refusal.arXiv preprint arXiv:2402.04249,","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2607.11997","last_updated":"2026-07-15T11:14:47Z","snapshot_observed_at":"2026-08-02T06:50:45.588647Z","submitted_at":"2026-07-13T16:08:30Z","title":"Are we Merging the Right Models? Impact of Expert Training Duration on Model Merging for LLMs","version":2},"reference_index":10,"source":"pdf_text","source_observed_at":"2026-08-02T06:50:46.683366Z"},"links":{"cited_paper":"/paper/2402.04249","citing_paper":"/paper/2607.11997"},"observation_digest":"sha256:49d3cec1373cb85b790baa8ee2d53c10cbf3634ea16f69571cff52c54c87bb90","observation_id":"1ad042f7-cf7b-46d1-a0e7-99a4b7bfdea5","resolution":{"observed_at":"2026-08-02T06:50:46.683366Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2404.07815","last_updated":"2024-10-04T00:11:05Z","snapshot_observed_at":"2026-08-05T06:40:25.340181Z","submitted_at":"2024-04-11T14:58:19Z","title":"Post-Hoc Reversal: Are We Selecting Models Prematurely?","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2404.07815","snapshot_observed_at":"2026-08-02T06:50:46.761806Z","title":null,"venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2607.11997","last_updated":"2026-07-15T11:14:47Z","snapshot_observed_at":"2026-08-02T06:50:45.588647Z","submitted_at":"2026-07-13T16:08:30Z","title":"Are we Merging the Right Models? Impact of Expert Training Duration on Model Merging for LLMs","version":2},"reference_index":11,"source":"pdf_text","source_observed_at":"2026-08-02T06:50:46.761806Z"},"links":{"cited_paper":"/paper/2404.07815","citing_paper":"/paper/2607.11997"},"observation_digest":"sha256:c8b85ed059ab57ee62d1e4e27b43879a1a6bef2432bccad3431965bf9faa6ccf","observation_id":"1b253a7f-37f8-451d-89b1-9d2b02ed5a58","resolution":{"observed_at":"2026-08-02T06:50:46.761806Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2604.13627","last_updated":"2026-08-03T12:13:35Z","snapshot_observed_at":"2026-08-06T23:31:01.769852Z","submitted_at":"2026-04-15T08:53:42Z","title":"(How) Learning Rates Regulate Catastrophic Overtraining","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2604.13627","snapshot_observed_at":"2026-08-02T06:50:46.840290Z","title":"(How) learning rates regulate catastrophic overtraining.arXiv preprint arXiv:2604.13627,","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2607.11997","last_updated":"2026-07-15T11:14:47Z","snapshot_observed_at":"2026-08-02T06:50:45.588647Z","submitted_at":"2026-07-13T16:08:30Z","title":"Are we Merging the Right Models? Impact of Expert Training Duration on Model Merging for LLMs","version":2},"reference_index":12,"source":"pdf_text","source_observed_at":"2026-08-02T06:50:46.840290Z"},"links":{"cited_paper":"/paper/2604.13627","citing_paper":"/paper/2607.11997"},"observation_digest":"sha256:f79c0953f01603c9b3722ddcd201bfd0e078af52be5407285dd67e061582478e","observation_id":"f1d29d39-3815-471e-8ad2-1ef6ea2147a0","resolution":{"observed_at":"2026-08-02T06:50:46.840290Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-02T06:50:47.120762Z","title":"Leveraging model soups to classify ICH images from the Mekong Delta","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2607.11997","last_updated":"2026-07-15T11:14:47Z","snapshot_observed_at":"2026-08-02T06:50:45.588647Z","submitted_at":"2026-07-13T16:08:30Z","title":"Are we Merging the Right Models? Impact of Expert Training Duration on Model Merging for LLMs","version":2},"reference_index":14,"source":"pdf_text","source_observed_at":"2026-08-02T06:50:47.120762Z"},"links":{"citing_paper":"/paper/2607.11997"},"observation_digest":"sha256:d47401d5a90f9fa33f1bf0b152b33a9e5c12c59dcb76e22694c4ebe7737b569a","observation_id":"06b57ed0-044b-42d7-806c-9cbbd301789c","resolution":{"observed_at":"2026-08-02T06:50:47.120762Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2407.13690","last_updated":"2024-12-23T17:32:21Z","snapshot_observed_at":"2026-08-05T19:05:54.552558Z","submitted_at":"2024-06-18T07:14:02Z","title":"DART-Math: Difficulty-Aware Rejection Tuning for Mathematical Problem-Solving","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2407.13690","snapshot_observed_at":"2026-08-02T06:50:47.198790Z","title":null,"venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2607.11997","last_updated":"2026-07-15T11:14:47Z","snapshot_observed_at":"2026-08-02T06:50:45.588647Z","submitted_at":"2026-07-13T16:08:30Z","title":"Are we Merging the Right Models? Impact of Expert Training Duration on Model Merging for LLMs","version":2},"reference_index":15,"source":"pdf_text","source_observed_at":"2026-08-02T06:50:47.198790Z"},"links":{"cited_paper":"/paper/2407.13690","citing_paper":"/paper/2607.11997"},"observation_digest":"sha256:d5428ed68cd3c6f98fba57f3d10edb35047340809d34e084dfda0f6984de8e4c","observation_id":"3de66465-e2a7-4063-a89f-5846ecb6dd93","resolution":{"observed_at":"2026-08-02T06:50:47.198790Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2311.07911","last_updated":"2023-11-14T05:13:55Z","snapshot_observed_at":"2026-07-06T16:47:08.877195Z","submitted_at":"2023-11-14T05:13:55Z","title":"Instruction-Following Evaluation for Large Language Models","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2311.07911","snapshot_observed_at":"2026-08-02T06:50:47.289943Z","title":"Instruction-following evaluation for large language models.arXiv preprint arXiv:2311.07911,","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2607.11997","last_updated":"2026-07-15T11:14:47Z","snapshot_observed_at":"2026-08-02T06:50:45.588647Z","submitted_at":"2026-07-13T16:08:30Z","title":"Are we Merging the Right Models? Impact of Expert Training Duration on Model Merging for LLMs","version":2},"reference_index":16,"source":"pdf_text","source_observed_at":"2026-08-02T06:50:47.289943Z"},"links":{"cited_paper":"/paper/2311.07911","citing_paper":"/paper/2607.11997"},"observation_digest":"sha256:1aafa07679abcdd28a4fdcd00d2e7c26af93a7a4e60590d0d07fcf3837759eb7","observation_id":"dde66761-ade0-403c-a483-fde4b3d1c713","resolution":{"observed_at":"2026-08-02T06:50:47.289943Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2107.03374","last_updated":"2021-07-14T17:16:02Z","snapshot_observed_at":"2026-08-08T11:58:24.516369Z","submitted_at":"2021-07-07T17:41:24Z","title":"Evaluating Large Language Models Trained on Code","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2107.03374","snapshot_observed_at":"2026-08-02T06:50:46.150392Z","title":null,"venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2607.11997","last_updated":"2026-07-15T11:14:47Z","snapshot_observed_at":"2026-08-02T06:50:45.588647Z","submitted_at":"2026-07-13T16:08:30Z","title":"Are we Merging the Right Models? Impact of Expert Training Duration on Model Merging for LLMs","version":2},"reference_index":2001,"source":"pdf_text","source_observed_at":"2026-08-02T06:50:46.150392Z"},"links":{"cited_paper":"/paper/2107.03374","citing_paper":"/paper/2607.11997"},"observation_digest":"sha256:56f20ba9db573276099aa0c55d6a544efdce3afc73ba97551c9dee104835b66e","observation_id":"e7acefc7-7e31-4ac9-95ac-98600eebf8fc","resolution":{"observed_at":"2026-08-02T06:50:46.150392Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2110.14168","last_updated":"2021-11-18T00:23:45Z","snapshot_observed_at":"2026-08-07T01:45:38.840969Z","submitted_at":"2021-10-27T04:49:45Z","title":"Training Verifiers to Solve Math Word Problems","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2110.14168","snapshot_observed_at":"2026-08-02T06:50:46.312944Z","title":"Training verifiers to solve math word problems","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2607.11997","last_updated":"2026-07-15T11:14:47Z","snapshot_observed_at":"2026-08-02T06:50:45.588647Z","submitted_at":"2026-07-13T16:08:30Z","title":"Are we Merging the Right Models? Impact of Expert Training Duration on Model Merging for LLMs","version":2},"reference_index":2018,"source":"pdf_text","source_observed_at":"2026-08-02T06:50:46.312944Z"},"links":{"cited_paper":"/paper/2110.14168","citing_paper":"/paper/2607.11997"},"observation_digest":"sha256:5cbb5662c9a159bf5b406002fa52eb21597d8f7a3650b1b1bde40e0c545392d3","observation_id":"03f676bb-151d-4718-a2ea-a6e969067fd2","resolution":{"observed_at":"2026-08-02T06:50:46.312944Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"1803.05457","last_updated":"2018-03-14T18:04:21Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2018-03-14T18:04:21Z","title":"Think you have Solved Question Answering? Try ARC, the AI2 Reasoning Challenge","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"1803.05457","snapshot_observed_at":"2026-08-02T06:50:46.228734Z","title":"Think you have solved question answering? try ARC, the AI2 reasoning chal- lenge.arXiv preprint arXiv:1803.05457,","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2607.11997","last_updated":"2026-07-15T11:14:47Z","snapshot_observed_at":"2026-08-02T06:50:45.588647Z","submitted_at":"2026-07-13T16:08:30Z","title":"Are we Merging the Right Models? Impact of Expert Training Duration on Model Merging for LLMs","version":2},"reference_index":2021,"source":"pdf_text","source_observed_at":"2026-08-02T06:50:46.228734Z"},"links":{"cited_paper":"/paper/1803.05457","citing_paper":"/paper/2607.11997"},"observation_digest":"sha256:b670f9b1a02a9ff89d7027beb51523cc6c4e90235a4227f2aa58bce099acd213","observation_id":"44d98251-a298-41ae-8a2f-6cb8f1d6fd65","resolution":{"observed_at":"2026-08-02T06:50:46.228734Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2411.15124","last_updated":"2025-04-14T22:39:09Z","snapshot_observed_at":"2026-08-08T12:58:42.430328Z","submitted_at":"2024-11-22T18:44:04Z","title":"Tulu 3: Pushing Frontiers in Open Language Model Post-Training","version":5},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2411.15124","snapshot_observed_at":"2026-08-02T06:50:46.535520Z","title":null,"venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2607.11997","last_updated":"2026-07-15T11:14:47Z","snapshot_observed_at":"2026-08-02T06:50:45.588647Z","submitted_at":"2026-07-13T16:08:30Z","title":"Are we Merging the Right Models? Impact of Expert Training Duration on Model Merging for LLMs","version":2},"reference_index":2022,"source":"pdf_text","source_observed_at":"2026-08-02T06:50:46.535520Z"},"links":{"cited_paper":"/paper/2411.15124","citing_paper":"/paper/2607.11997"},"observation_digest":"sha256:7fb94ca0ac5d421bcf4a55e3377cc2ea7993d15c31f49a1ed25b2a145e6eedc1","observation_id":"fa4101ca-5d6a-44de-a051-72d67ac4bfcb","resolution":{"observed_at":"2026-08-02T06:50:46.535520Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2411.03055","last_updated":"2025-08-08T14:13:09Z","snapshot_observed_at":"2026-07-06T19:45:30.325597Z","submitted_at":"2024-11-05T12:42:42Z","title":"ATM: Improving Model Merging by Alternating Tuning and Merging","version":4},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2411.03055","snapshot_observed_at":"2026-08-02T06:50:47.352443Z","title":"S., Silvestri, F., and Rodol`a, E","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2607.11997","last_updated":"2026-07-15T11:14:47Z","snapshot_observed_at":"2026-08-02T06:50:45.588647Z","submitted_at":"2026-07-13T16:08:30Z","title":"Are we Merging the Right Models? Impact of Expert Training Duration on Model Merging for LLMs","version":2},"reference_index":2023,"source":"pdf_text","source_observed_at":"2026-08-02T06:50:47.352443Z"},"links":{"cited_paper":"/paper/2411.03055","citing_paper":"/paper/2607.11997"},"observation_digest":"sha256:0fd853f479bca011eeb31d0399b5bbe565dd4f8e80e0a72811eebe37aa6c67aa","observation_id":"4133b9af-05ba-4708-8dc4-b794ed0739a3","resolution":{"observed_at":"2026-08-02T06:50:47.352443Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2506.02153","last_updated":"2025-09-15T22:15:00Z","snapshot_observed_at":"2026-08-05T18:56:11.132981Z","submitted_at":"2025-06-02T18:35:16Z","title":"Small Language Models are the Future of Agentic AI","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2506.02153","snapshot_observed_at":"2026-08-02T06:50:46.076212Z","title":"C., and Molchanov, P","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2607.11997","last_updated":"2026-07-15T11:14:47Z","snapshot_observed_at":"2026-08-02T06:50:45.588647Z","submitted_at":"2026-07-13T16:08:30Z","title":"Are we Merging the Right Models? Impact of Expert Training Duration on Model Merging for LLMs","version":2},"reference_index":2024,"source":"pdf_text","source_observed_at":"2026-08-02T06:50:46.076212Z"},"links":{"cited_paper":"/paper/2506.02153","citing_paper":"/paper/2607.11997"},"observation_digest":"sha256:ef7ffcd6f9b8faf05cd09264b4670e104117cf4d250a84cf7c52c9325731dc1e","observation_id":"ec3a490e-7255-444f-8403-9996ce34ccee","resolution":{"observed_at":"2026-08-02T06:50:46.076212Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2506.14126","last_updated":"2026-06-16T20:21:45Z","snapshot_observed_at":"2026-08-07T00:15:47.597756Z","submitted_at":"2025-06-17T02:42:10Z","title":"From Memorization to Parameter Interference: How Overtraining Experts Harms Model Merging","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2506.14126","snapshot_observed_at":"2026-08-02T06:50:46.438907Z","title":null,"venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2607.11997","last_updated":"2026-07-15T11:14:47Z","snapshot_observed_at":"2026-08-02T06:50:45.588647Z","submitted_at":"2026-07-13T16:08:30Z","title":"Are we Merging the Right Models? Impact of Expert Training Duration on Model Merging for LLMs","version":2},"reference_index":2025,"source":"pdf_text","source_observed_at":"2026-08-02T06:50:46.438907Z"},"links":{"cited_paper":"/paper/2506.14126","citing_paper":"/paper/2607.11997"},"observation_digest":"sha256:f66392c9f0b9828b1d7f1175edf73d6e10ed7a84ff57f3745ffc90b1c8a826a1","observation_id":"295c6de6-2825-41f3-a601-d3454334ba90","resolution":{"observed_at":"2026-08-02T06:50:46.438907Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2402.06619","last_updated":"2024-02-09T18:51:49Z","snapshot_observed_at":"2026-07-06T17:28:06.332906Z","submitted_at":"2024-02-09T18:51:49Z","title":"Aya Dataset: An Open-Access Collection for Multilingual Instruction Tuning","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2402.06619","snapshot_observed_at":"2026-08-02T06:50:46.973903Z","title":"Aya dataset: an open-access collection for multilingual instruction tuning","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2607.11997","last_updated":"2026-07-15T11:14:47Z","snapshot_observed_at":"2026-08-02T06:50:45.588647Z","submitted_at":"2026-07-13T16:08:30Z","title":"Are we Merging the Right Models? Impact of Expert Training Duration on Model Merging for LLMs","version":2},"reference_index":2026,"source":"pdf_text","source_observed_at":"2026-08-02T06:50:46.973903Z"},"links":{"cited_paper":"/paper/2402.06619","citing_paper":"/paper/2607.11997"},"observation_digest":"sha256:c28031786ff50218db672abe826baefcbfd7a64133f2e9f23d4122748c58e720","observation_id":"c0052134-c04b-455f-9403-6e5c415f3c8c","resolution":{"observed_at":"2026-08-02T06:50:46.973903Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}}],"paper":{"arxiv_id":"2607.11997","last_updated":"2026-07-15T11:14:47Z","latest_version":2,"primary_category":"cs.LG","snapshot_observed_at":"2026-08-02T06:50:45.588647Z","submitted_at":"2026-07-13T16:08:30Z","title":"Are we Merging the Right Models? Impact of Expert Training Duration on Model Merging for LLMs"},"reference_resolution":{"displayed":17,"state_counts":{"malformed_identifier":0,"metadata_mismatch":0,"parse_uncertain":0,"unresolved":17,"verified_exact":0,"verified_fuzzy":0},"total_outbound_references":17},"refusal":"A citation records a reference. It does not transfer a finding from one paper to another.","schema":"pith.paper-citation-record.v1","standing_sources":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"thesis":"As of 9 August 2026, this Paper Citation Record lists 17 of 17 outbound references and 0 inbound Pith citation observations for arXiv:2607.11997."}