{"as_of":"2026-08-14T12:24:00Z","caps":{"database_statements":6,"inbound":100,"outbound":100},"context_digest":"sha256:fe6709966dcfc2c017c0fc8ad43e516578852ef0f7db8d87d2cc1feab936bf9e","coverage":[{"denominator":0,"lane":"reference_resolution","note":"Typed states for the displayed outbound observations.","records_observed":0,"source":"paper_references, paper_reference_links","source_observed_at":null,"state":"measured"},{"denominator":18,"lane":"standing_notices","note":"One-hop event checks from named stored sources.","records_observed":18,"source":"scholarly_work_events, retraction_status_cache","source_observed_at":"2026-08-14T06:32:32.682623+00:00","state":"measured"},{"denominator":18,"lane":"inbound_itemization","note":"Pith citing papers itemized under the disclosed page cap.","records_observed":18,"source":"paper_references, paper_reference_links","source_observed_at":"2026-08-12T00:18:54.116775Z","state":"measured"},{"denominator":1,"lane":"external_citation_measurements","note":"A source-named dated measurement, never combined with another source.","records_observed":0,"source":"pith","source_observed_at":"2026-07-10T03:46:44.819123Z","state":"measured"}],"external_citation_measurements":[],"inbound":[{"citation":{"cited_paper":{"arxiv_id":"2503.10657","last_updated":"2025-05-20T14:59:21Z","snapshot_observed_at":"2026-08-13T09:25:43.676027Z","submitted_at":"2025-03-08T04:07:07Z","title":"RouterEval: A Comprehensive Benchmark for Routing LLMs to Explore Model-level Scaling Up in LLMs","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2503.10657","snapshot_observed_at":"2026-08-07T23:48:00.958174Z","title":"Liang, Yupei Lin, Yandong Chen, Shanshan Zhong, Hefeng Wu, and Liang Lin","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2502.08773","last_updated":"2025-07-22T15:27:33Z","snapshot_observed_at":"2026-08-14T02:29:31.486179Z","submitted_at":"2025-02-12T20:30:28Z","title":"Universal Model Routing for Efficient LLM Inference","version":2},"reference_index":43,"source":"arxiv_source","source_observed_at":"2026-08-07T23:48:00.958174Z"},"links":{"cited_paper":"/paper/2503.10657","citing_paper":"/paper/2502.08773"},"observation_digest":"sha256:486e57fad0b2fe65cc81799b39dc7336f1d23475ffaa01037d0c3d149cfd74e3","observation_id":"c8c53eb2-396d-4395-b2ca-5008e9421336","resolution":{"observed_at":"2026-08-07T23:48:00.958174Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2503.10657","last_updated":"2025-05-20T14:59:21Z","snapshot_observed_at":"2026-08-13T09:25:43.676027Z","submitted_at":"2025-03-08T04:07:07Z","title":"RouterEval: A Comprehensive Benchmark for Routing LLMs to Explore Model-level Scaling Up in LLMs","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2503.10657","snapshot_observed_at":"2026-08-07T15:11:17.431359Z","title":"Routereval: A comprehensive benchmark for routing llms to explore model-level scaling up in llms.arXiv preprint arXiv:2503.10657,","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2505.16221","last_updated":"2025-05-22T04:46:04Z","snapshot_observed_at":"2026-08-14T04:36:48.417359Z","submitted_at":"2025-05-22T04:46:04Z","title":"LightRouter: Towards Efficient LLM Collaboration with Minimal Overhead","version":1},"reference_index":13,"source":"pdf_text","source_observed_at":"2026-08-07T15:11:17.431359Z"},"links":{"cited_paper":"/paper/2503.10657","citing_paper":"/paper/2505.16221"},"observation_digest":"sha256:fbd2de06f842cea47b7d36ac044dd6581d3a7398ab51dea28c6f13b563fbfcc9","observation_id":"f9b76be3-78ac-49d5-82ed-63d56d03b83a","resolution":{"observed_at":"2026-08-07T15:11:17.431359Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2503.10657","last_updated":"2025-05-20T14:59:21Z","snapshot_observed_at":"2026-08-13T09:25:43.676027Z","submitted_at":"2025-03-08T04:07:07Z","title":"RouterEval: A Comprehensive Benchmark for Routing LLMs to Explore Model-level Scaling Up in LLMs","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2503.10657","snapshot_observed_at":"2026-08-07T14:12:37.520825Z","title":"Routereval: A comprehensive benchmark for routing llms to explore model-level scaling up in llms.arXiv preprint arXiv:2503.10657, 2025","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2505.19797","last_updated":"2025-06-18T09:47:20Z","snapshot_observed_at":"2026-08-13T15:45:35.029123Z","submitted_at":"2025-05-26T10:29:42Z","title":"The Avengers: A Simple Recipe for Uniting Smaller Language Models to Challenge Proprietary Giants","version":3},"reference_index":33,"source":"pdf_text","source_observed_at":"2026-08-07T14:12:37.520825Z"},"links":{"cited_paper":"/paper/2503.10657","citing_paper":"/paper/2505.19797"},"observation_digest":"sha256:1e53b6dc4089902eac3a63a3a495c274ffc974af9f57b83fa43e1ada29404b44","observation_id":"2deb8c0a-17a0-472d-bf21-777bba6b2774","resolution":{"observed_at":"2026-08-07T14:12:37.520825Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2503.10657","last_updated":"2025-05-20T14:59:21Z","snapshot_observed_at":"2026-08-13T09:25:43.676027Z","submitted_at":"2025-03-08T04:07:07Z","title":"RouterEval: A Comprehensive Benchmark for Routing LLMs to Explore Model-level Scaling Up in LLMs","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2503.10657","snapshot_observed_at":"2026-08-07T05:56:56.742261Z","title":"Routereval: A comprehensive benchmark for routing llms to explore model-level scaling up in llms,","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2506.06579","last_updated":"2025-06-06T23:13:08Z","snapshot_observed_at":"2026-08-13T00:55:48.528355Z","submitted_at":"2025-06-06T23:13:08Z","title":"Towards Efficient Multi-LLM Inference: Characterization and Analysis of LLM Routing and Hierarchical Techniques","version":1},"reference_index":17,"source":"pdf_text","source_observed_at":"2026-08-07T05:56:56.742261Z"},"links":{"cited_paper":"/paper/2503.10657","citing_paper":"/paper/2506.06579"},"observation_digest":"sha256:c29720c799115dda0d3f87c754fdf4edf40e9af6ee1b1526d336ca4b5342315f","observation_id":"e13a777d-61db-46c8-bd3e-3388bac034aa","resolution":{"observed_at":"2026-08-07T05:56:56.742261Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2503.10657","last_updated":"2025-05-20T14:59:21Z","snapshot_observed_at":"2026-08-13T09:25:43.676027Z","submitted_at":"2025-03-08T04:07:07Z","title":"RouterEval: A Comprehensive Benchmark for Routing LLMs to Explore Model-level Scaling Up in LLMs","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2503.10657","snapshot_observed_at":"2026-08-05T23:23:08.355033Z","title":"Liang, Yupei Lin, Yandong Chen, Shanshan Zhong, Hefeng Wu, and Liang Lin","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2508.05585","last_updated":"2025-08-07T17:22:33Z","snapshot_observed_at":"2026-08-08T05:46:23.473665Z","submitted_at":"2025-08-07T17:22:33Z","title":"DART: Dual Adaptive Refinement Transfer for Open-Vocabulary Multi-Label Recognition","version":1},"reference_index":21,"source":"pdf_text","source_observed_at":"2026-08-05T23:23:08.355033Z"},"links":{"cited_paper":"/paper/2503.10657","citing_paper":"/paper/2508.05585"},"observation_digest":"sha256:b85815e780e48df35cac7e3cd6956e7482d1d34b08236b0769b6fdd8142d2e45","observation_id":"d102ab14-2f5e-4e1f-b484-6d775d8326ee","resolution":{"observed_at":"2026-08-05T23:23:08.355033Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2503.10657","last_updated":"2025-05-20T14:59:21Z","snapshot_observed_at":"2026-08-13T09:25:43.676027Z","submitted_at":"2025-03-08T04:07:07Z","title":"RouterEval: A Comprehensive Benchmark for Routing LLMs to Explore Model-level Scaling Up in LLMs","version":2},"cited_work":{"arxiv_id":"2503.10657","doi":null,"metadata_source":"pith","pith_arxiv_id":"2503.10657","snapshot_observed_at":"2026-07-10T03:46:44.819123Z","title":"Routereval: A comprehensive benchmark for routing llms to explore model-level scaling up in llms.arXiv preprint arXiv:2503.10657","venue":"cs.CL","work_id":"7a1ee412-d726-449d-96d9-6253531a9aa6","year":2025},"citing_paper":{"arxiv_id":"2509.24814","last_updated":"2026-05-07T06:47:05Z","snapshot_observed_at":"2026-08-13T11:27:33.387352Z","submitted_at":"2025-09-29T14:02:27Z","title":"A Greedy PDE Router for Blending Neural Operators and Classical Methods","version":2},"reference_index":6,"source":"pdf_text","source_observed_at":"2026-05-18T12:32:13.577474Z"},"links":{"cited_paper":"/paper/2503.10657","citing_paper":"/paper/2509.24814"},"observation_digest":"sha256:0b1fde008caf7d06c162eff4d9f97d9b79c94e5192c3f3ce480c717a7a0cfc66","observation_id":"71724254-d1c4-43b2-b4fc-3e1f8ff3fd70","resolution":{"observed_at":"2026-05-18T12:32:36.156349Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-14T06:32:32.682623+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-14T06:32:32.682623+00:00","source":"crossref"},{"observed_at":"2026-08-14T06:32:18.44784+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2503.10657","last_updated":"2025-05-20T14:59:21Z","snapshot_observed_at":"2026-08-13T09:25:43.676027Z","submitted_at":"2025-03-08T04:07:07Z","title":"RouterEval: A Comprehensive Benchmark for Routing LLMs to Explore Model-level Scaling Up in LLMs","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2503.10657","snapshot_observed_at":"2026-08-03T05:17:29.774492Z","title":"Routereval: A comprehensive benchmark for routing llms to explore model-level scaling up in llms","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2602.02823","last_updated":"2026-05-29T19:48:02Z","snapshot_observed_at":"2026-08-03T05:17:28.452364Z","submitted_at":"2026-02-02T21:23:51Z","title":"R2-Router: A New Paradigm for LLM Routing with Reasoning","version":2},"reference_index":8,"source":"pdf_text","source_observed_at":"2026-08-03T05:17:29.774492Z"},"links":{"cited_paper":"/paper/2503.10657","citing_paper":"/paper/2602.02823"},"observation_digest":"sha256:90f196f0fe3d75f47981c985399327383646ed18e75f3bb7f24f13788c3f25b7","observation_id":"a8de10d9-19ca-4a24-bc9f-ac458b76624f","resolution":{"observed_at":"2026-08-03T05:17:29.774492Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2503.10657","last_updated":"2025-05-20T14:59:21Z","snapshot_observed_at":"2026-08-13T09:25:43.676027Z","submitted_at":"2025-03-08T04:07:07Z","title":"RouterEval: A Comprehensive Benchmark for Routing LLMs to Explore Model-level Scaling Up in LLMs","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2503.10657","snapshot_observed_at":"2026-08-02T20:32:11.860234Z","title":"arXiv preprint arXiv:2503.10657(2025)","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2602.23293","last_updated":"2026-06-18T18:11:36Z","snapshot_observed_at":"2026-08-14T03:17:05.370100Z","submitted_at":"2026-02-26T18:04:03Z","title":"Impacts of Aggregation on Model Diversity and Consumer Utility","version":2},"reference_index":2025,"source":"pdf_text","source_observed_at":"2026-08-02T20:32:11.860234Z"},"links":{"cited_paper":"/paper/2503.10657","citing_paper":"/paper/2602.23293"},"observation_digest":"sha256:038b6e7653356d93027517ae0fe9f3938f28bf0c3fe600543adc7ecfc2913f3b","observation_id":"e677c21e-75b1-4dde-a390-b216535a7791","resolution":{"observed_at":"2026-08-02T20:32:11.860234Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2503.10657","last_updated":"2025-05-20T14:59:21Z","snapshot_observed_at":"2026-08-13T09:25:43.676027Z","submitted_at":"2025-03-08T04:07:07Z","title":"RouterEval: A Comprehensive Benchmark for Routing LLMs to Explore Model-level Scaling Up in LLMs","version":2},"cited_work":{"arxiv_id":"2503.10657","doi":null,"metadata_source":"pith","pith_arxiv_id":"2503.10657","snapshot_observed_at":"2026-07-10T03:46:44.819123Z","title":"Routereval: A comprehensive benchmark for routing llms to explore model-level scaling up in llms.arXiv preprint arXiv:2503.10657","venue":"cs.CL","work_id":"7a1ee412-d726-449d-96d9-6253531a9aa6","year":2025},"citing_paper":{"arxiv_id":"2604.06296","last_updated":"2026-04-15T22:45:43Z","snapshot_observed_at":"2026-08-13T06:40:12.658137Z","submitted_at":"2026-04-07T17:13:47Z","title":"AgentOpt v0.1 Technical Report: Client-Side Optimization for LLM-Based Agent","version":2},"reference_index":12,"source":"pdf_text","source_observed_at":"2026-05-10T19:51:39.564680Z"},"links":{"cited_paper":"/paper/2503.10657","citing_paper":"/paper/2604.06296"},"observation_digest":"sha256:726a35718e1da15fd81153fcde8c9d3fa561559386415f5d282cc3934b2c84bf","observation_id":"a13b1713-3b64-44fb-9594-e0881eaa6342","resolution":{"observed_at":"2026-05-10T22:25:52.580766Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-14T06:32:32.682623+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-14T06:32:32.682623+00:00","source":"crossref"},{"observed_at":"2026-08-14T06:32:18.44784+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2503.10657","last_updated":"2025-05-20T14:59:21Z","snapshot_observed_at":"2026-08-13T09:25:43.676027Z","submitted_at":"2025-03-08T04:07:07Z","title":"RouterEval: A Comprehensive Benchmark for Routing LLMs to Explore Model-level Scaling Up in LLMs","version":2},"cited_work":{"arxiv_id":"2503.10657","doi":null,"metadata_source":"pith","pith_arxiv_id":"2503.10657","snapshot_observed_at":"2026-07-10T03:46:44.819123Z","title":"Routereval: A comprehensive benchmark for routing llms to explore model-level scaling up in llms.arXiv preprint arXiv:2503.10657","venue":"cs.CL","work_id":"7a1ee412-d726-449d-96d9-6253531a9aa6","year":2025},"citing_paper":{"arxiv_id":"2605.07075","last_updated":"2026-05-08T00:49:05Z","snapshot_observed_at":"2026-08-13T04:49:34.839848Z","submitted_at":"2026-05-08T00:49:05Z","title":"ModelLens: Finding the Best for Your Task from Myriads of Models","version":1},"reference_index":27,"source":"pdf_text","source_observed_at":"2026-05-11T01:55:57.223673Z"},"links":{"cited_paper":"/paper/2503.10657","citing_paper":"/paper/2605.07075"},"observation_digest":"sha256:c42c5e5315d0094195f761c1ca77586ef69c5c4f37a7a39e0a1459e909bad194","observation_id":"eb45bc2a-666a-4087-b0ea-097eaaaf8966","resolution":{"observed_at":"2026-05-11T04:05:59.512351Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-14T06:32:32.682623+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-14T06:32:32.682623+00:00","source":"crossref"},{"observed_at":"2026-08-14T06:32:18.44784+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2503.10657","last_updated":"2025-05-20T14:59:21Z","snapshot_observed_at":"2026-08-13T09:25:43.676027Z","submitted_at":"2025-03-08T04:07:07Z","title":"RouterEval: A Comprehensive Benchmark for Routing LLMs to Explore Model-level Scaling Up in LLMs","version":2},"cited_work":{"arxiv_id":"2503.10657","doi":null,"metadata_source":"pith","pith_arxiv_id":"2503.10657","snapshot_observed_at":"2026-07-10T03:46:44.819123Z","title":"Routereval: A comprehensive benchmark for routing llms to explore model-level scaling up in llms.arXiv preprint arXiv:2503.10657","venue":"cs.CL","work_id":"7a1ee412-d726-449d-96d9-6253531a9aa6","year":2025},"citing_paper":{"arxiv_id":"2605.25424","last_updated":"2026-05-25T04:52:10Z","snapshot_observed_at":"2026-08-14T02:29:08.641815Z","submitted_at":"2026-05-25T04:52:10Z","title":"SeqRoute: Global Budget-Aware Sequential LLM Routing via Offline Reinforcement Learning","version":1},"reference_index":10,"source":"pdf_text","source_observed_at":"2026-06-29T23:13:41.456647Z"},"links":{"cited_paper":"/paper/2503.10657","citing_paper":"/paper/2605.25424"},"observation_digest":"sha256:e94116ef84c3fca95d0995f4131509e23e8ebbda7bc0d34692fe9fcd0ed641d8","observation_id":"73e15b42-552a-4048-92c9-241bb1b828d1","resolution":{"observed_at":"2026-06-29T23:14:01.051611Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-14T06:32:32.682623+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-14T06:32:32.682623+00:00","source":"crossref"},{"observed_at":"2026-08-14T06:32:18.44784+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2503.10657","last_updated":"2025-05-20T14:59:21Z","snapshot_observed_at":"2026-08-13T09:25:43.676027Z","submitted_at":"2025-03-08T04:07:07Z","title":"RouterEval: A Comprehensive Benchmark for Routing LLMs to Explore Model-level Scaling Up in LLMs","version":2},"cited_work":{"arxiv_id":"2503.10657","doi":null,"metadata_source":"pith","pith_arxiv_id":"2503.10657","snapshot_observed_at":"2026-07-10T03:46:44.819123Z","title":"Routereval: A comprehensive benchmark for routing llms to explore model-level scaling up in llms.arXiv preprint arXiv:2503.10657","venue":"cs.CL","work_id":"7a1ee412-d726-449d-96d9-6253531a9aa6","year":2025},"citing_paper":{"arxiv_id":"2607.00053","last_updated":"2026-06-30T01:46:26Z","snapshot_observed_at":"2026-07-07T00:05:41.920778Z","submitted_at":"2026-06-30T01:46:26Z","title":"SWE-Router: Routing in Multi-turn Agentic Software Engineering Tasks","version":1},"reference_index":47,"source":"arxiv_source","source_observed_at":"2026-07-02T18:19:43.146102Z"},"links":{"cited_paper":"/paper/2503.10657","citing_paper":"/paper/2607.00053"},"observation_digest":"sha256:0b8046eb304c46d45d0ea566dde30cd60ccd39c2207057886d0bc481b47b3f64","observation_id":"f97faf56-7698-4646-a1b1-41303bc13c05","resolution":{"observed_at":"2026-07-02T18:37:16.528637Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-14T06:32:32.682623+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-14T06:32:32.682623+00:00","source":"crossref"},{"observed_at":"2026-08-14T06:32:18.44784+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2503.10657","last_updated":"2025-05-20T14:59:21Z","snapshot_observed_at":"2026-08-13T09:25:43.676027Z","submitted_at":"2025-03-08T04:07:07Z","title":"RouterEval: A Comprehensive Benchmark for Routing LLMs to Explore Model-level Scaling Up in LLMs","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2503.10657","snapshot_observed_at":"2026-07-12T02:29:58.886380Z","title":"RouterEval: A comprehensive benchmark for routing LLMs to explore model-level scaling up in LLMs,","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2607.03436","last_updated":"2026-07-07T05:35:41Z","snapshot_observed_at":"2026-08-13T02:54:32.190736Z","submitted_at":"2026-07-03T15:49:40Z","title":"How Much of the Routing Gap Is Real? Decomposing the Router-to-Oracle Gap into Reproducible Specialist Advantage and Single-Draw Label Noise","version":2},"reference_index":20,"source":"pdf_text","source_observed_at":"2026-07-12T02:29:58.886380Z"},"links":{"cited_paper":"/paper/2503.10657","citing_paper":"/paper/2607.03436"},"observation_digest":"sha256:82af576a3185f6926d8a996b7af42374d00a9e6e994e6c8088620b199b9decdf","observation_id":"7e2939c7-f61e-4730-8b58-d7796b92fc78","resolution":{"observed_at":"2026-07-12T02:29:58.886380Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2503.10657","last_updated":"2025-05-20T14:59:21Z","snapshot_observed_at":"2026-08-13T09:25:43.676027Z","submitted_at":"2025-03-08T04:07:07Z","title":"RouterEval: A Comprehensive Benchmark for Routing LLMs to Explore Model-level Scaling Up in LLMs","version":2},"cited_work":{"arxiv_id":"2503.10657","doi":null,"metadata_source":"pith","pith_arxiv_id":"2503.10657","snapshot_observed_at":"2026-07-10T03:46:44.819123Z","title":"Routereval: A comprehensive benchmark for routing llms to explore model-level scaling up in llms.arXiv preprint arXiv:2503.10657","venue":"cs.CL","work_id":"7a1ee412-d726-449d-96d9-6253531a9aa6","year":2025},"citing_paper":{"arxiv_id":"2607.08665","last_updated":"2026-07-10T17:35:50Z","snapshot_observed_at":"2026-08-02T17:01:44.356729Z","submitted_at":"2026-07-09T16:34:15Z","title":"Resample or Reroute? Budget-Aware Test-Time Model Selection for Large Language Models","version":1},"reference_index":20,"source":"pdf_text","source_observed_at":"2026-07-10T03:42:04.605214Z"},"links":{"cited_paper":"/paper/2503.10657","citing_paper":"/paper/2607.08665"},"observation_digest":"sha256:2b9817792acc5b2e3fc039d9f441a8c0b62574f1546913be9c4ebd3473c9bc1a","observation_id":"111bb04e-7c64-47e8-88e0-5a68f44b1bb2","resolution":{"observed_at":"2026-07-10T03:46:44.820435Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-14T06:32:32.682623+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-14T06:32:32.682623+00:00","source":"crossref"},{"observed_at":"2026-08-14T06:32:18.44784+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2503.10657","last_updated":"2025-05-20T14:59:21Z","snapshot_observed_at":"2026-08-13T09:25:43.676027Z","submitted_at":"2025-03-08T04:07:07Z","title":"RouterEval: A Comprehensive Benchmark for Routing LLMs to Explore Model-level Scaling Up in LLMs","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2503.10657","snapshot_observed_at":"2026-07-13T06:32:41.000897Z","title":"RouterEval: A comprehensive benchmark for routing LLMs to explore model-level scaling up in LLMs,","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2607.08665","last_updated":"2026-07-10T17:35:50Z","snapshot_observed_at":"2026-08-02T17:01:44.356729Z","submitted_at":"2026-07-09T16:34:15Z","title":"Resample or Reroute? Budget-Aware Test-Time Model Selection for Large Language Models","version":2},"reference_index":20,"source":"pdf_text","source_observed_at":"2026-07-13T06:32:41.000897Z"},"links":{"cited_paper":"/paper/2503.10657","citing_paper":"/paper/2607.08665"},"observation_digest":"sha256:ef7f77bfb73addc1734989b60ebc2cf21978b4b1a49802170e7c11f05516d352","observation_id":"11c8c7dd-b43e-4970-9362-f215dfb2c6e7","resolution":{"observed_at":"2026-07-13T06:32:41.000897Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2503.10657","last_updated":"2025-05-20T14:59:21Z","snapshot_observed_at":"2026-08-13T09:25:43.676027Z","submitted_at":"2025-03-08T04:07:07Z","title":"RouterEval: A Comprehensive Benchmark for Routing LLMs to Explore Model-level Scaling Up in LLMs","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2503.10657","snapshot_observed_at":"2026-08-01T18:19:38.026339Z","title":null,"venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2607.17347","last_updated":"2026-07-19T17:17:47Z","snapshot_observed_at":"2026-08-14T09:52:35.049575Z","submitted_at":"2026-07-19T17:17:47Z","title":"Adapting Embedding Models for Agent Capability Retrieval","version":1},"reference_index":4,"source":"pdf_text","source_observed_at":"2026-08-01T18:19:38.026339Z"},"links":{"cited_paper":"/paper/2503.10657","citing_paper":"/paper/2607.17347"},"observation_digest":"sha256:c58999f0bf2f44c5a874533d31c873fd175c3fbab086e32791ffbc7dd4e92b9e","observation_id":"a3809a90-6dbd-4c46-a8ad-aa7e746c43cc","resolution":{"observed_at":"2026-08-01T18:19:38.026339Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2503.10657","last_updated":"2025-05-20T14:59:21Z","snapshot_observed_at":"2026-08-13T09:25:43.676027Z","submitted_at":"2025-03-08T04:07:07Z","title":"RouterEval: A Comprehensive Benchmark for Routing LLMs to Explore Model-level Scaling Up in LLMs","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2503.10657","snapshot_observed_at":"2026-08-01T13:10:58.965763Z","title":"RouterEval: A comprehensive benchmark for routing LLMs to explore model- level scaling up in LLMs.arXiv preprint arXiv:2503.10657, 2025","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2607.19215","last_updated":"2026-07-21T15:47:35Z","snapshot_observed_at":"2026-08-14T07:37:50.679789Z","submitted_at":"2026-07-21T15:47:35Z","title":"HACO: Hedged Agent Computing for Reliable LLM Systems","version":1},"reference_index":26,"source":"pdf_text","source_observed_at":"2026-08-01T13:10:58.965763Z"},"links":{"cited_paper":"/paper/2503.10657","citing_paper":"/paper/2607.19215"},"observation_digest":"sha256:b736a2fc238a1ba873ce91a1c6ba9e835cbb0eeb1ba62c5192d3011c79861b67","observation_id":"84fe4562-b79c-4a38-a339-a776b79ced88","resolution":{"observed_at":"2026-08-01T13:10:58.965763Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2503.10657","last_updated":"2025-05-20T14:59:21Z","snapshot_observed_at":"2026-08-13T09:25:43.676027Z","submitted_at":"2025-03-08T04:07:07Z","title":"RouterEval: A Comprehensive Benchmark for Routing LLMs to Explore Model-level Scaling Up in LLMs","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2503.10657","snapshot_observed_at":"2026-08-12T00:18:54.116775Z","title":"Routereval: A comprehensive benchmark for routing llms to explore model-level scaling up in llms.arXiv preprint arXiv:2503.10657,","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2608.08239","last_updated":"2026-08-08T17:07:11Z","snapshot_observed_at":"2026-08-14T03:23:31.907906Z","submitted_at":"2026-08-08T17:07:11Z","title":"The Replay Gap: Static Evaluation of Model Switching in LLM Agents Scores the Wrong World","version":1},"reference_index":5,"source":"pdf_text","source_observed_at":"2026-08-12T00:18:54.116775Z"},"links":{"cited_paper":"/paper/2503.10657","citing_paper":"/paper/2608.08239"},"observation_digest":"sha256:db004f8c5c9d1de03d8e0ba607b4346d2a40efb92ebc47557b6510156cd3234d","observation_id":"acd73c12-76f2-472b-96cb-804c6265671b","resolution":{"observed_at":"2026-08-12T00:18:54.116775Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}}],"links":{"evidence":"/evidence","html":"/paper/2503.10657/citation-record","integrity":"/paper/2503.10657/integrity","json":"/paper/2503.10657/citation-record.json","paper":"/paper/2503.10657"},"outbound":[],"paper":{"arxiv_id":"2503.10657","last_updated":"2025-05-20T14:59:21Z","latest_version":2,"primary_category":"cs.CL","snapshot_observed_at":"2026-08-13T09:25:43.676027Z","submitted_at":"2025-03-08T04:07:07Z","title":"RouterEval: A Comprehensive Benchmark for Routing LLMs to Explore Model-level Scaling Up in LLMs"},"reference_resolution":{"displayed":0,"state_counts":{"malformed_identifier":0,"metadata_mismatch":0,"parse_uncertain":0,"unresolved":0,"verified_exact":0,"verified_fuzzy":0},"total_outbound_references":0},"refusal":"A citation records a reference. It does not transfer a finding from one paper to another.","schema":"pith.paper-citation-record.v1","standing_sources":[{"observed_at":"2026-08-14T06:32:32.682623+00:00","source":"crossref"},{"observed_at":"2026-08-14T06:32:18.44784+00:00","source":"retraction_watch"}],"thesis":"As of 14 August 2026, this Paper Citation Record lists 0 of 0 outbound references and 18 inbound Pith citation observations for arXiv:2503.10657."}