{"as_of":"2026-08-06T19:29:00Z","caps":{"database_statements":6,"inbound":100,"outbound":100},"context_digest":"sha256:11d56435411713a62155a17086235e792a80d0a931a068976ec768ca89cc9737","coverage":[{"denominator":51,"lane":"reference_resolution","note":"Typed states for the displayed outbound observations.","records_observed":51,"source":"paper_references, paper_reference_links","source_observed_at":"2026-08-06T04:46:39.935439Z","state":"measured"},{"denominator":51,"lane":"standing_notices","note":"One-hop event checks from named stored sources.","records_observed":51,"source":"scholarly_work_events, retraction_status_cache","source_observed_at":"2026-08-06T06:34:29.942622+00:00","state":"measured"},{"denominator":0,"lane":"inbound_itemization","note":"Pith citing papers itemized under the disclosed page cap.","records_observed":0,"source":"paper_references, paper_reference_links","source_observed_at":null,"state":"measured"},{"denominator":1,"lane":"external_citation_measurements","note":"A source-named dated measurement, never combined with another source.","records_observed":0,"source":"cited_works","source_observed_at":null,"state":"measured"}],"external_citation_measurements":[],"inbound":[],"links":{"evidence":"/evidence","html":"/paper/2608.05126/citation-record","integrity":"/paper/2608.05126/integrity","json":"/paper/2608.05126/citation-record.json","paper":"/paper/2608.05126"},"outbound":[{"citation":{"cited_paper":{"arxiv_id":"2508.10925","last_updated":"2025-08-08T19:24:38Z","snapshot_observed_at":"2026-08-01T16:27:35.664983Z","submitted_at":"2025-08-08T19:24:38Z","title":"gpt-oss-120b & gpt-oss-20b Model Card","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2508.10925","snapshot_observed_at":"2026-08-06T04:46:37.340915Z","title":null,"venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2608.05126","last_updated":"2026-08-05T17:50:31Z","snapshot_observed_at":"2026-08-06T19:07:00.208170Z","submitted_at":"2026-08-05T17:50:31Z","title":"Spoken Function Calling: A New Perspective on Spoken Language Understanding for Large Audio Language Models","version":1},"reference_index":1,"source":"pdf_text","source_observed_at":"2026-08-06T04:46:37.340915Z"},"links":{"cited_paper":"/paper/2508.10925","citing_paper":"/paper/2608.05126"},"observation_digest":"sha256:3879077e8c331f9c50ec937162df4ad2e17df28661bd50e5d1000c56a8412900","observation_id":"68c884cd-11ee-401d-831b-cce310f596cc","resolution":{"observed_at":"2026-08-06T04:46:37.340915Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T04:46:40.722156Z","title":null,"venue":null,"work_id":"97ab27a3-391c-48c3-821c-e0efcc33009c","year":null},"citing_paper":{"arxiv_id":"2608.05126","last_updated":"2026-08-05T17:50:31Z","snapshot_observed_at":"2026-08-06T19:07:00.208170Z","submitted_at":"2026-08-05T17:50:31Z","title":"Spoken Function Calling: A New Perspective on Spoken Language Understanding for Large Audio Language Models","version":1},"reference_index":2,"source":"pdf_text","source_observed_at":"2026-08-06T04:46:37.391520Z"},"links":{"citing_paper":"/paper/2608.05126"},"observation_digest":"sha256:b4256edd5a24d09831f1ac9039efc508f6d467cd5d41962486dce1127a2e672a","observation_id":"c78e825d-a7a0-4139-99ce-b8470f8f0b78","resolution":{"observed_at":"2026-08-06T04:46:40.726054Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-06T06:34:29.942622+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-06T06:34:29.942622+00:00","source":"crossref"},{"observed_at":"2026-08-06T06:34:23.284952+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T04:46:40.698651Z","title":null,"venue":null,"work_id":"58e4d142-28db-4aa8-8ccc-a69798f39ab1","year":2020},"citing_paper":{"arxiv_id":"2608.05126","last_updated":"2026-08-05T17:50:31Z","snapshot_observed_at":"2026-08-06T19:07:00.208170Z","submitted_at":"2026-08-05T17:50:31Z","title":"Spoken Function Calling: A New Perspective on Spoken Language Understanding for Large Audio Language Models","version":1},"reference_index":3,"source":"pdf_text","source_observed_at":"2026-08-06T04:46:37.503546Z"},"links":{"citing_paper":"/paper/2608.05126"},"observation_digest":"sha256:c540c744073bc22dff1d3f399ed2ea96be928e1319c70f401b719cd80fa11c01","observation_id":"5557c7eb-9be7-4197-ab52-78f7429cb924","resolution":{"observed_at":"2026-08-06T04:46:40.702510Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-06T06:34:29.942622+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-06T06:34:29.942622+00:00","source":"crossref"},{"observed_at":"2026-08-06T06:34:23.284952+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T04:46:40.687259Z","title":null,"venue":null,"work_id":"1ee9c2dc-7acd-4f82-9908-2a291a14a889","year":2020},"citing_paper":{"arxiv_id":"2608.05126","last_updated":"2026-08-05T17:50:31Z","snapshot_observed_at":"2026-08-06T19:07:00.208170Z","submitted_at":"2026-08-05T17:50:31Z","title":"Spoken Function Calling: A New Perspective on Spoken Language Understanding for Large Audio Language Models","version":1},"reference_index":4,"source":"pdf_text","source_observed_at":"2026-08-06T04:46:37.547561Z"},"links":{"citing_paper":"/paper/2608.05126"},"observation_digest":"sha256:36b0b6b161bb244f805dcfef171ac0a75266f7c0972b1cd24c53d6f98881cbe9","observation_id":"e8b31969-e2e8-45b4-89ba-c050f76011d6","resolution":{"observed_at":"2026-08-06T04:46:40.690532Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-06T06:34:29.942622+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-06T06:34:29.942622+00:00","source":"crossref"},{"observed_at":"2026-08-06T06:34:23.284952+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T04:46:40.676184Z","title":null,"venue":null,"work_id":"c7b42394-bbf9-40ff-8d0b-7af5f0835b80","year":2018},"citing_paper":{"arxiv_id":"2608.05126","last_updated":"2026-08-05T17:50:31Z","snapshot_observed_at":"2026-08-06T19:07:00.208170Z","submitted_at":"2026-08-05T17:50:31Z","title":"Spoken Function Calling: A New Perspective on Spoken Language Understanding for Large Audio Language Models","version":1},"reference_index":5,"source":"pdf_text","source_observed_at":"2026-08-06T04:46:37.602812Z"},"links":{"citing_paper":"/paper/2608.05126"},"observation_digest":"sha256:00c97f20cc35e565d41d2f33a54c5731308334e5654c122f4cf53f738073cc91","observation_id":"458f25a5-a2fe-4a26-9a9a-bdffa223d439","resolution":{"observed_at":"2026-08-06T04:46:40.679876Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-06T06:34:29.942622+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-06T06:34:29.942622+00:00","source":"crossref"},{"observed_at":"2026-08-06T06:34:23.284952+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2107.03374","last_updated":"2021-07-14T17:16:02Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2021-07-07T17:41:24Z","title":"Evaluating Large Language Models Trained on Code","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2107.03374","snapshot_observed_at":"2026-08-06T04:46:37.672658Z","title":null,"venue":null,"work_id":null,"year":2021},"citing_paper":{"arxiv_id":"2608.05126","last_updated":"2026-08-05T17:50:31Z","snapshot_observed_at":"2026-08-06T19:07:00.208170Z","submitted_at":"2026-08-05T17:50:31Z","title":"Spoken Function Calling: A New Perspective on Spoken Language Understanding for Large Audio Language Models","version":1},"reference_index":6,"source":"pdf_text","source_observed_at":"2026-08-06T04:46:37.672658Z"},"links":{"cited_paper":"/paper/2107.03374","citing_paper":"/paper/2608.05126"},"observation_digest":"sha256:05404b2f20c6ed5d4fa54dce6393c24b05a653355b85329d50ee27884066e38d","observation_id":"31389c0c-8fc3-4a6f-9318-bd65047390be","resolution":{"observed_at":"2026-08-06T04:46:37.672658Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"1805.10190","last_updated":"2018-12-06T16:34:25Z","snapshot_observed_at":"2026-08-05T10:58:51.399386Z","submitted_at":"2018-05-25T15:04:17Z","title":"Snips Voice Platform: an embedded Spoken Language Understanding system for private-by-design voice interfaces","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"1805.10190","snapshot_observed_at":"2026-08-06T04:46:37.724269Z","title":null,"venue":null,"work_id":null,"year":2018},"citing_paper":{"arxiv_id":"2608.05126","last_updated":"2026-08-05T17:50:31Z","snapshot_observed_at":"2026-08-06T19:07:00.208170Z","submitted_at":"2026-08-05T17:50:31Z","title":"Spoken Function Calling: A New Perspective on Spoken Language Understanding for Large Audio Language Models","version":1},"reference_index":7,"source":"pdf_text","source_observed_at":"2026-08-06T04:46:37.724269Z"},"links":{"cited_paper":"/paper/1805.10190","citing_paper":"/paper/2608.05126"},"observation_digest":"sha256:6b980a834a6494a52e51309673479cd55f148f4df9499ad1e151a161771ae4c1","observation_id":"9eb6d3c9-32cc-44b4-b4f9-2331d15b4c02","resolution":{"observed_at":"2026-08-06T04:46:37.724269Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T04:46:37.788400Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2608.05126","last_updated":"2026-08-05T17:50:31Z","snapshot_observed_at":"2026-08-06T19:07:00.208170Z","submitted_at":"2026-08-05T17:50:31Z","title":"Spoken Function Calling: A New Perspective on Spoken Language Understanding for Large Audio Language Models","version":1},"reference_index":8,"source":"pdf_text","source_observed_at":"2026-08-06T04:46:37.788400Z"},"links":{"citing_paper":"/paper/2608.05126"},"observation_digest":"sha256:b9ff0183b1ea55d606fc66eec02d86c5e934d550a2cce0f293c47ac67808227a","observation_id":"7aff2a80-1dc0-49ed-a832-152d8b538bf1","resolution":{"observed_at":"2026-08-06T04:46:37.788400Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T04:46:40.658498Z","title":null,"venue":null,"work_id":"3ee55727-524a-402e-af0e-17d00d861eab","year":2023},"citing_paper":{"arxiv_id":"2608.05126","last_updated":"2026-08-05T17:50:31Z","snapshot_observed_at":"2026-08-06T19:07:00.208170Z","submitted_at":"2026-08-05T17:50:31Z","title":"Spoken Function Calling: A New Perspective on Spoken Language Understanding for Large Audio Language Models","version":1},"reference_index":9,"source":"pdf_text","source_observed_at":"2026-08-06T04:46:37.817177Z"},"links":{"citing_paper":"/paper/2608.05126"},"observation_digest":"sha256:178be7758f1a448edc79c54148bd0d141c89fb0dca96f855dc9f34e7d2c96630","observation_id":"63a6155e-84d0-4671-b203-71d75ddbcf5e","resolution":{"observed_at":"2026-08-06T04:46:40.661567Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-06T06:34:29.942622+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-06T06:34:29.942622+00:00","source":"crossref"},{"observed_at":"2026-08-06T06:34:23.284952+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T04:46:40.647670Z","title":null,"venue":null,"work_id":"88cfc1e0-0f57-4df1-ab78-5b010c2bfd7c","year":1990},"citing_paper":{"arxiv_id":"2608.05126","last_updated":"2026-08-05T17:50:31Z","snapshot_observed_at":"2026-08-06T19:07:00.208170Z","submitted_at":"2026-08-05T17:50:31Z","title":"Spoken Function Calling: A New Perspective on Spoken Language Understanding for Large Audio Language Models","version":1},"reference_index":10,"source":"pdf_text","source_observed_at":"2026-08-06T04:46:37.871444Z"},"links":{"citing_paper":"/paper/2608.05126"},"observation_digest":"sha256:09feb513ce3a9d6fce56edc607df83b0f53d8c8580fc0fe64e28ae6aeb69083c","observation_id":"f72e710f-893a-45f6-b226-431fd323f308","resolution":{"observed_at":"2026-08-06T04:46:40.650909Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-06T06:34:29.942622+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-06T06:34:29.942622+00:00","source":"crossref"},{"observed_at":"2026-08-06T06:34:23.284952+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2503.23278","last_updated":"2025-10-07T07:13:32Z","snapshot_observed_at":"2026-07-06T21:00:55.979837Z","submitted_at":"2025-03-30T01:58:22Z","title":"Model Context Protocol (MCP): Landscape, Security Threats, and Future Research Directions","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2503.23278","snapshot_observed_at":"2026-08-06T04:46:37.952256Z","title":null,"venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2608.05126","last_updated":"2026-08-05T17:50:31Z","snapshot_observed_at":"2026-08-06T19:07:00.208170Z","submitted_at":"2026-08-05T17:50:31Z","title":"Spoken Function Calling: A New Perspective on Spoken Language Understanding for Large Audio Language Models","version":1},"reference_index":11,"source":"pdf_text","source_observed_at":"2026-08-06T04:46:37.952256Z"},"links":{"cited_paper":"/paper/2503.23278","citing_paper":"/paper/2608.05126"},"observation_digest":"sha256:f3a0b9c078682952c9593f3f384de9c920690c6c799140267a69d68cc39c2325","observation_id":"01bf1bb0-f601-4b2b-b3cb-9d487837f49a","resolution":{"observed_at":"2026-08-06T04:46:37.952256Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T04:46:40.636493Z","title":null,"venue":null,"work_id":"5546697f-f219-4290-a40a-720488c750ba","year":2022},"citing_paper":{"arxiv_id":"2608.05126","last_updated":"2026-08-05T17:50:31Z","snapshot_observed_at":"2026-08-06T19:07:00.208170Z","submitted_at":"2026-08-05T17:50:31Z","title":"Spoken Function Calling: A New Perspective on Spoken Language Understanding for Large Audio Language Models","version":1},"reference_index":12,"source":"pdf_text","source_observed_at":"2026-08-06T04:46:38.002010Z"},"links":{"citing_paper":"/paper/2608.05126"},"observation_digest":"sha256:9b572ef290cbff3e2158bde20977cf1d7b73d3853f42eb91d8537d10e50c8c27","observation_id":"ba28f402-8d13-4679-9981-a29766de2a49","resolution":{"observed_at":"2026-08-06T04:46:40.640326Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-06T06:34:29.942622+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-06T06:34:29.942622+00:00","source":"crossref"},{"observed_at":"2026-08-06T06:34:23.284952+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T04:46:38.049617Z","title":null,"venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2608.05126","last_updated":"2026-08-05T17:50:31Z","snapshot_observed_at":"2026-08-06T19:07:00.208170Z","submitted_at":"2026-08-05T17:50:31Z","title":"Spoken Function Calling: A New Perspective on Spoken Language Understanding for Large Audio Language Models","version":1},"reference_index":13,"source":"pdf_text","source_observed_at":"2026-08-06T04:46:38.049617Z"},"links":{"citing_paper":"/paper/2608.05126"},"observation_digest":"sha256:2db5ac70751ec26a49a52768a5ddc95a0a07fb3651c58f90d52afa6dd39fb7de","observation_id":"deabefdf-803c-4a8a-8a7e-8955f1363987","resolution":{"observed_at":"2026-08-06T04:46:38.049617Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2001.08361","last_updated":"2020-01-23T03:59:20Z","snapshot_observed_at":"2026-07-06T08:52:12.656082Z","submitted_at":"2020-01-23T03:59:20Z","title":"Scaling Laws for Neural Language Models","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2001.08361","snapshot_observed_at":"2026-08-06T04:46:38.170089Z","title":null,"venue":null,"work_id":null,"year":2020},"citing_paper":{"arxiv_id":"2608.05126","last_updated":"2026-08-05T17:50:31Z","snapshot_observed_at":"2026-08-06T19:07:00.208170Z","submitted_at":"2026-08-05T17:50:31Z","title":"Spoken Function Calling: A New Perspective on Spoken Language Understanding for Large Audio Language Models","version":1},"reference_index":14,"source":"pdf_text","source_observed_at":"2026-08-06T04:46:38.170089Z"},"links":{"cited_paper":"/paper/2001.08361","citing_paper":"/paper/2608.05126"},"observation_digest":"sha256:749461f3cd32fdbf6f8293a55a6668a9ca543f799f7f3caf3fa247d39dc58ec9","observation_id":"85d553fc-8658-4e88-8c81-93b61a9d1dad","resolution":{"observed_at":"2026-08-06T04:46:38.170089Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T04:46:40.610153Z","title":null,"venue":null,"work_id":"380ca5fc-14bc-403b-9c45-3c7d08ada54d","year":2023},"citing_paper":{"arxiv_id":"2608.05126","last_updated":"2026-08-05T17:50:31Z","snapshot_observed_at":"2026-08-06T19:07:00.208170Z","submitted_at":"2026-08-05T17:50:31Z","title":"Spoken Function Calling: A New Perspective on Spoken Language Understanding for Large Audio Language Models","version":1},"reference_index":15,"source":"pdf_text","source_observed_at":"2026-08-06T04:46:38.277118Z"},"links":{"citing_paper":"/paper/2608.05126"},"observation_digest":"sha256:522131cacb7e4f3e2ae534ef305f97c54adca65c485157fbc95fbb5f52c7a0a3","observation_id":"f8bafe28-69c9-4361-b7df-a6508fe6abce","resolution":{"observed_at":"2026-08-06T04:46:40.613198Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-06T06:34:29.942622+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-06T06:34:29.942622+00:00","source":"crossref"},{"observed_at":"2026-08-06T06:34:23.284952+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T04:46:40.600362Z","title":null,"venue":null,"work_id":"9a714657-de7d-4a78-84c4-40f585a271fd","year":2021},"citing_paper":{"arxiv_id":"2608.05126","last_updated":"2026-08-05T17:50:31Z","snapshot_observed_at":"2026-08-06T19:07:00.208170Z","submitted_at":"2026-08-05T17:50:31Z","title":"Spoken Function Calling: A New Perspective on Spoken Language Understanding for Large Audio Language Models","version":1},"reference_index":16,"source":"pdf_text","source_observed_at":"2026-08-06T04:46:38.387114Z"},"links":{"citing_paper":"/paper/2608.05126"},"observation_digest":"sha256:8db69a7e7cabf8c0c0c77367229d0383036d7082d04981f87b8d6d7776381d5b","observation_id":"5f0054d1-655c-4b56-aee4-7bdad37c3f57","resolution":{"observed_at":"2026-08-06T04:46:40.603627Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-06T06:34:29.942622+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-06T06:34:29.942622+00:00","source":"crossref"},{"observed_at":"2026-08-06T06:34:23.284952+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T04:46:40.590271Z","title":null,"venue":null,"work_id":"120196b9-0611-445f-bac8-4398599117c9","year":2024},"citing_paper":{"arxiv_id":"2608.05126","last_updated":"2026-08-05T17:50:31Z","snapshot_observed_at":"2026-08-06T19:07:00.208170Z","submitted_at":"2026-08-05T17:50:31Z","title":"Spoken Function Calling: A New Perspective on Spoken Language Understanding for Large Audio Language Models","version":1},"reference_index":17,"source":"pdf_text","source_observed_at":"2026-08-06T04:46:38.460511Z"},"links":{"citing_paper":"/paper/2608.05126"},"observation_digest":"sha256:fd453dc2588971f52f14f8f9576090c9b6c59813fc0f7b73fa4b556b02ce0560","observation_id":"06dc2690-b465-44e3-bf99-edc1c85d32da","resolution":{"observed_at":"2026-08-06T04:46:40.593465Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-06T06:34:29.942622+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-06T06:34:29.942622+00:00","source":"crossref"},{"observed_at":"2026-08-06T06:34:23.284952+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2304.08244","last_updated":"2023-10-25T06:54:12Z","snapshot_observed_at":"2026-08-02T00:07:12.855748Z","submitted_at":"2023-04-14T14:05:32Z","title":"API-Bank: A Comprehensive Benchmark for Tool-Augmented LLMs","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2304.08244","snapshot_observed_at":"2026-08-06T04:46:38.571002Z","title":null,"venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2608.05126","last_updated":"2026-08-05T17:50:31Z","snapshot_observed_at":"2026-08-06T19:07:00.208170Z","submitted_at":"2026-08-05T17:50:31Z","title":"Spoken Function Calling: A New Perspective on Spoken Language Understanding for Large Audio Language Models","version":1},"reference_index":18,"source":"pdf_text","source_observed_at":"2026-08-06T04:46:38.571002Z"},"links":{"cited_paper":"/paper/2304.08244","citing_paper":"/paper/2608.05126"},"observation_digest":"sha256:98455031479d712b08b066b2741410e4f0693bf84a2158d7cccb46cbf99f47d7","observation_id":"df2f28b3-d29a-4b28-a0b7-389c682f1a60","resolution":{"observed_at":"2026-08-06T04:46:38.571002Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2409.00920","last_updated":"2025-07-25T08:26:54Z","snapshot_observed_at":"2026-07-06T19:09:01.748727Z","submitted_at":"2024-09-02T03:19:56Z","title":"ToolACE: Winning the Points of LLM Function Calling","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2409.00920","snapshot_observed_at":"2026-08-06T04:46:38.644903Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2608.05126","last_updated":"2026-08-05T17:50:31Z","snapshot_observed_at":"2026-08-06T19:07:00.208170Z","submitted_at":"2026-08-05T17:50:31Z","title":"Spoken Function Calling: A New Perspective on Spoken Language Understanding for Large Audio Language Models","version":1},"reference_index":19,"source":"pdf_text","source_observed_at":"2026-08-06T04:46:38.644903Z"},"links":{"cited_paper":"/paper/2409.00920","citing_paper":"/paper/2608.05126"},"observation_digest":"sha256:85dcc308dd22e3fa4362c5ae6c14ec6d3126786c409a2da9a58080a3cf7aaac9","observation_id":"ac7b5666-6aea-4533-827d-6543ad4e0d3c","resolution":{"observed_at":"2026-08-06T04:46:38.644903Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T04:46:40.579940Z","title":null,"venue":null,"work_id":"b3893602-e668-456b-9d62-bca018de7a33","year":2024},"citing_paper":{"arxiv_id":"2608.05126","last_updated":"2026-08-05T17:50:31Z","snapshot_observed_at":"2026-08-06T19:07:00.208170Z","submitted_at":"2026-08-05T17:50:31Z","title":"Spoken Function Calling: A New Perspective on Spoken Language Understanding for Large Audio Language Models","version":1},"reference_index":20,"source":"pdf_text","source_observed_at":"2026-08-06T04:46:38.761255Z"},"links":{"citing_paper":"/paper/2608.05126"},"observation_digest":"sha256:a8682a057eede766622e1327f40fb864b6f2d399f266dc98d3c10b8a4819e5e7","observation_id":"eef6f6c9-0c87-41d7-b02b-8a75fc1e0051","resolution":{"observed_at":"2026-08-06T04:46:40.582959Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-06T06:34:29.942622+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-06T06:34:29.942622+00:00","source":"crossref"},{"observed_at":"2026-08-06T06:34:23.284952+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T04:46:40.569927Z","title":null,"venue":null,"work_id":"a99d8251-3a0b-4e93-ae3d-a4eb1aaaaa8b","year":2019},"citing_paper":{"arxiv_id":"2608.05126","last_updated":"2026-08-05T17:50:31Z","snapshot_observed_at":"2026-08-06T19:07:00.208170Z","submitted_at":"2026-08-05T17:50:31Z","title":"Spoken Function Calling: A New Perspective on Spoken Language Understanding for Large Audio Language Models","version":1},"reference_index":21,"source":"pdf_text","source_observed_at":"2026-08-06T04:46:38.878716Z"},"links":{"citing_paper":"/paper/2608.05126"},"observation_digest":"sha256:5e52da18d05b31847fbe19f1259d11f5c36458cc84022743a724d94ae313bb16","observation_id":"25b557ab-da37-4e4b-abe0-4db7e41d1184","resolution":{"observed_at":"2026-08-06T04:46:40.573344Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-06T06:34:29.942622+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-06T06:34:29.942622+00:00","source":"crossref"},{"observed_at":"2026-08-06T06:34:23.284952+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T04:46:40.559557Z","title":null,"venue":null,"work_id":"082b12ae-1380-4471-8fcf-14c369552a83","year":2015},"citing_paper":{"arxiv_id":"2608.05126","last_updated":"2026-08-05T17:50:31Z","snapshot_observed_at":"2026-08-06T19:07:00.208170Z","submitted_at":"2026-08-05T17:50:31Z","title":"Spoken Function Calling: A New Perspective on Spoken Language Understanding for Large Audio Language Models","version":1},"reference_index":22,"source":"pdf_text","source_observed_at":"2026-08-06T04:46:38.938967Z"},"links":{"citing_paper":"/paper/2608.05126"},"observation_digest":"sha256:a39c5d981b6f20da5a21205dce039d600c5fc91e20c23f05f62e78899daa6749","observation_id":"ca227e6e-6959-4e1b-8662-acb94a536bfe","resolution":{"observed_at":"2026-08-06T04:46:40.562940Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-06T06:34:29.942622+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-06T06:34:29.942622+00:00","source":"crossref"},{"observed_at":"2026-08-06T06:34:23.284952+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2303.09014","last_updated":"2023-03-16T01:04:45Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2023-03-16T01:04:45Z","title":"ART: Automatic multi-step reasoning and tool-use for large language models","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2303.09014","snapshot_observed_at":"2026-08-06T04:46:39.053319Z","title":null,"venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2608.05126","last_updated":"2026-08-05T17:50:31Z","snapshot_observed_at":"2026-08-06T19:07:00.208170Z","submitted_at":"2026-08-05T17:50:31Z","title":"Spoken Function Calling: A New Perspective on Spoken Language Understanding for Large Audio Language Models","version":1},"reference_index":23,"source":"pdf_text","source_observed_at":"2026-08-06T04:46:39.053319Z"},"links":{"cited_paper":"/paper/2303.09014","citing_paper":"/paper/2608.05126"},"observation_digest":"sha256:307bdc0b6ea2daf230a5a7421368ee2c27e8d0b028a000bb2a693fa72a5efda9","observation_id":"cd1d388a-d0b1-4ece-8de4-5e40215d3b49","resolution":{"observed_at":"2026-08-06T04:46:39.053319Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T04:46:40.549078Z","title":null,"venue":null,"work_id":"98adf8ca-53ff-459c-83fb-dfe000e9cbc4","year":2024},"citing_paper":{"arxiv_id":"2608.05126","last_updated":"2026-08-05T17:50:31Z","snapshot_observed_at":"2026-08-06T19:07:00.208170Z","submitted_at":"2026-08-05T17:50:31Z","title":"Spoken Function Calling: A New Perspective on Spoken Language Understanding for Large Audio Language Models","version":1},"reference_index":24,"source":"pdf_text","source_observed_at":"2026-08-06T04:46:39.100248Z"},"links":{"citing_paper":"/paper/2608.05126"},"observation_digest":"sha256:8b47eabc77d3c8a8aec5ac7f8714acb0f1a09dfd130059c6798145282141d77d","observation_id":"82474a83-d490-409c-aea6-05461f7cbde9","resolution":{"observed_at":"2026-08-06T04:46:40.552379Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-06T06:34:29.942622+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-06T06:34:29.942622+00:00","source":"crossref"},{"observed_at":"2026-08-06T06:34:23.284952+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T04:46:39.202981Z","title":null,"venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2608.05126","last_updated":"2026-08-05T17:50:31Z","snapshot_observed_at":"2026-08-06T19:07:00.208170Z","submitted_at":"2026-08-05T17:50:31Z","title":"Spoken Function Calling: A New Perspective on Spoken Language Understanding for Large Audio Language Models","version":1},"reference_index":25,"source":"pdf_text","source_observed_at":"2026-08-06T04:46:39.202981Z"},"links":{"citing_paper":"/paper/2608.05126"},"observation_digest":"sha256:5932d5d52ca07ba3bd82b5e30f940b141f2204b7887d5532532a405c5e774fd1","observation_id":"83355c5f-d210-4a60-a2be-6c58ea51fe01","resolution":{"observed_at":"2026-08-06T04:46:39.202981Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T04:46:39.328263Z","title":null,"venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2608.05126","last_updated":"2026-08-05T17:50:31Z","snapshot_observed_at":"2026-08-06T19:07:00.208170Z","submitted_at":"2026-08-05T17:50:31Z","title":"Spoken Function Calling: A New Perspective on Spoken Language Understanding for Large Audio Language Models","version":1},"reference_index":26,"source":"pdf_text","source_observed_at":"2026-08-06T04:46:39.328263Z"},"links":{"citing_paper":"/paper/2608.05126"},"observation_digest":"sha256:c549c6158848a3b87730e5db57b29e99533ec31dfebf2000c6957013d0c81cbf","observation_id":"4ca49e6a-6d8f-452d-83d9-cc0f03268979","resolution":{"observed_at":"2026-08-06T04:46:39.328263Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T04:46:40.537949Z","title":null,"venue":null,"work_id":"79b05501-504a-492d-8e51-e922f00d359a","year":2025},"citing_paper":{"arxiv_id":"2608.05126","last_updated":"2026-08-05T17:50:31Z","snapshot_observed_at":"2026-08-06T19:07:00.208170Z","submitted_at":"2026-08-05T17:50:31Z","title":"Spoken Function Calling: A New Perspective on Spoken Language Understanding for Large Audio Language Models","version":1},"reference_index":27,"source":"pdf_text","source_observed_at":"2026-08-06T04:46:39.437359Z"},"links":{"citing_paper":"/paper/2608.05126"},"observation_digest":"sha256:b0541114e04696713df9af9fb196ffffb713e352f3142ad24543a8bf5669ca44","observation_id":"e6fade8d-2742-498f-bac1-7c7015e0904e","resolution":{"observed_at":"2026-08-06T04:46:40.541299Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-06T06:34:29.942622+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-06T06:34:29.942622+00:00","source":"crossref"},{"observed_at":"2026-08-06T06:34:23.284952+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T04:46:40.527528Z","title":null,"venue":null,"work_id":"19e8431b-7db0-48e9-afce-ea8bae3c7e79","year":null},"citing_paper":{"arxiv_id":"2608.05126","last_updated":"2026-08-05T17:50:31Z","snapshot_observed_at":"2026-08-06T19:07:00.208170Z","submitted_at":"2026-08-05T17:50:31Z","title":"Spoken Function Calling: A New Perspective on Spoken Language Understanding for Large Audio Language Models","version":1},"reference_index":28,"source":"pdf_text","source_observed_at":"2026-08-06T04:46:39.503875Z"},"links":{"citing_paper":"/paper/2608.05126"},"observation_digest":"sha256:54f609c3a41fa9f91da59c55bd26bbe51e20527c3fc29ed671712c92de3c8185","observation_id":"756a5b05-f1cc-4f11-b4fa-ed34bf3a8932","resolution":{"observed_at":"2026-08-06T04:46:40.530752Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-06T06:34:29.942622+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-06T06:34:29.942622+00:00","source":"crossref"},{"observed_at":"2026-08-06T06:34:23.284952+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2004.10087","last_updated":"2020-10-17T04:28:29Z","snapshot_observed_at":"2026-08-04T15:04:40.800903Z","submitted_at":"2020-04-21T15:07:34Z","title":"AGIF: An Adaptive Graph-Interactive Framework for Joint Multiple Intent Detection and Slot Filling","version":4},"cited_work":{"arxiv_id":"2004.10087","doi":null,"metadata_source":"pith","pith_arxiv_id":"2004.10087","snapshot_observed_at":"2026-08-06T04:46:40.169869Z","title":"AGIF: An Adaptive Graph-Interactive Framework for Joint Multiple Intent Detection and Slot Filling","venue":"cs.CL","work_id":"0f62fb7b-3efd-480c-8926-4912808b55ac","year":2020},"citing_paper":{"arxiv_id":"2608.05126","last_updated":"2026-08-05T17:50:31Z","snapshot_observed_at":"2026-08-06T19:07:00.208170Z","submitted_at":"2026-08-05T17:50:31Z","title":"Spoken Function Calling: A New Perspective on Spoken Language Understanding for Large Audio Language Models","version":1},"reference_index":29,"source":"pdf_text","source_observed_at":"2026-08-06T04:46:39.666283Z"},"links":{"cited_paper":"/paper/2004.10087","citing_paper":"/paper/2608.05126"},"observation_digest":"sha256:c3d27a54969f1ed356b0f6c234f6bb78037ed997a85898633a5e03e4ca52b896","observation_id":"9fb7097d-4845-4a57-9791-e4eb8693a43a","resolution":{"observed_at":"2026-08-06T04:46:40.175397Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-06T06:34:29.942622+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-06T06:34:29.942622+00:00","source":"crossref"},{"observed_at":"2026-08-06T06:34:23.284952+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2307.16789","last_updated":"2023-10-03T14:45:48Z","snapshot_observed_at":"2026-07-06T16:00:46.542753Z","submitted_at":"2023-07-31T15:56:53Z","title":"ToolLLM: Facilitating Large Language Models to Master 16000+ Real-world APIs","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2307.16789","snapshot_observed_at":"2026-08-06T04:46:39.793942Z","title":null,"venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2608.05126","last_updated":"2026-08-05T17:50:31Z","snapshot_observed_at":"2026-08-06T19:07:00.208170Z","submitted_at":"2026-08-05T17:50:31Z","title":"Spoken Function Calling: A New Perspective on Spoken Language Understanding for Large Audio Language Models","version":1},"reference_index":30,"source":"pdf_text","source_observed_at":"2026-08-06T04:46:39.793942Z"},"links":{"cited_paper":"/paper/2307.16789","citing_paper":"/paper/2608.05126"},"observation_digest":"sha256:dcec1b5f4429e9e0c4f359de47447e7e698fe490ae0f448b654a6d194585c1e2","observation_id":"039249dd-a5e5-4468-b96b-d0cd39d38d43","resolution":{"observed_at":"2026-08-06T04:46:39.793942Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T04:46:40.504875Z","title":null,"venue":null,"work_id":"f18f0b19-b69a-411b-b4bf-c096a656f681","year":2023},"citing_paper":{"arxiv_id":"2608.05126","last_updated":"2026-08-05T17:50:31Z","snapshot_observed_at":"2026-08-06T19:07:00.208170Z","submitted_at":"2026-08-05T17:50:31Z","title":"Spoken Function Calling: A New Perspective on Spoken Language Understanding for Large Audio Language Models","version":1},"reference_index":31,"source":"pdf_text","source_observed_at":"2026-08-06T04:46:39.880243Z"},"links":{"citing_paper":"/paper/2608.05126"},"observation_digest":"sha256:75ebf2fc67ad2868fb9222ba60f086a96c20fb761d37c0312f3d91c55310bb2b","observation_id":"34945ba1-9765-4366-9a04-f1b063773316","resolution":{"observed_at":"2026-08-06T04:46:40.509063Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-06T06:34:29.942622+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-06T06:34:29.942622+00:00","source":"crossref"},{"observed_at":"2026-08-06T06:34:23.284952+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T04:46:40.494526Z","title":null,"venue":null,"work_id":"16df8829-3ebc-48bc-a2d5-10315a613278","year":2023},"citing_paper":{"arxiv_id":"2608.05126","last_updated":"2026-08-05T17:50:31Z","snapshot_observed_at":"2026-08-06T19:07:00.208170Z","submitted_at":"2026-08-05T17:50:31Z","title":"Spoken Function Calling: A New Perspective on Spoken Language Understanding for Large Audio Language Models","version":1},"reference_index":32,"source":"pdf_text","source_observed_at":"2026-08-06T04:46:39.883762Z"},"links":{"citing_paper":"/paper/2608.05126"},"observation_digest":"sha256:02d91397c40690757ff152d14cd0dcab9cd2f3b1a2b405d8e42dde4f865550d8","observation_id":"92e6dff4-ab88-4901-8457-2027f0f3ed01","resolution":{"observed_at":"2026-08-06T04:46:40.497702Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-06T06:34:29.942622+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-06T06:34:29.942622+00:00","source":"crossref"},{"observed_at":"2026-08-06T06:34:23.284952+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T04:46:39.887216Z","title":null,"venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2608.05126","last_updated":"2026-08-05T17:50:31Z","snapshot_observed_at":"2026-08-06T19:07:00.208170Z","submitted_at":"2026-08-05T17:50:31Z","title":"Spoken Function Calling: A New Perspective on Spoken Language Understanding for Large Audio Language Models","version":1},"reference_index":33,"source":"pdf_text","source_observed_at":"2026-08-06T04:46:39.887216Z"},"links":{"citing_paper":"/paper/2608.05126"},"observation_digest":"sha256:5e5c50719aab1b54acc7a4db9256d82f9fa7c8c3bf660bf397f1ac34b093d7bd","observation_id":"b0cdd6c2-f0e2-4df0-baaf-7344be6afaca","resolution":{"observed_at":"2026-08-06T04:46:39.887216Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2402.03300","last_updated":"2024-04-27T15:25:53Z","snapshot_observed_at":"2026-08-06T14:58:42.911363Z","submitted_at":"2024-02-05T18:55:32Z","title":"DeepSeekMath: Pushing the Limits of Mathematical Reasoning in Open Language Models","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2402.03300","snapshot_observed_at":"2026-08-06T04:46:39.893471Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2608.05126","last_updated":"2026-08-05T17:50:31Z","snapshot_observed_at":"2026-08-06T19:07:00.208170Z","submitted_at":"2026-08-05T17:50:31Z","title":"Spoken Function Calling: A New Perspective on Spoken Language Understanding for Large Audio Language Models","version":1},"reference_index":34,"source":"pdf_text","source_observed_at":"2026-08-06T04:46:39.893471Z"},"links":{"cited_paper":"/paper/2402.03300","citing_paper":"/paper/2608.05126"},"observation_digest":"sha256:0e522179fb283d7dc179cabfcb79a2807916bef917856d284639a383c6267f08","observation_id":"80f980d4-fe3a-4fcf-be54-cb1ce85a2083","resolution":{"observed_at":"2026-08-06T04:46:39.893471Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T04:46:39.896657Z","title":"2011.Spoken language understanding: Systems for extracting semantic information from speech","venue":null,"work_id":null,"year":2011},"citing_paper":{"arxiv_id":"2608.05126","last_updated":"2026-08-05T17:50:31Z","snapshot_observed_at":"2026-08-06T19:07:00.208170Z","submitted_at":"2026-08-05T17:50:31Z","title":"Spoken Function Calling: A New Perspective on Spoken Language Understanding for Large Audio Language Models","version":1},"reference_index":35,"source":"pdf_text","source_observed_at":"2026-08-06T04:46:39.896657Z"},"links":{"citing_paper":"/paper/2608.05126"},"observation_digest":"sha256:0fb0c105cc3a430214ea3f4e85bf6f0fc0336e18965e7a07f9c23beecd9f1667","observation_id":"4989de87-f25e-4662-9af8-021b05412024","resolution":{"observed_at":"2026-08-06T04:46:39.896657Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2506.04779","last_updated":"2026-03-16T09:24:19Z","snapshot_observed_at":"2026-08-02T22:58:18.776196Z","submitted_at":"2025-06-05T09:09:36Z","title":"MMSU: A Massive Multi-task Spoken Language Understanding and Reasoning Benchmark","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2506.04779","snapshot_observed_at":"2026-08-06T04:46:39.899707Z","title":null,"venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2608.05126","last_updated":"2026-08-05T17:50:31Z","snapshot_observed_at":"2026-08-06T19:07:00.208170Z","submitted_at":"2026-08-05T17:50:31Z","title":"Spoken Function Calling: A New Perspective on Spoken Language Understanding for Large Audio Language Models","version":1},"reference_index":36,"source":"pdf_text","source_observed_at":"2026-08-06T04:46:39.899707Z"},"links":{"cited_paper":"/paper/2506.04779","citing_paper":"/paper/2608.05126"},"observation_digest":"sha256:70e397294a9748e5b527386b2c51b6bb2a336fddb91c8bfa23c431c4f629bdb5","observation_id":"a6c5b19d-4c99-4a95-bd34-3365cf908893","resolution":{"observed_at":"2026-08-06T04:46:39.899707Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T04:46:40.469145Z","title":null,"venue":null,"work_id":"d8f1239e-86ab-432c-922a-71dd66cc75c9","year":2022},"citing_paper":{"arxiv_id":"2608.05126","last_updated":"2026-08-05T17:50:31Z","snapshot_observed_at":"2026-08-06T19:07:00.208170Z","submitted_at":"2026-08-05T17:50:31Z","title":"Spoken Function Calling: A New Perspective on Spoken Language Understanding for Large Audio Language Models","version":1},"reference_index":37,"source":"pdf_text","source_observed_at":"2026-08-06T04:46:39.903078Z"},"links":{"citing_paper":"/paper/2608.05126"},"observation_digest":"sha256:e8d445b3e9acce5a61782b41c62af88a72254017bc7946e3e215ea1322f4b9dc","observation_id":"35373006-b9fe-4226-b346-c936ca36fafc","resolution":{"observed_at":"2026-08-06T04:46:40.472812Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-06T06:34:29.942622+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-06T06:34:29.942622+00:00","source":"crossref"},{"observed_at":"2026-08-06T06:34:23.284952+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T04:46:40.457731Z","title":"Williams, Antoine Raux, and Matthew Henderson","venue":null,"work_id":"9d0551dc-1e4c-413e-829a-b60b3252fa15","year":2016},"citing_paper":{"arxiv_id":"2608.05126","last_updated":"2026-08-05T17:50:31Z","snapshot_observed_at":"2026-08-06T19:07:00.208170Z","submitted_at":"2026-08-05T17:50:31Z","title":"Spoken Function Calling: A New Perspective on Spoken Language Understanding for Large Audio Language Models","version":1},"reference_index":38,"source":"pdf_text","source_observed_at":"2026-08-06T04:46:39.906076Z"},"links":{"citing_paper":"/paper/2608.05126"},"observation_digest":"sha256:3d50a20016af2123f90921d9215db9d9377405858985ca44b3bad4ba4aa55d25","observation_id":"1d1b386c-f2cf-40d2-bc3e-d7c788abf714","resolution":{"observed_at":"2026-08-06T04:46:40.461459Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-06T06:34:29.942622+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-06T06:34:29.942622+00:00","source":"crossref"},{"observed_at":"2026-08-06T06:34:23.284952+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T04:46:40.445521Z","title":"Williams, Antoine Raux, Deepak Ramachandran, and Alan W","venue":null,"work_id":"6bc76f2b-3e04-4a49-b96f-172aaeacb1b9","year":2013},"citing_paper":{"arxiv_id":"2608.05126","last_updated":"2026-08-05T17:50:31Z","snapshot_observed_at":"2026-08-06T19:07:00.208170Z","submitted_at":"2026-08-05T17:50:31Z","title":"Spoken Function Calling: A New Perspective on Spoken Language Understanding for Large Audio Language Models","version":1},"reference_index":39,"source":"pdf_text","source_observed_at":"2026-08-06T04:46:39.909060Z"},"links":{"citing_paper":"/paper/2608.05126"},"observation_digest":"sha256:4f74d5c90338e0b20521595c395b4a3a4dad90e05ae212af9c89867a9b1548a0","observation_id":"8a7d6943-0bbd-43db-8a55-b7d8e772c7d7","resolution":{"observed_at":"2026-08-06T04:46:40.450010Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-06T06:34:29.942622+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-06T06:34:29.942622+00:00","source":"crossref"},{"observed_at":"2026-08-06T06:34:23.284952+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2507.16632","last_updated":"2025-08-27T16:42:11Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2025-07-22T14:23:55Z","title":"Step-Audio 2 Technical Report","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2507.16632","snapshot_observed_at":"2026-08-06T04:46:39.912304Z","title":null,"venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2608.05126","last_updated":"2026-08-05T17:50:31Z","snapshot_observed_at":"2026-08-06T19:07:00.208170Z","submitted_at":"2026-08-05T17:50:31Z","title":"Spoken Function Calling: A New Perspective on Spoken Language Understanding for Large Audio Language Models","version":1},"reference_index":40,"source":"pdf_text","source_observed_at":"2026-08-06T04:46:39.912304Z"},"links":{"cited_paper":"/paper/2507.16632","citing_paper":"/paper/2608.05126"},"observation_digest":"sha256:0c0f8b2688f0666ec168b99f954e79cade8d72b3494a3afd45b5dde593ea731b","observation_id":"7881e317-5737-4e76-b43e-7829ea46fb42","resolution":{"observed_at":"2026-08-06T04:46:39.912304Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2509.17765","last_updated":"2025-09-22T13:26:24Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2025-09-22T13:26:24Z","title":"Qwen3-Omni Technical Report","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2509.17765","snapshot_observed_at":"2026-08-06T04:46:39.915594Z","title":null,"venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2608.05126","last_updated":"2026-08-05T17:50:31Z","snapshot_observed_at":"2026-08-06T19:07:00.208170Z","submitted_at":"2026-08-05T17:50:31Z","title":"Spoken Function Calling: A New Perspective on Spoken Language Understanding for Large Audio Language Models","version":1},"reference_index":41,"source":"pdf_text","source_observed_at":"2026-08-06T04:46:39.915594Z"},"links":{"cited_paper":"/paper/2509.17765","citing_paper":"/paper/2608.05126"},"observation_digest":"sha256:4223ff76fcbcae7270236e8bd59697894d0d992160f5b67fb9f87ca04e6ea2a1","observation_id":"73abdf14-ba0f-4a94-a564-53744e0e8c6a","resolution":{"observed_at":"2026-08-06T04:46:39.915594Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2505.09388","last_updated":"2025-05-14T13:41:34Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2025-05-14T13:41:34Z","title":"Qwen3 Technical Report","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2505.09388","snapshot_observed_at":"2026-08-06T04:46:39.919197Z","title":null,"venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2608.05126","last_updated":"2026-08-05T17:50:31Z","snapshot_observed_at":"2026-08-06T19:07:00.208170Z","submitted_at":"2026-08-05T17:50:31Z","title":"Spoken Function Calling: A New Perspective on Spoken Language Understanding for Large Audio Language Models","version":1},"reference_index":42,"source":"pdf_text","source_observed_at":"2026-08-06T04:46:39.919197Z"},"links":{"cited_paper":"/paper/2505.09388","citing_paper":"/paper/2608.05126"},"observation_digest":"sha256:1c5c02ba3dea991a4c8d9c2e5a4c4f306dc9f6da913f2d657da509caea4ee4c1","observation_id":"7970618c-360d-4126-b987-5948f065eb9f","resolution":{"observed_at":"2026-08-06T04:46:39.919197Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T04:46:40.434830Z","title":null,"venue":null,"work_id":"3be5f060-323a-4a7f-922b-0b04c903fe46","year":2022},"citing_paper":{"arxiv_id":"2608.05126","last_updated":"2026-08-05T17:50:31Z","snapshot_observed_at":"2026-08-06T19:07:00.208170Z","submitted_at":"2026-08-05T17:50:31Z","title":"Spoken Function Calling: A New Perspective on Spoken Language Understanding for Large Audio Language Models","version":1},"reference_index":43,"source":"pdf_text","source_observed_at":"2026-08-06T04:46:39.922531Z"},"links":{"citing_paper":"/paper/2608.05126"},"observation_digest":"sha256:88bfe9e3b14fd30c9d9229dd9fe295e40f69160897b8b5964a3dbc0da2d86289","observation_id":"b58058e9-d409-41b3-98da-5830d8fde8c7","resolution":{"observed_at":"2026-08-06T04:46:40.438191Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-06T06:34:29.942622+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-06T06:34:29.942622+00:00","source":"crossref"},{"observed_at":"2026-08-06T06:34:23.284952+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T04:46:40.423702Z","title":null,"venue":null,"work_id":"4fa4f3fe-2842-4ee3-a887-fe81be351fb2","year":2025},"citing_paper":{"arxiv_id":"2608.05126","last_updated":"2026-08-05T17:50:31Z","snapshot_observed_at":"2026-08-06T19:07:00.208170Z","submitted_at":"2026-08-05T17:50:31Z","title":"Spoken Function Calling: A New Perspective on Spoken Language Understanding for Large Audio Language Models","version":1},"reference_index":44,"source":"pdf_text","source_observed_at":"2026-08-06T04:46:39.925456Z"},"links":{"citing_paper":"/paper/2608.05126"},"observation_digest":"sha256:f32a618a6059bc25144ab5a0fa8d6b7b59555ebf31fa788e6879e3912d0a14db","observation_id":"91434bf0-9ab8-47a9-9f65-a97ff655dba0","resolution":{"observed_at":"2026-08-06T04:46:40.427185Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-06T06:34:29.942622+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-06T06:34:29.942622+00:00","source":"crossref"},{"observed_at":"2026-08-06T06:34:23.284952+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2403.13372","last_updated":"2024-06-27T22:44:48Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2024-03-20T08:08:54Z","title":"LlamaFactory: Unified Efficient Fine-Tuning of 100+ Language Models","version":4},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2403.13372","snapshot_observed_at":"2026-08-06T04:46:39.928331Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2608.05126","last_updated":"2026-08-05T17:50:31Z","snapshot_observed_at":"2026-08-06T19:07:00.208170Z","submitted_at":"2026-08-05T17:50:31Z","title":"Spoken Function Calling: A New Perspective on Spoken Language Understanding for Large Audio Language Models","version":1},"reference_index":45,"source":"pdf_text","source_observed_at":"2026-08-06T04:46:39.928331Z"},"links":{"cited_paper":"/paper/2403.13372","citing_paper":"/paper/2608.05126"},"observation_digest":"sha256:a35971e19e55fd2c0a21249824f54fefe93b6ca8a7e30e8fa45a0767dccdcd57","observation_id":"3fc20ac7-a4d4-4069-a930-13f2fdba1d19","resolution":{"observed_at":"2026-08-06T04:46:39.928331Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2506.21619","last_updated":"2025-09-03T10:46:35Z","snapshot_observed_at":"2026-07-06T21:48:17.192656Z","submitted_at":"2025-06-23T08:33:40Z","title":"IndexTTS2: A Breakthrough in Emotionally Expressive and Duration-Controlled Auto-Regressive Zero-Shot Text-to-Speech","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2506.21619","snapshot_observed_at":"2026-08-06T04:46:39.932052Z","title":null,"venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2608.05126","last_updated":"2026-08-05T17:50:31Z","snapshot_observed_at":"2026-08-06T19:07:00.208170Z","submitted_at":"2026-08-05T17:50:31Z","title":"Spoken Function Calling: A New Perspective on Spoken Language Understanding for Large Audio Language Models","version":1},"reference_index":46,"source":"pdf_text","source_observed_at":"2026-08-06T04:46:39.932052Z"},"links":{"cited_paper":"/paper/2506.21619","citing_paper":"/paper/2608.05126"},"observation_digest":"sha256:4066c79f2e962278574165b186aedf10caecc17ed175199985c2f92eac996fff","observation_id":"f2a5988c-5344-4936-b561-59216c21ed4d","resolution":{"observed_at":"2026-08-06T04:46:39.932052Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T04:46:40.413278Z","title":null,"venue":null,"work_id":"79c64d94-b4b6-406f-bf66-dfcc06b761d7","year":2023},"citing_paper":{"arxiv_id":"2608.05126","last_updated":"2026-08-05T17:50:31Z","snapshot_observed_at":"2026-08-06T19:07:00.208170Z","submitted_at":"2026-08-05T17:50:31Z","title":"Spoken Function Calling: A New Perspective on Spoken Language Understanding for Large Audio Language Models","version":1},"reference_index":47,"source":"pdf_text","source_observed_at":"2026-08-06T04:46:39.935439Z"},"links":{"citing_paper":"/paper/2608.05126"},"observation_digest":"sha256:a4517791ab7e6d8cfb814b03ee2b53a1ef9c3714c34188f20607f7bc8e1bf3fd","observation_id":"c42c89f3-ab3b-4ea7-b902-ab61781e039f","resolution":{"observed_at":"2026-08-06T04:46:40.416426Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-06T06:34:29.942622+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-06T06:34:29.942622+00:00","source":"crossref"},{"observed_at":"2026-08-06T06:34:23.284952+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"1707.06347","last_updated":"2017-08-28T09:20:06Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2017-07-20T02:32:33Z","title":"Proximal Policy Optimization Algorithms","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"1707.06347","snapshot_observed_at":"2026-08-06T04:46:39.890464Z","title":null,"venue":null,"work_id":null,"year":2017},"citing_paper":{"arxiv_id":"2608.05126","last_updated":"2026-08-05T17:50:31Z","snapshot_observed_at":"2026-08-06T19:07:00.208170Z","submitted_at":"2026-08-05T17:50:31Z","title":"Spoken Function Calling: A New Perspective on Spoken Language Understanding for Large Audio Language Models","version":1},"reference_index":2017,"source":"pdf_text","source_observed_at":"2026-08-06T04:46:39.890464Z"},"links":{"cited_paper":"/paper/1707.06347","citing_paper":"/paper/2608.05126"},"observation_digest":"sha256:80d9fb0b7590f1808c796a6d8030d2b89afd57e55ea6958b4b979262f1c6fb2f","observation_id":"708e4085-8e2e-4f16-9509-9f0e895d3cbc","resolution":{"observed_at":"2026-08-06T04:46:39.890464Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T04:46:40.710692Z","title":null,"venue":null,"work_id":"ff11e4bf-cd15-404d-b61c-343a58ded0c4","year":null},"citing_paper":{"arxiv_id":"2608.05126","last_updated":"2026-08-05T17:50:31Z","snapshot_observed_at":"2026-08-06T19:07:00.208170Z","submitted_at":"2026-08-05T17:50:31Z","title":"Spoken Function Calling: A New Perspective on Spoken Language Understanding for Large Audio Language Models","version":1},"reference_index":2020,"source":"pdf_text","source_observed_at":"2026-08-06T04:46:37.475471Z"},"links":{"citing_paper":"/paper/2608.05126"},"observation_digest":"sha256:4d58ef43326ef85ea7cab52cb3feca11d37ddc7aeaabf44e11e193a18346be27","observation_id":"ba1fb315-acf5-4c05-9c16-be0a8fe839b5","resolution":{"observed_at":"2026-08-06T04:46:40.714262Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-06T06:34:29.942622+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-06T06:34:29.942622+00:00","source":"crossref"},{"observed_at":"2026-08-06T06:34:23.284952+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T04:46:40.516235Z","title":null,"venue":null,"work_id":"f41f660e-9005-4d28-907a-a67474574e6a","year":null},"citing_paper":{"arxiv_id":"2608.05126","last_updated":"2026-08-05T17:50:31Z","snapshot_observed_at":"2026-08-06T19:07:00.208170Z","submitted_at":"2026-08-05T17:50:31Z","title":"Spoken Function Calling: A New Perspective on Spoken Language Understanding for Large Audio Language Models","version":1},"reference_index":2021,"source":"pdf_text","source_observed_at":"2026-08-06T04:46:39.615862Z"},"links":{"citing_paper":"/paper/2608.05126"},"observation_digest":"sha256:49fe615edb8e587aaa91770a800f7901230db04cb4defb83bc44210ab90cf06a","observation_id":"8110ddcc-9896-4428-8428-62729b96b549","resolution":{"observed_at":"2026-08-06T04:46:40.519500Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-06T06:34:29.942622+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-06T06:34:29.942622+00:00","source":"crossref"},{"observed_at":"2026-08-06T06:34:23.284952+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T04:46:40.619774Z","title":null,"venue":null,"work_id":"c9bea1c5-a1f8-4c26-b58a-f231b995da2f","year":null},"citing_paper":{"arxiv_id":"2608.05126","last_updated":"2026-08-05T17:50:31Z","snapshot_observed_at":"2026-08-06T19:07:00.208170Z","submitted_at":"2026-08-05T17:50:31Z","title":"Spoken Function Calling: A New Perspective on Spoken Language Understanding for Large Audio Language Models","version":1},"reference_index":2023,"source":"pdf_text","source_observed_at":"2026-08-06T04:46:38.054926Z"},"links":{"citing_paper":"/paper/2608.05126"},"observation_digest":"sha256:571b2b771313f9fb4f0ee2b41316a0f17d5a15b49a10b131f6eb0716e7614950","observation_id":"7a241a35-f0c8-4233-b37f-328520f690ad","resolution":{"observed_at":"2026-08-06T04:46:40.622925Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-06T06:34:29.942622+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-06T06:34:29.942622+00:00","source":"crossref"},{"observed_at":"2026-08-06T06:34:23.284952+00:00","source":"retraction_watch"}],"state":"measured"}}],"paper":{"arxiv_id":"2608.05126","last_updated":"2026-08-05T17:50:31Z","latest_version":1,"primary_category":"cs.CL","snapshot_observed_at":"2026-08-06T19:07:00.208170Z","submitted_at":"2026-08-05T17:50:31Z","title":"Spoken Function Calling: A New Perspective on Spoken Language Understanding for Large Audio Language Models"},"reference_resolution":{"displayed":51,"state_counts":{"malformed_identifier":0,"metadata_mismatch":0,"parse_uncertain":0,"unresolved":48,"verified_exact":1,"verified_fuzzy":2},"total_outbound_references":51},"refusal":"A citation records a reference. It does not transfer a finding from one paper to another.","schema":"pith.paper-citation-record.v1","standing_sources":[{"observed_at":"2026-08-06T06:34:29.942622+00:00","source":"crossref"},{"observed_at":"2026-08-06T06:34:23.284952+00:00","source":"retraction_watch"}],"thesis":"As of 6 August 2026, this Paper Citation Record lists 51 of 51 outbound references and 0 inbound Pith citation observations for arXiv:2608.05126."}