{"as_of":"2026-08-10T09:10:00Z","caps":{"database_statements":6,"inbound":100,"outbound":100},"context_digest":"sha256:aa42c6c7307227711f66bd94886384215d02617597e97a72e87eac0521edfadb","coverage":[{"denominator":51,"lane":"reference_resolution","note":"Typed states for the displayed outbound observations.","records_observed":51,"source":"paper_references, paper_reference_links","source_observed_at":"2026-08-06T04:46:39.935439Z","state":"measured"},{"denominator":51,"lane":"standing_notices","note":"One-hop event checks from named stored sources.","records_observed":51,"source":"scholarly_work_events, retraction_status_cache","source_observed_at":"2026-08-10T06:31:04.303077+00:00","state":"measured"},{"denominator":0,"lane":"inbound_itemization","note":"Pith citing papers itemized under the disclosed page cap.","records_observed":0,"source":"paper_references, paper_reference_links","source_observed_at":null,"state":"measured"},{"denominator":1,"lane":"external_citation_measurements","note":"A source-named dated measurement, never combined with another source.","records_observed":0,"source":"cited_works","source_observed_at":null,"state":"measured"}],"external_citation_measurements":[],"inbound":[],"links":{"evidence":"/evidence","html":"/paper/2608.05126/citation-record","integrity":"/paper/2608.05126/integrity","json":"/paper/2608.05126/citation-record.json","paper":"/paper/2608.05126"},"outbound":[{"citation":{"cited_paper":{"arxiv_id":"2508.10925","last_updated":"2025-08-08T19:24:38Z","snapshot_observed_at":"2026-08-01T16:27:35.664983Z","submitted_at":"2025-08-08T19:24:38Z","title":"gpt-oss-120b & gpt-oss-20b Model Card","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2508.10925","snapshot_observed_at":"2026-08-06T04:46:37.340915Z","title":null,"venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2608.05126","last_updated":"2026-08-05T17:50:31Z","snapshot_observed_at":"2026-08-10T00:52:03.795266Z","submitted_at":"2026-08-05T17:50:31Z","title":"Spoken Function Calling: A New Perspective on Spoken Language Understanding for Large Audio Language Models","version":1},"reference_index":1,"source":"pdf_text","source_observed_at":"2026-08-06T04:46:37.340915Z"},"links":{"cited_paper":"/paper/2508.10925","citing_paper":"/paper/2608.05126"},"observation_digest":"sha256:e156ba23576b189f8e1d617f98c4fbac31b7cca203c51a5dd2bc9150dd9a5538","observation_id":"68c884cd-11ee-401d-831b-cce310f596cc","resolution":{"observed_at":"2026-08-06T04:46:37.340915Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T04:46:40.722156Z","title":null,"venue":null,"work_id":"97ab27a3-391c-48c3-821c-e0efcc33009c","year":null},"citing_paper":{"arxiv_id":"2608.05126","last_updated":"2026-08-05T17:50:31Z","snapshot_observed_at":"2026-08-10T00:52:03.795266Z","submitted_at":"2026-08-05T17:50:31Z","title":"Spoken Function Calling: A New Perspective on Spoken Language Understanding for Large Audio Language Models","version":1},"reference_index":2,"source":"pdf_text","source_observed_at":"2026-08-06T04:46:37.391520Z"},"links":{"citing_paper":"/paper/2608.05126"},"observation_digest":"sha256:2ada74ef95ee374c1bc6dfff5e881c2e75beeed9159d0f7a5211d763cf5cf705","observation_id":"c78e825d-a7a0-4139-99ce-b8470f8f0b78","resolution":{"observed_at":"2026-08-06T04:46:40.726054Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T04:46:40.698651Z","title":null,"venue":null,"work_id":"58e4d142-28db-4aa8-8ccc-a69798f39ab1","year":2020},"citing_paper":{"arxiv_id":"2608.05126","last_updated":"2026-08-05T17:50:31Z","snapshot_observed_at":"2026-08-10T00:52:03.795266Z","submitted_at":"2026-08-05T17:50:31Z","title":"Spoken Function Calling: A New Perspective on Spoken Language Understanding for Large Audio Language Models","version":1},"reference_index":3,"source":"pdf_text","source_observed_at":"2026-08-06T04:46:37.503546Z"},"links":{"citing_paper":"/paper/2608.05126"},"observation_digest":"sha256:8347e5592ed223b75c952b26cf8933f6af4e13f037f0ab7f5ff20d7d31ba9ed0","observation_id":"5557c7eb-9be7-4197-ab52-78f7429cb924","resolution":{"observed_at":"2026-08-06T04:46:40.702510Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T04:46:40.687259Z","title":null,"venue":null,"work_id":"1ee9c2dc-7acd-4f82-9908-2a291a14a889","year":2020},"citing_paper":{"arxiv_id":"2608.05126","last_updated":"2026-08-05T17:50:31Z","snapshot_observed_at":"2026-08-10T00:52:03.795266Z","submitted_at":"2026-08-05T17:50:31Z","title":"Spoken Function Calling: A New Perspective on Spoken Language Understanding for Large Audio Language Models","version":1},"reference_index":4,"source":"pdf_text","source_observed_at":"2026-08-06T04:46:37.547561Z"},"links":{"citing_paper":"/paper/2608.05126"},"observation_digest":"sha256:29e865cbd61059386414c3470db3b2136378a0d097b2021f59e333e3d1eb18c0","observation_id":"e8b31969-e2e8-45b4-89ba-c050f76011d6","resolution":{"observed_at":"2026-08-06T04:46:40.690532Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T04:46:40.676184Z","title":null,"venue":null,"work_id":"c7b42394-bbf9-40ff-8d0b-7af5f0835b80","year":2018},"citing_paper":{"arxiv_id":"2608.05126","last_updated":"2026-08-05T17:50:31Z","snapshot_observed_at":"2026-08-10T00:52:03.795266Z","submitted_at":"2026-08-05T17:50:31Z","title":"Spoken Function Calling: A New Perspective on Spoken Language Understanding for Large Audio Language Models","version":1},"reference_index":5,"source":"pdf_text","source_observed_at":"2026-08-06T04:46:37.602812Z"},"links":{"citing_paper":"/paper/2608.05126"},"observation_digest":"sha256:8abbad68f4bfc5b8bc480465352088cfad122fc7b3394c90d9b47f8dbc1ec494","observation_id":"458f25a5-a2fe-4a26-9a9a-bdffa223d439","resolution":{"observed_at":"2026-08-06T04:46:40.679876Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2107.03374","last_updated":"2021-07-14T17:16:02Z","snapshot_observed_at":"2026-08-08T11:58:24.516369Z","submitted_at":"2021-07-07T17:41:24Z","title":"Evaluating Large Language Models Trained on Code","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2107.03374","snapshot_observed_at":"2026-08-06T04:46:37.672658Z","title":null,"venue":null,"work_id":null,"year":2021},"citing_paper":{"arxiv_id":"2608.05126","last_updated":"2026-08-05T17:50:31Z","snapshot_observed_at":"2026-08-10T00:52:03.795266Z","submitted_at":"2026-08-05T17:50:31Z","title":"Spoken Function Calling: A New Perspective on Spoken Language Understanding for Large Audio Language Models","version":1},"reference_index":6,"source":"pdf_text","source_observed_at":"2026-08-06T04:46:37.672658Z"},"links":{"cited_paper":"/paper/2107.03374","citing_paper":"/paper/2608.05126"},"observation_digest":"sha256:66b4d7b6feada3eef792ad0ff2ddae7977dbe9954cbd1c372aa5d73afde9f697","observation_id":"31389c0c-8fc3-4a6f-9318-bd65047390be","resolution":{"observed_at":"2026-08-06T04:46:37.672658Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"1805.10190","last_updated":"2018-12-06T16:34:25Z","snapshot_observed_at":"2026-08-05T10:58:51.399386Z","submitted_at":"2018-05-25T15:04:17Z","title":"Snips Voice Platform: an embedded Spoken Language Understanding system for private-by-design voice interfaces","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"1805.10190","snapshot_observed_at":"2026-08-06T04:46:37.724269Z","title":null,"venue":null,"work_id":null,"year":2018},"citing_paper":{"arxiv_id":"2608.05126","last_updated":"2026-08-05T17:50:31Z","snapshot_observed_at":"2026-08-10T00:52:03.795266Z","submitted_at":"2026-08-05T17:50:31Z","title":"Spoken Function Calling: A New Perspective on Spoken Language Understanding for Large Audio Language Models","version":1},"reference_index":7,"source":"pdf_text","source_observed_at":"2026-08-06T04:46:37.724269Z"},"links":{"cited_paper":"/paper/1805.10190","citing_paper":"/paper/2608.05126"},"observation_digest":"sha256:5c9c0a8042064684f39cf104184d92465c8a1317494a840fb202f10c274a53db","observation_id":"9eb6d3c9-32cc-44b4-b4f9-2331d15b4c02","resolution":{"observed_at":"2026-08-06T04:46:37.724269Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T04:46:37.788400Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2608.05126","last_updated":"2026-08-05T17:50:31Z","snapshot_observed_at":"2026-08-10T00:52:03.795266Z","submitted_at":"2026-08-05T17:50:31Z","title":"Spoken Function Calling: A New Perspective on Spoken Language Understanding for Large Audio Language Models","version":1},"reference_index":8,"source":"pdf_text","source_observed_at":"2026-08-06T04:46:37.788400Z"},"links":{"citing_paper":"/paper/2608.05126"},"observation_digest":"sha256:44d3f68762cbc8042db6067a9d771e38a33a2b2512acafc62453e3795d034f90","observation_id":"7aff2a80-1dc0-49ed-a832-152d8b538bf1","resolution":{"observed_at":"2026-08-06T04:46:37.788400Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T04:46:40.658498Z","title":null,"venue":null,"work_id":"3ee55727-524a-402e-af0e-17d00d861eab","year":2023},"citing_paper":{"arxiv_id":"2608.05126","last_updated":"2026-08-05T17:50:31Z","snapshot_observed_at":"2026-08-10T00:52:03.795266Z","submitted_at":"2026-08-05T17:50:31Z","title":"Spoken Function Calling: A New Perspective on Spoken Language Understanding for Large Audio Language Models","version":1},"reference_index":9,"source":"pdf_text","source_observed_at":"2026-08-06T04:46:37.817177Z"},"links":{"citing_paper":"/paper/2608.05126"},"observation_digest":"sha256:1b945f9df175bc42977077814d09c45d837e99e8c9715e7b87c5ef3f30385440","observation_id":"63a6155e-84d0-4671-b203-71d75ddbcf5e","resolution":{"observed_at":"2026-08-06T04:46:40.661567Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T04:46:40.647670Z","title":null,"venue":null,"work_id":"88cfc1e0-0f57-4df1-ab78-5b010c2bfd7c","year":1990},"citing_paper":{"arxiv_id":"2608.05126","last_updated":"2026-08-05T17:50:31Z","snapshot_observed_at":"2026-08-10T00:52:03.795266Z","submitted_at":"2026-08-05T17:50:31Z","title":"Spoken Function Calling: A New Perspective on Spoken Language Understanding for Large Audio Language Models","version":1},"reference_index":10,"source":"pdf_text","source_observed_at":"2026-08-06T04:46:37.871444Z"},"links":{"citing_paper":"/paper/2608.05126"},"observation_digest":"sha256:66ab605aa49fe0bc72a5d86d094e060a63a166c8658f27af57e68f0c9a4ded15","observation_id":"f72e710f-893a-45f6-b226-431fd323f308","resolution":{"observed_at":"2026-08-06T04:46:40.650909Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2503.23278","last_updated":"2025-10-07T07:13:32Z","snapshot_observed_at":"2026-07-06T21:00:55.979837Z","submitted_at":"2025-03-30T01:58:22Z","title":"Model Context Protocol (MCP): Landscape, Security Threats, and Future Research Directions","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2503.23278","snapshot_observed_at":"2026-08-06T04:46:37.952256Z","title":null,"venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2608.05126","last_updated":"2026-08-05T17:50:31Z","snapshot_observed_at":"2026-08-10T00:52:03.795266Z","submitted_at":"2026-08-05T17:50:31Z","title":"Spoken Function Calling: A New Perspective on Spoken Language Understanding for Large Audio Language Models","version":1},"reference_index":11,"source":"pdf_text","source_observed_at":"2026-08-06T04:46:37.952256Z"},"links":{"cited_paper":"/paper/2503.23278","citing_paper":"/paper/2608.05126"},"observation_digest":"sha256:3300d52c881bde01f636a92a80d3358d1c85a14478cf0216c41cdc17bdc4805f","observation_id":"01bf1bb0-f601-4b2b-b3cb-9d487837f49a","resolution":{"observed_at":"2026-08-06T04:46:37.952256Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T04:46:40.636493Z","title":null,"venue":null,"work_id":"5546697f-f219-4290-a40a-720488c750ba","year":2022},"citing_paper":{"arxiv_id":"2608.05126","last_updated":"2026-08-05T17:50:31Z","snapshot_observed_at":"2026-08-10T00:52:03.795266Z","submitted_at":"2026-08-05T17:50:31Z","title":"Spoken Function Calling: A New Perspective on Spoken Language Understanding for Large Audio Language Models","version":1},"reference_index":12,"source":"pdf_text","source_observed_at":"2026-08-06T04:46:38.002010Z"},"links":{"citing_paper":"/paper/2608.05126"},"observation_digest":"sha256:c8adbd704628bb7f5e80cb938e2a991224509066baae890ea962d17abb911e17","observation_id":"ba28f402-8d13-4679-9981-a29766de2a49","resolution":{"observed_at":"2026-08-06T04:46:40.640326Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T04:46:38.049617Z","title":null,"venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2608.05126","last_updated":"2026-08-05T17:50:31Z","snapshot_observed_at":"2026-08-10T00:52:03.795266Z","submitted_at":"2026-08-05T17:50:31Z","title":"Spoken Function Calling: A New Perspective on Spoken Language Understanding for Large Audio Language Models","version":1},"reference_index":13,"source":"pdf_text","source_observed_at":"2026-08-06T04:46:38.049617Z"},"links":{"citing_paper":"/paper/2608.05126"},"observation_digest":"sha256:f7318f6939b0c9cc7bf6157bd9b44f4a8ff19af4f5267b67e4884d4ead180cec","observation_id":"deabefdf-803c-4a8a-8a7e-8955f1363987","resolution":{"observed_at":"2026-08-06T04:46:38.049617Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2001.08361","last_updated":"2020-01-23T03:59:20Z","snapshot_observed_at":"2026-07-06T08:52:12.656082Z","submitted_at":"2020-01-23T03:59:20Z","title":"Scaling Laws for Neural Language Models","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2001.08361","snapshot_observed_at":"2026-08-06T04:46:38.170089Z","title":null,"venue":null,"work_id":null,"year":2020},"citing_paper":{"arxiv_id":"2608.05126","last_updated":"2026-08-05T17:50:31Z","snapshot_observed_at":"2026-08-10T00:52:03.795266Z","submitted_at":"2026-08-05T17:50:31Z","title":"Spoken Function Calling: A New Perspective on Spoken Language Understanding for Large Audio Language Models","version":1},"reference_index":14,"source":"pdf_text","source_observed_at":"2026-08-06T04:46:38.170089Z"},"links":{"cited_paper":"/paper/2001.08361","citing_paper":"/paper/2608.05126"},"observation_digest":"sha256:bbaad4533fa4e69727d86acb9afe69f88e27fb723fbc9748dac3019ae41b95b1","observation_id":"85d553fc-8658-4e88-8c81-93b61a9d1dad","resolution":{"observed_at":"2026-08-06T04:46:38.170089Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T04:46:40.610153Z","title":null,"venue":null,"work_id":"380ca5fc-14bc-403b-9c45-3c7d08ada54d","year":2023},"citing_paper":{"arxiv_id":"2608.05126","last_updated":"2026-08-05T17:50:31Z","snapshot_observed_at":"2026-08-10T00:52:03.795266Z","submitted_at":"2026-08-05T17:50:31Z","title":"Spoken Function Calling: A New Perspective on Spoken Language Understanding for Large Audio Language Models","version":1},"reference_index":15,"source":"pdf_text","source_observed_at":"2026-08-06T04:46:38.277118Z"},"links":{"citing_paper":"/paper/2608.05126"},"observation_digest":"sha256:93594ac6d4989a08d840d126d80a69141ff31dc64af472708cfe5081333b85b2","observation_id":"f8bafe28-69c9-4361-b7df-a6508fe6abce","resolution":{"observed_at":"2026-08-06T04:46:40.613198Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T04:46:40.600362Z","title":null,"venue":null,"work_id":"9a714657-de7d-4a78-84c4-40f585a271fd","year":2021},"citing_paper":{"arxiv_id":"2608.05126","last_updated":"2026-08-05T17:50:31Z","snapshot_observed_at":"2026-08-10T00:52:03.795266Z","submitted_at":"2026-08-05T17:50:31Z","title":"Spoken Function Calling: A New Perspective on Spoken Language Understanding for Large Audio Language Models","version":1},"reference_index":16,"source":"pdf_text","source_observed_at":"2026-08-06T04:46:38.387114Z"},"links":{"citing_paper":"/paper/2608.05126"},"observation_digest":"sha256:7d4b2e812283b796b47de2d6bbedd8f678f7f0ab59ad6bf882e845f9a696a89e","observation_id":"5f0054d1-655c-4b56-aee4-7bdad37c3f57","resolution":{"observed_at":"2026-08-06T04:46:40.603627Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T04:46:40.590271Z","title":null,"venue":null,"work_id":"120196b9-0611-445f-bac8-4398599117c9","year":2024},"citing_paper":{"arxiv_id":"2608.05126","last_updated":"2026-08-05T17:50:31Z","snapshot_observed_at":"2026-08-10T00:52:03.795266Z","submitted_at":"2026-08-05T17:50:31Z","title":"Spoken Function Calling: A New Perspective on Spoken Language Understanding for Large Audio Language Models","version":1},"reference_index":17,"source":"pdf_text","source_observed_at":"2026-08-06T04:46:38.460511Z"},"links":{"citing_paper":"/paper/2608.05126"},"observation_digest":"sha256:fa300498e87ae89ce2d5b77933cc207f160ca5f31c1742f00bca86e771a1fc15","observation_id":"06dc2690-b465-44e3-bf99-edc1c85d32da","resolution":{"observed_at":"2026-08-06T04:46:40.593465Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2304.08244","last_updated":"2023-10-25T06:54:12Z","snapshot_observed_at":"2026-08-02T00:07:12.855748Z","submitted_at":"2023-04-14T14:05:32Z","title":"API-Bank: A Comprehensive Benchmark for Tool-Augmented LLMs","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2304.08244","snapshot_observed_at":"2026-08-06T04:46:38.571002Z","title":null,"venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2608.05126","last_updated":"2026-08-05T17:50:31Z","snapshot_observed_at":"2026-08-10T00:52:03.795266Z","submitted_at":"2026-08-05T17:50:31Z","title":"Spoken Function Calling: A New Perspective on Spoken Language Understanding for Large Audio Language Models","version":1},"reference_index":18,"source":"pdf_text","source_observed_at":"2026-08-06T04:46:38.571002Z"},"links":{"cited_paper":"/paper/2304.08244","citing_paper":"/paper/2608.05126"},"observation_digest":"sha256:8e661af81915d8c9b2ace96fa7e22604110eb161ae6ea5c4b85c38e92ca1463c","observation_id":"df2f28b3-d29a-4b28-a0b7-389c682f1a60","resolution":{"observed_at":"2026-08-06T04:46:38.571002Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2409.00920","last_updated":"2025-07-25T08:26:54Z","snapshot_observed_at":"2026-08-07T03:58:38.931628Z","submitted_at":"2024-09-02T03:19:56Z","title":"ToolACE: Winning the Points of LLM Function Calling","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2409.00920","snapshot_observed_at":"2026-08-06T04:46:38.644903Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2608.05126","last_updated":"2026-08-05T17:50:31Z","snapshot_observed_at":"2026-08-10T00:52:03.795266Z","submitted_at":"2026-08-05T17:50:31Z","title":"Spoken Function Calling: A New Perspective on Spoken Language Understanding for Large Audio Language Models","version":1},"reference_index":19,"source":"pdf_text","source_observed_at":"2026-08-06T04:46:38.644903Z"},"links":{"cited_paper":"/paper/2409.00920","citing_paper":"/paper/2608.05126"},"observation_digest":"sha256:68e38656f75773508f51d706555733427ec3e961d320f644889b74c772f522b8","observation_id":"ac7b5666-6aea-4533-827d-6543ad4e0d3c","resolution":{"observed_at":"2026-08-06T04:46:38.644903Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T04:46:40.579940Z","title":null,"venue":null,"work_id":"b3893602-e668-456b-9d62-bca018de7a33","year":2024},"citing_paper":{"arxiv_id":"2608.05126","last_updated":"2026-08-05T17:50:31Z","snapshot_observed_at":"2026-08-10T00:52:03.795266Z","submitted_at":"2026-08-05T17:50:31Z","title":"Spoken Function Calling: A New Perspective on Spoken Language Understanding for Large Audio Language Models","version":1},"reference_index":20,"source":"pdf_text","source_observed_at":"2026-08-06T04:46:38.761255Z"},"links":{"citing_paper":"/paper/2608.05126"},"observation_digest":"sha256:b24c4ddffc297951cf54a57eeb19f29926d02ab7f7b500528874fbef4e362405","observation_id":"eef6f6c9-0c87-41d7-b02b-8a75fc1e0051","resolution":{"observed_at":"2026-08-06T04:46:40.582959Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T04:46:40.569927Z","title":null,"venue":null,"work_id":"a99d8251-3a0b-4e93-ae3d-a4eb1aaaaa8b","year":2019},"citing_paper":{"arxiv_id":"2608.05126","last_updated":"2026-08-05T17:50:31Z","snapshot_observed_at":"2026-08-10T00:52:03.795266Z","submitted_at":"2026-08-05T17:50:31Z","title":"Spoken Function Calling: A New Perspective on Spoken Language Understanding for Large Audio Language Models","version":1},"reference_index":21,"source":"pdf_text","source_observed_at":"2026-08-06T04:46:38.878716Z"},"links":{"citing_paper":"/paper/2608.05126"},"observation_digest":"sha256:029d2905761c661b2b4b430672f79e1aaf2ec57f490e6b3a6aaef28b53e23bf3","observation_id":"25b557ab-da37-4e4b-abe0-4db7e41d1184","resolution":{"observed_at":"2026-08-06T04:46:40.573344Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T04:46:40.559557Z","title":null,"venue":null,"work_id":"082b12ae-1380-4471-8fcf-14c369552a83","year":2015},"citing_paper":{"arxiv_id":"2608.05126","last_updated":"2026-08-05T17:50:31Z","snapshot_observed_at":"2026-08-10T00:52:03.795266Z","submitted_at":"2026-08-05T17:50:31Z","title":"Spoken Function Calling: A New Perspective on Spoken Language Understanding for Large Audio Language Models","version":1},"reference_index":22,"source":"pdf_text","source_observed_at":"2026-08-06T04:46:38.938967Z"},"links":{"citing_paper":"/paper/2608.05126"},"observation_digest":"sha256:587bb3bc5edb630820a737215c66d527ef3b1943e3d55a7cdf23acab20d9b1c1","observation_id":"ca227e6e-6959-4e1b-8662-acb94a536bfe","resolution":{"observed_at":"2026-08-06T04:46:40.562940Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2303.09014","last_updated":"2023-03-16T01:04:45Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2023-03-16T01:04:45Z","title":"ART: Automatic multi-step reasoning and tool-use for large language models","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2303.09014","snapshot_observed_at":"2026-08-06T04:46:39.053319Z","title":null,"venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2608.05126","last_updated":"2026-08-05T17:50:31Z","snapshot_observed_at":"2026-08-10T00:52:03.795266Z","submitted_at":"2026-08-05T17:50:31Z","title":"Spoken Function Calling: A New Perspective on Spoken Language Understanding for Large Audio Language Models","version":1},"reference_index":23,"source":"pdf_text","source_observed_at":"2026-08-06T04:46:39.053319Z"},"links":{"cited_paper":"/paper/2303.09014","citing_paper":"/paper/2608.05126"},"observation_digest":"sha256:5558d156645ebee15687ffa6660b665470d296830dddffc6b44ec2f48fbf2a29","observation_id":"cd1d388a-d0b1-4ece-8de4-5e40215d3b49","resolution":{"observed_at":"2026-08-06T04:46:39.053319Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T04:46:40.549078Z","title":null,"venue":null,"work_id":"98adf8ca-53ff-459c-83fb-dfe000e9cbc4","year":2024},"citing_paper":{"arxiv_id":"2608.05126","last_updated":"2026-08-05T17:50:31Z","snapshot_observed_at":"2026-08-10T00:52:03.795266Z","submitted_at":"2026-08-05T17:50:31Z","title":"Spoken Function Calling: A New Perspective on Spoken Language Understanding for Large Audio Language Models","version":1},"reference_index":24,"source":"pdf_text","source_observed_at":"2026-08-06T04:46:39.100248Z"},"links":{"citing_paper":"/paper/2608.05126"},"observation_digest":"sha256:01c2ccb9f1f8178886624faa819a8acd392255f2b25ea51883165b4c2a8028ed","observation_id":"82474a83-d490-409c-aea6-05461f7cbde9","resolution":{"observed_at":"2026-08-06T04:46:40.552379Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T04:46:39.202981Z","title":null,"venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2608.05126","last_updated":"2026-08-05T17:50:31Z","snapshot_observed_at":"2026-08-10T00:52:03.795266Z","submitted_at":"2026-08-05T17:50:31Z","title":"Spoken Function Calling: A New Perspective on Spoken Language Understanding for Large Audio Language Models","version":1},"reference_index":25,"source":"pdf_text","source_observed_at":"2026-08-06T04:46:39.202981Z"},"links":{"citing_paper":"/paper/2608.05126"},"observation_digest":"sha256:8e2bbb70323d909d4b09e1cfef6bbea94f72b9a45398ef39c07e009df162a230","observation_id":"83355c5f-d210-4a60-a2be-6c58ea51fe01","resolution":{"observed_at":"2026-08-06T04:46:39.202981Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T04:46:39.328263Z","title":null,"venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2608.05126","last_updated":"2026-08-05T17:50:31Z","snapshot_observed_at":"2026-08-10T00:52:03.795266Z","submitted_at":"2026-08-05T17:50:31Z","title":"Spoken Function Calling: A New Perspective on Spoken Language Understanding for Large Audio Language Models","version":1},"reference_index":26,"source":"pdf_text","source_observed_at":"2026-08-06T04:46:39.328263Z"},"links":{"citing_paper":"/paper/2608.05126"},"observation_digest":"sha256:c564bd1394e75098508d6eb74131d2beb1523262aed69eb35973523bdeb5fdb5","observation_id":"4ca49e6a-6d8f-452d-83d9-cc0f03268979","resolution":{"observed_at":"2026-08-06T04:46:39.328263Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T04:46:40.537949Z","title":null,"venue":null,"work_id":"79b05501-504a-492d-8e51-e922f00d359a","year":2025},"citing_paper":{"arxiv_id":"2608.05126","last_updated":"2026-08-05T17:50:31Z","snapshot_observed_at":"2026-08-10T00:52:03.795266Z","submitted_at":"2026-08-05T17:50:31Z","title":"Spoken Function Calling: A New Perspective on Spoken Language Understanding for Large Audio Language Models","version":1},"reference_index":27,"source":"pdf_text","source_observed_at":"2026-08-06T04:46:39.437359Z"},"links":{"citing_paper":"/paper/2608.05126"},"observation_digest":"sha256:1368ca505fae04cce977c689537a886d08ccb75e829df84d7d7a090ad208eb40","observation_id":"e6fade8d-2742-498f-bac1-7c7015e0904e","resolution":{"observed_at":"2026-08-06T04:46:40.541299Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T04:46:40.527528Z","title":null,"venue":null,"work_id":"19e8431b-7db0-48e9-afce-ea8bae3c7e79","year":null},"citing_paper":{"arxiv_id":"2608.05126","last_updated":"2026-08-05T17:50:31Z","snapshot_observed_at":"2026-08-10T00:52:03.795266Z","submitted_at":"2026-08-05T17:50:31Z","title":"Spoken Function Calling: A New Perspective on Spoken Language Understanding for Large Audio Language Models","version":1},"reference_index":28,"source":"pdf_text","source_observed_at":"2026-08-06T04:46:39.503875Z"},"links":{"citing_paper":"/paper/2608.05126"},"observation_digest":"sha256:f1bb44ef28592789710835ddb6525603909bbe36811e5addf3a64b8c64470255","observation_id":"756a5b05-f1cc-4f11-b4fa-ed34bf3a8932","resolution":{"observed_at":"2026-08-06T04:46:40.530752Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2004.10087","last_updated":"2020-10-17T04:28:29Z","snapshot_observed_at":"2026-08-10T02:25:30.909313Z","submitted_at":"2020-04-21T15:07:34Z","title":"AGIF: An Adaptive Graph-Interactive Framework for Joint Multiple Intent Detection and Slot Filling","version":4},"cited_work":{"arxiv_id":"2004.10087","doi":null,"metadata_source":"pith","pith_arxiv_id":"2004.10087","snapshot_observed_at":"2026-08-06T04:46:40.169869Z","title":"AGIF: An Adaptive Graph-Interactive Framework for Joint Multiple Intent Detection and Slot Filling","venue":"cs.CL","work_id":"0f62fb7b-3efd-480c-8926-4912808b55ac","year":2020},"citing_paper":{"arxiv_id":"2608.05126","last_updated":"2026-08-05T17:50:31Z","snapshot_observed_at":"2026-08-10T00:52:03.795266Z","submitted_at":"2026-08-05T17:50:31Z","title":"Spoken Function Calling: A New Perspective on Spoken Language Understanding for Large Audio Language Models","version":1},"reference_index":29,"source":"pdf_text","source_observed_at":"2026-08-06T04:46:39.666283Z"},"links":{"cited_paper":"/paper/2004.10087","citing_paper":"/paper/2608.05126"},"observation_digest":"sha256:9549637f713acaab30a25a9f3e89e0ff55dc460abf899aab0f061282ee91171c","observation_id":"9fb7097d-4845-4a57-9791-e4eb8693a43a","resolution":{"observed_at":"2026-08-06T04:46:40.175397Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2307.16789","last_updated":"2023-10-03T14:45:48Z","snapshot_observed_at":"2026-07-06T16:00:46.542753Z","submitted_at":"2023-07-31T15:56:53Z","title":"ToolLLM: Facilitating Large Language Models to Master 16000+ Real-world APIs","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2307.16789","snapshot_observed_at":"2026-08-06T04:46:39.793942Z","title":null,"venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2608.05126","last_updated":"2026-08-05T17:50:31Z","snapshot_observed_at":"2026-08-10T00:52:03.795266Z","submitted_at":"2026-08-05T17:50:31Z","title":"Spoken Function Calling: A New Perspective on Spoken Language Understanding for Large Audio Language Models","version":1},"reference_index":30,"source":"pdf_text","source_observed_at":"2026-08-06T04:46:39.793942Z"},"links":{"cited_paper":"/paper/2307.16789","citing_paper":"/paper/2608.05126"},"observation_digest":"sha256:a56c9a9d5628eddd19b1f9ebe30959bdbda72c454fd0d739513761c402218f09","observation_id":"039249dd-a5e5-4468-b96b-d0cd39d38d43","resolution":{"observed_at":"2026-08-06T04:46:39.793942Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T04:46:40.504875Z","title":null,"venue":null,"work_id":"f18f0b19-b69a-411b-b4bf-c096a656f681","year":2023},"citing_paper":{"arxiv_id":"2608.05126","last_updated":"2026-08-05T17:50:31Z","snapshot_observed_at":"2026-08-10T00:52:03.795266Z","submitted_at":"2026-08-05T17:50:31Z","title":"Spoken Function Calling: A New Perspective on Spoken Language Understanding for Large Audio Language Models","version":1},"reference_index":31,"source":"pdf_text","source_observed_at":"2026-08-06T04:46:39.880243Z"},"links":{"citing_paper":"/paper/2608.05126"},"observation_digest":"sha256:9028c5352ad4c81d30c017d741a66815278ddb5896339e55b7724a294664fd69","observation_id":"34945ba1-9765-4366-9a04-f1b063773316","resolution":{"observed_at":"2026-08-06T04:46:40.509063Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T04:46:40.494526Z","title":null,"venue":null,"work_id":"16df8829-3ebc-48bc-a2d5-10315a613278","year":2023},"citing_paper":{"arxiv_id":"2608.05126","last_updated":"2026-08-05T17:50:31Z","snapshot_observed_at":"2026-08-10T00:52:03.795266Z","submitted_at":"2026-08-05T17:50:31Z","title":"Spoken Function Calling: A New Perspective on Spoken Language Understanding for Large Audio Language Models","version":1},"reference_index":32,"source":"pdf_text","source_observed_at":"2026-08-06T04:46:39.883762Z"},"links":{"citing_paper":"/paper/2608.05126"},"observation_digest":"sha256:4f9adcaef5a8836132d5d00f7d80c32c6efd0bcf1adaadaed5005ce37e451c80","observation_id":"92e6dff4-ab88-4901-8457-2027f0f3ed01","resolution":{"observed_at":"2026-08-06T04:46:40.497702Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T04:46:39.887216Z","title":null,"venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2608.05126","last_updated":"2026-08-05T17:50:31Z","snapshot_observed_at":"2026-08-10T00:52:03.795266Z","submitted_at":"2026-08-05T17:50:31Z","title":"Spoken Function Calling: A New Perspective on Spoken Language Understanding for Large Audio Language Models","version":1},"reference_index":33,"source":"pdf_text","source_observed_at":"2026-08-06T04:46:39.887216Z"},"links":{"citing_paper":"/paper/2608.05126"},"observation_digest":"sha256:84a65d4d40f44057ecf56c8c776b3c87c7b4973859d4e2752d37792dece7e4c5","observation_id":"b0cdd6c2-f0e2-4df0-baaf-7344be6afaca","resolution":{"observed_at":"2026-08-06T04:46:39.887216Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2402.03300","last_updated":"2024-04-27T15:25:53Z","snapshot_observed_at":"2026-08-06T14:58:42.911363Z","submitted_at":"2024-02-05T18:55:32Z","title":"DeepSeekMath: Pushing the Limits of Mathematical Reasoning in Open Language Models","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2402.03300","snapshot_observed_at":"2026-08-06T04:46:39.893471Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2608.05126","last_updated":"2026-08-05T17:50:31Z","snapshot_observed_at":"2026-08-10T00:52:03.795266Z","submitted_at":"2026-08-05T17:50:31Z","title":"Spoken Function Calling: A New Perspective on Spoken Language Understanding for Large Audio Language Models","version":1},"reference_index":34,"source":"pdf_text","source_observed_at":"2026-08-06T04:46:39.893471Z"},"links":{"cited_paper":"/paper/2402.03300","citing_paper":"/paper/2608.05126"},"observation_digest":"sha256:521c6cc4fa27a0897d338c09ef4234142554d9963badd140f37b05e3b317006e","observation_id":"80f980d4-fe3a-4fcf-be54-cb1ce85a2083","resolution":{"observed_at":"2026-08-06T04:46:39.893471Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T04:46:39.896657Z","title":"2011.Spoken language understanding: Systems for extracting semantic information from speech","venue":null,"work_id":null,"year":2011},"citing_paper":{"arxiv_id":"2608.05126","last_updated":"2026-08-05T17:50:31Z","snapshot_observed_at":"2026-08-10T00:52:03.795266Z","submitted_at":"2026-08-05T17:50:31Z","title":"Spoken Function Calling: A New Perspective on Spoken Language Understanding for Large Audio Language Models","version":1},"reference_index":35,"source":"pdf_text","source_observed_at":"2026-08-06T04:46:39.896657Z"},"links":{"citing_paper":"/paper/2608.05126"},"observation_digest":"sha256:183bcf091df8a14168167351bcafad4e4c61d2a5dacd269dcf31dc5e7545d683","observation_id":"4989de87-f25e-4662-9af8-021b05412024","resolution":{"observed_at":"2026-08-06T04:46:39.896657Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2506.04779","last_updated":"2026-03-16T09:24:19Z","snapshot_observed_at":"2026-08-02T22:58:18.776196Z","submitted_at":"2025-06-05T09:09:36Z","title":"MMSU: A Massive Multi-task Spoken Language Understanding and Reasoning Benchmark","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2506.04779","snapshot_observed_at":"2026-08-06T04:46:39.899707Z","title":null,"venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2608.05126","last_updated":"2026-08-05T17:50:31Z","snapshot_observed_at":"2026-08-10T00:52:03.795266Z","submitted_at":"2026-08-05T17:50:31Z","title":"Spoken Function Calling: A New Perspective on Spoken Language Understanding for Large Audio Language Models","version":1},"reference_index":36,"source":"pdf_text","source_observed_at":"2026-08-06T04:46:39.899707Z"},"links":{"cited_paper":"/paper/2506.04779","citing_paper":"/paper/2608.05126"},"observation_digest":"sha256:faacb26115686f3518774079f5f1c746a4b6fcfe7f79cbb5f0265833d697eae2","observation_id":"a6c5b19d-4c99-4a95-bd34-3365cf908893","resolution":{"observed_at":"2026-08-06T04:46:39.899707Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T04:46:40.469145Z","title":null,"venue":null,"work_id":"d8f1239e-86ab-432c-922a-71dd66cc75c9","year":2022},"citing_paper":{"arxiv_id":"2608.05126","last_updated":"2026-08-05T17:50:31Z","snapshot_observed_at":"2026-08-10T00:52:03.795266Z","submitted_at":"2026-08-05T17:50:31Z","title":"Spoken Function Calling: A New Perspective on Spoken Language Understanding for Large Audio Language Models","version":1},"reference_index":37,"source":"pdf_text","source_observed_at":"2026-08-06T04:46:39.903078Z"},"links":{"citing_paper":"/paper/2608.05126"},"observation_digest":"sha256:c109a130d3001c7a760907d948f1b4b857550fd00baf9aa17ca6a2ea4c4558b2","observation_id":"35373006-b9fe-4226-b346-c936ca36fafc","resolution":{"observed_at":"2026-08-06T04:46:40.472812Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T04:46:40.457731Z","title":"Williams, Antoine Raux, and Matthew Henderson","venue":null,"work_id":"9d0551dc-1e4c-413e-829a-b60b3252fa15","year":2016},"citing_paper":{"arxiv_id":"2608.05126","last_updated":"2026-08-05T17:50:31Z","snapshot_observed_at":"2026-08-10T00:52:03.795266Z","submitted_at":"2026-08-05T17:50:31Z","title":"Spoken Function Calling: A New Perspective on Spoken Language Understanding for Large Audio Language Models","version":1},"reference_index":38,"source":"pdf_text","source_observed_at":"2026-08-06T04:46:39.906076Z"},"links":{"citing_paper":"/paper/2608.05126"},"observation_digest":"sha256:931d4b12ddc5d5a83d7ba5f8c0b997e192573e4e44261dd8efd1e83f038132a4","observation_id":"1d1b386c-f2cf-40d2-bc3e-d7c788abf714","resolution":{"observed_at":"2026-08-06T04:46:40.461459Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T04:46:40.445521Z","title":"Williams, Antoine Raux, Deepak Ramachandran, and Alan W","venue":null,"work_id":"6bc76f2b-3e04-4a49-b96f-172aaeacb1b9","year":2013},"citing_paper":{"arxiv_id":"2608.05126","last_updated":"2026-08-05T17:50:31Z","snapshot_observed_at":"2026-08-10T00:52:03.795266Z","submitted_at":"2026-08-05T17:50:31Z","title":"Spoken Function Calling: A New Perspective on Spoken Language Understanding for Large Audio Language Models","version":1},"reference_index":39,"source":"pdf_text","source_observed_at":"2026-08-06T04:46:39.909060Z"},"links":{"citing_paper":"/paper/2608.05126"},"observation_digest":"sha256:0f867ba8c631237f36d394e7ea03425cacab1f6830ad0604846c18a37acb47dd","observation_id":"8a7d6943-0bbd-43db-8a55-b7d8e772c7d7","resolution":{"observed_at":"2026-08-06T04:46:40.450010Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2507.16632","last_updated":"2025-08-27T16:42:11Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2025-07-22T14:23:55Z","title":"Step-Audio 2 Technical Report","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2507.16632","snapshot_observed_at":"2026-08-06T04:46:39.912304Z","title":null,"venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2608.05126","last_updated":"2026-08-05T17:50:31Z","snapshot_observed_at":"2026-08-10T00:52:03.795266Z","submitted_at":"2026-08-05T17:50:31Z","title":"Spoken Function Calling: A New Perspective on Spoken Language Understanding for Large Audio Language Models","version":1},"reference_index":40,"source":"pdf_text","source_observed_at":"2026-08-06T04:46:39.912304Z"},"links":{"cited_paper":"/paper/2507.16632","citing_paper":"/paper/2608.05126"},"observation_digest":"sha256:1079b1d4f8b464d99ca992d7980e870068295b91173d4a6ee8dec44be307289f","observation_id":"7881e317-5737-4e76-b43e-7829ea46fb42","resolution":{"observed_at":"2026-08-06T04:46:39.912304Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2509.17765","last_updated":"2025-09-22T13:26:24Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2025-09-22T13:26:24Z","title":"Qwen3-Omni Technical Report","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2509.17765","snapshot_observed_at":"2026-08-06T04:46:39.915594Z","title":null,"venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2608.05126","last_updated":"2026-08-05T17:50:31Z","snapshot_observed_at":"2026-08-10T00:52:03.795266Z","submitted_at":"2026-08-05T17:50:31Z","title":"Spoken Function Calling: A New Perspective on Spoken Language Understanding for Large Audio Language Models","version":1},"reference_index":41,"source":"pdf_text","source_observed_at":"2026-08-06T04:46:39.915594Z"},"links":{"cited_paper":"/paper/2509.17765","citing_paper":"/paper/2608.05126"},"observation_digest":"sha256:0830f388ff521a5435f14754ff0f06b61ba7c519c739304332a73138c18a9e85","observation_id":"73abdf14-ba0f-4a94-a564-53744e0e8c6a","resolution":{"observed_at":"2026-08-06T04:46:39.915594Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2505.09388","last_updated":"2025-05-14T13:41:34Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2025-05-14T13:41:34Z","title":"Qwen3 Technical Report","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2505.09388","snapshot_observed_at":"2026-08-06T04:46:39.919197Z","title":null,"venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2608.05126","last_updated":"2026-08-05T17:50:31Z","snapshot_observed_at":"2026-08-10T00:52:03.795266Z","submitted_at":"2026-08-05T17:50:31Z","title":"Spoken Function Calling: A New Perspective on Spoken Language Understanding for Large Audio Language Models","version":1},"reference_index":42,"source":"pdf_text","source_observed_at":"2026-08-06T04:46:39.919197Z"},"links":{"cited_paper":"/paper/2505.09388","citing_paper":"/paper/2608.05126"},"observation_digest":"sha256:72a23698cf2e8e68f3b03f9c82cd6fb2988d0038e6aaf4454814b90bad13aa88","observation_id":"7970618c-360d-4126-b987-5948f065eb9f","resolution":{"observed_at":"2026-08-06T04:46:39.919197Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T04:46:40.434830Z","title":null,"venue":null,"work_id":"3be5f060-323a-4a7f-922b-0b04c903fe46","year":2022},"citing_paper":{"arxiv_id":"2608.05126","last_updated":"2026-08-05T17:50:31Z","snapshot_observed_at":"2026-08-10T00:52:03.795266Z","submitted_at":"2026-08-05T17:50:31Z","title":"Spoken Function Calling: A New Perspective on Spoken Language Understanding for Large Audio Language Models","version":1},"reference_index":43,"source":"pdf_text","source_observed_at":"2026-08-06T04:46:39.922531Z"},"links":{"citing_paper":"/paper/2608.05126"},"observation_digest":"sha256:3ea04cbc3045ce2dde430c0c7886c9010e9009bec9b5702865e798046075e9cf","observation_id":"b58058e9-d409-41b3-98da-5830d8fde8c7","resolution":{"observed_at":"2026-08-06T04:46:40.438191Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T04:46:40.423702Z","title":null,"venue":null,"work_id":"4fa4f3fe-2842-4ee3-a887-fe81be351fb2","year":2025},"citing_paper":{"arxiv_id":"2608.05126","last_updated":"2026-08-05T17:50:31Z","snapshot_observed_at":"2026-08-10T00:52:03.795266Z","submitted_at":"2026-08-05T17:50:31Z","title":"Spoken Function Calling: A New Perspective on Spoken Language Understanding for Large Audio Language Models","version":1},"reference_index":44,"source":"pdf_text","source_observed_at":"2026-08-06T04:46:39.925456Z"},"links":{"citing_paper":"/paper/2608.05126"},"observation_digest":"sha256:cceaf72c8fd7854e9031cb7dfedbbd3df3605ad45974451490281c05c046894e","observation_id":"91434bf0-9ab8-47a9-9f65-a97ff655dba0","resolution":{"observed_at":"2026-08-06T04:46:40.427185Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2403.13372","last_updated":"2024-06-27T22:44:48Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2024-03-20T08:08:54Z","title":"LlamaFactory: Unified Efficient Fine-Tuning of 100+ Language Models","version":4},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2403.13372","snapshot_observed_at":"2026-08-06T04:46:39.928331Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2608.05126","last_updated":"2026-08-05T17:50:31Z","snapshot_observed_at":"2026-08-10T00:52:03.795266Z","submitted_at":"2026-08-05T17:50:31Z","title":"Spoken Function Calling: A New Perspective on Spoken Language Understanding for Large Audio Language Models","version":1},"reference_index":45,"source":"pdf_text","source_observed_at":"2026-08-06T04:46:39.928331Z"},"links":{"cited_paper":"/paper/2403.13372","citing_paper":"/paper/2608.05126"},"observation_digest":"sha256:cb75b07ac91ab813aa63f0b0132b0d5f21f558b1e3a09c8a82837bbeec1ba5cd","observation_id":"3fc20ac7-a4d4-4069-a930-13f2fdba1d19","resolution":{"observed_at":"2026-08-06T04:46:39.928331Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2506.21619","last_updated":"2025-09-03T10:46:35Z","snapshot_observed_at":"2026-08-08T06:46:54.528137Z","submitted_at":"2025-06-23T08:33:40Z","title":"IndexTTS2: A Breakthrough in Emotionally Expressive and Duration-Controlled Auto-Regressive Zero-Shot Text-to-Speech","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2506.21619","snapshot_observed_at":"2026-08-06T04:46:39.932052Z","title":null,"venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2608.05126","last_updated":"2026-08-05T17:50:31Z","snapshot_observed_at":"2026-08-10T00:52:03.795266Z","submitted_at":"2026-08-05T17:50:31Z","title":"Spoken Function Calling: A New Perspective on Spoken Language Understanding for Large Audio Language Models","version":1},"reference_index":46,"source":"pdf_text","source_observed_at":"2026-08-06T04:46:39.932052Z"},"links":{"cited_paper":"/paper/2506.21619","citing_paper":"/paper/2608.05126"},"observation_digest":"sha256:1787743e106a4085785a058b26d5726f6fc4c68408512436da8d5380a9cdd90c","observation_id":"f2a5988c-5344-4936-b561-59216c21ed4d","resolution":{"observed_at":"2026-08-06T04:46:39.932052Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T04:46:40.413278Z","title":null,"venue":null,"work_id":"79c64d94-b4b6-406f-bf66-dfcc06b761d7","year":2023},"citing_paper":{"arxiv_id":"2608.05126","last_updated":"2026-08-05T17:50:31Z","snapshot_observed_at":"2026-08-10T00:52:03.795266Z","submitted_at":"2026-08-05T17:50:31Z","title":"Spoken Function Calling: A New Perspective on Spoken Language Understanding for Large Audio Language Models","version":1},"reference_index":47,"source":"pdf_text","source_observed_at":"2026-08-06T04:46:39.935439Z"},"links":{"citing_paper":"/paper/2608.05126"},"observation_digest":"sha256:75c04502feb8ccab4c2e7222175d97583ad1471d53f08eb9dad363767fc3acbf","observation_id":"c42c89f3-ab3b-4ea7-b902-ab61781e039f","resolution":{"observed_at":"2026-08-06T04:46:40.416426Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"1707.06347","last_updated":"2017-08-28T09:20:06Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2017-07-20T02:32:33Z","title":"Proximal Policy Optimization Algorithms","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"1707.06347","snapshot_observed_at":"2026-08-06T04:46:39.890464Z","title":null,"venue":null,"work_id":null,"year":2017},"citing_paper":{"arxiv_id":"2608.05126","last_updated":"2026-08-05T17:50:31Z","snapshot_observed_at":"2026-08-10T00:52:03.795266Z","submitted_at":"2026-08-05T17:50:31Z","title":"Spoken Function Calling: A New Perspective on Spoken Language Understanding for Large Audio Language Models","version":1},"reference_index":2017,"source":"pdf_text","source_observed_at":"2026-08-06T04:46:39.890464Z"},"links":{"cited_paper":"/paper/1707.06347","citing_paper":"/paper/2608.05126"},"observation_digest":"sha256:1887f397afbdc9773699cd15359bad610fc0af20b66be5ccdc188174de69d6f2","observation_id":"708e4085-8e2e-4f16-9509-9f0e895d3cbc","resolution":{"observed_at":"2026-08-06T04:46:39.890464Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T04:46:40.710692Z","title":null,"venue":null,"work_id":"ff11e4bf-cd15-404d-b61c-343a58ded0c4","year":null},"citing_paper":{"arxiv_id":"2608.05126","last_updated":"2026-08-05T17:50:31Z","snapshot_observed_at":"2026-08-10T00:52:03.795266Z","submitted_at":"2026-08-05T17:50:31Z","title":"Spoken Function Calling: A New Perspective on Spoken Language Understanding for Large Audio Language Models","version":1},"reference_index":2020,"source":"pdf_text","source_observed_at":"2026-08-06T04:46:37.475471Z"},"links":{"citing_paper":"/paper/2608.05126"},"observation_digest":"sha256:e911adfb32246fbc4e1959aa666f07126347685a1606f822d13c2e08462e056d","observation_id":"ba1fb315-acf5-4c05-9c16-be0a8fe839b5","resolution":{"observed_at":"2026-08-06T04:46:40.714262Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T04:46:40.516235Z","title":null,"venue":null,"work_id":"f41f660e-9005-4d28-907a-a67474574e6a","year":null},"citing_paper":{"arxiv_id":"2608.05126","last_updated":"2026-08-05T17:50:31Z","snapshot_observed_at":"2026-08-10T00:52:03.795266Z","submitted_at":"2026-08-05T17:50:31Z","title":"Spoken Function Calling: A New Perspective on Spoken Language Understanding for Large Audio Language Models","version":1},"reference_index":2021,"source":"pdf_text","source_observed_at":"2026-08-06T04:46:39.615862Z"},"links":{"citing_paper":"/paper/2608.05126"},"observation_digest":"sha256:61928a45cbab987bd2061127e36e34c224d6e6a5bc72708c1b741c3e10f2c557","observation_id":"8110ddcc-9896-4428-8428-62729b96b549","resolution":{"observed_at":"2026-08-06T04:46:40.519500Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T04:46:40.619774Z","title":null,"venue":null,"work_id":"c9bea1c5-a1f8-4c26-b58a-f231b995da2f","year":null},"citing_paper":{"arxiv_id":"2608.05126","last_updated":"2026-08-05T17:50:31Z","snapshot_observed_at":"2026-08-10T00:52:03.795266Z","submitted_at":"2026-08-05T17:50:31Z","title":"Spoken Function Calling: A New Perspective on Spoken Language Understanding for Large Audio Language Models","version":1},"reference_index":2023,"source":"pdf_text","source_observed_at":"2026-08-06T04:46:38.054926Z"},"links":{"citing_paper":"/paper/2608.05126"},"observation_digest":"sha256:d42ec653cf7492b04c1e9d50b5114eaa5ae4be1c05c579cf48e309c2aff545a7","observation_id":"7a241a35-f0c8-4233-b37f-328520f690ad","resolution":{"observed_at":"2026-08-06T04:46:40.622925Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}}],"paper":{"arxiv_id":"2608.05126","last_updated":"2026-08-05T17:50:31Z","latest_version":1,"primary_category":"cs.CL","snapshot_observed_at":"2026-08-10T00:52:03.795266Z","submitted_at":"2026-08-05T17:50:31Z","title":"Spoken Function Calling: A New Perspective on Spoken Language Understanding for Large Audio Language Models"},"reference_resolution":{"displayed":51,"state_counts":{"malformed_identifier":0,"metadata_mismatch":0,"parse_uncertain":0,"unresolved":48,"verified_exact":1,"verified_fuzzy":2},"total_outbound_references":51},"refusal":"A citation records a reference. It does not transfer a finding from one paper to another.","schema":"pith.paper-citation-record.v1","standing_sources":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"thesis":"As of 10 August 2026, this Paper Citation Record lists 51 of 51 outbound references and 0 inbound Pith citation observations for arXiv:2608.05126."}