{"as_of":"2026-08-11T00:10:00Z","caps":{"database_statements":6,"inbound":100,"outbound":100},"context_digest":"sha256:3a4bb60e45cf1c916797427fb8fb549d024b2dd939e284613999eebf2a0ed060","coverage":[{"denominator":0,"lane":"reference_resolution","note":"Typed states for the displayed outbound observations.","records_observed":0,"source":"paper_references, paper_reference_links","source_observed_at":null,"state":"measured"},{"denominator":14,"lane":"standing_notices","note":"One-hop event checks from named stored sources.","records_observed":14,"source":"scholarly_work_events, retraction_status_cache","source_observed_at":"2026-08-10T06:31:04.303077+00:00","state":"measured"},{"denominator":14,"lane":"inbound_itemization","note":"Pith citing papers itemized under the disclosed page cap.","records_observed":14,"source":"paper_references, paper_reference_links","source_observed_at":"2026-08-07T00:33:19.577773Z","state":"measured"},{"denominator":1,"lane":"external_citation_measurements","note":"A source-named dated measurement, never combined with another source.","records_observed":1,"source":"arxiv_reference","source_observed_at":"2026-08-05T02:28:24.338817Z","state":"measured"}],"external_citation_measurements":[{"count":1,"observed_at":"2026-08-05T02:28:24.338817Z","source":"arxiv_reference"}],"inbound":[{"citation":{"cited_paper":{"arxiv_id":"2505.02625","last_updated":"2025-05-05T12:53:09Z","snapshot_observed_at":"2026-08-07T15:56:21.924963Z","submitted_at":"2025-05-05T12:53:09Z","title":"LLaMA-Omni2: LLM-based Real-time Spoken Chatbot with Autoregressive Streaming Speech Synthesis","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2505.02625","snapshot_observed_at":"2026-08-07T00:33:19.577773Z","title":"Llama-omni2: Llm- based real-time spoken chatbot with autoregressive streaming speech synthesis, 2025","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2506.13642","last_updated":"2025-06-22T07:56:58Z","snapshot_observed_at":"2026-08-10T10:57:00.100394Z","submitted_at":"2025-06-16T16:06:45Z","title":"Stream-Omni: Simultaneous Multimodal Interactions with Large Language-Vision-Speech Model","version":2},"reference_index":72,"source":"pdf_text","source_observed_at":"2026-08-07T00:33:19.577773Z"},"links":{"cited_paper":"/paper/2505.02625","citing_paper":"/paper/2506.13642"},"observation_digest":"sha256:725ce5b87fedba904165fc3cfe0db9141921b51744205348969eeb86ebac1c0e","observation_id":"f5e67540-2a28-4247-9b5c-61fabd5ceea5","resolution":{"observed_at":"2026-08-07T00:33:19.577773Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2505.02625","last_updated":"2025-05-05T12:53:09Z","snapshot_observed_at":"2026-08-07T15:56:21.924963Z","submitted_at":"2025-05-05T12:53:09Z","title":"LLaMA-Omni2: LLM-based Real-time Spoken Chatbot with Autoregressive Streaming Speech Synthesis","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2505.02625","snapshot_observed_at":"2026-08-05T15:52:43.998471Z","title":"Llama-omni2: Llm- based real-time spoken chatbot with autoregressive streaming speech synthesis,","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2509.00078","last_updated":"2025-08-26T20:40:24Z","snapshot_observed_at":"2026-08-05T15:52:43.760838Z","submitted_at":"2025-08-26T20:40:24Z","title":"ChipChat: Low-Latency Cascaded Conversational Agent in MLX","version":1},"reference_index":8,"source":"pdf_text","source_observed_at":"2026-08-05T15:52:43.998471Z"},"links":{"cited_paper":"/paper/2505.02625","citing_paper":"/paper/2509.00078"},"observation_digest":"sha256:4f66400e7182d8bcbff969c5ee21a0e71d3037f3c6fba07fa49e856b6900fff0","observation_id":"dc1e401d-6272-4516-afe5-6481c35deecb","resolution":{"observed_at":"2026-08-05T15:52:43.998471Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2505.02625","last_updated":"2025-05-05T12:53:09Z","snapshot_observed_at":"2026-08-07T15:56:21.924963Z","submitted_at":"2025-05-05T12:53:09Z","title":"LLaMA-Omni2: LLM-based Real-time Spoken Chatbot with Autoregressive Streaming Speech Synthesis","version":1},"cited_work":{"arxiv_id":"2505.02625","doi":"10.48550/arxiv.2505.02625","metadata_source":"arxiv_reference","pith_arxiv_id":"2505.02625","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Llama-omni2: Llm-based real-time spoken chatbot with autoregressive streaming speech synthesis","venue":"ArXiv.org","work_id":"ebfa628e-d823-44f5-a686-0ac0a7d5f4ba","year":2025},"citing_paper":{"arxiv_id":"2509.22220","last_updated":"2026-04-13T11:56:11Z","snapshot_observed_at":"2026-08-02T12:47:50.534468Z","submitted_at":"2025-09-26T11:32:51Z","title":"StableToken: A Noise-Robust Semantic Speech Tokenizer for Resilient SpeechLLMs","version":2},"reference_index":27,"source":"arxiv_source","source_observed_at":"2026-05-18T12:57:04.450462Z"},"links":{"cited_paper":"/paper/2505.02625","citing_paper":"/paper/2509.22220"},"observation_digest":"sha256:e2b8f7355da1902420eefa91c08b1101f6b8f09ea90c1cfb2b0eaee37c5d0809","observation_id":"1497d69c-0b13-4a0c-8c71-a5d4043ee17c","resolution":{"observed_at":"2026-05-18T13:01:24.365147Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2505.02625","last_updated":"2025-05-05T12:53:09Z","snapshot_observed_at":"2026-08-07T15:56:21.924963Z","submitted_at":"2025-05-05T12:53:09Z","title":"LLaMA-Omni2: LLM-based Real-time Spoken Chatbot with Autoregressive Streaming Speech Synthesis","version":1},"cited_work":{"arxiv_id":"2505.02625","doi":"10.48550/arxiv.2505.02625","metadata_source":"arxiv_reference","pith_arxiv_id":"2505.02625","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Llama-omni2: Llm-based real-time spoken chatbot with autoregressive streaming speech synthesis","venue":"ArXiv.org","work_id":"ebfa628e-d823-44f5-a686-0ac0a7d5f4ba","year":2025},"citing_paper":{"arxiv_id":"2510.09592","last_updated":"2026-05-10T15:21:35Z","snapshot_observed_at":"2026-07-06T22:32:22.953630Z","submitted_at":"2025-10-10T17:50:59Z","title":"Mind-Paced Speaking: A Dual-Brain Approach to Real-Time Reasoning in Spoken Language Models","version":2},"reference_index":18,"source":"pdf_text","source_observed_at":"2026-05-18T07:43:23.913399Z"},"links":{"cited_paper":"/paper/2505.02625","citing_paper":"/paper/2510.09592"},"observation_digest":"sha256:26d0cd86e5dbd7d7d09f7b2a39bd179c6dbc370d79859771e9072652bf0ff447","observation_id":"d503e0da-a3f1-44fe-a927-eed086ae7943","resolution":{"observed_at":"2026-05-18T07:46:03.589866Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2505.02625","last_updated":"2025-05-05T12:53:09Z","snapshot_observed_at":"2026-08-07T15:56:21.924963Z","submitted_at":"2025-05-05T12:53:09Z","title":"LLaMA-Omni2: LLM-based Real-time Spoken Chatbot with Autoregressive Streaming Speech Synthesis","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2505.02625","snapshot_observed_at":"2026-08-03T12:25:01.710490Z","title":null,"venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2601.03190","last_updated":"2026-07-23T19:26:16Z","snapshot_observed_at":"2026-08-09T14:57:40.708124Z","submitted_at":"2026-01-06T17:10:48Z","title":"Maximizing Local Entropy Where It Matters: Prefix-Aware Localized LLM Unlearning","version":4},"reference_index":12,"source":"arxiv_source","source_observed_at":"2026-08-03T12:25:01.710490Z"},"links":{"cited_paper":"/paper/2505.02625","citing_paper":"/paper/2601.03190"},"observation_digest":"sha256:192178cf2e9522f7d05b5b7316ff8649513358ed5b122ff268af62d9c8ad5bad","observation_id":"0b2bf6c6-73b0-46cb-9d81-f1357ac89c8f","resolution":{"observed_at":"2026-08-03T12:25:01.710490Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2505.02625","last_updated":"2025-05-05T12:53:09Z","snapshot_observed_at":"2026-08-07T15:56:21.924963Z","submitted_at":"2025-05-05T12:53:09Z","title":"LLaMA-Omni2: LLM-based Real-time Spoken Chatbot with Autoregressive Streaming Speech Synthesis","version":1},"cited_work":{"arxiv_id":"2505.02625","doi":"10.48550/arxiv.2505.02625","metadata_source":"arxiv_reference","pith_arxiv_id":"2505.02625","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Llama-omni2: Llm-based real-time spoken chatbot with autoregressive streaming speech synthesis","venue":"ArXiv.org","work_id":"ebfa628e-d823-44f5-a686-0ac0a7d5f4ba","year":2025},"citing_paper":{"arxiv_id":"2604.13804","last_updated":"2026-04-15T12:39:03Z","snapshot_observed_at":"2026-07-06T23:01:41.643335Z","submitted_at":"2026-04-15T12:39:03Z","title":"Character Beyond Speech: Leveraging Role-Playing Evaluation in Audio Large Language Models via Reinforcement Learning","version":1},"reference_index":12,"source":"pdf_text","source_observed_at":"2026-05-10T14:22:25.660785Z"},"links":{"cited_paper":"/paper/2505.02625","citing_paper":"/paper/2604.13804"},"observation_digest":"sha256:f032a0f6b08b95791069c4e5c7981afb7cc3782507c1a86fd7b4242c94dddea1","observation_id":"2c54e787-bd4b-49e5-8bd5-f205b4baa4a4","resolution":{"observed_at":"2026-05-10T14:25:30.047240Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2505.02625","last_updated":"2025-05-05T12:53:09Z","snapshot_observed_at":"2026-08-07T15:56:21.924963Z","submitted_at":"2025-05-05T12:53:09Z","title":"LLaMA-Omni2: LLM-based Real-time Spoken Chatbot with Autoregressive Streaming Speech Synthesis","version":1},"cited_work":{"arxiv_id":"2505.02625","doi":"10.48550/arxiv.2505.02625","metadata_source":"arxiv_reference","pith_arxiv_id":"2505.02625","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Llama-omni2: Llm-based real-time spoken chatbot with autoregressive streaming speech synthesis","venue":"ArXiv.org","work_id":"ebfa628e-d823-44f5-a686-0ac0a7d5f4ba","year":2025},"citing_paper":{"arxiv_id":"2604.19300","last_updated":"2026-04-21T10:05:28Z","snapshot_observed_at":"2026-08-03T21:49:52.344628Z","submitted_at":"2026-04-21T10:05:28Z","title":"HalluAudio: A Comprehensive Benchmark for Hallucination Detection in Large Audio-Language Models","version":1},"reference_index":8,"source":"pdf_text","source_observed_at":"2026-05-10T01:34:54.375266Z"},"links":{"cited_paper":"/paper/2505.02625","citing_paper":"/paper/2604.19300"},"observation_digest":"sha256:312a4fa8cb9310a71bfa1b42cfee7116a40ca4cccc12816913a9f5017a265e28","observation_id":"5ff634e1-5961-4926-8e3d-76138e3aa974","resolution":{"observed_at":"2026-05-11T13:31:03.386871Z","resolver_source":"arxiv_id","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2505.02625","last_updated":"2025-05-05T12:53:09Z","snapshot_observed_at":"2026-08-07T15:56:21.924963Z","submitted_at":"2025-05-05T12:53:09Z","title":"LLaMA-Omni2: LLM-based Real-time Spoken Chatbot with Autoregressive Streaming Speech Synthesis","version":1},"cited_work":{"arxiv_id":"2505.02625","doi":"10.48550/arxiv.2505.02625","metadata_source":"arxiv_reference","pith_arxiv_id":"2505.02625","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Llama-omni2: Llm-based real-time spoken chatbot with autoregressive streaming speech synthesis","venue":"ArXiv.org","work_id":"ebfa628e-d823-44f5-a686-0ac0a7d5f4ba","year":2025},"citing_paper":{"arxiv_id":"2604.20842","last_updated":"2026-04-22T17:59:58Z","snapshot_observed_at":"2026-08-04T16:34:40.989674Z","submitted_at":"2026-04-22T17:59:58Z","title":"SpeechParaling-Bench: A Comprehensive Benchmark for Paralinguistic-Aware Speech Generation","version":1},"reference_index":11,"source":"pdf_text","source_observed_at":"2026-05-10T00:39:04.303837Z"},"links":{"cited_paper":"/paper/2505.02625","citing_paper":"/paper/2604.20842"},"observation_digest":"sha256:26a1116f7d5222f8bd27ddd34ed5c57432e63c24b88ffe705d67de938d089393","observation_id":"6c489d98-1bb1-4ec1-8352-cbfaa2e5f95b","resolution":{"observed_at":"2026-05-10T00:39:48.480242Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2505.02625","last_updated":"2025-05-05T12:53:09Z","snapshot_observed_at":"2026-08-07T15:56:21.924963Z","submitted_at":"2025-05-05T12:53:09Z","title":"LLaMA-Omni2: LLM-based Real-time Spoken Chatbot with Autoregressive Streaming Speech Synthesis","version":1},"cited_work":{"arxiv_id":"2505.02625","doi":"10.48550/arxiv.2505.02625","metadata_source":"arxiv_reference","pith_arxiv_id":"2505.02625","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Llama-omni2: Llm-based real-time spoken chatbot with autoregressive streaming speech synthesis","venue":"ArXiv.org","work_id":"ebfa628e-d823-44f5-a686-0ac0a7d5f4ba","year":2025},"citing_paper":{"arxiv_id":"2606.01016","last_updated":"2026-05-31T05:13:32Z","snapshot_observed_at":"2026-07-06T23:41:39.172310Z","submitted_at":"2026-05-31T05:13:32Z","title":"PolySpeech-100: A Large-Scale Benchmark for Speech Understanding Across 100+ Languages and Dialects","version":1},"reference_index":35,"source":"pdf_text","source_observed_at":"2026-06-28T17:44:07.669223Z"},"links":{"cited_paper":"/paper/2505.02625","citing_paper":"/paper/2606.01016"},"observation_digest":"sha256:2bf3a75a324e356abc77b107cb2ec7961e52d48d5114ff0673a3c0177dec606a","observation_id":"9d99baf3-d54d-4dcc-af2b-3f8950b489ac","resolution":{"observed_at":"2026-06-28T17:52:26.993555Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2505.02625","last_updated":"2025-05-05T12:53:09Z","snapshot_observed_at":"2026-08-07T15:56:21.924963Z","submitted_at":"2025-05-05T12:53:09Z","title":"LLaMA-Omni2: LLM-based Real-time Spoken Chatbot with Autoregressive Streaming Speech Synthesis","version":1},"cited_work":{"arxiv_id":"2505.02625","doi":"10.48550/arxiv.2505.02625","metadata_source":"arxiv_reference","pith_arxiv_id":"2505.02625","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Llama-omni2: Llm-based real-time spoken chatbot with autoregressive streaming speech synthesis","venue":"ArXiv.org","work_id":"ebfa628e-d823-44f5-a686-0ac0a7d5f4ba","year":2025},"citing_paper":{"arxiv_id":"2606.05121","last_updated":"2026-06-03T17:26:11Z","snapshot_observed_at":"2026-08-02T17:56:34.317060Z","submitted_at":"2026-06-03T17:26:11Z","title":"Audio Interaction Model","version":1},"reference_index":12,"source":"pdf_text","source_observed_at":"2026-06-28T04:57:05.062465Z"},"links":{"cited_paper":"/paper/2505.02625","citing_paper":"/paper/2606.05121"},"observation_digest":"sha256:264cb38c17cb0de98d573509762bddafd819b8963ed5313c4da89f8f0152cfff","observation_id":"023d51a6-739e-4ff4-b00a-dabae30b93a1","resolution":{"observed_at":"2026-07-02T10:46:52.382243Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2505.02625","last_updated":"2025-05-05T12:53:09Z","snapshot_observed_at":"2026-08-07T15:56:21.924963Z","submitted_at":"2025-05-05T12:53:09Z","title":"LLaMA-Omni2: LLM-based Real-time Spoken Chatbot with Autoregressive Streaming Speech Synthesis","version":1},"cited_work":{"arxiv_id":"2505.02625","doi":"10.48550/arxiv.2505.02625","metadata_source":"arxiv_reference","pith_arxiv_id":"2505.02625","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Llama-omni2: Llm-based real-time spoken chatbot with autoregressive streaming speech synthesis","venue":"ArXiv.org","work_id":"ebfa628e-d823-44f5-a686-0ac0a7d5f4ba","year":2025},"citing_paper":{"arxiv_id":"2606.07433","last_updated":"2026-06-05T16:29:13Z","snapshot_observed_at":"2026-08-01T21:05:06.439607Z","submitted_at":"2026-06-05T16:29:13Z","title":"Watch, Remember, Reason: Human-View Video Understanding with MLLMs","version":1},"reference_index":127,"source":"pdf_text","source_observed_at":"2026-06-27T22:00:28.350003Z"},"links":{"cited_paper":"/paper/2505.02625","citing_paper":"/paper/2606.07433"},"observation_digest":"sha256:59ce2c94ba18b11f8ffc60c37dfe4dfc4eef91e3efeb92fbc0210d5ad87849db","observation_id":"0f7c03a4-2614-492e-96a3-577d3af0b60c","resolution":{"observed_at":"2026-07-02T17:27:14.747442Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2505.02625","last_updated":"2025-05-05T12:53:09Z","snapshot_observed_at":"2026-08-07T15:56:21.924963Z","submitted_at":"2025-05-05T12:53:09Z","title":"LLaMA-Omni2: LLM-based Real-time Spoken Chatbot with Autoregressive Streaming Speech Synthesis","version":1},"cited_work":{"arxiv_id":"2505.02625","doi":"10.48550/arxiv.2505.02625","metadata_source":"arxiv_reference","pith_arxiv_id":"2505.02625","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Llama-omni2: Llm-based real-time spoken chatbot with autoregressive streaming speech synthesis","venue":"ArXiv.org","work_id":"ebfa628e-d823-44f5-a686-0ac0a7d5f4ba","year":2025},"citing_paper":{"arxiv_id":"2606.25444","last_updated":"2026-06-24T06:15:18Z","snapshot_observed_at":"2026-08-02T19:50:32.355214Z","submitted_at":"2026-06-24T06:15:18Z","title":"Does Translation-Enhanced Speech Encoder Pre-training Affect Speech LLMs?","version":1},"reference_index":26,"source":"pdf_text","source_observed_at":"2026-06-25T20:03:10.858349Z"},"links":{"cited_paper":"/paper/2505.02625","citing_paper":"/paper/2606.25444"},"observation_digest":"sha256:56d06100484255ed161f6929d2fefd0b3deff043c247af092d994805d2813b73","observation_id":"1347e7a1-5d37-48c1-9d6b-1b492fd9f58d","resolution":{"observed_at":"2026-07-04T20:30:08.189345Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2505.02625","last_updated":"2025-05-05T12:53:09Z","snapshot_observed_at":"2026-08-07T15:56:21.924963Z","submitted_at":"2025-05-05T12:53:09Z","title":"LLaMA-Omni2: LLM-based Real-time Spoken Chatbot with Autoregressive Streaming Speech Synthesis","version":1},"cited_work":{"arxiv_id":"2505.02625","doi":"10.48550/arxiv.2505.02625","metadata_source":"arxiv_reference","pith_arxiv_id":"2505.02625","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Llama-omni2: Llm-based real-time spoken chatbot with autoregressive streaming speech synthesis","venue":"ArXiv.org","work_id":"ebfa628e-d823-44f5-a686-0ac0a7d5f4ba","year":2025},"citing_paper":{"arxiv_id":"2606.30944","last_updated":"2026-06-29T21:55:18Z","snapshot_observed_at":"2026-08-07T21:37:48.909123Z","submitted_at":"2026-06-29T21:55:18Z","title":"Preserving Speech-to-Text LLM Capabilities in Speech-to-Speech Generation","version":1},"reference_index":29,"source":"pdf_text","source_observed_at":"2026-07-01T01:01:24.536821Z"},"links":{"cited_paper":"/paper/2505.02625","citing_paper":"/paper/2606.30944"},"observation_digest":"sha256:82efc7b2ccfdc08b0d3c2dfcf28c42d7af8cdcac45c7a66ea909779489b2002c","observation_id":"b9d06f2e-9fcc-459b-87e8-153712bfeec4","resolution":{"observed_at":"2026-07-01T13:05:45.816779Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2505.02625","last_updated":"2025-05-05T12:53:09Z","snapshot_observed_at":"2026-08-07T15:56:21.924963Z","submitted_at":"2025-05-05T12:53:09Z","title":"LLaMA-Omni2: LLM-based Real-time Spoken Chatbot with Autoregressive Streaming Speech Synthesis","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2505.02625","snapshot_observed_at":"2026-08-01T07:12:17.669270Z","title":"arXiv preprint arXiv:2505.02625 , year=","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2607.21550","last_updated":"2026-07-23T17:35:20Z","snapshot_observed_at":"2026-08-08T08:48:36.880078Z","submitted_at":"2026-07-23T17:35:20Z","title":"X$^3$-OPD: Distilling Reasoning into Large Audio-Language Models via On-Policy Alignment","version":1},"reference_index":139,"source":"arxiv_source","source_observed_at":"2026-08-01T07:12:17.669270Z"},"links":{"cited_paper":"/paper/2505.02625","citing_paper":"/paper/2607.21550"},"observation_digest":"sha256:d5e77ea54262b55369f553366d7e9e9e60de2884d394de3ee034d451262e1571","observation_id":"93186e2e-187a-4d75-a08c-e04e9e56982b","resolution":{"observed_at":"2026-08-01T07:12:17.669270Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}}],"links":{"evidence":"/evidence","html":"/paper/2505.02625/citation-record","integrity":"/paper/2505.02625/integrity","json":"/paper/2505.02625/citation-record.json","paper":"/paper/2505.02625"},"outbound":[],"paper":{"arxiv_id":"2505.02625","last_updated":"2025-05-05T12:53:09Z","latest_version":1,"primary_category":"cs.CL","snapshot_observed_at":"2026-08-07T15:56:21.924963Z","submitted_at":"2025-05-05T12:53:09Z","title":"LLaMA-Omni2: LLM-based Real-time Spoken Chatbot with Autoregressive Streaming Speech Synthesis"},"reference_resolution":{"displayed":0,"state_counts":{"malformed_identifier":0,"metadata_mismatch":0,"parse_uncertain":0,"unresolved":0,"verified_exact":0,"verified_fuzzy":0},"total_outbound_references":0},"refusal":"A citation records a reference. It does not transfer a finding from one paper to another.","schema":"pith.paper-citation-record.v1","standing_sources":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"thesis":"As of 11 August 2026, this Paper Citation Record lists 0 of 0 outbound references and 14 inbound Pith citation observations for arXiv:2505.02625."}