{"as_of":"2026-08-09T07:47:00Z","caps":{"database_statements":6,"inbound":100,"outbound":100},"context_digest":"sha256:2579f05b0bfef2e114542efd00e1077dc69d3cb7db75c3bbef1136347882d5ea","coverage":[{"denominator":0,"lane":"reference_resolution","note":"Typed states for the displayed outbound observations.","records_observed":0,"source":"paper_references, paper_reference_links","source_observed_at":null,"state":"measured"},{"denominator":22,"lane":"standing_notices","note":"One-hop event checks from named stored sources.","records_observed":22,"source":"scholarly_work_events, retraction_status_cache","source_observed_at":"2026-08-09T06:31:02.800959+00:00","state":"measured"},{"denominator":22,"lane":"inbound_itemization","note":"Pith citing papers itemized under the disclosed page cap.","records_observed":22,"source":"paper_references, paper_reference_links","source_observed_at":"2026-08-06T19:20:50.875290Z","state":"measured"},{"denominator":1,"lane":"external_citation_measurements","note":"A source-named dated measurement, never combined with another source.","records_observed":0,"source":"arxiv_reference","source_observed_at":"2026-07-04T08:29:41.902212Z","state":"measured"}],"external_citation_measurements":[],"inbound":[{"citation":{"cited_paper":{"arxiv_id":"2107.03312","last_updated":"2021-07-07T15:45:42Z","snapshot_observed_at":"2026-08-09T07:05:32.492915Z","submitted_at":"2021-07-07T15:45:42Z","title":"SoundStream: An End-to-End Neural Audio Codec","version":1},"cited_work":{"arxiv_id":"2107.03312","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2107.03312","snapshot_observed_at":"2026-07-04T08:29:41.902212Z","title":"Soundstream: An end-to-end neural audio codec","venue":null,"work_id":"e2f34c6d-cf73-4f48-a80a-732b812cc299","year":2021},"citing_paper":{"arxiv_id":"2501.09747","last_updated":"2025-01-16T18:57:04Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2025-01-16T18:57:04Z","title":"FAST: Efficient Action Tokenization for Vision-Language-Action Models","version":1},"reference_index":69,"source":"pdf_text","source_observed_at":"2026-05-11T08:52:31.686474Z"},"links":{"cited_paper":"/paper/2107.03312","citing_paper":"/paper/2501.09747"},"observation_digest":"sha256:3d2db63fe9f8267be2b73d7c9b770ef7ff580363c7919a09153aa59f249df1cc","observation_id":"79955a72-0acc-47f7-a02a-8a740e4e8185","resolution":{"observed_at":"2026-05-11T08:52:32.118854Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2107.03312","last_updated":"2021-07-07T15:45:42Z","snapshot_observed_at":"2026-08-09T07:05:32.492915Z","submitted_at":"2021-07-07T15:45:42Z","title":"SoundStream: An End-to-End Neural Audio Codec","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2107.03312","snapshot_observed_at":"2026-08-06T19:20:50.875290Z","title":"Soundstream: An end-to-end neural audio codec,","venue":null,"work_id":null,"year":2021},"citing_paper":{"arxiv_id":"2507.06040","last_updated":"2025-07-08T14:41:42Z","snapshot_observed_at":"2026-08-09T07:05:59.481970Z","submitted_at":"2025-07-08T14:41:42Z","title":"EdgeCodec: Onboard Lightweight High Fidelity Neural Compressor with Residual Vector Quantization","version":1},"reference_index":4,"source":"pdf_text","source_observed_at":"2026-08-06T19:20:50.875290Z"},"links":{"cited_paper":"/paper/2107.03312","citing_paper":"/paper/2507.06040"},"observation_digest":"sha256:f5d93a0e8e11257ac20de44368102d8b747ac1a99018c655a4ee0d082c4207ff","observation_id":"47a07caf-a475-4e70-85fe-b4cde9d82706","resolution":{"observed_at":"2026-08-06T19:20:50.875290Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2107.03312","last_updated":"2021-07-07T15:45:42Z","snapshot_observed_at":"2026-08-09T07:05:32.492915Z","submitted_at":"2021-07-07T15:45:42Z","title":"SoundStream: An End-to-End Neural Audio Codec","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2107.03312","snapshot_observed_at":"2026-08-06T18:28:30.286712Z","title":"Zeghidour, A","venue":null,"work_id":null,"year":2021},"citing_paper":{"arxiv_id":"2507.08236","last_updated":"2025-07-11T00:47:08Z","snapshot_observed_at":"2026-08-09T07:05:55.440219Z","submitted_at":"2025-07-11T00:47:08Z","title":"Distilling Spectrograms into Tokens: Fast and Lightweight Bioacoustic Classification for BirdCLEF+ 2025","version":1},"reference_index":7,"source":"pdf_text","source_observed_at":"2026-08-06T18:28:30.286712Z"},"links":{"cited_paper":"/paper/2107.03312","citing_paper":"/paper/2507.08236"},"observation_digest":"sha256:fca513e19db7d878372ec8b629e35616e17422946c90dbe13c07aaa121deaee8","observation_id":"c2fefea6-c362-4b77-a9d6-e5da99af7fa5","resolution":{"observed_at":"2026-08-06T18:28:30.286712Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2107.03312","last_updated":"2021-07-07T15:45:42Z","snapshot_observed_at":"2026-08-09T07:05:32.492915Z","submitted_at":"2021-07-07T15:45:42Z","title":"SoundStream: An End-to-End Neural Audio Codec","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2107.03312","snapshot_observed_at":"2026-08-05T14:36:37.718921Z","title":"SoundStream: An End-to-End Neural Audio Codec,","venue":null,"work_id":null,"year":2021},"citing_paper":{"arxiv_id":"2508.21153","last_updated":"2025-08-28T18:38:42Z","snapshot_observed_at":"2026-08-07T14:49:09.772290Z","submitted_at":"2025-08-28T18:38:42Z","title":"WaveLLDM: Design and Development of a Lightweight Latent Diffusion Model for Speech Enhancement and Restoration","version":1},"reference_index":59,"source":"pdf_text","source_observed_at":"2026-08-05T14:36:37.718921Z"},"links":{"cited_paper":"/paper/2107.03312","citing_paper":"/paper/2508.21153"},"observation_digest":"sha256:dc5ee79d6a55f6515025a5f67c2ad51d07929dd49bba714c293e88e7b0356347","observation_id":"f8aa1a5a-b4ef-4032-9bd8-cfdf23931ef9","resolution":{"observed_at":"2026-08-05T14:36:37.718921Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2107.03312","last_updated":"2021-07-07T15:45:42Z","snapshot_observed_at":"2026-08-09T07:05:32.492915Z","submitted_at":"2021-07-07T15:45:42Z","title":"SoundStream: An End-to-End Neural Audio Codec","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2107.03312","snapshot_observed_at":"2026-08-05T11:41:03.517998Z","title":null,"venue":null,"work_id":null,"year":2021},"citing_paper":{"arxiv_id":"2509.02349","last_updated":"2025-09-04T14:25:57Z","snapshot_observed_at":"2026-08-09T00:24:32.254551Z","submitted_at":"2025-09-02T14:15:22Z","title":"AudioCodecBench: A Comprehensive Benchmark for Audio Codec Evaluation","version":2},"reference_index":49,"source":"arxiv_source","source_observed_at":"2026-08-05T11:41:03.517998Z"},"links":{"cited_paper":"/paper/2107.03312","citing_paper":"/paper/2509.02349"},"observation_digest":"sha256:f99eba64eb2683225f73089d2f8813cd89755400ebc50d7a449a8b6090f0b34e","observation_id":"f7a8ed48-8e2f-4b77-985e-251722f8d049","resolution":{"observed_at":"2026-08-05T11:41:03.517998Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2107.03312","last_updated":"2021-07-07T15:45:42Z","snapshot_observed_at":"2026-08-09T07:05:32.492915Z","submitted_at":"2021-07-07T15:45:42Z","title":"SoundStream: An End-to-End Neural Audio Codec","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2107.03312","snapshot_observed_at":"2026-08-05T11:29:01.660239Z","title":"Zeghidour, A","venue":null,"work_id":null,"year":2021},"citing_paper":{"arxiv_id":"2509.02771","last_updated":"2025-09-02T19:20:06Z","snapshot_observed_at":"2026-08-09T07:32:23.956379Z","submitted_at":"2025-09-02T19:20:06Z","title":"Analysis of Speaker Verification Performance Trade-offs with Neural Audio Codec Transmission","version":1},"reference_index":3,"source":"pdf_text","source_observed_at":"2026-08-05T11:29:01.660239Z"},"links":{"cited_paper":"/paper/2107.03312","citing_paper":"/paper/2509.02771"},"observation_digest":"sha256:f517783ac2110578dcb71090574253ff74a218bedf55c08109b6b330435b747e","observation_id":"a76137b1-16cc-4475-9fc3-80d93e4a3337","resolution":{"observed_at":"2026-08-05T11:29:01.660239Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2107.03312","last_updated":"2021-07-07T15:45:42Z","snapshot_observed_at":"2026-08-09T07:05:32.492915Z","submitted_at":"2021-07-07T15:45:42Z","title":"SoundStream: An End-to-End Neural Audio Codec","version":1},"cited_work":{"arxiv_id":"2107.03312","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2107.03312","snapshot_observed_at":"2026-07-04T08:29:41.902212Z","title":"Soundstream: An end-to-end neural audio codec","venue":null,"work_id":"e2f34c6d-cf73-4f48-a80a-732b812cc299","year":2021},"citing_paper":{"arxiv_id":"2604.01929","last_updated":"2026-04-29T10:54:09Z","snapshot_observed_at":"2026-07-06T22:51:40.181565Z","submitted_at":"2026-04-02T11:49:00Z","title":"Woosh: A Sound Effects Foundation Model","version":3},"reference_index":21,"source":"pdf_text","source_observed_at":"2026-05-13T20:51:08.144573Z"},"links":{"cited_paper":"/paper/2107.03312","citing_paper":"/paper/2604.01929"},"observation_digest":"sha256:5cb3277c788048e9a3c1bb97ba8ba0a9611805b2af2bfdfeb9a18600a8df9f59","observation_id":"797383dd-4533-42af-9127-d130d7dcafc7","resolution":{"observed_at":"2026-05-13T20:53:15.789605Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2107.03312","last_updated":"2021-07-07T15:45:42Z","snapshot_observed_at":"2026-08-09T07:05:32.492915Z","submitted_at":"2021-07-07T15:45:42Z","title":"SoundStream: An End-to-End Neural Audio Codec","version":1},"cited_work":{"arxiv_id":"2107.03312","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2107.03312","snapshot_observed_at":"2026-07-04T08:29:41.902212Z","title":"Soundstream: An end-to-end neural audio codec","venue":null,"work_id":"e2f34c6d-cf73-4f48-a80a-732b812cc299","year":2021},"citing_paper":{"arxiv_id":"2604.12965","last_updated":"2026-04-14T16:59:03Z","snapshot_observed_at":"2026-08-06T12:38:28.688700Z","submitted_at":"2026-04-14T16:59:03Z","title":"Efficient Retrieval Scaling with Hierarchical Indexing for Large Scale Recommendation","version":1},"reference_index":68,"source":"pdf_text","source_observed_at":"2026-05-10T14:25:53.909672Z"},"links":{"cited_paper":"/paper/2107.03312","citing_paper":"/paper/2604.12965"},"observation_digest":"sha256:4c9cf5f50d6b0f026dc76fd9f3f64a0ccbdfee8782e2a3a4a4a4b670caf96d95","observation_id":"2b0019ce-021a-4bee-a712-34ed89a03939","resolution":{"observed_at":"2026-05-11T11:36:01.029236Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2107.03312","last_updated":"2021-07-07T15:45:42Z","snapshot_observed_at":"2026-08-09T07:05:32.492915Z","submitted_at":"2021-07-07T15:45:42Z","title":"SoundStream: An End-to-End Neural Audio Codec","version":1},"cited_work":{"arxiv_id":"2107.03312","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2107.03312","snapshot_observed_at":"2026-07-04T08:29:41.902212Z","title":"Soundstream: An end-to-end neural audio codec","venue":null,"work_id":"e2f34c6d-cf73-4f48-a80a-732b812cc299","year":2021},"citing_paper":{"arxiv_id":"2605.13789","last_updated":"2026-05-14T01:44:35Z","snapshot_observed_at":"2026-07-30T00:30:13.289148Z","submitted_at":"2026-05-13T17:08:41Z","title":"ENSEMBITS: an alphabet of protein conformational ensembles","version":1},"reference_index":28,"source":"pdf_text","source_observed_at":"2026-05-14T19:11:29.935218Z"},"links":{"cited_paper":"/paper/2107.03312","citing_paper":"/paper/2605.13789"},"observation_digest":"sha256:2e0a7a922625c723db159774ab8e2b11027b9bdec78f04bc1e1ae6410dddc2c9","observation_id":"cd7b13a5-b159-4f4f-b0c1-75ee70cb3f54","resolution":{"observed_at":"2026-05-14T19:12:50.731283Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2107.03312","last_updated":"2021-07-07T15:45:42Z","snapshot_observed_at":"2026-08-09T07:05:32.492915Z","submitted_at":"2021-07-07T15:45:42Z","title":"SoundStream: An End-to-End Neural Audio Codec","version":1},"cited_work":{"arxiv_id":"2107.03312","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2107.03312","snapshot_observed_at":"2026-07-04T08:29:41.902212Z","title":"Soundstream: An end-to-end neural audio codec","venue":null,"work_id":"e2f34c6d-cf73-4f48-a80a-732b812cc299","year":2021},"citing_paper":{"arxiv_id":"2605.13789","last_updated":"2026-05-14T01:44:35Z","snapshot_observed_at":"2026-07-30T00:30:13.289148Z","submitted_at":"2026-05-13T17:08:41Z","title":"ENSEMBITS: an alphabet of protein conformational ensembles","version":2},"reference_index":28,"source":"pdf_text","source_observed_at":"2026-05-15T04:51:32.821791Z"},"links":{"cited_paper":"/paper/2107.03312","citing_paper":"/paper/2605.13789"},"observation_digest":"sha256:a9edfc4dc7e1bc434acf72a025f56c59b9ba0cc9b821a05d781f29a0cad0ec40","observation_id":"732d25d3-0863-4f1e-a2f9-a5f8ba34b17e","resolution":{"observed_at":"2026-05-15T04:55:03.987291Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2107.03312","last_updated":"2021-07-07T15:45:42Z","snapshot_observed_at":"2026-08-09T07:05:32.492915Z","submitted_at":"2021-07-07T15:45:42Z","title":"SoundStream: An End-to-End Neural Audio Codec","version":1},"cited_work":{"arxiv_id":"2107.03312","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2107.03312","snapshot_observed_at":"2026-07-04T08:29:41.902212Z","title":"Soundstream: An end-to-end neural audio codec","venue":null,"work_id":"e2f34c6d-cf73-4f48-a80a-732b812cc299","year":2021},"citing_paper":{"arxiv_id":"2605.20519","last_updated":"2026-05-21T17:49:53Z","snapshot_observed_at":"2026-07-06T23:31:03.878152Z","submitted_at":"2026-05-19T21:39:52Z","title":"Codec-Robust Attacks on Audio LLMs","version":1},"reference_index":55,"source":"pdf_text","source_observed_at":"2026-05-21T06:43:52.735211Z"},"links":{"cited_paper":"/paper/2107.03312","citing_paper":"/paper/2605.20519"},"observation_digest":"sha256:a6193ee28fe945aa7152b8c4a9374dce3db7c83426ffd7d19e6a22e12aaba0cb","observation_id":"854eea0f-0743-49e4-8e56-6a0053862c56","resolution":{"observed_at":"2026-05-21T06:44:00.663100Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2107.03312","last_updated":"2021-07-07T15:45:42Z","snapshot_observed_at":"2026-08-09T07:05:32.492915Z","submitted_at":"2021-07-07T15:45:42Z","title":"SoundStream: An End-to-End Neural Audio Codec","version":1},"cited_work":{"arxiv_id":"2107.03312","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2107.03312","snapshot_observed_at":"2026-07-04T08:29:41.902212Z","title":"Soundstream: An end-to-end neural audio codec","venue":null,"work_id":"e2f34c6d-cf73-4f48-a80a-732b812cc299","year":2021},"citing_paper":{"arxiv_id":"2605.20519","last_updated":"2026-05-21T17:49:53Z","snapshot_observed_at":"2026-07-06T23:31:03.878152Z","submitted_at":"2026-05-19T21:39:52Z","title":"Codec-Robust Attacks on Audio LLMs","version":2},"reference_index":55,"source":"pdf_text","source_observed_at":"2026-05-25T05:44:46.831360Z"},"links":{"cited_paper":"/paper/2107.03312","citing_paper":"/paper/2605.20519"},"observation_digest":"sha256:9bcf280e52ad97aa5f83dda61064c3520a36cb31033b38ba01d227bd853f703d","observation_id":"592f59bb-4ce9-4807-b40e-97ec1e5107ac","resolution":{"observed_at":"2026-05-25T05:45:23.656482Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2107.03312","last_updated":"2021-07-07T15:45:42Z","snapshot_observed_at":"2026-08-09T07:05:32.492915Z","submitted_at":"2021-07-07T15:45:42Z","title":"SoundStream: An End-to-End Neural Audio Codec","version":1},"cited_work":{"arxiv_id":"2107.03312","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2107.03312","snapshot_observed_at":"2026-07-04T08:29:41.902212Z","title":"Soundstream: An end-to-end neural audio codec","venue":null,"work_id":"e2f34c6d-cf73-4f48-a80a-732b812cc299","year":2021},"citing_paper":{"arxiv_id":"2605.20649","last_updated":"2026-05-20T03:09:45Z","snapshot_observed_at":"2026-08-04T03:27:22.737171Z","submitted_at":"2026-05-20T03:09:45Z","title":"AMAR: Lightweight Attention-Based Multi-User Activity Recognition from Wi-Fi CSI","version":1},"reference_index":29,"source":"pdf_text","source_observed_at":"2026-05-21T03:11:44.222945Z"},"links":{"cited_paper":"/paper/2107.03312","citing_paper":"/paper/2605.20649"},"observation_digest":"sha256:714e594ff6bfd7ff82d531359e6be79f0fccc22a2f1794808e1a17e205acf45d","observation_id":"1c163d7f-68ac-47d0-8275-0921e3cd588b","resolution":{"observed_at":"2026-05-21T03:13:56.115057Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2107.03312","last_updated":"2021-07-07T15:45:42Z","snapshot_observed_at":"2026-08-09T07:05:32.492915Z","submitted_at":"2021-07-07T15:45:42Z","title":"SoundStream: An End-to-End Neural Audio Codec","version":1},"cited_work":{"arxiv_id":"2107.03312","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2107.03312","snapshot_observed_at":"2026-07-04T08:29:41.902212Z","title":"Soundstream: An end-to-end neural audio codec","venue":null,"work_id":"e2f34c6d-cf73-4f48-a80a-732b812cc299","year":2021},"citing_paper":{"arxiv_id":"2605.21081","last_updated":"2026-05-20T12:16:28Z","snapshot_observed_at":"2026-08-06T01:21:51.969876Z","submitted_at":"2026-05-20T12:16:28Z","title":"Musical Attention Transformer: Music Generation Using a Music-Specific Attention Model","version":1},"reference_index":23,"source":"pdf_text","source_observed_at":"2026-05-21T01:44:20.282896Z"},"links":{"cited_paper":"/paper/2107.03312","citing_paper":"/paper/2605.21081"},"observation_digest":"sha256:821a0c106dd7bcea9def7a54e2271dfb56d52a345bf682c75ce7dd26965d5ced","observation_id":"57f43a6d-5d8e-4b41-9c21-f261757f3f0c","resolution":{"observed_at":"2026-05-21T01:44:22.778590Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2107.03312","last_updated":"2021-07-07T15:45:42Z","snapshot_observed_at":"2026-08-09T07:05:32.492915Z","submitted_at":"2021-07-07T15:45:42Z","title":"SoundStream: An End-to-End Neural Audio Codec","version":1},"cited_work":{"arxiv_id":"2107.03312","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2107.03312","snapshot_observed_at":"2026-07-04T08:29:41.902212Z","title":"Soundstream: An end-to-end neural audio codec","venue":null,"work_id":"e2f34c6d-cf73-4f48-a80a-732b812cc299","year":2021},"citing_paper":{"arxiv_id":"2606.02631","last_updated":"2026-05-30T14:59:57Z","snapshot_observed_at":"2026-08-08T06:42:18.072725Z","submitted_at":"2026-05-30T14:59:57Z","title":"Wavelet as Tokenizer: Preliminary Results on a Shared Wavelet Token Schema for Natural Signals","version":1},"reference_index":10,"source":"pdf_text","source_observed_at":"2026-06-28T18:05:11.392504Z"},"links":{"cited_paper":"/paper/2107.03312","citing_paper":"/paper/2606.02631"},"observation_digest":"sha256:24ba9edcdd72d5b7324c148860e4886277e36dfbd1d4cdd2dc99ae057094037e","observation_id":"a7b0c50a-f784-47ab-a015-258f2a3c6361","resolution":{"observed_at":"2026-07-01T20:46:13.380021Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2107.03312","last_updated":"2021-07-07T15:45:42Z","snapshot_observed_at":"2026-08-09T07:05:32.492915Z","submitted_at":"2021-07-07T15:45:42Z","title":"SoundStream: An End-to-End Neural Audio Codec","version":1},"cited_work":{"arxiv_id":"2107.03312","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2107.03312","snapshot_observed_at":"2026-07-04T08:29:41.902212Z","title":"Soundstream: An end-to-end neural audio codec","venue":null,"work_id":"e2f34c6d-cf73-4f48-a80a-732b812cc299","year":2021},"citing_paper":{"arxiv_id":"2606.02739","last_updated":"2026-06-01T18:05:18Z","snapshot_observed_at":"2026-07-06T23:43:07.940839Z","submitted_at":"2026-06-01T18:05:18Z","title":"EntangleCodec: A Unified Discrete Audio Tokenizer via Semantic-Acoustic Entanglement","version":1},"reference_index":36,"source":"arxiv_source","source_observed_at":"2026-06-28T12:34:06.024192Z"},"links":{"cited_paper":"/paper/2107.03312","citing_paper":"/paper/2606.02739"},"observation_digest":"sha256:24268e172f408611bff2474d8521d2ac642e9aca1784390f93b5b07c8bac9927","observation_id":"51248e43-eef7-4c61-a155-310d1a5acc03","resolution":{"observed_at":"2026-07-02T01:06:24.820994Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2107.03312","last_updated":"2021-07-07T15:45:42Z","snapshot_observed_at":"2026-08-09T07:05:32.492915Z","submitted_at":"2021-07-07T15:45:42Z","title":"SoundStream: An End-to-End Neural Audio Codec","version":1},"cited_work":{"arxiv_id":"2107.03312","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2107.03312","snapshot_observed_at":"2026-07-04T08:29:41.902212Z","title":"Soundstream: An end-to-end neural audio codec","venue":null,"work_id":"e2f34c6d-cf73-4f48-a80a-732b812cc299","year":2021},"citing_paper":{"arxiv_id":"2606.06357","last_updated":"2026-06-04T16:25:07Z","snapshot_observed_at":"2026-07-06T23:46:10.612537Z","submitted_at":"2026-06-04T16:25:07Z","title":"F3-Tokenizer: Taming Audio Autoencoder Latents for Understanding and Generation","version":1},"reference_index":22,"source":"pdf_text","source_observed_at":"2026-06-27T23:36:27.369551Z"},"links":{"cited_paper":"/paper/2107.03312","citing_paper":"/paper/2606.06357"},"observation_digest":"sha256:b62c9707533493881e012dd533a9b30f94c0a77a33e050f2e860ddd930180b35","observation_id":"169498f5-45de-4c48-b16c-6ba7b734649d","resolution":{"observed_at":"2026-07-02T15:47:06.019471Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2107.03312","last_updated":"2021-07-07T15:45:42Z","snapshot_observed_at":"2026-08-09T07:05:32.492915Z","submitted_at":"2021-07-07T15:45:42Z","title":"SoundStream: An End-to-End Neural Audio Codec","version":1},"cited_work":{"arxiv_id":"2107.03312","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2107.03312","snapshot_observed_at":"2026-07-04T08:29:41.902212Z","title":"Soundstream: An end-to-end neural audio codec","venue":null,"work_id":"e2f34c6d-cf73-4f48-a80a-732b812cc299","year":2021},"citing_paper":{"arxiv_id":"2606.21893","last_updated":"2026-06-20T05:58:50Z","snapshot_observed_at":"2026-08-06T08:56:11.950592Z","submitted_at":"2026-06-20T05:58:50Z","title":"AugCodec: A Low-Bitrate Disentangled Neural Speech Codec via Data Augmentation","version":1},"reference_index":6,"source":"pdf_text","source_observed_at":"2026-06-26T11:37:39.960212Z"},"links":{"cited_paper":"/paper/2107.03312","citing_paper":"/paper/2606.21893"},"observation_digest":"sha256:1e565314ff54683397e78bcb52544e733bb300cf73e1a029eeb31ed42e47c581","observation_id":"c54b9722-9099-4269-a7bd-ec2b71ab558d","resolution":{"observed_at":"2026-07-04T08:29:41.903636Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2107.03312","last_updated":"2021-07-07T15:45:42Z","snapshot_observed_at":"2026-08-09T07:05:32.492915Z","submitted_at":"2021-07-07T15:45:42Z","title":"SoundStream: An End-to-End Neural Audio Codec","version":1},"cited_work":{"arxiv_id":"2107.03312","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2107.03312","snapshot_observed_at":"2026-07-04T08:29:41.902212Z","title":"Soundstream: An end-to-end neural audio codec","venue":null,"work_id":"e2f34c6d-cf73-4f48-a80a-732b812cc299","year":2021},"citing_paper":{"arxiv_id":"2606.28779","last_updated":"2026-06-27T07:14:04Z","snapshot_observed_at":"2026-08-09T07:06:34.532654Z","submitted_at":"2026-06-27T07:14:04Z","title":"Telephony Voice Agent for Banking Services","version":1},"reference_index":22,"source":"pdf_text","source_observed_at":"2026-06-30T08:57:32.212566Z"},"links":{"cited_paper":"/paper/2107.03312","citing_paper":"/paper/2606.28779"},"observation_digest":"sha256:05daa572ef9f20b2603268af04387126f9af9b2b04391b72ca558a20a8d1d0c3","observation_id":"04b19d85-4832-4674-a219-a630ceffafca","resolution":{"observed_at":"2026-06-30T09:04:32.871308Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2107.03312","last_updated":"2021-07-07T15:45:42Z","snapshot_observed_at":"2026-08-09T07:05:32.492915Z","submitted_at":"2021-07-07T15:45:42Z","title":"SoundStream: An End-to-End Neural Audio Codec","version":1},"cited_work":{"arxiv_id":"2107.03312","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2107.03312","snapshot_observed_at":"2026-07-04T08:29:41.902212Z","title":"Soundstream: An end-to-end neural audio codec","venue":null,"work_id":"e2f34c6d-cf73-4f48-a80a-732b812cc299","year":2021},"citing_paper":{"arxiv_id":"2607.02266","last_updated":"2026-07-02T14:51:42Z","snapshot_observed_at":"2026-08-08T15:31:37.930459Z","submitted_at":"2026-07-02T14:51:42Z","title":"HERMES: A Multi-Granularity Labeling Substrate for Pre-training Data Mixtures","version":1},"reference_index":18,"source":"arxiv_source","source_observed_at":"2026-07-03T16:44:41.720388Z"},"links":{"cited_paper":"/paper/2107.03312","citing_paper":"/paper/2607.02266"},"observation_digest":"sha256:5e875062ea89eda503b485da256982189a1f2b838e8b8a342cd650ac74040738","observation_id":"fb7580f2-7310-4645-9a55-bba81d679027","resolution":{"observed_at":"2026-07-03T16:48:39.426776Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2107.03312","last_updated":"2021-07-07T15:45:42Z","snapshot_observed_at":"2026-08-09T07:05:32.492915Z","submitted_at":"2021-07-07T15:45:42Z","title":"SoundStream: An End-to-End Neural Audio Codec","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2107.03312","snapshot_observed_at":"2026-07-11T21:55:00.514210Z","title":null,"venue":null,"work_id":null,"year":2021},"citing_paper":{"arxiv_id":"2607.04068","last_updated":"2026-07-05T01:00:58Z","snapshot_observed_at":"2026-08-09T07:05:55.860168Z","submitted_at":"2026-07-05T01:00:58Z","title":"UniSGR: Unified Framework for Semantic ID Generation and Ranking","version":1},"reference_index":24,"source":"pdf_text","source_observed_at":"2026-07-11T21:55:00.514210Z"},"links":{"cited_paper":"/paper/2107.03312","citing_paper":"/paper/2607.04068"},"observation_digest":"sha256:04374b230dc0d70993ef056d42dbb9f00585003013d699fc89abb4100e7f9159","observation_id":"9cf854e5-57ba-441c-b304-1deda6cfeccb","resolution":{"observed_at":"2026-07-11T21:55:00.514210Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2107.03312","last_updated":"2021-07-07T15:45:42Z","snapshot_observed_at":"2026-08-09T07:05:32.492915Z","submitted_at":"2021-07-07T15:45:42Z","title":"SoundStream: An End-to-End Neural Audio Codec","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2107.03312","snapshot_observed_at":"2026-08-02T06:23:32.601187Z","title":null,"venue":null,"work_id":null,"year":2021},"citing_paper":{"arxiv_id":"2607.14148","last_updated":"2026-07-14T13:05:53Z","snapshot_observed_at":"2026-08-07T05:42:24.503168Z","submitted_at":"2026-07-14T13:05:53Z","title":"ITGPT: A Transformer Based Architecture for the Generation of Dance Dance Revolution and In the Groove Charts","version":1},"reference_index":21,"source":"pdf_text","source_observed_at":"2026-08-02T06:23:32.601187Z"},"links":{"cited_paper":"/paper/2107.03312","citing_paper":"/paper/2607.14148"},"observation_digest":"sha256:5a8b68a30dfd5af08d3d1af14d6007b82f47f2d4cf13720f5e60380020517492","observation_id":"7c1e7c73-db9d-4ad5-8761-1ec6f0521e5f","resolution":{"observed_at":"2026-08-02T06:23:32.601187Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}}],"links":{"evidence":"/evidence","html":"/paper/2107.03312/citation-record","integrity":"/paper/2107.03312/integrity","json":"/paper/2107.03312/citation-record.json","paper":"/paper/2107.03312"},"outbound":[],"paper":{"arxiv_id":"2107.03312","last_updated":"2021-07-07T15:45:42Z","latest_version":1,"primary_category":"cs.SD","snapshot_observed_at":"2026-08-09T07:05:32.492915Z","submitted_at":"2021-07-07T15:45:42Z","title":"SoundStream: An End-to-End Neural Audio Codec"},"reference_resolution":{"displayed":0,"state_counts":{"malformed_identifier":0,"metadata_mismatch":0,"parse_uncertain":0,"unresolved":0,"verified_exact":0,"verified_fuzzy":0},"total_outbound_references":0},"refusal":"A citation records a reference. It does not transfer a finding from one paper to another.","schema":"pith.paper-citation-record.v1","standing_sources":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"thesis":"As of 9 August 2026, this Paper Citation Record lists 0 of 0 outbound references and 22 inbound Pith citation observations for arXiv:2107.03312."}