{"as_of":"2026-08-23T16:17:00Z","caps":{"database_statements":6,"inbound":100,"outbound":100},"context_digest":"sha256:753b0ee8dbfb952098a02ebcea4f89601ab8d938d0303c0e5efee450dd21805a","coverage":[{"denominator":14,"lane":"reference_resolution","note":"Typed states for the displayed outbound observations.","records_observed":14,"source":"paper_references, paper_reference_links","source_observed_at":"2026-08-15T20:25:19.483315Z","state":"measured"},{"denominator":16,"lane":"standing_notices","note":"One-hop event checks from named stored sources.","records_observed":16,"source":"scholarly_work_events, retraction_status_cache","source_observed_at":"2026-08-23T06:30:58.430688+00:00","state":"measured"},{"denominator":2,"lane":"inbound_itemization","note":"Pith citing papers itemized under the disclosed page cap.","records_observed":2,"source":"paper_references, paper_reference_links","source_observed_at":"2026-07-15T12:21:48.698333Z","state":"measured"},{"denominator":1,"lane":"external_citation_measurements","note":"A source-named dated measurement, never combined with another source.","records_observed":0,"source":"arxiv_reference","source_observed_at":"2026-07-02T12:36:56.317033Z","state":"measured"}],"external_citation_measurements":[],"inbound":[{"citation":{"cited_paper":{"arxiv_id":"2505.13085","last_updated":"2025-05-20T10:22:17Z","snapshot_observed_at":"2026-08-18T17:55:17.132101Z","submitted_at":"2025-05-19T13:19:49Z","title":"Universal Semantic Disentangled Privacy-preserving Speech Representation Learning","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2505.13085","snapshot_observed_at":"2026-07-15T12:21:48.698333Z","title":"Available: https://arxiv.org/abs/2505.13085","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2603.08977","last_updated":"2026-06-07T05:11:25Z","snapshot_observed_at":"2026-08-17T06:01:11.312584Z","submitted_at":"2026-03-09T22:11:40Z","title":"Universal Speech Content Factorization","version":2},"reference_index":37,"source":"pdf_text","source_observed_at":"2026-07-15T12:21:48.698333Z"},"links":{"cited_paper":"/paper/2505.13085","citing_paper":"/paper/2603.08977"},"observation_digest":"sha256:4a0d6105e1aee60c6c3865ecb0a5031a51ec0c65266e60d9cb51e1d74384e190","observation_id":"d57897c5-718f-44f5-8623-670205c2ead3","resolution":{"observed_at":"2026-07-15T12:21:48.698333Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2505.13085","last_updated":"2025-05-20T10:22:17Z","snapshot_observed_at":"2026-08-18T17:55:17.132101Z","submitted_at":"2025-05-19T13:19:49Z","title":"Universal Semantic Disentangled Privacy-preserving Speech Representation Learning","version":2},"cited_work":{"arxiv_id":"2505.13085","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2505.13085","snapshot_observed_at":"2026-07-02T12:36:56.317033Z","title":"Univer- sal semantic disentangled privacy-preserving speech representa- tion learning,","venue":null,"work_id":"7ec2aa56-2617-457b-a167-49b036860dad","year":2025},"citing_paper":{"arxiv_id":"2606.05561","last_updated":"2026-06-04T01:16:01Z","snapshot_observed_at":"2026-08-02T18:07:49.784049Z","submitted_at":"2026-06-04T01:16:01Z","title":"InfoShield: Privacy-Preserving Speech Representations for Mental Health Screening via Information-Theoretic Optimization","version":1},"reference_index":25,"source":"pdf_text","source_observed_at":"2026-06-28T02:03:42.928266Z"},"links":{"cited_paper":"/paper/2505.13085","citing_paper":"/paper/2606.05561"},"observation_digest":"sha256:0bc6ff4affcf78ee2aafef3adf0a10af329c0b33418136742c5ea2aace41e9b8","observation_id":"dabfbd56-6831-4555-9d23-563724d449ac","resolution":{"observed_at":"2026-07-02T12:36:56.318538Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-23T06:30:58.430688+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-23T06:30:58.430688+00:00","source":"crossref"},{"observed_at":"2026-08-23T06:30:53.778098+00:00","source":"retraction_watch"}],"state":"measured"}}],"links":{"evidence":"/evidence","html":"/paper/2505.13085/citation-record","integrity":"/paper/2505.13085/integrity","json":"/paper/2505.13085/citation-record.json","paper":"/paper/2505.13085"},"outbound":[{"citation":{"cited_paper":{"arxiv_id":"2303.08774","last_updated":"2024-03-04T06:01:33Z","snapshot_observed_at":"2026-08-17T09:58:46.058102Z","submitted_at":"2023-03-15T17:15:04Z","title":"GPT-4 Technical Report","version":6},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2303.08774","snapshot_observed_at":"2026-08-15T20:25:19.402775Z","title":"GPT-4 technical report","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2505.13085","last_updated":"2025-05-20T10:22:17Z","snapshot_observed_at":"2026-08-18T17:55:17.132101Z","submitted_at":"2025-05-19T13:19:49Z","title":"Universal Semantic Disentangled Privacy-preserving Speech Representation Learning","version":2},"reference_index":1,"source":"pdf_text","source_observed_at":"2026-08-15T20:25:19.402775Z"},"links":{"cited_paper":"/paper/2303.08774","citing_paper":"/paper/2505.13085"},"observation_digest":"sha256:45710d408a66b43ca7957794976e15a6d0d9bd0297de62e4145593d6e0cf081b","observation_id":"9149a5a0-73de-42ef-9eeb-1aaed13c5678","resolution":{"observed_at":"2026-08-15T20:25:19.402775Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2403.03100","last_updated":"2024-04-23T08:38:03Z","snapshot_observed_at":"2026-08-18T17:50:58.009731Z","submitted_at":"2024-03-05T16:35:25Z","title":"NaturalSpeech 3: Zero-Shot Speech Synthesis with Factorized Codec and Diffusion Models","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2403.03100","snapshot_observed_at":"2026-08-15T20:25:19.432858Z","title":"Naturalspeech 3: Zero-shot speech synthesis with factor- ized codec and diffusion models","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2505.13085","last_updated":"2025-05-20T10:22:17Z","snapshot_observed_at":"2026-08-18T17:55:17.132101Z","submitted_at":"2025-05-19T13:19:49Z","title":"Universal Semantic Disentangled Privacy-preserving Speech Representation Learning","version":2},"reference_index":6,"source":"pdf_text","source_observed_at":"2026-08-15T20:25:19.432858Z"},"links":{"cited_paper":"/paper/2403.03100","citing_paper":"/paper/2505.13085"},"observation_digest":"sha256:b9a9affd11af2b4a7c88c9453ffa2244f412adc49c1d7242934912b7e05d3753","observation_id":"7a9a1220-3137-4e4e-a788-ecbce25fe5ca","resolution":{"observed_at":"2026-08-15T20:25:19.432858Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2402.08093","last_updated":"2024-02-15T18:57:26Z","snapshot_observed_at":"2026-08-16T20:07:00.407637Z","submitted_at":"2024-02-12T22:21:30Z","title":"BASE TTS: Lessons from building a billion-parameter Text-to-Speech model on 100K hours of data","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2402.08093","snapshot_observed_at":"2026-08-15T20:25:19.438932Z","title":"BASE TTS: Lessons from building a billion-parameter text-to-speech model on 100k hours of data","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2505.13085","last_updated":"2025-05-20T10:22:17Z","snapshot_observed_at":"2026-08-18T17:55:17.132101Z","submitted_at":"2025-05-19T13:19:49Z","title":"Universal Semantic Disentangled Privacy-preserving Speech Representation Learning","version":2},"reference_index":7,"source":"pdf_text","source_observed_at":"2026-08-15T20:25:19.438932Z"},"links":{"cited_paper":"/paper/2402.08093","citing_paper":"/paper/2505.13085"},"observation_digest":"sha256:82b7312557df3c2ce9ec9157cadf784372b757ff0586f350b95bc41265eab97d","observation_id":"b50488e6-70aa-4576-9500-a20364b4eec2","resolution":{"observed_at":"2026-08-15T20:25:19.438932Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-15T20:25:19.739935Z","title":"Npu-ntu system for voice privacy 2024 challenge","venue":null,"work_id":"be4a0098-3ed4-48ff-9768-809b346c651a","year":2024},"citing_paper":{"arxiv_id":"2505.13085","last_updated":"2025-05-20T10:22:17Z","snapshot_observed_at":"2026-08-18T17:55:17.132101Z","submitted_at":"2025-05-19T13:19:49Z","title":"Universal Semantic Disentangled Privacy-preserving Speech Representation Learning","version":2},"reference_index":11,"source":"pdf_text","source_observed_at":"2026-08-15T20:25:19.463734Z"},"links":{"citing_paper":"/paper/2505.13085"},"observation_digest":"sha256:86e3f11e146819ead58c3d3f0e294f4093b6262b73378967a68eb849af91bf47","observation_id":"426a4006-3f02-43bf-bf7e-ba2e38f22435","resolution":{"observed_at":"2026-08-15T20:25:19.746832Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-23T06:30:58.430688+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-23T06:30:58.430688+00:00","source":"crossref"},{"observed_at":"2026-08-23T06:30:53.778098+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-15T20:25:19.682764Z","title":null,"venue":null,"work_id":"89bf0eaa-ffa5-4337-87cc-f203e33de6bb","year":2023},"citing_paper":{"arxiv_id":"2505.13085","last_updated":"2025-05-20T10:22:17Z","snapshot_observed_at":"2026-08-18T17:55:17.132101Z","submitted_at":"2025-05-19T13:19:49Z","title":"Universal Semantic Disentangled Privacy-preserving Speech Representation Learning","version":2},"reference_index":14,"source":"pdf_text","source_observed_at":"2026-08-15T20:25:19.483315Z"},"links":{"citing_paper":"/paper/2505.13085"},"observation_digest":"sha256:e28a49b0a1bba313ae7dba6631d68d08d47e327466efe00f6ee5dffe2e2948c5","observation_id":"b2f01ec4-e5a9-42c4-9cd0-2fe08793b803","resolution":{"observed_at":"2026-08-15T20:25:19.688928Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-23T06:30:58.430688+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-23T06:30:58.430688+00:00","source":"crossref"},{"observed_at":"2026-08-23T06:30:53.778098+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-15T20:25:19.759105Z","title":null,"venue":null,"work_id":"abb870c8-dcbe-4aee-aa16-bc1e84e46552","year":2014},"citing_paper":{"arxiv_id":"2505.13085","last_updated":"2025-05-20T10:22:17Z","snapshot_observed_at":"2026-08-18T17:55:17.132101Z","submitted_at":"2025-05-19T13:19:49Z","title":"Universal Semantic Disentangled Privacy-preserving Speech Representation Learning","version":2},"reference_index":29,"source":"pdf_text","source_observed_at":"2026-08-15T20:25:19.421317Z"},"links":{"citing_paper":"/paper/2505.13085"},"observation_digest":"sha256:4c5f8fcaa068d484cf9a48936354635536ba6e513dd8413f7cbf988d154abd77","observation_id":"78090832-6b96-41e2-974f-790cfae10af3","resolution":{"observed_at":"2026-08-15T20:25:19.764348Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-23T06:30:58.430688+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-23T06:30:58.430688+00:00","source":"crossref"},{"observed_at":"2026-08-23T06:30:53.778098+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-15T20:25:19.702224Z","title":"For each resolution discriminator, we set n_fft = (2048, 1024,","venue":null,"work_id":"d4ce7e70-1b0e-46dc-8a17-639b11d92bb5","year":null},"citing_paper":{"arxiv_id":"2505.13085","last_updated":"2025-05-20T10:22:17Z","snapshot_observed_at":"2026-08-18T17:55:17.132101Z","submitted_at":"2025-05-19T13:19:49Z","title":"Universal Semantic Disentangled Privacy-preserving Speech Representation Learning","version":2},"reference_index":1024,"source":"pdf_text","source_observed_at":"2026-08-15T20:25:19.476480Z"},"links":{"citing_paper":"/paper/2505.13085"},"observation_digest":"sha256:f2a762b828373c5f87edbb4792ebb5c4cba0947a2a592e80c7356eb7e7dc7646","observation_id":"c43604bf-89e2-48fa-bf87-e0839939f7ed","resolution":{"observed_at":"2026-08-15T20:25:19.708599Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-23T06:30:58.430688+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-23T06:30:58.430688+00:00","source":"crossref"},{"observed_at":"2026-08-23T06:30:53.778098+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2306.12925","last_updated":"2023-06-22T14:37:54Z","snapshot_observed_at":"2026-08-21T19:19:25.551780Z","submitted_at":"2023-06-22T14:37:54Z","title":"AudioPaLM: A Large Language Model That Can Speak and Listen","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2306.12925","snapshot_observed_at":"2026-08-15T20:25:19.445069Z","title":"AudiopaLM: A large language model that can speak and listen.arXiv preprint arXiv:2306.12925,","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2505.13085","last_updated":"2025-05-20T10:22:17Z","snapshot_observed_at":"2026-08-18T17:55:17.132101Z","submitted_at":"2025-05-19T13:19:49Z","title":"Universal Semantic Disentangled Privacy-preserving Speech Representation Learning","version":2},"reference_index":2001,"source":"pdf_text","source_observed_at":"2026-08-15T20:25:19.445069Z"},"links":{"cited_paper":"/paper/2306.12925","citing_paper":"/paper/2505.13085"},"observation_digest":"sha256:ef8044e9182ea05200a01d91b489686a0acde5a53388a1c593a8286eb6497836","observation_id":"b45e9446-7196-4692-a401-20a613632a79","resolution":{"observed_at":"2026-08-15T20:25:19.445069Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2407.21783","last_updated":"2024-11-23T23:27:33Z","snapshot_observed_at":"2026-08-13T17:20:44.002518Z","submitted_at":"2024-07-31T17:54:27Z","title":"The Llama 3 Herd of Models","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2407.21783","snapshot_observed_at":"2026-08-15T20:25:19.415027Z","title":"The Llama 3 Herd of Models","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2505.13085","last_updated":"2025-05-20T10:22:17Z","snapshot_observed_at":"2026-08-18T17:55:17.132101Z","submitted_at":"2025-05-19T13:19:49Z","title":"Universal Semantic Disentangled Privacy-preserving Speech Representation Learning","version":2},"reference_index":2010,"source":"pdf_text","source_observed_at":"2026-08-15T20:25:19.415027Z"},"links":{"cited_paper":"/paper/2407.21783","citing_paper":"/paper/2505.13085"},"observation_digest":"sha256:522ca6ac12802a44a4b8153e7ab2cb787382d95d7e3efebb9275a2fbaf421312","observation_id":"4aae48df-3b07-4016-9589-36316f9151a8","resolution":{"observed_at":"2026-08-15T20:25:19.415027Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2404.10301","last_updated":"2024-07-29T14:52:26Z","snapshot_observed_at":"2026-08-17T11:23:59.787786Z","submitted_at":"2024-04-16T06:09:33Z","title":"Long-form music generation with latent diffusion","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2404.10301","snapshot_observed_at":"2026-08-15T20:25:19.426727Z","title":"Zach Evans, Julian D Parker, CJ Carr, Zack Zukowski, Josiah Taylor, and Jordi Pons","venue":null,"work_id":null,"year":2014},"citing_paper":{"arxiv_id":"2505.13085","last_updated":"2025-05-20T10:22:17Z","snapshot_observed_at":"2026-08-18T17:55:17.132101Z","submitted_at":"2025-05-19T13:19:49Z","title":"Universal Semantic Disentangled Privacy-preserving Speech Representation Learning","version":2},"reference_index":2014,"source":"pdf_text","source_observed_at":"2026-08-15T20:25:19.426727Z"},"links":{"cited_paper":"/paper/2404.10301","citing_paper":"/paper/2505.13085"},"observation_digest":"sha256:99c0543de8bd351f2ec146d71681809af8960a02e3a5b04d3ff69fff3369d13b","observation_id":"18e83ac3-3a05-4a89-a806-c1f937e01239","resolution":{"observed_at":"2026-08-15T20:25:19.426727Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-15T20:25:19.721919Z","title":null,"venue":null,"work_id":"e2ec2153-c7fe-4662-8417-627f3e71b898","year":2024},"citing_paper":{"arxiv_id":"2505.13085","last_updated":"2025-05-20T10:22:17Z","snapshot_observed_at":"2026-08-18T17:55:17.132101Z","submitted_at":"2025-05-19T13:19:49Z","title":"Universal Semantic Disentangled Privacy-preserving Speech Representation Learning","version":2},"reference_index":2020,"source":"pdf_text","source_observed_at":"2026-08-15T20:25:19.470372Z"},"links":{"citing_paper":"/paper/2505.13085"},"observation_digest":"sha256:b755466e84498b492399f4f0feafc011a47f42624d3dca47ab220b426969dac7","observation_id":"b60e147d-6440-45b5-aba7-804d042a88fe","resolution":{"observed_at":"2026-08-15T20:25:19.727335Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-23T06:30:58.430688+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-23T06:30:58.430688+00:00","source":"crossref"},{"observed_at":"2026-08-23T06:30:53.778098+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2312.11805","last_updated":"2025-05-09T21:04:06Z","snapshot_observed_at":"2026-08-20T18:27:04.837880Z","submitted_at":"2023-12-19T02:39:27Z","title":"Gemini: A Family of Highly Capable Multimodal Models","version":5},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2312.11805","snapshot_observed_at":"2026-08-15T20:25:19.409174Z","title":"Gemini: a family of highly capable multimodal models","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2505.13085","last_updated":"2025-05-20T10:22:17Z","snapshot_observed_at":"2026-08-18T17:55:17.132101Z","submitted_at":"2025-05-19T13:19:49Z","title":"Universal Semantic Disentangled Privacy-preserving Speech Representation Learning","version":2},"reference_index":2023,"source":"pdf_text","source_observed_at":"2026-08-15T20:25:19.409174Z"},"links":{"cited_paper":"/paper/2312.11805","citing_paper":"/paper/2505.13085"},"observation_digest":"sha256:db0609c2b782ae5ea55e2f0b382614c4eb3298b411cf1087e922df1e201fe839","observation_id":"3cb18b70-92b3-41ed-96e7-b2b51220cf4e","resolution":{"observed_at":"2026-08-15T20:25:19.409174Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2307.09288","last_updated":"2023-07-19T17:08:59Z","snapshot_observed_at":"2026-08-07T12:56:43.323460Z","submitted_at":"2023-07-18T14:31:57Z","title":"Llama 2: Open Foundation and Fine-Tuned Chat Models","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2307.09288","snapshot_observed_at":"2026-08-15T20:25:19.457780Z","title":"Llama 2: Open founda- tion and fine-tuned chat models","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2505.13085","last_updated":"2025-05-20T10:22:17Z","snapshot_observed_at":"2026-08-18T17:55:17.132101Z","submitted_at":"2025-05-19T13:19:49Z","title":"Universal Semantic Disentangled Privacy-preserving Speech Representation Learning","version":2},"reference_index":2024,"source":"pdf_text","source_observed_at":"2026-08-15T20:25:19.457780Z"},"links":{"cited_paper":"/paper/2307.09288","citing_paper":"/paper/2505.13085"},"observation_digest":"sha256:19fcd4e5fb23def0fd05ccbbea13843e7eb301e7ebe05019597fc0dcc0514bb7","observation_id":"fb785f51-a93d-4c06-9944-31b8c2bca1fd","resolution":{"observed_at":"2026-08-15T20:25:19.457780Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2404.02677","last_updated":"2024-06-12T14:05:43Z","snapshot_observed_at":"2026-08-19T22:00:02.604229Z","submitted_at":"2024-04-03T12:20:51Z","title":"The VoicePrivacy 2024 Challenge Evaluation Plan","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2404.02677","snapshot_observed_at":"2026-08-15T20:25:19.451836Z","title":"The voiceprivacy 2024 challenge evaluation plan","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2505.13085","last_updated":"2025-05-20T10:22:17Z","snapshot_observed_at":"2026-08-18T17:55:17.132101Z","submitted_at":"2025-05-19T13:19:49Z","title":"Universal Semantic Disentangled Privacy-preserving Speech Representation Learning","version":2},"reference_index":2025,"source":"pdf_text","source_observed_at":"2026-08-15T20:25:19.451836Z"},"links":{"cited_paper":"/paper/2404.02677","citing_paper":"/paper/2505.13085"},"observation_digest":"sha256:c6d0d56eed997f981b42326ec7f2f6f2317b9def056e81f331595f3394f545fa","observation_id":"c6fc5a32-62bd-40d0-8f66-4dec280795b2","resolution":{"observed_at":"2026-08-15T20:25:19.451836Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}}],"paper":{"arxiv_id":"2505.13085","last_updated":"2025-05-20T10:22:17Z","latest_version":2,"primary_category":"eess.AS","snapshot_observed_at":"2026-08-18T17:55:17.132101Z","submitted_at":"2025-05-19T13:19:49Z","title":"Universal Semantic Disentangled Privacy-preserving Speech Representation Learning"},"reference_resolution":{"displayed":14,"state_counts":{"malformed_identifier":0,"metadata_mismatch":0,"parse_uncertain":0,"unresolved":12,"verified_exact":0,"verified_fuzzy":2},"total_outbound_references":14},"refusal":"A citation records a reference. It does not transfer a finding from one paper to another.","schema":"pith.paper-citation-record.v1","standing_sources":[{"observed_at":"2026-08-23T06:30:58.430688+00:00","source":"crossref"},{"observed_at":"2026-08-23T06:30:53.778098+00:00","source":"retraction_watch"}],"thesis":"As of 23 August 2026, this Paper Citation Record lists 14 of 14 outbound references and 2 inbound Pith citation observations for arXiv:2505.13085."}