{"as_of":"2026-08-12T08:02:00Z","caps":{"database_statements":6,"inbound":100,"outbound":100},"context_digest":"sha256:bd4d76b5030ce08a5bb675c144347bc333ccf93d6913af6f1e8d436ca145efc6","coverage":[{"denominator":54,"lane":"reference_resolution","note":"Typed states for the displayed outbound observations.","records_observed":54,"source":"paper_references, paper_reference_links","source_observed_at":"2026-08-07T00:59:56.303718Z","state":"measured"},{"denominator":55,"lane":"standing_notices","note":"One-hop event checks from named stored sources.","records_observed":55,"source":"scholarly_work_events, retraction_status_cache","source_observed_at":"2026-08-12T06:34:41.77262+00:00","state":"measured"},{"denominator":1,"lane":"inbound_itemization","note":"Pith citing papers itemized under the disclosed page cap.","records_observed":1,"source":"paper_references, paper_reference_links","source_observed_at":"2026-08-06T21:34:07.924156Z","state":"measured"},{"denominator":1,"lane":"external_citation_measurements","note":"A source-named dated measurement, never combined with another source.","records_observed":0,"source":"pith","source_observed_at":"2026-08-06T21:34:11.220887Z","state":"measured"}],"external_citation_measurements":[],"inbound":[{"citation":{"cited_paper":{"arxiv_id":"2506.12260","last_updated":"2025-08-22T06:43:05Z","snapshot_observed_at":"2026-08-11T15:08:26.236040Z","submitted_at":"2025-06-13T22:22:19Z","title":"Improving Speech Enhancement with Multi-Metric Supervision from Learned Quality Assessment","version":2},"cited_work":{"arxiv_id":"2506.12260","doi":null,"metadata_source":"pith","pith_arxiv_id":"2506.12260","snapshot_observed_at":"2026-08-06T21:34:11.220887Z","title":"Improving Speech Enhancement with Multi-Metric Supervision from Learned Quality Assessment","venue":"cs.SD","work_id":"066ac648-864a-4e98-ab50-bc1232e0ab66","year":2025},"citing_paper":{"arxiv_id":"2506.23859","last_updated":"2025-08-19T09:17:36Z","snapshot_observed_at":"2026-08-12T03:52:58.429144Z","submitted_at":"2025-06-30T13:55:10Z","title":"Less is More: Data Curation Matters in Scaling Speech Enhancement","version":2},"reference_index":7,"source":"pdf_text","source_observed_at":"2026-08-06T21:34:07.924156Z"},"links":{"cited_paper":"/paper/2506.12260","citing_paper":"/paper/2506.23859"},"observation_digest":"sha256:9d504b40476cbaf7f2b3e60f7876959aee41420a6724061a97f6aa54a88c70dd","observation_id":"3fbac6c8-1f69-4e11-b181-9c5c61e0a827","resolution":{"observed_at":"2026-08-06T21:34:11.310798Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-12T06:34:41.77262+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-12T06:34:41.77262+00:00","source":"crossref"},{"observed_at":"2026-08-12T06:34:36.333875+00:00","source":"retraction_watch"}],"state":"measured"}}],"links":{"evidence":"/evidence","html":"/paper/2506.12260/citation-record","integrity":"/paper/2506.12260/integrity","json":"/paper/2506.12260/citation-record.json","paper":"/paper/2506.12260"},"outbound":[{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T00:59:51.318929Z","title":null,"venue":null,"work_id":null,"year":2013},"citing_paper":{"arxiv_id":"2506.12260","last_updated":"2025-08-22T06:43:05Z","snapshot_observed_at":"2026-08-11T15:08:26.236040Z","submitted_at":"2025-06-13T22:22:19Z","title":"Improving Speech Enhancement with Multi-Metric Supervision from Learned Quality Assessment","version":2},"reference_index":1,"source":"pdf_text","source_observed_at":"2026-08-07T00:59:51.318929Z"},"links":{"citing_paper":"/paper/2506.12260"},"observation_digest":"sha256:3799558db2909fbaef3f60075692584c303dbaa34f7fc03b35ed46321dd1c7c6","observation_id":"b28c77b7-c622-455f-8923-0c0497d98042","resolution":{"observed_at":"2026-08-07T00:59:51.318929Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T01:00:02.784094Z","title":"SDR–half- baked or well done?","venue":null,"work_id":"9f7557ad-4788-4b1f-9edb-119c7d8afdfe","year":2019},"citing_paper":{"arxiv_id":"2506.12260","last_updated":"2025-08-22T06:43:05Z","snapshot_observed_at":"2026-08-11T15:08:26.236040Z","submitted_at":"2025-06-13T22:22:19Z","title":"Improving Speech Enhancement with Multi-Metric Supervision from Learned Quality Assessment","version":2},"reference_index":2,"source":"pdf_text","source_observed_at":"2026-08-07T00:59:51.393421Z"},"links":{"citing_paper":"/paper/2506.12260"},"observation_digest":"sha256:7325b2a6cca935415b5e661e88b0aac15bc19e614fdb653f361d8d2c83adc4de","observation_id":"cf9aca71-7c68-471f-8309-c9ab7386e0d9","resolution":{"observed_at":"2026-08-07T01:00:02.810913Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-12T06:34:41.77262+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-12T06:34:41.77262+00:00","source":"crossref"},{"observed_at":"2026-08-12T06:34:36.333875+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T01:00:02.704834Z","title":"How bad are artifacts?: Analyzing the impact of speech enhancement errors on asr,","venue":null,"work_id":"b6c36e56-46ad-4996-9651-e9ec5090d19a","year":2022},"citing_paper":{"arxiv_id":"2506.12260","last_updated":"2025-08-22T06:43:05Z","snapshot_observed_at":"2026-08-11T15:08:26.236040Z","submitted_at":"2025-06-13T22:22:19Z","title":"Improving Speech Enhancement with Multi-Metric Supervision from Learned Quality Assessment","version":2},"reference_index":3,"source":"pdf_text","source_observed_at":"2026-08-07T00:59:51.494513Z"},"links":{"citing_paper":"/paper/2506.12260"},"observation_digest":"sha256:f154df865eabb173ffc1e68298447bbcc17263a6d0e1276a4841b208d19505ea","observation_id":"21a0fe25-e3cb-480e-8759-75a7212852d2","resolution":{"observed_at":"2026-08-07T01:00:02.732807Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-12T06:34:41.77262+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-12T06:34:41.77262+00:00","source":"crossref"},{"observed_at":"2026-08-12T06:34:36.333875+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T01:00:02.640903Z","title":"Bridging the gap between monaural speech enhancement and recognition with distortion-independent acous- tic modeling,","venue":null,"work_id":"3398b4f9-a5df-4386-8960-afdcd2b74e42","year":2020},"citing_paper":{"arxiv_id":"2506.12260","last_updated":"2025-08-22T06:43:05Z","snapshot_observed_at":"2026-08-11T15:08:26.236040Z","submitted_at":"2025-06-13T22:22:19Z","title":"Improving Speech Enhancement with Multi-Metric Supervision from Learned Quality Assessment","version":2},"reference_index":4,"source":"pdf_text","source_observed_at":"2026-08-07T00:59:51.592392Z"},"links":{"citing_paper":"/paper/2506.12260"},"observation_digest":"sha256:84b27568ba3ccd0824749da488299083eab7f6a27e6a1ae32cc071e31a2039ea","observation_id":"462bd6d8-9579-48ab-8a2b-b97a2c17bb4b","resolution":{"observed_at":"2026-08-07T01:00:02.663269Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-12T06:34:41.77262+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-12T06:34:41.77262+00:00","source":"crossref"},{"observed_at":"2026-08-12T06:34:36.333875+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T01:00:02.583223Z","title":"Advancing non-intrusive suppression on enhancement distortion for noise robust asr,","venue":null,"work_id":"dcf59929-2c15-4006-b0db-39ab04d20bbb","year":2025},"citing_paper":{"arxiv_id":"2506.12260","last_updated":"2025-08-22T06:43:05Z","snapshot_observed_at":"2026-08-11T15:08:26.236040Z","submitted_at":"2025-06-13T22:22:19Z","title":"Improving Speech Enhancement with Multi-Metric Supervision from Learned Quality Assessment","version":2},"reference_index":5,"source":"pdf_text","source_observed_at":"2026-08-07T00:59:51.731425Z"},"links":{"citing_paper":"/paper/2506.12260"},"observation_digest":"sha256:a7e65e8edc24d10fe573ae705023e729f42ede1a28c58dd76b1147f26ef95b2b","observation_id":"18d0c539-84f8-47f7-a0f3-de04eb4b5ebe","resolution":{"observed_at":"2026-08-07T01:00:02.606075Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-12T06:34:41.77262+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-12T06:34:41.77262+00:00","source":"crossref"},{"observed_at":"2026-08-12T06:34:36.333875+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T01:00:02.529131Z","title":"Fat-hubert: Front-end adaptive training of hidden-unit bert for distortion-invariant robust speech recognition,","venue":null,"work_id":"3a771e6d-b5d6-4e0e-bb71-6c5500820466","year":2023},"citing_paper":{"arxiv_id":"2506.12260","last_updated":"2025-08-22T06:43:05Z","snapshot_observed_at":"2026-08-11T15:08:26.236040Z","submitted_at":"2025-06-13T22:22:19Z","title":"Improving Speech Enhancement with Multi-Metric Supervision from Learned Quality Assessment","version":2},"reference_index":6,"source":"pdf_text","source_observed_at":"2026-08-07T00:59:51.863593Z"},"links":{"citing_paper":"/paper/2506.12260"},"observation_digest":"sha256:db06164bdb605b8685b826b551afd53313e89a5cdebce390cd2f7d91566cbce7","observation_id":"2c90d128-7dfd-4885-bdaa-c7e8af5ba188","resolution":{"observed_at":"2026-08-07T01:00:02.555239Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-12T06:34:41.77262+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-12T06:34:41.77262+00:00","source":"crossref"},{"observed_at":"2026-08-12T06:34:36.333875+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T01:00:02.467084Z","title":"Closing the gap between time-domain multi-channel speech enhancement on real and simulation conditions,","venue":null,"work_id":"b247b830-0711-4d64-8760-d15643d95cdd","year":2021},"citing_paper":{"arxiv_id":"2506.12260","last_updated":"2025-08-22T06:43:05Z","snapshot_observed_at":"2026-08-11T15:08:26.236040Z","submitted_at":"2025-06-13T22:22:19Z","title":"Improving Speech Enhancement with Multi-Metric Supervision from Learned Quality Assessment","version":2},"reference_index":7,"source":"pdf_text","source_observed_at":"2026-08-07T00:59:52.007591Z"},"links":{"citing_paper":"/paper/2506.12260"},"observation_digest":"sha256:4c0439256290a1ef4026a4f3867847286398d22b30dcc3dc2217b4e38a3a6bd5","observation_id":"aa56e567-42db-46b9-8ad8-67614693cccb","resolution":{"observed_at":"2026-08-07T01:00:02.496086Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-12T06:34:41.77262+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-12T06:34:41.77262+00:00","source":"crossref"},{"observed_at":"2026-08-12T06:34:36.333875+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2506.23859","last_updated":"2025-08-19T09:17:36Z","snapshot_observed_at":"2026-08-12T03:52:58.429144Z","submitted_at":"2025-06-30T13:55:10Z","title":"Less is More: Data Curation Matters in Scaling Speech Enhancement","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2506.23859","snapshot_observed_at":"2026-08-07T00:59:52.138592Z","title":"Less is more: Data curation matters in scaling speech enhancement,","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2506.12260","last_updated":"2025-08-22T06:43:05Z","snapshot_observed_at":"2026-08-11T15:08:26.236040Z","submitted_at":"2025-06-13T22:22:19Z","title":"Improving Speech Enhancement with Multi-Metric Supervision from Learned Quality Assessment","version":2},"reference_index":8,"source":"pdf_text","source_observed_at":"2026-08-07T00:59:52.138592Z"},"links":{"cited_paper":"/paper/2506.23859","citing_paper":"/paper/2506.12260"},"observation_digest":"sha256:3d41a61cdd2e974867f9213285b0b7c86dc3304a55aea81d197cf28164abb5eb","observation_id":"799acbfa-1e7f-4781-8821-d5544ec02d31","resolution":{"observed_at":"2026-08-07T00:59:52.138592Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T01:00:02.307482Z","title":"Lightweight Front-end Enhancement for Robust ASR via Frame Resampling and Sub-Band Pruning,","venue":null,"work_id":"325baf87-6147-436d-94f0-e98b12a63166","year":2025},"citing_paper":{"arxiv_id":"2506.12260","last_updated":"2025-08-22T06:43:05Z","snapshot_observed_at":"2026-08-11T15:08:26.236040Z","submitted_at":"2025-06-13T22:22:19Z","title":"Improving Speech Enhancement with Multi-Metric Supervision from Learned Quality Assessment","version":2},"reference_index":9,"source":"pdf_text","source_observed_at":"2026-08-07T00:59:52.248087Z"},"links":{"citing_paper":"/paper/2506.12260"},"observation_digest":"sha256:d9866ea20bdf5f29333715f66d57ae165a5a8460d569d81c0b0decb623a1fd95","observation_id":"96a43799-7db9-4a3f-874e-205fe6af9ebb","resolution":{"observed_at":"2026-08-07T01:00:02.437242Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-12T06:34:41.77262+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-12T06:34:41.77262+00:00","source":"crossref"},{"observed_at":"2026-08-12T06:34:36.333875+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T01:00:02.044149Z","title":"A review on subjective and objective evaluation of synthetic speech,","venue":null,"work_id":"de423ff6-9a2e-4ce0-ad35-3a026f12c318","year":2024},"citing_paper":{"arxiv_id":"2506.12260","last_updated":"2025-08-22T06:43:05Z","snapshot_observed_at":"2026-08-11T15:08:26.236040Z","submitted_at":"2025-06-13T22:22:19Z","title":"Improving Speech Enhancement with Multi-Metric Supervision from Learned Quality Assessment","version":2},"reference_index":10,"source":"pdf_text","source_observed_at":"2026-08-07T00:59:52.384331Z"},"links":{"citing_paper":"/paper/2506.12260"},"observation_digest":"sha256:e074dc145c4204660e2d57b83630f1b534a681c0db2773544e8f28a3501431a8","observation_id":"71ce7602-6300-4836-a3d3-a682f31ac579","resolution":{"observed_at":"2026-08-07T01:00:02.156596Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-12T06:34:41.77262+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-12T06:34:41.77262+00:00","source":"crossref"},{"observed_at":"2026-08-12T06:34:36.333875+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T01:00:01.855739Z","title":"Objective measures of perceptual audio quality reviewed: An evaluation of their application domain dependence,","venue":null,"work_id":"17f0a796-ef09-4539-b24e-163d09a606d9","year":2021},"citing_paper":{"arxiv_id":"2506.12260","last_updated":"2025-08-22T06:43:05Z","snapshot_observed_at":"2026-08-11T15:08:26.236040Z","submitted_at":"2025-06-13T22:22:19Z","title":"Improving Speech Enhancement with Multi-Metric Supervision from Learned Quality Assessment","version":2},"reference_index":11,"source":"pdf_text","source_observed_at":"2026-08-07T00:59:52.526177Z"},"links":{"citing_paper":"/paper/2506.12260"},"observation_digest":"sha256:096a07ec0430ac958fe111854807fdbff0ad0afcd34d7238760a21003fe7171a","observation_id":"9b077439-c7e1-4fdf-9589-6e8474f368cc","resolution":{"observed_at":"2026-08-07T01:00:01.982525Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-12T06:34:41.77262+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-12T06:34:41.77262+00:00","source":"crossref"},{"observed_at":"2026-08-12T06:34:36.333875+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T01:00:01.638559Z","title":"Versa: A versatile evaluation toolkit for speech, audio, and music,","venue":null,"work_id":"98636bb1-9154-4349-ad48-33e3ac52555e","year":2024},"citing_paper":{"arxiv_id":"2506.12260","last_updated":"2025-08-22T06:43:05Z","snapshot_observed_at":"2026-08-11T15:08:26.236040Z","submitted_at":"2025-06-13T22:22:19Z","title":"Improving Speech Enhancement with Multi-Metric Supervision from Learned Quality Assessment","version":2},"reference_index":12,"source":"pdf_text","source_observed_at":"2026-08-07T00:59:52.638340Z"},"links":{"citing_paper":"/paper/2506.12260"},"observation_digest":"sha256:06231f14477ee0445d53de8a00a1fc7837075157626b16ec0152eca39329c904","observation_id":"3e5aa4c9-92f2-4de0-b13b-3d3faf4c71ad","resolution":{"observed_at":"2026-08-07T01:00:01.716700Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-12T06:34:41.77262+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-12T06:34:41.77262+00:00","source":"crossref"},{"observed_at":"2026-08-12T06:34:36.333875+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2506.01611","last_updated":"2025-06-02T12:50:37Z","snapshot_observed_at":"2026-08-07T11:35:03.897104Z","submitted_at":"2025-06-02T12:50:37Z","title":"Lessons Learned from the URGENT 2024 Speech Enhancement Challenge","version":1},"cited_work":{"arxiv_id":"2506.01611","doi":null,"metadata_source":"pith","pith_arxiv_id":"2506.01611","snapshot_observed_at":"2026-08-07T00:59:56.760465Z","title":"Lessons Learned from the URGENT 2024 Speech Enhancement Challenge","venue":"eess.AS","work_id":"af144f5f-e926-40a2-b4f7-00508ad4f919","year":2025},"citing_paper":{"arxiv_id":"2506.12260","last_updated":"2025-08-22T06:43:05Z","snapshot_observed_at":"2026-08-11T15:08:26.236040Z","submitted_at":"2025-06-13T22:22:19Z","title":"Improving Speech Enhancement with Multi-Metric Supervision from Learned Quality Assessment","version":2},"reference_index":13,"source":"pdf_text","source_observed_at":"2026-08-07T00:59:52.754052Z"},"links":{"cited_paper":"/paper/2506.01611","citing_paper":"/paper/2506.12260"},"observation_digest":"sha256:1d5753aa73a9f4e67d08b7bc49ccf290b6c58f2d6e5403ed3b7bf53602629e2f","observation_id":"5d88ab18-9866-414a-93e4-188c6102559e","resolution":{"observed_at":"2026-08-07T00:59:56.812979Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-12T06:34:41.77262+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-12T06:34:41.77262+00:00","source":"crossref"},{"observed_at":"2026-08-12T06:34:36.333875+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2502.05139","last_updated":"2025-02-07T18:15:57Z","snapshot_observed_at":"2026-07-06T20:33:01.960703Z","submitted_at":"2025-02-07T18:15:57Z","title":"Meta Audiobox Aesthetics: Unified Automatic Quality Assessment for Speech, Music, and Sound","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2502.05139","snapshot_observed_at":"2026-08-07T00:59:52.863711Z","title":"Meta audiobox aesthetics: Unified automatic quality assessment for speech, music, and sound,","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2506.12260","last_updated":"2025-08-22T06:43:05Z","snapshot_observed_at":"2026-08-11T15:08:26.236040Z","submitted_at":"2025-06-13T22:22:19Z","title":"Improving Speech Enhancement with Multi-Metric Supervision from Learned Quality Assessment","version":2},"reference_index":14,"source":"pdf_text","source_observed_at":"2026-08-07T00:59:52.863711Z"},"links":{"cited_paper":"/paper/2502.05139","citing_paper":"/paper/2506.12260"},"observation_digest":"sha256:3ece1b0f0c20ab118717c1a3b452c5f8a4cace1278837768734704761c14112b","observation_id":"9fda4650-5519-426c-ad6a-23f20fcf6493","resolution":{"observed_at":"2026-08-07T00:59:52.863711Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T01:00:01.429502Z","title":"DNSMOS P.835: A non-intrusive perceptual objective speech quality metric to evaluate noise suppressors,","venue":null,"work_id":"3b547d78-8c3a-43d9-a20c-147c668147fa","year":2022},"citing_paper":{"arxiv_id":"2506.12260","last_updated":"2025-08-22T06:43:05Z","snapshot_observed_at":"2026-08-11T15:08:26.236040Z","submitted_at":"2025-06-13T22:22:19Z","title":"Improving Speech Enhancement with Multi-Metric Supervision from Learned Quality Assessment","version":2},"reference_index":15,"source":"pdf_text","source_observed_at":"2026-08-07T00:59:52.978436Z"},"links":{"citing_paper":"/paper/2506.12260"},"observation_digest":"sha256:33278712ac4869d8dc367d774f8f85d157bfa1055506eb6045cb3e1f00dfed0b","observation_id":"80f9f217-b210-4834-896a-ec540a5dcbd5","resolution":{"observed_at":"2026-08-07T01:00:01.527396Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-12T06:34:41.77262+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-12T06:34:41.77262+00:00","source":"crossref"},{"observed_at":"2026-08-12T06:34:36.333875+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T01:00:01.174139Z","title":"UTMOS: UTokyo-SaruLab system for V oiceMOS chal- lenge 2022,","venue":null,"work_id":"cdd69636-ff91-4670-9693-449637bb36c6","year":2022},"citing_paper":{"arxiv_id":"2506.12260","last_updated":"2025-08-22T06:43:05Z","snapshot_observed_at":"2026-08-11T15:08:26.236040Z","submitted_at":"2025-06-13T22:22:19Z","title":"Improving Speech Enhancement with Multi-Metric Supervision from Learned Quality Assessment","version":2},"reference_index":16,"source":"pdf_text","source_observed_at":"2026-08-07T00:59:53.142609Z"},"links":{"citing_paper":"/paper/2506.12260"},"observation_digest":"sha256:865373fa1e25e43e2e635c519b29c3b9aab76608f1c8b1b438c2bbc62b76dfa5","observation_id":"27c49265-0fba-42b5-bd78-2e6f52b627d0","resolution":{"observed_at":"2026-08-07T01:00:01.324060Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-12T06:34:41.77262+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-12T06:34:41.77262+00:00","source":"crossref"},{"observed_at":"2026-08-12T06:34:36.333875+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T01:00:00.804801Z","title":"The t05 system for the voicemos challenge 2024: Transfer learning from deep image classifier to naturalness mos prediction of high-quality synthetic speech,","venue":null,"work_id":"ffe6b44a-93ea-49b8-a93a-1516272dcabb","year":2024},"citing_paper":{"arxiv_id":"2506.12260","last_updated":"2025-08-22T06:43:05Z","snapshot_observed_at":"2026-08-11T15:08:26.236040Z","submitted_at":"2025-06-13T22:22:19Z","title":"Improving Speech Enhancement with Multi-Metric Supervision from Learned Quality Assessment","version":2},"reference_index":17,"source":"pdf_text","source_observed_at":"2026-08-07T00:59:53.243207Z"},"links":{"citing_paper":"/paper/2506.12260"},"observation_digest":"sha256:54ce326d944496edb7172209825aeb66ffd224595896d387af3ae917ae839bd4","observation_id":"97c161ad-e9fe-4bc3-9e0e-38510567681c","resolution":{"observed_at":"2026-08-07T01:00:01.010390Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-12T06:34:41.77262+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-12T06:34:41.77262+00:00","source":"crossref"},{"observed_at":"2026-08-12T06:34:36.333875+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T01:00:00.497999Z","title":"The voicemos challenge 2022,","venue":null,"work_id":"09deaa18-07c6-4bbf-958e-bf640fd0fc64","year":2022},"citing_paper":{"arxiv_id":"2506.12260","last_updated":"2025-08-22T06:43:05Z","snapshot_observed_at":"2026-08-11T15:08:26.236040Z","submitted_at":"2025-06-13T22:22:19Z","title":"Improving Speech Enhancement with Multi-Metric Supervision from Learned Quality Assessment","version":2},"reference_index":18,"source":"pdf_text","source_observed_at":"2026-08-07T00:59:53.351630Z"},"links":{"citing_paper":"/paper/2506.12260"},"observation_digest":"sha256:ace3850ecaeb23cc89fbe8755b514556fa793f3ba47695b37c7024858ddb4ed0","observation_id":"526ae2a3-fbfd-4f51-bcc8-6f7dc5585c70","resolution":{"observed_at":"2026-08-07T01:00:00.617091Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-12T06:34:41.77262+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-12T06:34:41.77262+00:00","source":"crossref"},{"observed_at":"2026-08-12T06:34:36.333875+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T01:00:00.330324Z","title":"The voicemos challenge 2023: Zero-shot subjective speech quality prediction for multiple domains,","venue":null,"work_id":"fef8a573-0a63-4e72-a682-660955e926da","year":2023},"citing_paper":{"arxiv_id":"2506.12260","last_updated":"2025-08-22T06:43:05Z","snapshot_observed_at":"2026-08-11T15:08:26.236040Z","submitted_at":"2025-06-13T22:22:19Z","title":"Improving Speech Enhancement with Multi-Metric Supervision from Learned Quality Assessment","version":2},"reference_index":19,"source":"pdf_text","source_observed_at":"2026-08-07T00:59:53.468606Z"},"links":{"citing_paper":"/paper/2506.12260"},"observation_digest":"sha256:bdd48a710695d4b959f8bbab29ad7ed046e03bd63f1095ec331abadf3b86321d","observation_id":"73aad329-0b7a-4d60-9102-eb326a1451c0","resolution":{"observed_at":"2026-08-07T01:00:00.399408Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-12T06:34:41.77262+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-12T06:34:41.77262+00:00","source":"crossref"},{"observed_at":"2026-08-12T06:34:36.333875+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T01:00:00.174318Z","title":"The voicemos challenge 2024: Beyond speech quality prediction,","venue":null,"work_id":"2b12063a-8cfa-4c9c-8925-7420d3388de3","year":2024},"citing_paper":{"arxiv_id":"2506.12260","last_updated":"2025-08-22T06:43:05Z","snapshot_observed_at":"2026-08-11T15:08:26.236040Z","submitted_at":"2025-06-13T22:22:19Z","title":"Improving Speech Enhancement with Multi-Metric Supervision from Learned Quality Assessment","version":2},"reference_index":20,"source":"pdf_text","source_observed_at":"2026-08-07T00:59:53.604460Z"},"links":{"citing_paper":"/paper/2506.12260"},"observation_digest":"sha256:26bd4358a2af92a6f4e73ffd7df9f0b5ee57ef39bf7712839f474bea3b0f856c","observation_id":"b1911d94-4095-45af-a45c-c3fd70cc2894","resolution":{"observed_at":"2026-08-07T01:00:00.262689Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-12T06:34:41.77262+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-12T06:34:41.77262+00:00","source":"crossref"},{"observed_at":"2026-08-12T06:34:36.333875+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2506.23874","last_updated":"2025-06-30T14:05:17Z","snapshot_observed_at":"2026-08-10T06:02:30.411704Z","submitted_at":"2025-06-30T14:05:17Z","title":"URGENT-PK: Perceptually-Aligned Ranking Model Designed for Speech Enhancement Competition","version":1},"cited_work":{"arxiv_id":"2506.23874","doi":null,"metadata_source":"pith","pith_arxiv_id":"2506.23874","snapshot_observed_at":"2026-08-07T00:59:56.633705Z","title":"URGENT-PK: Perceptually-Aligned Ranking Model Designed for Speech Enhancement Competition","venue":"eess.AS","work_id":"95cfa651-bc11-405c-852c-9e1c71f5a631","year":2025},"citing_paper":{"arxiv_id":"2506.12260","last_updated":"2025-08-22T06:43:05Z","snapshot_observed_at":"2026-08-11T15:08:26.236040Z","submitted_at":"2025-06-13T22:22:19Z","title":"Improving Speech Enhancement with Multi-Metric Supervision from Learned Quality Assessment","version":2},"reference_index":21,"source":"pdf_text","source_observed_at":"2026-08-07T00:59:53.700804Z"},"links":{"cited_paper":"/paper/2506.23874","citing_paper":"/paper/2506.12260"},"observation_digest":"sha256:1d2e20ba3d03ff5129100f3722bf9ff76719bebc066610af7b5685cdc9b0bfb8","observation_id":"828a28ec-a6a1-4453-be17-03e14ef6c3ef","resolution":{"observed_at":"2026-08-07T00:59:56.682451Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-12T06:34:41.77262+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-12T06:34:41.77262+00:00","source":"crossref"},{"observed_at":"2026-08-12T06:34:36.333875+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T01:00:00.050869Z","title":"ICASSP 2024 speech signal improvement challenge,","venue":null,"work_id":"a6d1335e-8fc6-48bc-82af-a212d322883c","year":2024},"citing_paper":{"arxiv_id":"2506.12260","last_updated":"2025-08-22T06:43:05Z","snapshot_observed_at":"2026-08-11T15:08:26.236040Z","submitted_at":"2025-06-13T22:22:19Z","title":"Improving Speech Enhancement with Multi-Metric Supervision from Learned Quality Assessment","version":2},"reference_index":22,"source":"pdf_text","source_observed_at":"2026-08-07T00:59:53.768549Z"},"links":{"citing_paper":"/paper/2506.12260"},"observation_digest":"sha256:e18947eb0e0bfe137edf88e92f1c5b0fea8aca4c5c2b2a328d6cebde0ae6afb8","observation_id":"1610e23b-d9f4-4649-8891-e3c20ab1fbde","resolution":{"observed_at":"2026-08-07T01:00:00.101072Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-12T06:34:41.77262+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-12T06:34:41.77262+00:00","source":"crossref"},{"observed_at":"2026-08-12T06:34:36.333875+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2505.20741","last_updated":"2025-05-27T05:31:19Z","snapshot_observed_at":"2026-08-11T15:07:59.957029Z","submitted_at":"2025-05-27T05:31:19Z","title":"Uni-VERSA: Versatile Speech Assessment with a Unified Network","version":1},"cited_work":{"arxiv_id":"2505.20741","doi":null,"metadata_source":"pith","pith_arxiv_id":"2505.20741","snapshot_observed_at":"2026-08-07T00:59:56.523759Z","title":"Uni-VERSA: Versatile Speech Assessment with a Unified Network","venue":"cs.SD","work_id":"cdb00c1b-7340-444c-b69a-2e1b50468391","year":2025},"citing_paper":{"arxiv_id":"2506.12260","last_updated":"2025-08-22T06:43:05Z","snapshot_observed_at":"2026-08-11T15:08:26.236040Z","submitted_at":"2025-06-13T22:22:19Z","title":"Improving Speech Enhancement with Multi-Metric Supervision from Learned Quality Assessment","version":2},"reference_index":23,"source":"pdf_text","source_observed_at":"2026-08-07T00:59:53.899052Z"},"links":{"cited_paper":"/paper/2505.20741","citing_paper":"/paper/2506.12260"},"observation_digest":"sha256:6f2a9594a2868add881459c8a3a73d6b163ba8030a27eee95688af2ccacb70aa","observation_id":"64a7a598-988f-45c9-8b04-891bd755a587","resolution":{"observed_at":"2026-08-07T00:59:56.557082Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-12T06:34:41.77262+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-12T06:34:41.77262+00:00","source":"crossref"},{"observed_at":"2026-08-12T06:34:36.333875+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T00:59:59.874908Z","title":"Perceptual evaluation of speech quality (PESQ)—a new method for speech quality assessment of telephone networks and codecs,","venue":null,"work_id":"b3987799-dfb0-4cf5-95bb-198afe09c905","year":2001},"citing_paper":{"arxiv_id":"2506.12260","last_updated":"2025-08-22T06:43:05Z","snapshot_observed_at":"2026-08-11T15:08:26.236040Z","submitted_at":"2025-06-13T22:22:19Z","title":"Improving Speech Enhancement with Multi-Metric Supervision from Learned Quality Assessment","version":2},"reference_index":24,"source":"pdf_text","source_observed_at":"2026-08-07T00:59:54.030231Z"},"links":{"citing_paper":"/paper/2506.12260"},"observation_digest":"sha256:5b57e6173878447bd9d97608c4d1889dc9885bec62d60af859f996f9b3cdcb48","observation_id":"c6de50a5-1516-4526-a048-fd359071ccdd","resolution":{"observed_at":"2026-08-07T00:59:59.954336Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-12T06:34:41.77262+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-12T06:34:41.77262+00:00","source":"crossref"},{"observed_at":"2026-08-12T06:34:36.333875+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T00:59:59.690046Z","title":"Perceptual objective listening quality assess- ment (POLQA), the third generation ITU-T standard for end-to-end speech quality measurement part I–—temporal alignment,","venue":null,"work_id":"5330cf9a-34cb-4af8-bade-9bef379c524c","year":2013},"citing_paper":{"arxiv_id":"2506.12260","last_updated":"2025-08-22T06:43:05Z","snapshot_observed_at":"2026-08-11T15:08:26.236040Z","submitted_at":"2025-06-13T22:22:19Z","title":"Improving Speech Enhancement with Multi-Metric Supervision from Learned Quality Assessment","version":2},"reference_index":25,"source":"pdf_text","source_observed_at":"2026-08-07T00:59:54.165739Z"},"links":{"citing_paper":"/paper/2506.12260"},"observation_digest":"sha256:7af12f15b132691b93d2a4e71496c336c464318bdc395a2c4b39e98586784ef1","observation_id":"717af5b2-ce14-497d-a98e-e3e55a746213","resolution":{"observed_at":"2026-08-07T00:59:59.760642Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-12T06:34:41.77262+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-12T06:34:41.77262+00:00","source":"crossref"},{"observed_at":"2026-08-12T06:34:36.333875+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T00:59:59.547647Z","title":"URGENT challenge: Universality, robustness, and generalizability for speech en- hancement,","venue":null,"work_id":"9ccdd92f-2989-4421-916a-89274503ee6d","year":2024},"citing_paper":{"arxiv_id":"2506.12260","last_updated":"2025-08-22T06:43:05Z","snapshot_observed_at":"2026-08-11T15:08:26.236040Z","submitted_at":"2025-06-13T22:22:19Z","title":"Improving Speech Enhancement with Multi-Metric Supervision from Learned Quality Assessment","version":2},"reference_index":26,"source":"pdf_text","source_observed_at":"2026-08-07T00:59:54.277950Z"},"links":{"citing_paper":"/paper/2506.12260"},"observation_digest":"sha256:4039250c727fdef2389d2ddd59cab0eefe5a9003a5c69390470d7e6c9d757501","observation_id":"286b5014-978a-4f2d-a3a5-d389b67e1ec7","resolution":{"observed_at":"2026-08-07T00:59:59.624900Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-12T06:34:41.77262+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-12T06:34:41.77262+00:00","source":"crossref"},{"observed_at":"2026-08-12T06:34:36.333875+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T00:59:59.422542Z","title":"Inter- speech 2025 URGENT speech enhancement challenge,","venue":null,"work_id":"5146e0bd-e107-4c78-b74c-af42e651b18b","year":2025},"citing_paper":{"arxiv_id":"2506.12260","last_updated":"2025-08-22T06:43:05Z","snapshot_observed_at":"2026-08-11T15:08:26.236040Z","submitted_at":"2025-06-13T22:22:19Z","title":"Improving Speech Enhancement with Multi-Metric Supervision from Learned Quality Assessment","version":2},"reference_index":27,"source":"pdf_text","source_observed_at":"2026-08-07T00:59:54.413312Z"},"links":{"citing_paper":"/paper/2506.12260"},"observation_digest":"sha256:a96cddf1c4f3cd84646ce57c2fd06f9db6eae09b6f37100f24c4a7b45a70baf1","observation_id":"b3bbf6c2-0e84-4dc6-8791-bbab95095f28","resolution":{"observed_at":"2026-08-07T00:59:59.472832Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-12T06:34:41.77262+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-12T06:34:41.77262+00:00","source":"crossref"},{"observed_at":"2026-08-12T06:34:36.333875+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T00:59:59.295446Z","title":"Performance measurement in blind audio source separation,","venue":null,"work_id":"ec524388-313e-4252-9f57-c8c9092336c2","year":2006},"citing_paper":{"arxiv_id":"2506.12260","last_updated":"2025-08-22T06:43:05Z","snapshot_observed_at":"2026-08-11T15:08:26.236040Z","submitted_at":"2025-06-13T22:22:19Z","title":"Improving Speech Enhancement with Multi-Metric Supervision from Learned Quality Assessment","version":2},"reference_index":28,"source":"pdf_text","source_observed_at":"2026-08-07T00:59:54.511671Z"},"links":{"citing_paper":"/paper/2506.12260"},"observation_digest":"sha256:57bbdc94d92f60d6c9fccffe2dac09e857de7c9ef7a7db8fc6b5228393e0ee8b","observation_id":"13f11994-095c-42dd-ac5b-d32284d9a7c2","resolution":{"observed_at":"2026-08-07T00:59:59.343918Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-12T06:34:41.77262+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-12T06:34:41.77262+00:00","source":"crossref"},{"observed_at":"2026-08-12T06:34:36.333875+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T00:59:59.124792Z","title":"Distillation and pruning for scalable self- supervised representation-based speech quality assessment,","venue":null,"work_id":"5f7cdd23-6c0a-486f-8143-5a14f83d273f","year":2025},"citing_paper":{"arxiv_id":"2506.12260","last_updated":"2025-08-22T06:43:05Z","snapshot_observed_at":"2026-08-11T15:08:26.236040Z","submitted_at":"2025-06-13T22:22:19Z","title":"Improving Speech Enhancement with Multi-Metric Supervision from Learned Quality Assessment","version":2},"reference_index":29,"source":"pdf_text","source_observed_at":"2026-08-07T00:59:54.648344Z"},"links":{"citing_paper":"/paper/2506.12260"},"observation_digest":"sha256:4ccfbb3f950c69c47c019c10dc67d11c7f74ef94cf27fda52624836885454bf8","observation_id":"78437f9c-7a8c-442d-8745-7c857bbeb1ab","resolution":{"observed_at":"2026-08-07T00:59:59.213900Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-12T06:34:41.77262+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-12T06:34:41.77262+00:00","source":"crossref"},{"observed_at":"2026-08-12T06:34:36.333875+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T00:59:58.975232Z","title":"NISQA: A deep CNN- self-attention model for multidimensional speech quality prediction with crowdsourced datasets,","venue":null,"work_id":"876db6ed-0e11-4aef-bc7a-3615bb9d661d","year":2021},"citing_paper":{"arxiv_id":"2506.12260","last_updated":"2025-08-22T06:43:05Z","snapshot_observed_at":"2026-08-11T15:08:26.236040Z","submitted_at":"2025-06-13T22:22:19Z","title":"Improving Speech Enhancement with Multi-Metric Supervision from Learned Quality Assessment","version":2},"reference_index":30,"source":"pdf_text","source_observed_at":"2026-08-07T00:59:54.745525Z"},"links":{"citing_paper":"/paper/2506.12260"},"observation_digest":"sha256:768bdea62a100d4c0c4454f9c77e0f06709d81ddfbeaff9e520c1d321a61c2a2","observation_id":"daf7cdf0-5a0b-4cd1-a709-c00337bfc8d7","resolution":{"observed_at":"2026-08-07T00:59:59.041433Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-12T06:34:41.77262+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-12T06:34:41.77262+00:00","source":"crossref"},{"observed_at":"2026-08-12T06:34:36.333875+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T00:59:58.785907Z","title":"SCOREQ: Speech quality assessment with contrastive regression,","venue":null,"work_id":"1c8651ea-37b3-4e22-88ea-ee89d37c2bd6","year":2024},"citing_paper":{"arxiv_id":"2506.12260","last_updated":"2025-08-22T06:43:05Z","snapshot_observed_at":"2026-08-11T15:08:26.236040Z","submitted_at":"2025-06-13T22:22:19Z","title":"Improving Speech Enhancement with Multi-Metric Supervision from Learned Quality Assessment","version":2},"reference_index":31,"source":"pdf_text","source_observed_at":"2026-08-07T00:59:54.832101Z"},"links":{"citing_paper":"/paper/2506.12260"},"observation_digest":"sha256:dda227730b3f60ae461803bdc46efcb8b370728ff509fd7f7ba0152052408904","observation_id":"963dbb25-bd7e-4808-9049-20944ecea082","resolution":{"observed_at":"2026-08-07T00:59:58.879269Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-12T06:34:41.77262+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-12T06:34:41.77262+00:00","source":"crossref"},{"observed_at":"2026-08-12T06:34:36.333875+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T00:59:58.643919Z","title":"Owsm v3. 1: Better and faster open whisper-style speech models based on e-branchformer,","venue":null,"work_id":"978089c7-27c4-4ed9-ae5b-b3987dba6e03","year":2024},"citing_paper":{"arxiv_id":"2506.12260","last_updated":"2025-08-22T06:43:05Z","snapshot_observed_at":"2026-08-11T15:08:26.236040Z","submitted_at":"2025-06-13T22:22:19Z","title":"Improving Speech Enhancement with Multi-Metric Supervision from Learned Quality Assessment","version":2},"reference_index":32,"source":"pdf_text","source_observed_at":"2026-08-07T00:59:54.923054Z"},"links":{"citing_paper":"/paper/2506.12260"},"observation_digest":"sha256:c1bef2be0f38e1857ca5e5f66e1cb6f5b080b0721fab56c3c55a0e795663578b","observation_id":"0faae484-a6a3-4fa9-b9cf-a0852255993b","resolution":{"observed_at":"2026-08-07T00:59:58.708717Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-12T06:34:41.77262+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-12T06:34:41.77262+00:00","source":"crossref"},{"observed_at":"2026-08-12T06:34:36.333875+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T00:59:58.541765Z","title":"An algorithm for predicting the intelligibility of speech masked by modulated noise maskers,","venue":null,"work_id":"23b29705-df92-4b92-bbdb-ae5aa90310a2","year":2009},"citing_paper":{"arxiv_id":"2506.12260","last_updated":"2025-08-22T06:43:05Z","snapshot_observed_at":"2026-08-11T15:08:26.236040Z","submitted_at":"2025-06-13T22:22:19Z","title":"Improving Speech Enhancement with Multi-Metric Supervision from Learned Quality Assessment","version":2},"reference_index":33,"source":"pdf_text","source_observed_at":"2026-08-07T00:59:55.025151Z"},"links":{"citing_paper":"/paper/2506.12260"},"observation_digest":"sha256:26ff436a8a60257685d06e5fd7f0d0a5d53feed1c54a7b40284a3e85b9cf172f","observation_id":"39fe393c-1152-4e33-bc00-37f251892288","resolution":{"observed_at":"2026-08-07T00:59:58.586149Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-12T06:34:41.77262+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-12T06:34:41.77262+00:00","source":"crossref"},{"observed_at":"2026-08-12T06:34:36.333875+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T00:59:58.458714Z","title":"SpeechBERTScore: Reference-aware automatic evaluation of speech generation leveraging NLP evaluation metrics,","venue":null,"work_id":"de2e3d8e-0d68-4609-92d6-f0d734ba5e1d","year":2024},"citing_paper":{"arxiv_id":"2506.12260","last_updated":"2025-08-22T06:43:05Z","snapshot_observed_at":"2026-08-11T15:08:26.236040Z","submitted_at":"2025-06-13T22:22:19Z","title":"Improving Speech Enhancement with Multi-Metric Supervision from Learned Quality Assessment","version":2},"reference_index":34,"source":"pdf_text","source_observed_at":"2026-08-07T00:59:55.155079Z"},"links":{"citing_paper":"/paper/2506.12260"},"observation_digest":"sha256:995f36c7620857425265287560bbd1be3ac5be6cad8a794ad4ad3e4fac3c9281","observation_id":"7be65e64-7583-49de-98a1-d3ab0c68ef69","resolution":{"observed_at":"2026-08-07T00:59:58.488141Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-12T06:34:41.77262+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-12T06:34:41.77262+00:00","source":"crossref"},{"observed_at":"2026-08-12T06:34:36.333875+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T00:59:55.245748Z","title":"Hubert: Self-supervised speech representation learning by masked prediction of hidden units,","venue":null,"work_id":null,"year":2021},"citing_paper":{"arxiv_id":"2506.12260","last_updated":"2025-08-22T06:43:05Z","snapshot_observed_at":"2026-08-11T15:08:26.236040Z","submitted_at":"2025-06-13T22:22:19Z","title":"Improving Speech Enhancement with Multi-Metric Supervision from Learned Quality Assessment","version":2},"reference_index":35,"source":"pdf_text","source_observed_at":"2026-08-07T00:59:55.245748Z"},"links":{"citing_paper":"/paper/2506.12260"},"observation_digest":"sha256:eef05a6a82ba63b5de721358d18905dce85b928646fa13e52415ea8ed814600f","observation_id":"5acf0661-06e7-448f-80d5-79851f3403fe","resolution":{"observed_at":"2026-08-07T00:59:55.245748Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T00:59:58.299665Z","title":"Evaluation metrics for generative speech enhancement methods: Issues and perspectives,","venue":null,"work_id":"8981e35c-4a5b-4938-9b1f-576c9a764840","year":2023},"citing_paper":{"arxiv_id":"2506.12260","last_updated":"2025-08-22T06:43:05Z","snapshot_observed_at":"2026-08-11T15:08:26.236040Z","submitted_at":"2025-06-13T22:22:19Z","title":"Improving Speech Enhancement with Multi-Metric Supervision from Learned Quality Assessment","version":2},"reference_index":36,"source":"pdf_text","source_observed_at":"2026-08-07T00:59:55.351874Z"},"links":{"citing_paper":"/paper/2506.12260"},"observation_digest":"sha256:522cbc63379d5dd3e578773371c22e3866b798d987208973450f79e94418e3d1","observation_id":"4abb97a3-c6e6-43af-b5f2-a7c70b786a03","resolution":{"observed_at":"2026-08-07T00:59:58.367515Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-12T06:34:41.77262+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-12T06:34:41.77262+00:00","source":"crossref"},{"observed_at":"2026-08-12T06:34:36.333875+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T00:59:58.132328Z","title":"Espnet-spk: full pipeline speaker embedding toolkit with reproducible recipes, self- supervised front-ends, and off-the-shelf models,","venue":null,"work_id":"55998d1e-1192-4674-848b-a475f0f23bf9","year":2024},"citing_paper":{"arxiv_id":"2506.12260","last_updated":"2025-08-22T06:43:05Z","snapshot_observed_at":"2026-08-11T15:08:26.236040Z","submitted_at":"2025-06-13T22:22:19Z","title":"Improving Speech Enhancement with Multi-Metric Supervision from Learned Quality Assessment","version":2},"reference_index":37,"source":"pdf_text","source_observed_at":"2026-08-07T00:59:55.433303Z"},"links":{"citing_paper":"/paper/2506.12260"},"observation_digest":"sha256:ef69a144afa431f9225d3d8eb56260542aa62d58a086cfa07a62f22fd92a1983","observation_id":"ba36eb9c-7f0d-4ea1-bdc6-2d178ed17c46","resolution":{"observed_at":"2026-08-07T00:59:58.188012Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-12T06:34:41.77262+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-12T06:34:41.77262+00:00","source":"crossref"},{"observed_at":"2026-08-12T06:34:36.333875+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T00:59:58.004147Z","title":"Mel-cepstral distance measure for objective speech qual- ity assessment,","venue":null,"work_id":"2d4f79b3-61c4-47ba-ad33-041b8f1e630b","year":1993},"citing_paper":{"arxiv_id":"2506.12260","last_updated":"2025-08-22T06:43:05Z","snapshot_observed_at":"2026-08-11T15:08:26.236040Z","submitted_at":"2025-06-13T22:22:19Z","title":"Improving Speech Enhancement with Multi-Metric Supervision from Learned Quality Assessment","version":2},"reference_index":38,"source":"pdf_text","source_observed_at":"2026-08-07T00:59:55.488912Z"},"links":{"citing_paper":"/paper/2506.12260"},"observation_digest":"sha256:22284875abb0eb6b63c8a85fb002f3ac11d45dcbbdadcec2c7eff3ac0693f242","observation_id":"825bc888-e9c2-43d5-bc37-41bf4fb89d03","resolution":{"observed_at":"2026-08-07T00:59:58.069927Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-12T06:34:41.77262+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-12T06:34:41.77262+00:00","source":"crossref"},{"observed_at":"2026-08-12T06:34:36.333875+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T00:59:57.838552Z","title":"Lessons learned from the URGENT 2024 speech enhancement chal- lenge,","venue":null,"work_id":"3c2ad75a-da88-4281-ae34-bdf61ca08349","year":2024},"citing_paper":{"arxiv_id":"2506.12260","last_updated":"2025-08-22T06:43:05Z","snapshot_observed_at":"2026-08-11T15:08:26.236040Z","submitted_at":"2025-06-13T22:22:19Z","title":"Improving Speech Enhancement with Multi-Metric Supervision from Learned Quality Assessment","version":2},"reference_index":39,"source":"pdf_text","source_observed_at":"2026-08-07T00:59:55.528778Z"},"links":{"citing_paper":"/paper/2506.12260"},"observation_digest":"sha256:c8b452a3cd041197fbfb204262a4e2d1c7b8bb8c2ebd340f0e6808f7f3e0e156","observation_id":"f53e7362-7fc5-4143-8137-8d17d957a590","resolution":{"observed_at":"2026-08-07T00:59:57.928233Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-12T06:34:41.77262+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-12T06:34:41.77262+00:00","source":"crossref"},{"observed_at":"2026-08-12T06:34:36.333875+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T00:59:57.729300Z","title":"Distance measures for speech processing,","venue":null,"work_id":"5bcd6ee4-fbed-4184-b9dc-2a8d3392f005","year":1976},"citing_paper":{"arxiv_id":"2506.12260","last_updated":"2025-08-22T06:43:05Z","snapshot_observed_at":"2026-08-11T15:08:26.236040Z","submitted_at":"2025-06-13T22:22:19Z","title":"Improving Speech Enhancement with Multi-Metric Supervision from Learned Quality Assessment","version":2},"reference_index":40,"source":"pdf_text","source_observed_at":"2026-08-07T00:59:55.577652Z"},"links":{"citing_paper":"/paper/2506.12260"},"observation_digest":"sha256:1985dc940bd52ffc7bcad6b46c86d25fe9592858d704b41e06fc9afe40b1a9f6","observation_id":"f8d7e008-305d-4c47-a548-798d160bd59d","resolution":{"observed_at":"2026-08-07T00:59:57.771171Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-12T06:34:41.77262+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-12T06:34:41.77262+00:00","source":"crossref"},{"observed_at":"2026-08-12T06:34:36.333875+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2505.15061","last_updated":"2025-05-21T03:30:23Z","snapshot_observed_at":"2026-08-11T15:08:24.695331Z","submitted_at":"2025-05-21T03:30:23Z","title":"SHEET: A Multi-purpose Open-source Speech Human Evaluation Estimation Toolkit","version":1},"cited_work":{"arxiv_id":"2505.15061","doi":null,"metadata_source":"pith","pith_arxiv_id":"2505.15061","snapshot_observed_at":"2026-08-07T00:59:56.411210Z","title":"SHEET: A Multi-purpose Open-source Speech Human Evaluation Estimation Toolkit","venue":"cs.SD","work_id":"7c98b9b2-c521-450b-8def-4d2fc50711bd","year":2025},"citing_paper":{"arxiv_id":"2506.12260","last_updated":"2025-08-22T06:43:05Z","snapshot_observed_at":"2026-08-11T15:08:26.236040Z","submitted_at":"2025-06-13T22:22:19Z","title":"Improving Speech Enhancement with Multi-Metric Supervision from Learned Quality Assessment","version":2},"reference_index":41,"source":"pdf_text","source_observed_at":"2026-08-07T00:59:55.643986Z"},"links":{"cited_paper":"/paper/2505.15061","citing_paper":"/paper/2506.12260"},"observation_digest":"sha256:a31cbd4d2e18302989e283150ec9f507f0ac7ec6a8aefff6bca6851e24a0c6f3","observation_id":"7b40b310-2311-414f-955d-0e1c5e80c18d","resolution":{"observed_at":"2026-08-07T00:59:56.456405Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-12T06:34:41.77262+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-12T06:34:41.77262+00:00","source":"crossref"},{"observed_at":"2026-08-12T06:34:36.333875+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T00:59:57.575877Z","title":"The chime-7 udase task: Unsupervised domain adaptation for conversational speech enhancement,","venue":null,"work_id":"a38b05a9-1ba1-4551-9541-40834ab4f18e","year":2023},"citing_paper":{"arxiv_id":"2506.12260","last_updated":"2025-08-22T06:43:05Z","snapshot_observed_at":"2026-08-11T15:08:26.236040Z","submitted_at":"2025-06-13T22:22:19Z","title":"Improving Speech Enhancement with Multi-Metric Supervision from Learned Quality Assessment","version":2},"reference_index":42,"source":"pdf_text","source_observed_at":"2026-08-07T00:59:55.698323Z"},"links":{"citing_paper":"/paper/2506.12260"},"observation_digest":"sha256:471fc9890465d112b399fdaaee04fa834a54050203ee7602a1b7a30cb8e7c7db","observation_id":"3013255a-d3d3-4827-bc4f-2c9f25c1cac1","resolution":{"observed_at":"2026-08-07T00:59:57.667117Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-12T06:34:41.77262+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-12T06:34:41.77262+00:00","source":"crossref"},{"observed_at":"2026-08-12T06:34:36.333875+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T00:59:57.465390Z","title":"Generalization ability of mos prediction networks,","venue":null,"work_id":"ceb1e584-5ce5-4a81-ad66-237c986c7d93","year":2022},"citing_paper":{"arxiv_id":"2506.12260","last_updated":"2025-08-22T06:43:05Z","snapshot_observed_at":"2026-08-11T15:08:26.236040Z","submitted_at":"2025-06-13T22:22:19Z","title":"Improving Speech Enhancement with Multi-Metric Supervision from Learned Quality Assessment","version":2},"reference_index":43,"source":"pdf_text","source_observed_at":"2026-08-07T00:59:55.749707Z"},"links":{"citing_paper":"/paper/2506.12260"},"observation_digest":"sha256:464febc6f183a80c992e164eeb505e0aa9e30809f4013e55f5904ad8a14578de","observation_id":"fba4b467-5ee4-4b4f-8a1b-3e52edcded0e","resolution":{"observed_at":"2026-08-07T00:59:57.516889Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-12T06:34:41.77262+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-12T06:34:41.77262+00:00","source":"crossref"},{"observed_at":"2026-08-12T06:34:36.333875+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T00:59:57.331298Z","title":"The blizzard challenge 2019,","venue":null,"work_id":"e5dec4ca-b4cb-4c74-94c1-76a778f3dcfe","year":2019},"citing_paper":{"arxiv_id":"2506.12260","last_updated":"2025-08-22T06:43:05Z","snapshot_observed_at":"2026-08-11T15:08:26.236040Z","submitted_at":"2025-06-13T22:22:19Z","title":"Improving Speech Enhancement with Multi-Metric Supervision from Learned Quality Assessment","version":2},"reference_index":44,"source":"pdf_text","source_observed_at":"2026-08-07T00:59:55.796894Z"},"links":{"citing_paper":"/paper/2506.12260"},"observation_digest":"sha256:aa3eb8bd0e3effab73f1c80f86a778612604e459e5dd273ba100140c0dd16de6","observation_id":"dcf98de6-a306-46dd-a800-20bd7f478278","resolution":{"observed_at":"2026-08-07T00:59:57.400223Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-12T06:34:41.77262+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-12T06:34:41.77262+00:00","source":"crossref"},{"observed_at":"2026-08-12T06:34:36.333875+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T00:59:55.841645Z","title":"Mos-bench: Benchmarking generalization abilities of subjective speech quality assessment models,","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2506.12260","last_updated":"2025-08-22T06:43:05Z","snapshot_observed_at":"2026-08-11T15:08:26.236040Z","submitted_at":"2025-06-13T22:22:19Z","title":"Improving Speech Enhancement with Multi-Metric Supervision from Learned Quality Assessment","version":2},"reference_index":45,"source":"pdf_text","source_observed_at":"2026-08-07T00:59:55.841645Z"},"links":{"citing_paper":"/paper/2506.12260"},"observation_digest":"sha256:5562c36b0fd0d8611c43304b2a187797a4065a301d185883896e08d34fb6b2e6","observation_id":"6cc6fd5e-383a-4eef-a2d3-8f238e7618a1","resolution":{"observed_at":"2026-08-07T00:59:55.841645Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T00:59:57.212467Z","title":"The INTERSPEECH 2020 deep noise suppression challenge: Datasets, subjective testing framework, and challenge results,","venue":null,"work_id":"1ec21386-60ea-413a-82a3-be45f60fa5eb","year":2020},"citing_paper":{"arxiv_id":"2506.12260","last_updated":"2025-08-22T06:43:05Z","snapshot_observed_at":"2026-08-11T15:08:26.236040Z","submitted_at":"2025-06-13T22:22:19Z","title":"Improving Speech Enhancement with Multi-Metric Supervision from Learned Quality Assessment","version":2},"reference_index":46,"source":"pdf_text","source_observed_at":"2026-08-07T00:59:55.967591Z"},"links":{"citing_paper":"/paper/2506.12260"},"observation_digest":"sha256:796667200d28b184b5110f55af340286bf2d12d462cae6e667686b972f0c7053","observation_id":"57d83909-e1ff-4e8e-9e5e-0ce9d9f0bbdd","resolution":{"observed_at":"2026-08-07T00:59:57.250392Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-12T06:34:41.77262+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-12T06:34:41.77262+00:00","source":"crossref"},{"observed_at":"2026-08-12T06:34:36.333875+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T00:59:57.124243Z","title":"An analysis of environment, microphone and data simulation mismatches in robust speech recognition,","venue":null,"work_id":"25b64b03-9ba5-4d58-a57b-a43b43cb50c6","year":2017},"citing_paper":{"arxiv_id":"2506.12260","last_updated":"2025-08-22T06:43:05Z","snapshot_observed_at":"2026-08-11T15:08:26.236040Z","submitted_at":"2025-06-13T22:22:19Z","title":"Improving Speech Enhancement with Multi-Metric Supervision from Learned Quality Assessment","version":2},"reference_index":47,"source":"pdf_text","source_observed_at":"2026-08-07T00:59:56.011864Z"},"links":{"citing_paper":"/paper/2506.12260"},"observation_digest":"sha256:06ba92003012ecc7fed3de3ba9c5a39ba392614a3eb6a2c73fe5de7830a93c6a","observation_id":"26c706da-74de-42f9-9042-ce318aa2c9fb","resolution":{"observed_at":"2026-08-07T00:59:57.155828Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-12T06:34:41.77262+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-12T06:34:41.77262+00:00","source":"crossref"},{"observed_at":"2026-08-12T06:34:36.333875+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"1804.00015","last_updated":"2018-03-30T18:09:39Z","snapshot_observed_at":"2026-08-10T10:28:25.759884Z","submitted_at":"2018-03-30T18:09:39Z","title":"ESPnet: End-to-End Speech Processing Toolkit","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"1804.00015","snapshot_observed_at":"2026-08-07T00:59:56.058833Z","title":"Espnet: End-to-end speech processing toolkit,","venue":null,"work_id":null,"year":2018},"citing_paper":{"arxiv_id":"2506.12260","last_updated":"2025-08-22T06:43:05Z","snapshot_observed_at":"2026-08-11T15:08:26.236040Z","submitted_at":"2025-06-13T22:22:19Z","title":"Improving Speech Enhancement with Multi-Metric Supervision from Learned Quality Assessment","version":2},"reference_index":48,"source":"pdf_text","source_observed_at":"2026-08-07T00:59:56.058833Z"},"links":{"cited_paper":"/paper/1804.00015","citing_paper":"/paper/2506.12260"},"observation_digest":"sha256:17b23e487e376a1d2a3a09fc00b0e7a3dca6439c90bb7cd2cd3532a6ef50abcf","observation_id":"7dd16f2b-a7b8-4757-87ff-b1e5e21a8ad7","resolution":{"observed_at":"2026-08-07T00:59:56.058833Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T00:59:56.102596Z","title":"Wavlm: Large-scale self-supervised pre- training for full stack speech processing,","venue":null,"work_id":null,"year":2022},"citing_paper":{"arxiv_id":"2506.12260","last_updated":"2025-08-22T06:43:05Z","snapshot_observed_at":"2026-08-11T15:08:26.236040Z","submitted_at":"2025-06-13T22:22:19Z","title":"Improving Speech Enhancement with Multi-Metric Supervision from Learned Quality Assessment","version":2},"reference_index":49,"source":"pdf_text","source_observed_at":"2026-08-07T00:59:56.102596Z"},"links":{"citing_paper":"/paper/2506.12260"},"observation_digest":"sha256:db0092b3d7ce6e9a251ce34cf8035adeb9b53d7c40e8799abed9727913907136","observation_id":"b99c0718-0b12-4817-8cb4-94b26172514d","resolution":{"observed_at":"2026-08-07T00:59:56.102596Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T00:59:56.156431Z","title":"SUPERB: Speech Processing Universal PERformance Benchmark,","venue":null,"work_id":null,"year":2021},"citing_paper":{"arxiv_id":"2506.12260","last_updated":"2025-08-22T06:43:05Z","snapshot_observed_at":"2026-08-11T15:08:26.236040Z","submitted_at":"2025-06-13T22:22:19Z","title":"Improving Speech Enhancement with Multi-Metric Supervision from Learned Quality Assessment","version":2},"reference_index":50,"source":"pdf_text","source_observed_at":"2026-08-07T00:59:56.156431Z"},"links":{"citing_paper":"/paper/2506.12260"},"observation_digest":"sha256:a2ea273ff2b92f71d8b2bc304df1c7fc64c36b2e8c89a689a7dac66e50be7527","observation_id":"336466e3-0897-4f1e-822f-4945ea164596","resolution":{"observed_at":"2026-08-07T00:59:56.156431Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T00:59:57.005769Z","title":"Music source separation with band-split rnn,","venue":null,"work_id":"ca8f2049-8772-4c04-8805-11c2508454f1","year":1901},"citing_paper":{"arxiv_id":"2506.12260","last_updated":"2025-08-22T06:43:05Z","snapshot_observed_at":"2026-08-11T15:08:26.236040Z","submitted_at":"2025-06-13T22:22:19Z","title":"Improving Speech Enhancement with Multi-Metric Supervision from Learned Quality Assessment","version":2},"reference_index":51,"source":"pdf_text","source_observed_at":"2026-08-07T00:59:56.215030Z"},"links":{"citing_paper":"/paper/2506.12260"},"observation_digest":"sha256:e4f384778a4b43d43ca0047f389e0230eccda030642f4a3aa0fd7b2d14b87c8c","observation_id":"5a8b846c-7665-4d28-965c-27e0a789b056","resolution":{"observed_at":"2026-08-07T00:59:57.051037Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-12T06:34:41.77262+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-12T06:34:41.77262+00:00","source":"crossref"},{"observed_at":"2026-08-12T06:34:36.333875+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T00:59:56.261770Z","title":"Towards deep learning models resistant to adversarial attacks,","venue":null,"work_id":null,"year":2018},"citing_paper":{"arxiv_id":"2506.12260","last_updated":"2025-08-22T06:43:05Z","snapshot_observed_at":"2026-08-11T15:08:26.236040Z","submitted_at":"2025-06-13T22:22:19Z","title":"Improving Speech Enhancement with Multi-Metric Supervision from Learned Quality Assessment","version":2},"reference_index":52,"source":"pdf_text","source_observed_at":"2026-08-07T00:59:56.261770Z"},"links":{"citing_paper":"/paper/2506.12260"},"observation_digest":"sha256:e151e32ec2e8ac65e27fdfc0a03703d1367962ea743acc9a9a17fd139a72e976","observation_id":"6fdeea55-fb72-4c5b-b4ea-702eb9a77ec6","resolution":{"observed_at":"2026-08-07T00:59:56.261770Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T00:59:56.909646Z","title":"Adversarial attacks on automatic speech recognition (asr): A survey,","venue":null,"work_id":"5c3d8848-0861-487b-8991-7bb36f0b08c9","year":2024},"citing_paper":{"arxiv_id":"2506.12260","last_updated":"2025-08-22T06:43:05Z","snapshot_observed_at":"2026-08-11T15:08:26.236040Z","submitted_at":"2025-06-13T22:22:19Z","title":"Improving Speech Enhancement with Multi-Metric Supervision from Learned Quality Assessment","version":2},"reference_index":53,"source":"pdf_text","source_observed_at":"2026-08-07T00:59:56.303718Z"},"links":{"citing_paper":"/paper/2506.12260"},"observation_digest":"sha256:312d4e8b2772b3bf296ca99c5a88a7f7de6057d193986b37ac51e754a793629c","observation_id":"7e6c78d8-7051-49ec-9026-26acd3bf90bd","resolution":{"observed_at":"2026-08-07T00:59:56.935708Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-12T06:34:41.77262+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-12T06:34:41.77262+00:00","source":"crossref"},{"observed_at":"2026-08-12T06:34:36.333875+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2411.03715","last_updated":"2026-04-24T06:25:24Z","snapshot_observed_at":"2026-07-06T19:46:02.181982Z","submitted_at":"2024-11-06T07:29:28Z","title":"MOS-Bench: Benchmarking Generalization Abilities of Subjective Speech Quality Assessment Models","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2411.03715","snapshot_observed_at":"2026-08-07T00:59:55.909879Z","title":"Available: https://arxiv.org/abs/2411.03715","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2506.12260","last_updated":"2025-08-22T06:43:05Z","snapshot_observed_at":"2026-08-11T15:08:26.236040Z","submitted_at":"2025-06-13T22:22:19Z","title":"Improving Speech Enhancement with Multi-Metric Supervision from Learned Quality Assessment","version":2},"reference_index":2024,"source":"pdf_text","source_observed_at":"2026-08-07T00:59:55.909879Z"},"links":{"cited_paper":"/paper/2411.03715","citing_paper":"/paper/2506.12260"},"observation_digest":"sha256:27e98217fb2334a8a91723178f31e6276855756b7ddd3bd3165fd4fa143de648","observation_id":"0ec3ff65-a3be-4a88-863f-a0c4a7c59ef8","resolution":{"observed_at":"2026-08-07T00:59:55.909879Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}}],"paper":{"arxiv_id":"2506.12260","last_updated":"2025-08-22T06:43:05Z","latest_version":2,"primary_category":"cs.SD","snapshot_observed_at":"2026-08-11T15:08:26.236040Z","submitted_at":"2025-06-13T22:22:19Z","title":"Improving Speech Enhancement with Multi-Metric Supervision from Learned Quality Assessment"},"reference_resolution":{"displayed":54,"state_counts":{"malformed_identifier":0,"metadata_mismatch":0,"parse_uncertain":0,"unresolved":10,"verified_exact":4,"verified_fuzzy":40},"total_outbound_references":54},"refusal":"A citation records a reference. It does not transfer a finding from one paper to another.","schema":"pith.paper-citation-record.v1","standing_sources":[{"observed_at":"2026-08-12T06:34:41.77262+00:00","source":"crossref"},{"observed_at":"2026-08-12T06:34:36.333875+00:00","source":"retraction_watch"}],"thesis":"As of 12 August 2026, this Paper Citation Record lists 54 of 54 outbound references and 1 inbound Pith citation observation for arXiv:2506.12260."}