{"as_of":"2026-08-13T11:10:00Z","caps":{"database_statements":6,"inbound":100,"outbound":100},"context_digest":"sha256:a7ac5995e5b9a8bfaa155b4f91a2152cc639abdbef958d639bf097a909a385a8","coverage":[{"denominator":39,"lane":"reference_resolution","note":"Typed states for the displayed outbound observations.","records_observed":39,"source":"paper_references, paper_reference_links","source_observed_at":"2026-08-06T14:59:53.498359Z","state":"measured"},{"denominator":40,"lane":"standing_notices","note":"One-hop event checks from named stored sources.","records_observed":40,"source":"scholarly_work_events, retraction_status_cache","source_observed_at":"2026-08-13T06:32:02.005865+00:00","state":"measured"},{"denominator":1,"lane":"inbound_itemization","note":"Pith citing papers itemized under the disclosed page cap.","records_observed":1,"source":"paper_references, paper_reference_links","source_observed_at":"2026-08-06T14:59:53.349227Z","state":"measured"},{"denominator":1,"lane":"external_citation_measurements","note":"A source-named dated measurement, never combined with another source.","records_observed":0,"source":"pith","source_observed_at":"2026-08-06T14:59:53.578636Z","state":"measured"}],"external_citation_measurements":[],"inbound":[{"citation":{"cited_paper":{"arxiv_id":"2507.17208","last_updated":"2025-07-23T04:55:28Z","snapshot_observed_at":"2026-08-13T08:38:34.334066Z","submitted_at":"2025-07-23T04:55:28Z","title":"SLASH: Self-Supervised Speech Pitch Estimation Leveraging DSP-derived Absolute Pitch","version":1},"cited_work":{"arxiv_id":"2507.17208","doi":null,"metadata_source":"pith","pith_arxiv_id":"2507.17208","snapshot_observed_at":"2026-08-06T14:59:53.578636Z","title":"SLASH: Self-Supervised Speech Pitch Estimation Leveraging DSP-derived Absolute Pitch","venue":"eess.AS","work_id":"a6b2eb6a-137a-471a-aa82-c3e09ee5d53d","year":2025},"citing_paper":{"arxiv_id":"2507.17208","last_updated":"2025-07-23T04:55:28Z","snapshot_observed_at":"2026-08-13T08:38:34.334066Z","submitted_at":"2025-07-23T04:55:28Z","title":"SLASH: Self-Supervised Speech Pitch Estimation Leveraging DSP-derived Absolute Pitch","version":1},"reference_index":2,"source":"pdf_text","source_observed_at":"2026-08-06T14:59:53.349227Z"},"links":{"cited_paper":"/paper/2507.17208","citing_paper":"/paper/2507.17208"},"observation_digest":"sha256:bf862dac6c60829fcd7ec2ef57e0b5cfe9a15c991cbf111c07eacc44a050cdc1","observation_id":"0557ffbd-8a97-4897-80f3-28dc7d2cd753","resolution":{"observed_at":"2026-08-06T14:59:53.582393Z","resolver_source":"local_arxiv","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-13T06:32:02.005865+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-13T06:32:02.005865+00:00","source":"crossref"},{"observed_at":"2026-08-13T06:31:53.387327+00:00","source":"retraction_watch"}],"state":"measured"}}],"links":{"evidence":"/evidence","html":"/paper/2507.17208/citation-record","integrity":"/paper/2507.17208/integrity","json":"/paper/2507.17208/citation-record.json","paper":"/paper/2507.17208"},"outbound":[{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T14:59:53.944207Z","title":"It has been applied to many kinds of applica- tions, such as text-to-speech and emotion recognition, among others [1, 2, 3]","venue":null,"work_id":"b838fd75-5e63-4cff-834f-69169195591c","year":null},"citing_paper":{"arxiv_id":"2507.17208","last_updated":"2025-07-23T04:55:28Z","snapshot_observed_at":"2026-08-13T08:38:34.334066Z","submitted_at":"2025-07-23T04:55:28Z","title":"SLASH: Self-Supervised Speech Pitch Estimation Leveraging DSP-derived Absolute Pitch","version":1},"reference_index":1,"source":"pdf_text","source_observed_at":"2026-08-06T14:59:53.345146Z"},"links":{"citing_paper":"/paper/2507.17208"},"observation_digest":"sha256:bf9da5735b9fe052db5767d37cc570349241f909aa5c96fcf1021c8ef2547e8c","observation_id":"c4132dcd-19d7-4bcc-9573-d40fb7f739e9","resolution":{"observed_at":"2026-08-06T14:59:53.948258Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-13T06:32:02.005865+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-13T06:32:02.005865+00:00","source":"crossref"},{"observed_at":"2026-08-13T06:31:53.387327+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2507.17208","last_updated":"2025-07-23T04:55:28Z","snapshot_observed_at":"2026-08-13T08:38:34.334066Z","submitted_at":"2025-07-23T04:55:28Z","title":"SLASH: Self-Supervised Speech Pitch Estimation Leveraging DSP-derived Absolute Pitch","version":1},"cited_work":{"arxiv_id":"2507.17208","doi":null,"metadata_source":"pith","pith_arxiv_id":"2507.17208","snapshot_observed_at":"2026-08-06T14:59:53.578636Z","title":"SLASH: Self-Supervised Speech Pitch Estimation Leveraging DSP-derived Absolute Pitch","venue":"eess.AS","work_id":"a6b2eb6a-137a-471a-aa82-c3e09ee5d53d","year":2025},"citing_paper":{"arxiv_id":"2507.17208","last_updated":"2025-07-23T04:55:28Z","snapshot_observed_at":"2026-08-13T08:38:34.334066Z","submitted_at":"2025-07-23T04:55:28Z","title":"SLASH: Self-Supervised Speech Pitch Estimation Leveraging DSP-derived Absolute Pitch","version":1},"reference_index":2,"source":"pdf_text","source_observed_at":"2026-08-06T14:59:53.349227Z"},"links":{"cited_paper":"/paper/2507.17208","citing_paper":"/paper/2507.17208"},"observation_digest":"sha256:bf862dac6c60829fcd7ec2ef57e0b5cfe9a15c991cbf111c07eacc44a050cdc1","observation_id":"0557ffbd-8a97-4897-80f3-28dc7d2cd753","resolution":{"observed_at":"2026-08-06T14:59:53.582393Z","resolver_source":"local_arxiv","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-13T06:32:02.005865+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-13T06:32:02.005865+00:00","source":"crossref"},{"observed_at":"2026-08-13T06:31:53.387327+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T14:59:53.933404Z","title":"Experimental setup Datasets: In our experiments, we used two datasets","venue":null,"work_id":"75959f6d-b3ae-470e-a92d-bfba04184693","year":null},"citing_paper":{"arxiv_id":"2507.17208","last_updated":"2025-07-23T04:55:28Z","snapshot_observed_at":"2026-08-13T08:38:34.334066Z","submitted_at":"2025-07-23T04:55:28Z","title":"SLASH: Self-Supervised Speech Pitch Estimation Leveraging DSP-derived Absolute Pitch","version":1},"reference_index":3,"source":"pdf_text","source_observed_at":"2026-08-06T14:59:53.353247Z"},"links":{"citing_paper":"/paper/2507.17208"},"observation_digest":"sha256:c99a1cad946eb828cf956821750205c15db3dbec7a7440cfbbf1a706671e2932","observation_id":"840a7197-cdaa-4ac2-a62a-854b663a06d4","resolution":{"observed_at":"2026-08-06T14:59:53.937831Z","resolver_source":"raw_fallback","status":"malformed_identifier"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-13T06:32:02.005865+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-13T06:32:02.005865+00:00","source":"crossref"},{"observed_at":"2026-08-13T06:31:53.387327+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T14:59:53.922311Z","title":"By incorporating the absolute pitch into the model, SLASH enhanced the pitch prediction accuracy of conventional SSL-based methods, which depend on relative pitch objectives","venue":null,"work_id":"c786aac7-92a0-4afd-a68d-3e05dc494027","year":null},"citing_paper":{"arxiv_id":"2507.17208","last_updated":"2025-07-23T04:55:28Z","snapshot_observed_at":"2026-08-13T08:38:34.334066Z","submitted_at":"2025-07-23T04:55:28Z","title":"SLASH: Self-Supervised Speech Pitch Estimation Leveraging DSP-derived Absolute Pitch","version":1},"reference_index":4,"source":"pdf_text","source_observed_at":"2026-08-06T14:59:53.356902Z"},"links":{"citing_paper":"/paper/2507.17208"},"observation_digest":"sha256:c6b5265edf30a624000aabf85d1fbbd9d0c8fd98f49697da8ac1bf31285e6ae7","observation_id":"1324bbdf-067c-4ad8-8505-30a905589631","resolution":{"observed_at":"2026-08-06T14:59:53.926436Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-13T06:32:02.005865+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-13T06:32:02.005865+00:00","source":"crossref"},{"observed_at":"2026-08-13T06:31:53.387327+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2106.15561","last_updated":"2021-07-23T12:32:52Z","snapshot_observed_at":"2026-08-08T12:28:23.489422Z","submitted_at":"2021-06-29T16:50:51Z","title":"A Survey on Neural Speech Synthesis","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2106.15561","snapshot_observed_at":"2026-08-06T14:59:53.360356Z","title":"A survey on neural speech synthesis,","venue":null,"work_id":null,"year":2021},"citing_paper":{"arxiv_id":"2507.17208","last_updated":"2025-07-23T04:55:28Z","snapshot_observed_at":"2026-08-13T08:38:34.334066Z","submitted_at":"2025-07-23T04:55:28Z","title":"SLASH: Self-Supervised Speech Pitch Estimation Leveraging DSP-derived Absolute Pitch","version":1},"reference_index":5,"source":"pdf_text","source_observed_at":"2026-08-06T14:59:53.360356Z"},"links":{"cited_paper":"/paper/2106.15561","citing_paper":"/paper/2507.17208"},"observation_digest":"sha256:96bcdbff66903ec8b6cc3b9ee7d5c41bb65a6a46b4814e9bff70d0fcb45e59aa","observation_id":"bbec28ee-caa4-468a-ba38-b175895f5f7d","resolution":{"observed_at":"2026-08-06T14:59:53.360356Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T14:59:53.911655Z","title":"Speech emotion recognition using deep learning techniques: A review,","venue":null,"work_id":"d750e9a5-acfd-42fc-bef2-016a52653b3b","year":2019},"citing_paper":{"arxiv_id":"2507.17208","last_updated":"2025-07-23T04:55:28Z","snapshot_observed_at":"2026-08-13T08:38:34.334066Z","submitted_at":"2025-07-23T04:55:28Z","title":"SLASH: Self-Supervised Speech Pitch Estimation Leveraging DSP-derived Absolute Pitch","version":1},"reference_index":6,"source":"pdf_text","source_observed_at":"2026-08-06T14:59:53.365170Z"},"links":{"citing_paper":"/paper/2507.17208"},"observation_digest":"sha256:64faa1f17e318ddfbd85c334e3a737665350708a77872ea92c960ee029898622","observation_id":"31749ecb-f8b6-4cc9-8fbd-2fb868272c4d","resolution":{"observed_at":"2026-08-06T14:59:53.915930Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-13T06:32:02.005865+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-13T06:32:02.005865+00:00","source":"crossref"},{"observed_at":"2026-08-13T06:31:53.387327+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T14:59:53.900438Z","title":"An overview of voice conversion and its challenges: From statistical modeling to deep learning,","venue":null,"work_id":"bc3145ec-c910-4a97-b3a6-a46189f35e32","year":2020},"citing_paper":{"arxiv_id":"2507.17208","last_updated":"2025-07-23T04:55:28Z","snapshot_observed_at":"2026-08-13T08:38:34.334066Z","submitted_at":"2025-07-23T04:55:28Z","title":"SLASH: Self-Supervised Speech Pitch Estimation Leveraging DSP-derived Absolute Pitch","version":1},"reference_index":7,"source":"pdf_text","source_observed_at":"2026-08-06T14:59:53.368977Z"},"links":{"citing_paper":"/paper/2507.17208"},"observation_digest":"sha256:2eff638a1a522866af27151153be64ad4c86769069b33964c9760d626cce30df","observation_id":"d5712d0c-c578-4dd4-9d7d-ab9487760f8a","resolution":{"observed_at":"2026-08-06T14:59:53.904766Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-13T06:32:02.005865+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-13T06:32:02.005865+00:00","source":"crossref"},{"observed_at":"2026-08-13T06:31:53.387327+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T14:59:53.888466Z","title":"PYIN: A fundamental frequency estimator using probabilistic threshold distributions,","venue":null,"work_id":"992a6f6b-ed7a-4966-9b3a-7bfa9d83b330","year":2014},"citing_paper":{"arxiv_id":"2507.17208","last_updated":"2025-07-23T04:55:28Z","snapshot_observed_at":"2026-08-13T08:38:34.334066Z","submitted_at":"2025-07-23T04:55:28Z","title":"SLASH: Self-Supervised Speech Pitch Estimation Leveraging DSP-derived Absolute Pitch","version":1},"reference_index":8,"source":"pdf_text","source_observed_at":"2026-08-06T14:59:53.373190Z"},"links":{"citing_paper":"/paper/2507.17208"},"observation_digest":"sha256:22470b7dfebd7dc2d5e2967f70616348c1ecbbe8c01100e629b3c6fe597d75f2","observation_id":"30d970ca-d21d-4527-b327-a87b1f4fb527","resolution":{"observed_at":"2026-08-06T14:59:53.893264Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-13T06:32:02.005865+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-13T06:32:02.005865+00:00","source":"crossref"},{"observed_at":"2026-08-13T06:31:53.387327+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T14:59:53.876961Z","title":"A sawtooth waveform inspired pitch estimator for speech and music,","venue":null,"work_id":"5a141d85-bebb-4db1-b8e5-408d736d6899","year":2008},"citing_paper":{"arxiv_id":"2507.17208","last_updated":"2025-07-23T04:55:28Z","snapshot_observed_at":"2026-08-13T08:38:34.334066Z","submitted_at":"2025-07-23T04:55:28Z","title":"SLASH: Self-Supervised Speech Pitch Estimation Leveraging DSP-derived Absolute Pitch","version":1},"reference_index":9,"source":"pdf_text","source_observed_at":"2026-08-06T14:59:53.377359Z"},"links":{"citing_paper":"/paper/2507.17208"},"observation_digest":"sha256:3b2dd2d867c6a32bd1bf2a013a3d875646602ad9f3314c0cd25bfd9a7cdab053","observation_id":"161e625f-55e5-4574-a0c0-545b3059ed36","resolution":{"observed_at":"2026-08-06T14:59:53.881585Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-13T06:32:02.005865+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-13T06:32:02.005865+00:00","source":"crossref"},{"observed_at":"2026-08-13T06:31:53.387327+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T14:59:53.865048Z","title":"A robust algorithm for pitch tracking (RAPT),","venue":null,"work_id":"13b8c058-4b95-4df8-9e35-890aa9f70a5b","year":1995},"citing_paper":{"arxiv_id":"2507.17208","last_updated":"2025-07-23T04:55:28Z","snapshot_observed_at":"2026-08-13T08:38:34.334066Z","submitted_at":"2025-07-23T04:55:28Z","title":"SLASH: Self-Supervised Speech Pitch Estimation Leveraging DSP-derived Absolute Pitch","version":1},"reference_index":10,"source":"pdf_text","source_observed_at":"2026-08-06T14:59:53.380682Z"},"links":{"citing_paper":"/paper/2507.17208"},"observation_digest":"sha256:835b182506525fbc351d832a1d10c639b684a96939caafc14099e4e73bd60f1a","observation_id":"a623cd81-325e-41a4-b19f-ea2be75783e5","resolution":{"observed_at":"2026-08-06T14:59:53.869194Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-13T06:32:02.005865+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-13T06:32:02.005865+00:00","source":"crossref"},{"observed_at":"2026-08-13T06:31:53.387327+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T14:59:53.853911Z","title":"PESTO: Pitch estimation with self-supervised transposition-equivariant objec- tive,","venue":null,"work_id":"c49bb5f3-69a7-48b8-8de6-952c0c757ae1","year":2023},"citing_paper":{"arxiv_id":"2507.17208","last_updated":"2025-07-23T04:55:28Z","snapshot_observed_at":"2026-08-13T08:38:34.334066Z","submitted_at":"2025-07-23T04:55:28Z","title":"SLASH: Self-Supervised Speech Pitch Estimation Leveraging DSP-derived Absolute Pitch","version":1},"reference_index":11,"source":"pdf_text","source_observed_at":"2026-08-06T14:59:53.384838Z"},"links":{"citing_paper":"/paper/2507.17208"},"observation_digest":"sha256:c2203122e3ec65cd34779c15888bffada06c4d7b7dc7ad803d39941804f44c9d","observation_id":"34a6407b-58d7-4bab-89f1-7dc329742893","resolution":{"observed_at":"2026-08-06T14:59:53.858147Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-13T06:32:02.005865+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-13T06:32:02.005865+00:00","source":"crossref"},{"observed_at":"2026-08-13T06:31:53.387327+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T14:59:53.842429Z","title":"SPICE: Self-supervised pitch estimation,","venue":null,"work_id":"851b098d-a294-4fba-adb3-8ccaa9856895","year":2020},"citing_paper":{"arxiv_id":"2507.17208","last_updated":"2025-07-23T04:55:28Z","snapshot_observed_at":"2026-08-13T08:38:34.334066Z","submitted_at":"2025-07-23T04:55:28Z","title":"SLASH: Self-Supervised Speech Pitch Estimation Leveraging DSP-derived Absolute Pitch","version":1},"reference_index":12,"source":"pdf_text","source_observed_at":"2026-08-06T14:59:53.388406Z"},"links":{"citing_paper":"/paper/2507.17208"},"observation_digest":"sha256:2775187c9bfbebe2e438090faade5e0c267cb7ff5ff14820237f2e469f09a7a3","observation_id":"01f4518f-a093-4662-84f1-f4c571122f28","resolution":{"observed_at":"2026-08-06T14:59:53.846610Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-13T06:32:02.005865+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-13T06:32:02.005865+00:00","source":"crossref"},{"observed_at":"2026-08-13T06:31:53.387327+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T14:59:53.832044Z","title":"Crepe: A convolu- tional representation for pitch estimation,","venue":null,"work_id":"7b9b74b2-096f-4089-b3cb-6c127d3c78ba","year":2018},"citing_paper":{"arxiv_id":"2507.17208","last_updated":"2025-07-23T04:55:28Z","snapshot_observed_at":"2026-08-13T08:38:34.334066Z","submitted_at":"2025-07-23T04:55:28Z","title":"SLASH: Self-Supervised Speech Pitch Estimation Leveraging DSP-derived Absolute Pitch","version":1},"reference_index":13,"source":"pdf_text","source_observed_at":"2026-08-06T14:59:53.391880Z"},"links":{"citing_paper":"/paper/2507.17208"},"observation_digest":"sha256:fd741220ffa7c8459b922adf3c600b7794985df87cf85378645c3102a0d23448","observation_id":"f9b7e45e-f122-4d58-8f57-1d0c5ff4cc65","resolution":{"observed_at":"2026-08-06T14:59:53.835702Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-13T06:32:02.005865+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-13T06:32:02.005865+00:00","source":"crossref"},{"observed_at":"2026-08-13T06:31:53.387327+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T14:59:53.821699Z","title":"DeepF0: End-to-end fundamen- tal frequency estimation for music and speech signals,","venue":null,"work_id":"20bc0609-a829-4f82-a250-b391d8590ded","year":2021},"citing_paper":{"arxiv_id":"2507.17208","last_updated":"2025-07-23T04:55:28Z","snapshot_observed_at":"2026-08-13T08:38:34.334066Z","submitted_at":"2025-07-23T04:55:28Z","title":"SLASH: Self-Supervised Speech Pitch Estimation Leveraging DSP-derived Absolute Pitch","version":1},"reference_index":14,"source":"pdf_text","source_observed_at":"2026-08-06T14:59:53.395560Z"},"links":{"citing_paper":"/paper/2507.17208"},"observation_digest":"sha256:b63ae7cbf04992a13db88bf0d10b531f4ad88dd163de9c66897b25e1b8540e66","observation_id":"38625b17-a927-43b2-afee-3205498d9c63","resolution":{"observed_at":"2026-08-06T14:59:53.825369Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-13T06:32:02.005865+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-13T06:32:02.005865+00:00","source":"crossref"},{"observed_at":"2026-08-13T06:31:53.387327+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T14:59:53.810151Z","title":"Noise-robust DSP-assisted neural pitch estimation with very low complexity,","venue":null,"work_id":"0fada1dd-3fe3-4878-9874-b67fc9bdb100","year":2024},"citing_paper":{"arxiv_id":"2507.17208","last_updated":"2025-07-23T04:55:28Z","snapshot_observed_at":"2026-08-13T08:38:34.334066Z","submitted_at":"2025-07-23T04:55:28Z","title":"SLASH: Self-Supervised Speech Pitch Estimation Leveraging DSP-derived Absolute Pitch","version":1},"reference_index":15,"source":"pdf_text","source_observed_at":"2026-08-06T14:59:53.400094Z"},"links":{"citing_paper":"/paper/2507.17208"},"observation_digest":"sha256:b143ca51eb51ded9dbdb554ca8bae0946fe2af43aeaaf1351ac4b1165a13b6a6","observation_id":"089d3406-ee7f-406d-af5a-62ea0deb3f46","resolution":{"observed_at":"2026-08-06T14:59:53.814824Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-13T06:32:02.005865+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-13T06:32:02.005865+00:00","source":"crossref"},{"observed_at":"2026-08-13T06:31:53.387327+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2211.09407","last_updated":"2022-11-17T08:29:57Z","snapshot_observed_at":"2026-08-13T00:32:24.100018Z","submitted_at":"2022-11-17T08:29:57Z","title":"NANSY++: Unified Voice Synthesis with Neural Analysis and Synthesis","version":1},"cited_work":{"arxiv_id":"2211.09407","doi":null,"metadata_source":"pith","pith_arxiv_id":"2211.09407","snapshot_observed_at":"2026-08-06T14:59:53.556325Z","title":"NANSY++: Unified Voice Synthesis with Neural Analysis and Synthesis","venue":"cs.SD","work_id":"5a25fe9f-32d1-4d40-83d0-1c3539621cfc","year":2022},"citing_paper":{"arxiv_id":"2507.17208","last_updated":"2025-07-23T04:55:28Z","snapshot_observed_at":"2026-08-13T08:38:34.334066Z","submitted_at":"2025-07-23T04:55:28Z","title":"SLASH: Self-Supervised Speech Pitch Estimation Leveraging DSP-derived Absolute Pitch","version":1},"reference_index":16,"source":"pdf_text","source_observed_at":"2026-08-06T14:59:53.403572Z"},"links":{"cited_paper":"/paper/2211.09407","citing_paper":"/paper/2507.17208"},"observation_digest":"sha256:1945813d7861aa526ef62e2fe33ec9b3502c5b94aadbf3169e004d7d0d97a7e2","observation_id":"cdf6c810-7c59-415b-ab9e-ddd0f8336852","resolution":{"observed_at":"2026-08-06T14:59:53.560739Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-13T06:32:02.005865+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-13T06:32:02.005865+00:00","source":"crossref"},{"observed_at":"2026-08-13T06:31:53.387327+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T14:59:53.798840Z","title":"Singing voice separa- tion and vocal F0 estimation based on mutual combination of ro- bust principal component analysis and subharmonic summation,","venue":null,"work_id":"e6f2d341-c6e0-4857-a7b7-a9165bec4312","year":2016},"citing_paper":{"arxiv_id":"2507.17208","last_updated":"2025-07-23T04:55:28Z","snapshot_observed_at":"2026-08-13T08:38:34.334066Z","submitted_at":"2025-07-23T04:55:28Z","title":"SLASH: Self-Supervised Speech Pitch Estimation Leveraging DSP-derived Absolute Pitch","version":1},"reference_index":17,"source":"pdf_text","source_observed_at":"2026-08-06T14:59:53.407328Z"},"links":{"citing_paper":"/paper/2507.17208"},"observation_digest":"sha256:f57baf7fa6cb137f51a1859066749e9d1bdb3aadf9731f141bc2000cf5742d5f","observation_id":"193cd4ad-c9b8-45ee-8ddf-66446c1f1741","resolution":{"observed_at":"2026-08-06T14:59:53.802898Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-13T06:32:02.005865+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-13T06:32:02.005865+00:00","source":"crossref"},{"observed_at":"2026-08-13T06:31:53.387327+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T14:59:53.788698Z","title":"Unsupervised harmonic parameter estimation using differentiable DSP and spectral opti- mal transport,","venue":null,"work_id":"0155c5e5-0e9b-46e4-94f8-50a4e0436582","year":2024},"citing_paper":{"arxiv_id":"2507.17208","last_updated":"2025-07-23T04:55:28Z","snapshot_observed_at":"2026-08-13T08:38:34.334066Z","submitted_at":"2025-07-23T04:55:28Z","title":"SLASH: Self-Supervised Speech Pitch Estimation Leveraging DSP-derived Absolute Pitch","version":1},"reference_index":18,"source":"pdf_text","source_observed_at":"2026-08-06T14:59:53.411442Z"},"links":{"citing_paper":"/paper/2507.17208"},"observation_digest":"sha256:b8a10b8faebab310375c72119c36935daa9996f66730dcdcbf1f4989410af6b6","observation_id":"50fbb57d-84cb-41cf-b93c-8f534433d212","resolution":{"observed_at":"2026-08-06T14:59:53.792416Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-13T06:32:02.005865+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-13T06:32:02.005865+00:00","source":"crossref"},{"observed_at":"2026-08-13T06:31:53.387327+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2001.04643","last_updated":"2020-01-14T06:49:37Z","snapshot_observed_at":"2026-08-12T04:40:34.468424Z","submitted_at":"2020-01-14T06:49:37Z","title":"DDSP: Differentiable Digital Signal Processing","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2001.04643","snapshot_observed_at":"2026-08-06T14:59:53.415326Z","title":"DDSP: Differentiable digital signal processing,","venue":null,"work_id":null,"year":2001},"citing_paper":{"arxiv_id":"2507.17208","last_updated":"2025-07-23T04:55:28Z","snapshot_observed_at":"2026-08-13T08:38:34.334066Z","submitted_at":"2025-07-23T04:55:28Z","title":"SLASH: Self-Supervised Speech Pitch Estimation Leveraging DSP-derived Absolute Pitch","version":1},"reference_index":19,"source":"pdf_text","source_observed_at":"2026-08-06T14:59:53.415326Z"},"links":{"cited_paper":"/paper/2001.04643","citing_paper":"/paper/2507.17208"},"observation_digest":"sha256:20685e01ac9048c931b30ea1a4425a012431346a6c9e072dc7d636c0305db38d","observation_id":"59b00769-c895-4bf5-a692-58d5202f61b8","resolution":{"observed_at":"2026-08-06T14:59:53.415326Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T14:59:53.777197Z","title":"Technical foundations of tandem- straight, a speech analysis, modification and synthesis frame- work,","venue":null,"work_id":"2ddbe840-0a35-4f93-b8f7-228f4f595e36","year":2011},"citing_paper":{"arxiv_id":"2507.17208","last_updated":"2025-07-23T04:55:28Z","snapshot_observed_at":"2026-08-13T08:38:34.334066Z","submitted_at":"2025-07-23T04:55:28Z","title":"SLASH: Self-Supervised Speech Pitch Estimation Leveraging DSP-derived Absolute Pitch","version":1},"reference_index":20,"source":"pdf_text","source_observed_at":"2026-08-06T14:59:53.419347Z"},"links":{"citing_paper":"/paper/2507.17208"},"observation_digest":"sha256:9eb18e52af0b32665b8bfc08731d888a6fe8924cb48f21ab1c636c465f1f2fa0","observation_id":"e98f068b-29ce-4a61-8492-5ffad18ea928","resolution":{"observed_at":"2026-08-06T14:59:53.781101Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-13T06:32:02.005865+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-13T06:32:02.005865+00:00","source":"crossref"},{"observed_at":"2026-08-13T06:31:53.387327+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T14:59:53.767488Z","title":"WORLD: A vocoder- based high-quality speech synthesis system for real-time applica- tions,","venue":null,"work_id":"da6284c2-c7b5-477d-b381-056935007779","year":2016},"citing_paper":{"arxiv_id":"2507.17208","last_updated":"2025-07-23T04:55:28Z","snapshot_observed_at":"2026-08-13T08:38:34.334066Z","submitted_at":"2025-07-23T04:55:28Z","title":"SLASH: Self-Supervised Speech Pitch Estimation Leveraging DSP-derived Absolute Pitch","version":1},"reference_index":21,"source":"pdf_text","source_observed_at":"2026-08-06T14:59:53.423455Z"},"links":{"citing_paper":"/paper/2507.17208"},"observation_digest":"sha256:7febb22c645bc8840e0ff087ffd32c10c1f1bf66eb9d5c2e21aeee79a10b98a6","observation_id":"7a51550f-7c5e-40e4-a1c4-715c2c9e511c","resolution":{"observed_at":"2026-08-06T14:59:53.770781Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-13T06:32:02.005865+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-13T06:32:02.005865+00:00","source":"crossref"},{"observed_at":"2026-08-13T06:31:53.387327+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T14:59:53.756931Z","title":"A spectral envelope estimation method based on F0-adaptive multi-frame integration analysis,","venue":null,"work_id":"81365add-6c7a-4aaa-b2b1-414b6e134f00","year":2012},"citing_paper":{"arxiv_id":"2507.17208","last_updated":"2025-07-23T04:55:28Z","snapshot_observed_at":"2026-08-13T08:38:34.334066Z","submitted_at":"2025-07-23T04:55:28Z","title":"SLASH: Self-Supervised Speech Pitch Estimation Leveraging DSP-derived Absolute Pitch","version":1},"reference_index":22,"source":"pdf_text","source_observed_at":"2026-08-06T14:59:53.428007Z"},"links":{"citing_paper":"/paper/2507.17208"},"observation_digest":"sha256:71b9f3acc35ce11c94f842fa46701df062f2dce9a4bc9f534b6b23ef0e7aaafe","observation_id":"beac2734-bf82-41da-914b-31c8840608b5","resolution":{"observed_at":"2026-08-06T14:59:53.760809Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-13T06:32:02.005865+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-13T06:32:02.005865+00:00","source":"crossref"},{"observed_at":"2026-08-13T06:31:53.387327+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2208.07282","last_updated":"2023-05-08T13:45:05Z","snapshot_observed_at":"2026-08-02T18:02:32.792463Z","submitted_at":"2022-08-15T15:48:36Z","title":"Differentiable WORLD Synthesizer-based Neural Vocoder With Application To End-To-End Audio Style Transfer","version":5},"cited_work":{"arxiv_id":"2208.07282","doi":null,"metadata_source":"pith","pith_arxiv_id":"2208.07282","snapshot_observed_at":"2026-08-06T14:59:53.529368Z","title":"Differentiable WORLD Synthesizer-based Neural Vocoder With Application To End-To-End Audio Style Transfer","venue":"eess.AS","work_id":"1a15d6dd-1f9a-4e7a-8537-5d9c199b9d50","year":2022},"citing_paper":{"arxiv_id":"2507.17208","last_updated":"2025-07-23T04:55:28Z","snapshot_observed_at":"2026-08-13T08:38:34.334066Z","submitted_at":"2025-07-23T04:55:28Z","title":"SLASH: Self-Supervised Speech Pitch Estimation Leveraging DSP-derived Absolute Pitch","version":1},"reference_index":23,"source":"pdf_text","source_observed_at":"2026-08-06T14:59:53.431793Z"},"links":{"cited_paper":"/paper/2208.07282","citing_paper":"/paper/2507.17208"},"observation_digest":"sha256:be2af7a20fabe34db293234077e9e72dad339546ed1019709ebb700282b297d8","observation_id":"c91b2f6b-5bec-4449-a631-c908ac175f8e","resolution":{"observed_at":"2026-08-06T14:59:53.535942Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-13T06:32:02.005865+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-13T06:32:02.005865+00:00","source":"crossref"},{"observed_at":"2026-08-13T06:31:53.387327+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T14:59:53.744589Z","title":"Calculation of a constant Q spectral transform,","venue":null,"work_id":"af601a2e-2176-43a8-932f-ad7afc7065eb","year":1991},"citing_paper":{"arxiv_id":"2507.17208","last_updated":"2025-07-23T04:55:28Z","snapshot_observed_at":"2026-08-13T08:38:34.334066Z","submitted_at":"2025-07-23T04:55:28Z","title":"SLASH: Self-Supervised Speech Pitch Estimation Leveraging DSP-derived Absolute Pitch","version":1},"reference_index":24,"source":"pdf_text","source_observed_at":"2026-08-06T14:59:53.436314Z"},"links":{"citing_paper":"/paper/2507.17208"},"observation_digest":"sha256:b1847e1933e684e6a3cd8ceb2b9298907aeb98ccd76eb5523b4ca17d279c395d","observation_id":"759dd120-e030-4022-a8f1-d2911e2ba727","resolution":{"observed_at":"2026-08-06T14:59:53.748476Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-13T06:32:02.005865+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-13T06:32:02.005865+00:00","source":"crossref"},{"observed_at":"2026-08-13T06:31:53.387327+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T14:59:53.733477Z","title":"Robust estimation of a location parameter,","venue":null,"work_id":"e5dc1766-8c02-4201-a486-4b802d64ddfc","year":1964},"citing_paper":{"arxiv_id":"2507.17208","last_updated":"2025-07-23T04:55:28Z","snapshot_observed_at":"2026-08-13T08:38:34.334066Z","submitted_at":"2025-07-23T04:55:28Z","title":"SLASH: Self-Supervised Speech Pitch Estimation Leveraging DSP-derived Absolute Pitch","version":1},"reference_index":25,"source":"pdf_text","source_observed_at":"2026-08-06T14:59:53.441294Z"},"links":{"citing_paper":"/paper/2507.17208"},"observation_digest":"sha256:509a6be997256a7540d581640511c5c2700a2e014eb93124aa3541d3e1b5f9a9","observation_id":"e3645cea-de11-4d8f-b690-c84aeede71fd","resolution":{"observed_at":"2026-08-06T14:59:53.737799Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-13T06:32:02.005865+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-13T06:32:02.005865+00:00","source":"crossref"},{"observed_at":"2026-08-13T06:31:53.387327+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T14:59:53.721051Z","title":"Spectral smoothing technique in PARCOR speech analysis-synthesis,","venue":null,"work_id":"bf8a5b0e-dfad-4be8-b120-e8ca163ac9d1","year":1978},"citing_paper":{"arxiv_id":"2507.17208","last_updated":"2025-07-23T04:55:28Z","snapshot_observed_at":"2026-08-13T08:38:34.334066Z","submitted_at":"2025-07-23T04:55:28Z","title":"SLASH: Self-Supervised Speech Pitch Estimation Leveraging DSP-derived Absolute Pitch","version":1},"reference_index":26,"source":"pdf_text","source_observed_at":"2026-08-06T14:59:53.444820Z"},"links":{"citing_paper":"/paper/2507.17208"},"observation_digest":"sha256:bc9568ca6cf8df423cc8b3f9fd6a0bdff74018ab479dd40f0e4e9855afb04588","observation_id":"283b9773-82ae-4bab-8818-dc977acb93fa","resolution":{"observed_at":"2026-08-06T14:59:53.725470Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-13T06:32:02.005865+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-13T06:32:02.005865+00:00","source":"crossref"},{"observed_at":"2026-08-13T06:31:53.387327+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T14:59:53.710958Z","title":"Learning to learn by gradient descent by gradient descent,","venue":null,"work_id":"a7701b6c-6630-4d46-b1cf-5c6c87c5e42b","year":2016},"citing_paper":{"arxiv_id":"2507.17208","last_updated":"2025-07-23T04:55:28Z","snapshot_observed_at":"2026-08-13T08:38:34.334066Z","submitted_at":"2025-07-23T04:55:28Z","title":"SLASH: Self-Supervised Speech Pitch Estimation Leveraging DSP-derived Absolute Pitch","version":1},"reference_index":27,"source":"pdf_text","source_observed_at":"2026-08-06T14:59:53.448371Z"},"links":{"citing_paper":"/paper/2507.17208"},"observation_digest":"sha256:f89a657a88fbeae3c0cc34e95ec77641514e398e15cc7fa8fd06003f1940cd8b","observation_id":"15a6d693-2c07-4c86-b055-b3414fcc257d","resolution":{"observed_at":"2026-08-06T14:59:53.715232Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-13T06:32:02.005865+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-13T06:32:02.005865+00:00","source":"crossref"},{"observed_at":"2026-08-13T06:31:53.387327+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T14:59:53.700460Z","title":"A spectral energy distance for parallel speech synthesis,","venue":null,"work_id":"90637fe3-3411-4e77-abce-a91d47a487c2","year":2020},"citing_paper":{"arxiv_id":"2507.17208","last_updated":"2025-07-23T04:55:28Z","snapshot_observed_at":"2026-08-13T08:38:34.334066Z","submitted_at":"2025-07-23T04:55:28Z","title":"SLASH: Self-Supervised Speech Pitch Estimation Leveraging DSP-derived Absolute Pitch","version":1},"reference_index":28,"source":"pdf_text","source_observed_at":"2026-08-06T14:59:53.451733Z"},"links":{"citing_paper":"/paper/2507.17208"},"observation_digest":"sha256:6d3cccf7183479dd4ea277132f1dd65d0e7c69e932b2dfb6524cc56565506c0f","observation_id":"4eb387b1-fdcd-4ebf-aa7b-eb7a1cbe1ba1","resolution":{"observed_at":"2026-08-06T14:59:53.704764Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-13T06:32:02.005865+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-13T06:32:02.005865+00:00","source":"crossref"},{"observed_at":"2026-08-13T06:31:53.387327+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T14:59:53.690582Z","title":"LibriTTS-R: A restored multi-speaker text-to-speech corpus,","venue":null,"work_id":"dbac2cd5-8557-4778-a6c9-7335f11a35f2","year":2023},"citing_paper":{"arxiv_id":"2507.17208","last_updated":"2025-07-23T04:55:28Z","snapshot_observed_at":"2026-08-13T08:38:34.334066Z","submitted_at":"2025-07-23T04:55:28Z","title":"SLASH: Self-Supervised Speech Pitch Estimation Leveraging DSP-derived Absolute Pitch","version":1},"reference_index":29,"source":"pdf_text","source_observed_at":"2026-08-06T14:59:53.455624Z"},"links":{"citing_paper":"/paper/2507.17208"},"observation_digest":"sha256:78af564ac0dc82528ed755b4647d99d5dd4c0a5ad32ce90b5c2d24ce08d720c5","observation_id":"69281a0f-f797-4506-a81f-5913d60816bb","resolution":{"observed_at":"2026-08-06T14:59:53.694773Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-13T06:32:02.005865+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-13T06:32:02.005865+00:00","source":"crossref"},{"observed_at":"2026-08-13T06:31:53.387327+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T14:59:53.681206Z","title":"On the improvement of singing voice separation for monaural recordings using the MIR-1K dataset,","venue":null,"work_id":"52a39441-bb49-42ea-a9cc-4c97c1d15422","year":2009},"citing_paper":{"arxiv_id":"2507.17208","last_updated":"2025-07-23T04:55:28Z","snapshot_observed_at":"2026-08-13T08:38:34.334066Z","submitted_at":"2025-07-23T04:55:28Z","title":"SLASH: Self-Supervised Speech Pitch Estimation Leveraging DSP-derived Absolute Pitch","version":1},"reference_index":30,"source":"pdf_text","source_observed_at":"2026-08-06T14:59:53.460278Z"},"links":{"citing_paper":"/paper/2507.17208"},"observation_digest":"sha256:b0d153d5e9ee776b2a27459608a64ff6f6c0a91b7e6cd84722434ddd45fa89ad","observation_id":"78b63fb6-8ae4-4dae-a47c-c64f5978276e","resolution":{"observed_at":"2026-08-06T14:59:53.684669Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-13T06:32:02.005865+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-13T06:32:02.005865+00:00","source":"crossref"},{"observed_at":"2026-08-13T06:31:53.387327+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T14:59:53.671583Z","title":"ESPnet-TTS: Uni- fied, reproducible, and integratable open source end-to-end text- to-speech toolkit,","venue":null,"work_id":"96e5d04e-f4dd-47a4-93ad-ad2706bd8363","year":2020},"citing_paper":{"arxiv_id":"2507.17208","last_updated":"2025-07-23T04:55:28Z","snapshot_observed_at":"2026-08-13T08:38:34.334066Z","submitted_at":"2025-07-23T04:55:28Z","title":"SLASH: Self-Supervised Speech Pitch Estimation Leveraging DSP-derived Absolute Pitch","version":1},"reference_index":31,"source":"pdf_text","source_observed_at":"2026-08-06T14:59:53.464870Z"},"links":{"citing_paper":"/paper/2507.17208"},"observation_digest":"sha256:fd863b4846908617cfae976ea1a881de93c6d7461ecfd2308c0c19c40af225bd","observation_id":"d69b9152-61c8-4b7c-9688-33bf9aa99174","resolution":{"observed_at":"2026-08-06T14:59:53.675128Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-13T06:32:02.005865+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-13T06:32:02.005865+00:00","source":"crossref"},{"observed_at":"2026-08-13T06:31:53.387327+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T14:59:53.661634Z","title":"Decoupled weight decay regulariza- tion,","venue":null,"work_id":"eca28dd4-193f-45e5-be12-83e4ae3f8092","year":2019},"citing_paper":{"arxiv_id":"2507.17208","last_updated":"2025-07-23T04:55:28Z","snapshot_observed_at":"2026-08-13T08:38:34.334066Z","submitted_at":"2025-07-23T04:55:28Z","title":"SLASH: Self-Supervised Speech Pitch Estimation Leveraging DSP-derived Absolute Pitch","version":1},"reference_index":32,"source":"pdf_text","source_observed_at":"2026-08-06T14:59:53.469735Z"},"links":{"citing_paper":"/paper/2507.17208"},"observation_digest":"sha256:092dc1d7b02b13313d18c9a15b54cae795e11d7c163825cad95d53be975e7a4c","observation_id":"f934d510-c05d-4a77-96f8-ddee088ea706","resolution":{"observed_at":"2026-08-06T14:59:53.665163Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-13T06:32:02.005865+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-13T06:32:02.005865+00:00","source":"crossref"},{"observed_at":"2026-08-13T06:31:53.387327+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T14:59:53.652150Z","title":"Fast and reliable F0 estimation method based on the period extraction of vocal fold vibration of singing voice and speech,","venue":null,"work_id":"66bbaa64-7a06-479a-9b22-42ddc4f7fc0f","year":2009},"citing_paper":{"arxiv_id":"2507.17208","last_updated":"2025-07-23T04:55:28Z","snapshot_observed_at":"2026-08-13T08:38:34.334066Z","submitted_at":"2025-07-23T04:55:28Z","title":"SLASH: Self-Supervised Speech Pitch Estimation Leveraging DSP-derived Absolute Pitch","version":1},"reference_index":33,"source":"pdf_text","source_observed_at":"2026-08-06T14:59:53.473562Z"},"links":{"citing_paper":"/paper/2507.17208"},"observation_digest":"sha256:2d1991f1d015e00ac2dd0a25a8985eb862791bb8a849542febbd0b8d0636c706","observation_id":"1312bd19-4574-4fd7-9daf-d4fb8ad39220","resolution":{"observed_at":"2026-08-06T14:59:53.655731Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-13T06:32:02.005865+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-13T06:32:02.005865+00:00","source":"crossref"},{"observed_at":"2026-08-13T06:31:53.387327+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T14:59:53.642825Z","title":"Harvest: A high-performance fundamental frequency estimator from speech signals,","venue":null,"work_id":"645a5d7d-aac2-4c9f-b265-ec9d7cbf9d8f","year":2017},"citing_paper":{"arxiv_id":"2507.17208","last_updated":"2025-07-23T04:55:28Z","snapshot_observed_at":"2026-08-13T08:38:34.334066Z","submitted_at":"2025-07-23T04:55:28Z","title":"SLASH: Self-Supervised Speech Pitch Estimation Leveraging DSP-derived Absolute Pitch","version":1},"reference_index":34,"source":"pdf_text","source_observed_at":"2026-08-06T14:59:53.477496Z"},"links":{"citing_paper":"/paper/2507.17208"},"observation_digest":"sha256:2bcf155f7c482aa98d42204a2f2ef4c21be5e924afc6e02852970b6ea41e5932","observation_id":"8d711df3-e4ec-4798-8aea-95d783e9333d","resolution":{"observed_at":"2026-08-06T14:59:53.646507Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-13T06:32:02.005865+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-13T06:32:02.005865+00:00","source":"crossref"},{"observed_at":"2026-08-13T06:31:53.387327+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T14:59:53.632033Z","title":"Multiple fundamental fre- quency estimation by modeling spectral peaks and non-peak re- gions,","venue":null,"work_id":"6f4ea88a-162f-409a-ab95-af9e2840b125","year":2010},"citing_paper":{"arxiv_id":"2507.17208","last_updated":"2025-07-23T04:55:28Z","snapshot_observed_at":"2026-08-13T08:38:34.334066Z","submitted_at":"2025-07-23T04:55:28Z","title":"SLASH: Self-Supervised Speech Pitch Estimation Leveraging DSP-derived Absolute Pitch","version":1},"reference_index":35,"source":"pdf_text","source_observed_at":"2026-08-06T14:59:53.481576Z"},"links":{"citing_paper":"/paper/2507.17208"},"observation_digest":"sha256:6e26890a312c1f9669dffc8af620ef32cb102fd51d1ccf341d7aad11eeabb5a9","observation_id":"af257b4e-b97d-4928-97be-372c9c0fe827","resolution":{"observed_at":"2026-08-06T14:59:53.635556Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-13T06:32:02.005865+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-13T06:32:02.005865+00:00","source":"crossref"},{"observed_at":"2026-08-13T06:31:53.387327+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T14:59:53.621327Z","title":"MedleyDB: A multitrack dataset for annotation- intensive MIR research","venue":null,"work_id":"1b8acd38-63a6-44ee-880a-9b1f126a6c85","year":2014},"citing_paper":{"arxiv_id":"2507.17208","last_updated":"2025-07-23T04:55:28Z","snapshot_observed_at":"2026-08-13T08:38:34.334066Z","submitted_at":"2025-07-23T04:55:28Z","title":"SLASH: Self-Supervised Speech Pitch Estimation Leveraging DSP-derived Absolute Pitch","version":1},"reference_index":36,"source":"pdf_text","source_observed_at":"2026-08-06T14:59:53.485441Z"},"links":{"citing_paper":"/paper/2507.17208"},"observation_digest":"sha256:97b95cea4511bf85b2186ad1ab32533f3537d489fc109aa3c287edffe7d3761d","observation_id":"4a1de76e-c04a-4485-abf5-98f53fdb42a3","resolution":{"observed_at":"2026-08-06T14:59:53.625833Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-13T06:32:02.005865+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-13T06:32:02.005865+00:00","source":"crossref"},{"observed_at":"2026-08-13T06:31:53.387327+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T14:59:53.610062Z","title":"An analysis/synthesis framework for automatic F0 annotation of multitrack datasets,","venue":null,"work_id":"fafabc49-095f-47d2-bf43-c2facee5c7ce","year":2017},"citing_paper":{"arxiv_id":"2507.17208","last_updated":"2025-07-23T04:55:28Z","snapshot_observed_at":"2026-08-13T08:38:34.334066Z","submitted_at":"2025-07-23T04:55:28Z","title":"SLASH: Self-Supervised Speech Pitch Estimation Leveraging DSP-derived Absolute Pitch","version":1},"reference_index":37,"source":"pdf_text","source_observed_at":"2026-08-06T14:59:53.489549Z"},"links":{"citing_paper":"/paper/2507.17208"},"observation_digest":"sha256:b4948f2cbccdc6be8d581d00f72dac904c1b2acd10af065a05d6694261578bb7","observation_id":"597ffbe7-1f90-46a6-a1cb-184ba3471167","resolution":{"observed_at":"2026-08-06T14:59:53.613854Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-13T06:32:02.005865+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-13T06:32:02.005865+00:00","source":"crossref"},{"observed_at":"2026-08-13T06:31:53.387327+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T14:59:53.599803Z","title":"Neural audio synthesis of musi- cal notes with wavenet autoencoders,","venue":null,"work_id":"0af480d3-1b60-46e3-8bf4-1d51e681bbed","year":2017},"citing_paper":{"arxiv_id":"2507.17208","last_updated":"2025-07-23T04:55:28Z","snapshot_observed_at":"2026-08-13T08:38:34.334066Z","submitted_at":"2025-07-23T04:55:28Z","title":"SLASH: Self-Supervised Speech Pitch Estimation Leveraging DSP-derived Absolute Pitch","version":1},"reference_index":38,"source":"pdf_text","source_observed_at":"2026-08-06T14:59:53.494244Z"},"links":{"citing_paper":"/paper/2507.17208"},"observation_digest":"sha256:64c9037103b20628b495db3594d690b50cc1afc2fbed9bcd9f1c8aab007e0571","observation_id":"2f2ca08e-d900-46f3-a801-4e72f78c26a6","resolution":{"observed_at":"2026-08-06T14:59:53.603344Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-13T06:32:02.005865+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-13T06:32:02.005865+00:00","source":"crossref"},{"observed_at":"2026-08-13T06:31:53.387327+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T14:59:53.589408Z","title":"Melody extraction from polyphonic music signals: Approaches, applica- tions, and challenges,","venue":null,"work_id":"c69ed8ea-72a2-4899-ac63-4366448e755d","year":2014},"citing_paper":{"arxiv_id":"2507.17208","last_updated":"2025-07-23T04:55:28Z","snapshot_observed_at":"2026-08-13T08:38:34.334066Z","submitted_at":"2025-07-23T04:55:28Z","title":"SLASH: Self-Supervised Speech Pitch Estimation Leveraging DSP-derived Absolute Pitch","version":1},"reference_index":39,"source":"pdf_text","source_observed_at":"2026-08-06T14:59:53.498359Z"},"links":{"citing_paper":"/paper/2507.17208"},"observation_digest":"sha256:4888c19cd0dd7e9a10539bbf4bef368a364e25e6e4585b70ab149bbee09df6c4","observation_id":"cf7fdb41-2f9f-4183-8f90-7ad8f384b377","resolution":{"observed_at":"2026-08-06T14:59:53.593284Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-13T06:32:02.005865+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-13T06:32:02.005865+00:00","source":"crossref"},{"observed_at":"2026-08-13T06:31:53.387327+00:00","source":"retraction_watch"}],"state":"measured"}}],"paper":{"arxiv_id":"2507.17208","last_updated":"2025-07-23T04:55:28Z","latest_version":1,"primary_category":"eess.AS","snapshot_observed_at":"2026-08-13T08:38:34.334066Z","submitted_at":"2025-07-23T04:55:28Z","title":"SLASH: Self-Supervised Speech Pitch Estimation Leveraging DSP-derived Absolute Pitch"},"reference_resolution":{"displayed":39,"state_counts":{"malformed_identifier":1,"metadata_mismatch":1,"parse_uncertain":0,"unresolved":2,"verified_exact":2,"verified_fuzzy":33},"total_outbound_references":39},"refusal":"A citation records a reference. It does not transfer a finding from one paper to another.","schema":"pith.paper-citation-record.v1","standing_sources":[{"observed_at":"2026-08-13T06:32:02.005865+00:00","source":"crossref"},{"observed_at":"2026-08-13T06:31:53.387327+00:00","source":"retraction_watch"}],"thesis":"As of 13 August 2026, this Paper Citation Record lists 39 of 39 outbound references and 1 inbound Pith citation observation for arXiv:2507.17208."}