{"as_of":"2026-08-18T08:43:00Z","caps":{"database_statements":6,"inbound":100,"outbound":100},"context_digest":"sha256:c1bbb243a3c42834bc3c09f44e2460fc7e24f01a91f36580107330b0cb9bc308","coverage":[{"denominator":46,"lane":"reference_resolution","note":"Typed states for the displayed outbound observations.","records_observed":46,"source":"paper_references, paper_reference_links","source_observed_at":"2026-08-07T11:32:17.610645Z","state":"measured"},{"denominator":47,"lane":"standing_notices","note":"One-hop event checks from named stored sources.","records_observed":47,"source":"scholarly_work_events, retraction_status_cache","source_observed_at":"2026-08-18T06:34:40.430872+00:00","state":"measured"},{"denominator":1,"lane":"inbound_itemization","note":"Pith citing papers itemized under the disclosed page cap.","records_observed":1,"source":"paper_references, paper_reference_links","source_observed_at":"2026-08-07T11:32:13.765343Z","state":"measured"},{"denominator":1,"lane":"external_citation_measurements","note":"A source-named dated measurement, never combined with another source.","records_observed":0,"source":"pith","source_observed_at":"2026-08-07T11:32:17.963370Z","state":"measured"}],"external_citation_measurements":[],"inbound":[{"citation":{"cited_paper":{"arxiv_id":"2506.02181","last_updated":"2025-06-02T19:11:16Z","snapshot_observed_at":"2026-08-08T05:50:50.489478Z","submitted_at":"2025-06-02T19:11:16Z","title":"Echoes of Phonetics: Unveiling Relevant Acoustic Cues for ASR via Feature Attribution","version":1},"cited_work":{"arxiv_id":"2506.02181","doi":null,"metadata_source":"pith","pith_arxiv_id":"2506.02181","snapshot_observed_at":"2026-08-07T11:32:17.963370Z","title":"Echoes of Phonetics: Unveiling Relevant Acoustic Cues for ASR via Feature Attribution","venue":"cs.CL","work_id":"3e06c9d4-24cb-4dce-aa67-06b109b9e287","year":2025},"citing_paper":{"arxiv_id":"2506.02181","last_updated":"2025-06-02T19:11:16Z","snapshot_observed_at":"2026-08-08T05:50:50.489478Z","submitted_at":"2025-06-02T19:11:16Z","title":"Echoes of Phonetics: Unveiling Relevant Acoustic Cues for ASR via Feature Attribution","version":1},"reference_index":2,"source":"pdf_text","source_observed_at":"2026-08-07T11:32:13.765343Z"},"links":{"cited_paper":"/paper/2506.02181","citing_paper":"/paper/2506.02181"},"observation_digest":"sha256:68040d3ca1988724e85727e1c6e3ac07ca434b0f213aef2855705fa3bf78ac21","observation_id":"0b075e5a-482f-47e0-ab8a-9b01b6920469","resolution":{"observed_at":"2026-08-07T11:32:18.040121Z","resolver_source":"local_arxiv","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-18T06:34:40.430872+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-18T06:34:40.430872+00:00","source":"crossref"},{"observed_at":"2026-08-18T06:34:34.496301+00:00","source":"retraction_watch"}],"state":"measured"}}],"links":{"evidence":"/evidence","html":"/paper/2506.02181/citation-record","integrity":"/paper/2506.02181/integrity","json":"/paper/2506.02181/citation-record.json","paper":"/paper/2506.02181"},"outbound":[{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T11:32:25.602331Z","title":null,"venue":null,"work_id":"cae9b231-54ee-4ae3-ae62-4e677f9b4f73","year":null},"citing_paper":{"arxiv_id":"2506.02181","last_updated":"2025-06-02T19:11:16Z","snapshot_observed_at":"2026-08-08T05:50:50.489478Z","submitted_at":"2025-06-02T19:11:16Z","title":"Echoes of Phonetics: Unveiling Relevant Acoustic Cues for ASR via Feature Attribution","version":1},"reference_index":1,"source":"pdf_text","source_observed_at":"2026-08-07T11:32:13.696269Z"},"links":{"citing_paper":"/paper/2506.02181"},"observation_digest":"sha256:04642064e9a392abcdf5bd6a01dd8a449c450fbd60388e3b7817fa64c02d5fc9","observation_id":"32be8668-a1d0-4b36-9bae-41ad920028b8","resolution":{"observed_at":"2026-08-07T11:32:25.650868Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-18T06:34:40.430872+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-18T06:34:40.430872+00:00","source":"crossref"},{"observed_at":"2026-08-18T06:34:34.496301+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2506.02181","last_updated":"2025-06-02T19:11:16Z","snapshot_observed_at":"2026-08-08T05:50:50.489478Z","submitted_at":"2025-06-02T19:11:16Z","title":"Echoes of Phonetics: Unveiling Relevant Acoustic Cues for ASR via Feature Attribution","version":1},"cited_work":{"arxiv_id":"2506.02181","doi":null,"metadata_source":"pith","pith_arxiv_id":"2506.02181","snapshot_observed_at":"2026-08-07T11:32:17.963370Z","title":"Echoes of Phonetics: Unveiling Relevant Acoustic Cues for ASR via Feature Attribution","venue":"cs.CL","work_id":"3e06c9d4-24cb-4dce-aa67-06b109b9e287","year":2025},"citing_paper":{"arxiv_id":"2506.02181","last_updated":"2025-06-02T19:11:16Z","snapshot_observed_at":"2026-08-08T05:50:50.489478Z","submitted_at":"2025-06-02T19:11:16Z","title":"Echoes of Phonetics: Unveiling Relevant Acoustic Cues for ASR via Feature Attribution","version":1},"reference_index":2,"source":"pdf_text","source_observed_at":"2026-08-07T11:32:13.765343Z"},"links":{"cited_paper":"/paper/2506.02181","citing_paper":"/paper/2506.02181"},"observation_digest":"sha256:68040d3ca1988724e85727e1c6e3ac07ca434b0f213aef2855705fa3bf78ac21","observation_id":"0b075e5a-482f-47e0-ab8a-9b01b6920469","resolution":{"observed_at":"2026-08-07T11:32:18.040121Z","resolver_source":"local_arxiv","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-18T06:34:40.430872+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-18T06:34:40.430872+00:00","source":"crossref"},{"observed_at":"2026-08-18T06:34:34.496301+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T11:32:25.433599Z","title":"Output text is encoded into BPE","venue":null,"work_id":"1e751472-8ab0-4f23-914c-39df773c0231","year":null},"citing_paper":{"arxiv_id":"2506.02181","last_updated":"2025-06-02T19:11:16Z","snapshot_observed_at":"2026-08-08T05:50:50.489478Z","submitted_at":"2025-06-02T19:11:16Z","title":"Echoes of Phonetics: Unveiling Relevant Acoustic Cues for ASR via Feature Attribution","version":1},"reference_index":3,"source":"pdf_text","source_observed_at":"2026-08-07T11:32:13.851341Z"},"links":{"citing_paper":"/paper/2506.02181"},"observation_digest":"sha256:e3fffee0bef94b11d6d7fba98686962c025ea97817a6247239222b8d36833eb1","observation_id":"b0cc402e-91c3-4731-bb11-86d9ab1b1021","resolution":{"observed_at":"2026-08-07T11:32:25.519045Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-18T06:34:40.430872+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-18T06:34:40.430872+00:00","source":"crossref"},{"observed_at":"2026-08-18T06:34:34.496301+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T11:32:25.056826Z","title":"Time Fig","venue":null,"work_id":"0a125716-7ad1-4ab8-b730-fdd6d8bbe35f","year":null},"citing_paper":{"arxiv_id":"2506.02181","last_updated":"2025-06-02T19:11:16Z","snapshot_observed_at":"2026-08-08T05:50:50.489478Z","submitted_at":"2025-06-02T19:11:16Z","title":"Echoes of Phonetics: Unveiling Relevant Acoustic Cues for ASR via Feature Attribution","version":1},"reference_index":4,"source":"pdf_text","source_observed_at":"2026-08-07T11:32:14.077263Z"},"links":{"citing_paper":"/paper/2506.02181"},"observation_digest":"sha256:e885e77c7f0fc23f422ff645e99490a4069e90695af407a7c554711734198f01","observation_id":"55b8654c-edf4-4196-a818-749efbd5c139","resolution":{"observed_at":"2026-08-07T11:32:25.090547Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-18T06:34:40.430872+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-18T06:34:40.430872+00:00","source":"crossref"},{"observed_at":"2026-08-18T06:34:34.496301+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T11:32:25.161548Z","title":"The encoder layers use a convolution kernel size of31, an em- bedding size of512, and a linear layer hidden size of 2 048","venue":null,"work_id":"e49f7d6e-34d3-4295-9ee5-9ab185d8c6a8","year":null},"citing_paper":{"arxiv_id":"2506.02181","last_updated":"2025-06-02T19:11:16Z","snapshot_observed_at":"2026-08-08T05:50:50.489478Z","submitted_at":"2025-06-02T19:11:16Z","title":"Echoes of Phonetics: Unveiling Relevant Acoustic Cues for ASR via Feature Attribution","version":1},"reference_index":5,"source":"pdf_text","source_observed_at":"2026-08-07T11:32:13.957612Z"},"links":{"citing_paper":"/paper/2506.02181"},"observation_digest":"sha256:08dac3e5419a7a3c34f68e6520141b56388c7025025db5b90d632a5b32d7bcbc","observation_id":"df277d58-ba0c-4348-87c0-c04a91f3b2e3","resolution":{"observed_at":"2026-08-07T11:32:25.248711Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-18T06:34:40.430872+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-18T06:34:40.430872+00:00","source":"crossref"},{"observed_at":"2026-08-18T06:34:34.496301+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T11:32:24.841459Z","title":null,"venue":null,"work_id":"2f086064-7106-4396-aa77-60c4e25927a0","year":null},"citing_paper":{"arxiv_id":"2506.02181","last_updated":"2025-06-02T19:11:16Z","snapshot_observed_at":"2026-08-08T05:50:50.489478Z","submitted_at":"2025-06-02T19:11:16Z","title":"Echoes of Phonetics: Unveiling Relevant Acoustic Cues for ASR via Feature Attribution","version":1},"reference_index":6,"source":"pdf_text","source_observed_at":"2026-08-07T11:32:14.179728Z"},"links":{"citing_paper":"/paper/2506.02181"},"observation_digest":"sha256:47f55da8e1f3b0a93a4f4d4c11bbb3f0635f811a6de729f8e1023cf5181ce1f4","observation_id":"237cf734-eaba-43a1-a1a0-261ae9493e39","resolution":{"observed_at":"2026-08-07T11:32:24.899323Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-18T06:34:40.430872+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-18T06:34:40.430872+00:00","source":"crossref"},{"observed_at":"2026-08-18T06:34:34.496301+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T11:32:24.964832Z","title":"This is the first in-depth analysis of saliency maps in relation to fine-grained acoustic patterns across three phoneme classes","venue":null,"work_id":"54a9dd8b-77cc-4b9a-bc13-0c8b1758f063","year":null},"citing_paper":{"arxiv_id":"2506.02181","last_updated":"2025-06-02T19:11:16Z","snapshot_observed_at":"2026-08-08T05:50:50.489478Z","submitted_at":"2025-06-02T19:11:16Z","title":"Echoes of Phonetics: Unveiling Relevant Acoustic Cues for ASR via Feature Attribution","version":1},"reference_index":7,"source":"pdf_text","source_observed_at":"2026-08-07T11:32:14.125900Z"},"links":{"citing_paper":"/paper/2506.02181"},"observation_digest":"sha256:ab1564643ff4b143f5c63cd02a1c37166fd4ac6d6b60d4c64593dcaf93671857","observation_id":"192b2f02-aa97-4c3d-82e6-7deffcdda717","resolution":{"observed_at":"2026-08-07T11:32:25.001122Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-18T06:34:40.430872+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-18T06:34:40.430872+00:00","source":"crossref"},{"observed_at":"2026-08-18T06:34:34.496301+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T11:32:23.299291Z","title":"Understanding the representa- tion and computation of multilayer perceptrons: A case study in speech recognition,","venue":null,"work_id":"a2ed62e7-cf3a-42d3-82ba-80c7ff1417e0","year":2017},"citing_paper":{"arxiv_id":"2506.02181","last_updated":"2025-06-02T19:11:16Z","snapshot_observed_at":"2026-08-08T05:50:50.489478Z","submitted_at":"2025-06-02T19:11:16Z","title":"Echoes of Phonetics: Unveiling Relevant Acoustic Cues for ASR via Feature Attribution","version":1},"reference_index":8,"source":"pdf_text","source_observed_at":"2026-08-07T11:32:14.927096Z"},"links":{"citing_paper":"/paper/2506.02181"},"observation_digest":"sha256:2f6783b5fe41450ffe0e8929b322953e9a8e39ebd9565402e2bc72b654567e7f","observation_id":"bfa9bd22-277f-41f7-99d1-40563c5d85ee","resolution":{"observed_at":"2026-08-07T11:32:23.407824Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-18T06:34:40.430872+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-18T06:34:40.430872+00:00","source":"crossref"},{"observed_at":"2026-08-18T06:34:34.496301+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T11:32:24.618317Z","title":"Analyzing phonetic and graphemic representations in end-to-end automatic speech recog- nition,","venue":null,"work_id":"34c9f2a4-4ee3-4c41-a263-e0781ee395d0","year":2019},"citing_paper":{"arxiv_id":"2506.02181","last_updated":"2025-06-02T19:11:16Z","snapshot_observed_at":"2026-08-08T05:50:50.489478Z","submitted_at":"2025-06-02T19:11:16Z","title":"Echoes of Phonetics: Unveiling Relevant Acoustic Cues for ASR via Feature Attribution","version":1},"reference_index":9,"source":"pdf_text","source_observed_at":"2026-08-07T11:32:14.217408Z"},"links":{"citing_paper":"/paper/2506.02181"},"observation_digest":"sha256:4147a3ec302cd90c021c662249f9cb6f3f7c9f0bf7db6d49e1a3c53e96467c55","observation_id":"8427c8e7-6bed-4d11-9105-189ad95a3a56","resolution":{"observed_at":"2026-08-07T11:32:24.714209Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-18T06:34:40.430872+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-18T06:34:40.430872+00:00","source":"crossref"},{"observed_at":"2026-08-18T06:34:34.496301+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T11:32:24.388354Z","title":"Probing phoneme, language and speaker information in unsupervised speech representations,","venue":null,"work_id":"fb4e2295-1d74-4ab1-95c1-9abdf225543f","year":2022},"citing_paper":{"arxiv_id":"2506.02181","last_updated":"2025-06-02T19:11:16Z","snapshot_observed_at":"2026-08-08T05:50:50.489478Z","submitted_at":"2025-06-02T19:11:16Z","title":"Echoes of Phonetics: Unveiling Relevant Acoustic Cues for ASR via Feature Attribution","version":1},"reference_index":10,"source":"pdf_text","source_observed_at":"2026-08-07T11:32:14.303707Z"},"links":{"citing_paper":"/paper/2506.02181"},"observation_digest":"sha256:c27786b9fd715a517969f8ee2b8510a325c67793c435f2a41f8b59a550324b89","observation_id":"29201fdd-5aa4-4714-9f74-777854121a46","resolution":{"observed_at":"2026-08-07T11:32:24.485514Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-18T06:34:40.430872+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-18T06:34:40.430872+00:00","source":"crossref"},{"observed_at":"2026-08-18T06:34:34.496301+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T11:32:24.257066Z","title":"Domain-Informed Probing of wav2vec 2.0 Embeddings for Pho- netic Features,","venue":null,"work_id":"f12da747-6440-4d37-b135-f64f3d0bb77c","year":2022},"citing_paper":{"arxiv_id":"2506.02181","last_updated":"2025-06-02T19:11:16Z","snapshot_observed_at":"2026-08-08T05:50:50.489478Z","submitted_at":"2025-06-02T19:11:16Z","title":"Echoes of Phonetics: Unveiling Relevant Acoustic Cues for ASR via Feature Attribution","version":1},"reference_index":11,"source":"pdf_text","source_observed_at":"2026-08-07T11:32:14.412565Z"},"links":{"citing_paper":"/paper/2506.02181"},"observation_digest":"sha256:d5543471fd8e3a48bc89c12c6237f7f507e36a6d2d179b8a719b899d0a1a1594","observation_id":"dbe9d202-7fd3-4c78-aeab-be3dba4606a2","resolution":{"observed_at":"2026-08-07T11:32:24.312303Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-18T06:34:40.430872+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-18T06:34:40.430872+00:00","source":"crossref"},{"observed_at":"2026-08-18T06:34:34.496301+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T11:32:24.043243Z","title":"Probing self- supervised speech models for phonetic and phonemic informa- tion: A case study in aspiration,","venue":null,"work_id":"c9ed01ff-0c6b-456a-b890-9d939a3dfb39","year":2023},"citing_paper":{"arxiv_id":"2506.02181","last_updated":"2025-06-02T19:11:16Z","snapshot_observed_at":"2026-08-08T05:50:50.489478Z","submitted_at":"2025-06-02T19:11:16Z","title":"Echoes of Phonetics: Unveiling Relevant Acoustic Cues for ASR via Feature Attribution","version":1},"reference_index":12,"source":"pdf_text","source_observed_at":"2026-08-07T11:32:14.556191Z"},"links":{"citing_paper":"/paper/2506.02181"},"observation_digest":"sha256:d919432eca97b892ac2852e940afe0f3b59e73e98692580428c04616dc995901","observation_id":"d1e2d11f-e127-44a7-88c1-9123d62fd077","resolution":{"observed_at":"2026-08-07T11:32:24.127112Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-18T06:34:40.430872+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-18T06:34:40.430872+00:00","source":"crossref"},{"observed_at":"2026-08-18T06:34:34.496301+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T11:32:23.818388Z","title":"Self-supervised speech representations are more phonetic than semantic,","venue":null,"work_id":"c5114d3e-d08e-4522-9f8f-853e009e6c1d","year":2024},"citing_paper":{"arxiv_id":"2506.02181","last_updated":"2025-06-02T19:11:16Z","snapshot_observed_at":"2026-08-08T05:50:50.489478Z","submitted_at":"2025-06-02T19:11:16Z","title":"Echoes of Phonetics: Unveiling Relevant Acoustic Cues for ASR via Feature Attribution","version":1},"reference_index":13,"source":"pdf_text","source_observed_at":"2026-08-07T11:32:14.636029Z"},"links":{"citing_paper":"/paper/2506.02181"},"observation_digest":"sha256:b745bf76a171f69898f68f17c374150a41f3da10a74512e5a2798f4b4ea4311f","observation_id":"c469bcd5-9759-4445-90a0-6495cc0b9107","resolution":{"observed_at":"2026-08-07T11:32:23.912683Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-18T06:34:40.430872+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-18T06:34:40.430872+00:00","source":"crossref"},{"observed_at":"2026-08-18T06:34:34.496301+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T11:32:23.676261Z","title":"Encoding of phonol- ogy in a recurrent neural model of grounded speech,","venue":null,"work_id":"51f65c0a-43b0-4715-8bd3-2731eca2e522","year":2017},"citing_paper":{"arxiv_id":"2506.02181","last_updated":"2025-06-02T19:11:16Z","snapshot_observed_at":"2026-08-08T05:50:50.489478Z","submitted_at":"2025-06-02T19:11:16Z","title":"Echoes of Phonetics: Unveiling Relevant Acoustic Cues for ASR via Feature Attribution","version":1},"reference_index":14,"source":"pdf_text","source_observed_at":"2026-08-07T11:32:14.725407Z"},"links":{"citing_paper":"/paper/2506.02181"},"observation_digest":"sha256:c0937b0aa6f8826436964b966121c2e7c31c43cd3c63d47204cc28be66a0d49d","observation_id":"a9a232c3-1c1d-478b-b710-8fb9d5be4099","resolution":{"observed_at":"2026-08-07T11:32:23.745518Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-18T06:34:40.430872+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-18T06:34:40.430872+00:00","source":"crossref"},{"observed_at":"2026-08-18T06:34:34.496301+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T11:32:23.500349Z","title":"Learn- ing weakly supervised multimodal phoneme embeddings,","venue":null,"work_id":"e4d7c9e5-77d6-43b4-a1c3-df45af143dc1","year":2017},"citing_paper":{"arxiv_id":"2506.02181","last_updated":"2025-06-02T19:11:16Z","snapshot_observed_at":"2026-08-08T05:50:50.489478Z","submitted_at":"2025-06-02T19:11:16Z","title":"Echoes of Phonetics: Unveiling Relevant Acoustic Cues for ASR via Feature Attribution","version":1},"reference_index":15,"source":"pdf_text","source_observed_at":"2026-08-07T11:32:14.829948Z"},"links":{"citing_paper":"/paper/2506.02181"},"observation_digest":"sha256:b543ef908c24fcbc2f2fcbf90d42b30381d7b1b7d26043c057515a99c1f0a604","observation_id":"cbaf82b0-f548-412b-b390-a1ef98ef0803","resolution":{"observed_at":"2026-08-07T11:32:23.605363Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-18T06:34:40.430872+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-18T06:34:40.430872+00:00","source":"crossref"},{"observed_at":"2026-08-18T06:34:34.496301+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T11:32:22.253522Z","title":"Acous- tic characteristics of American English vowels,","venue":null,"work_id":"8b287103-9ba2-471a-bc08-259c94f6fe20","year":1995},"citing_paper":{"arxiv_id":"2506.02181","last_updated":"2025-06-02T19:11:16Z","snapshot_observed_at":"2026-08-08T05:50:50.489478Z","submitted_at":"2025-06-02T19:11:16Z","title":"Echoes of Phonetics: Unveiling Relevant Acoustic Cues for ASR via Feature Attribution","version":1},"reference_index":16,"source":"pdf_text","source_observed_at":"2026-08-07T11:32:15.511358Z"},"links":{"citing_paper":"/paper/2506.02181"},"observation_digest":"sha256:69f8f357556ac1c53c9363c663f8a0f8cedbcb7f5e5797e444e34842b894ce41","observation_id":"4da04087-fc0b-499c-a2b8-d111cefd052a","resolution":{"observed_at":"2026-08-07T11:32:22.306585Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-18T06:34:40.430872+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-18T06:34:40.430872+00:00","source":"crossref"},{"observed_at":"2026-08-18T06:34:34.496301+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T11:32:23.142122Z","title":"Neuron Activation Profiles for Interpreting Convolutional Speech Recognition Models,","venue":null,"work_id":"5ad1a56b-fb7d-4c7e-af11-4a685a97126a","year":2018},"citing_paper":{"arxiv_id":"2506.02181","last_updated":"2025-06-02T19:11:16Z","snapshot_observed_at":"2026-08-08T05:50:50.489478Z","submitted_at":"2025-06-02T19:11:16Z","title":"Echoes of Phonetics: Unveiling Relevant Acoustic Cues for ASR via Feature Attribution","version":1},"reference_index":17,"source":"pdf_text","source_observed_at":"2026-08-07T11:32:15.004400Z"},"links":{"citing_paper":"/paper/2506.02181"},"observation_digest":"sha256:49de18c20acf55f2cfe9523109668b2d91fbdfb11c981dfbb71271372e4e233c","observation_id":"235f25a8-d890-43e1-8166-70c1b049d5af","resolution":{"observed_at":"2026-08-07T11:32:23.215208Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-18T06:34:40.430872+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-18T06:34:40.430872+00:00","source":"crossref"},{"observed_at":"2026-08-18T06:34:34.496301+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T11:32:22.962369Z","title":"Interpretable Convolutional Filters with SincNet,","venue":null,"work_id":"cdc748b7-2b60-4685-9b81-aa2d4aea9756","year":2018},"citing_paper":{"arxiv_id":"2506.02181","last_updated":"2025-06-02T19:11:16Z","snapshot_observed_at":"2026-08-08T05:50:50.489478Z","submitted_at":"2025-06-02T19:11:16Z","title":"Echoes of Phonetics: Unveiling Relevant Acoustic Cues for ASR via Feature Attribution","version":1},"reference_index":18,"source":"pdf_text","source_observed_at":"2026-08-07T11:32:15.077490Z"},"links":{"citing_paper":"/paper/2506.02181"},"observation_digest":"sha256:78d0ddfdd557a6807a145131d62579c7d2c73ad9ab5701bbbeeb021d493ea147","observation_id":"af0cc813-9177-45ea-841e-5d938a432c13","resolution":{"observed_at":"2026-08-07T11:32:23.047254Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-18T06:34:40.430872+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-18T06:34:40.430872+00:00","source":"crossref"},{"observed_at":"2026-08-18T06:34:34.496301+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T11:32:22.780102Z","title":"Introspection for convolutional automatic speech recognition,","venue":null,"work_id":"7e4b1e69-715d-4732-b6ae-c7dcb2f64827","year":2018},"citing_paper":{"arxiv_id":"2506.02181","last_updated":"2025-06-02T19:11:16Z","snapshot_observed_at":"2026-08-08T05:50:50.489478Z","submitted_at":"2025-06-02T19:11:16Z","title":"Echoes of Phonetics: Unveiling Relevant Acoustic Cues for ASR via Feature Attribution","version":1},"reference_index":19,"source":"pdf_text","source_observed_at":"2026-08-07T11:32:15.146016Z"},"links":{"citing_paper":"/paper/2506.02181"},"observation_digest":"sha256:1ff5132354c1f95d1ccabab5d9bdf4f32bca90b01a425100a7b7b6a828b33557","observation_id":"1c7f36dc-b57b-4146-8590-64f887c27d60","resolution":{"observed_at":"2026-08-07T11:32:22.858990Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-18T06:34:40.430872+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-18T06:34:40.430872+00:00","source":"crossref"},{"observed_at":"2026-08-18T06:34:34.496301+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T11:32:22.577865Z","title":"End-to-end acoustic modeling using convolutional neural networks for HMM- based automatic speech recognition,","venue":null,"work_id":"c2b6f17a-665c-4cce-9f0d-faf48b718f4b","year":2019},"citing_paper":{"arxiv_id":"2506.02181","last_updated":"2025-06-02T19:11:16Z","snapshot_observed_at":"2026-08-08T05:50:50.489478Z","submitted_at":"2025-06-02T19:11:16Z","title":"Echoes of Phonetics: Unveiling Relevant Acoustic Cues for ASR via Feature Attribution","version":1},"reference_index":20,"source":"pdf_text","source_observed_at":"2026-08-07T11:32:15.227371Z"},"links":{"citing_paper":"/paper/2506.02181"},"observation_digest":"sha256:c6ad3e1c00703711b8fbee152bfd77775b67607aedbb571ca551dd325a0406d3","observation_id":"cbad7af6-f918-4f29-b034-15bd8714eaef","resolution":{"observed_at":"2026-08-07T11:32:22.652588Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-18T06:34:40.430872+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-18T06:34:40.430872+00:00","source":"crossref"},{"observed_at":"2026-08-18T06:34:34.496301+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2002.08125","last_updated":"2020-02-19T11:59:36Z","snapshot_observed_at":"2026-08-08T18:22:23.033620Z","submitted_at":"2020-02-19T11:59:36Z","title":"Gradient-Adjusted Neuron Activation Profiles for Comprehensive Introspection of Convolutional Speech Recognition Models","version":1},"cited_work":{"arxiv_id":"2002.08125","doi":null,"metadata_source":"pith","pith_arxiv_id":"2002.08125","snapshot_observed_at":"2026-08-07T11:32:17.809189Z","title":"Gradient-Adjusted Neuron Activation Profiles for Comprehensive Introspection of Convolutional Speech Recognition Models","venue":"cs.LG","work_id":"d4a415d6-65ba-431c-bb01-e4df23020b4b","year":2020},"citing_paper":{"arxiv_id":"2506.02181","last_updated":"2025-06-02T19:11:16Z","snapshot_observed_at":"2026-08-08T05:50:50.489478Z","submitted_at":"2025-06-02T19:11:16Z","title":"Echoes of Phonetics: Unveiling Relevant Acoustic Cues for ASR via Feature Attribution","version":1},"reference_index":21,"source":"pdf_text","source_observed_at":"2026-08-07T11:32:15.281363Z"},"links":{"cited_paper":"/paper/2002.08125","citing_paper":"/paper/2506.02181"},"observation_digest":"sha256:b4c8682a48ecd1c9270c1c59f18d5325ddb04f7b0ae81bebfdafcd23a1b54f81","observation_id":"2c80d404-16f8-4f7f-876e-d7acd6fc5683","resolution":{"observed_at":"2026-08-07T11:32:17.857757Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-18T06:34:40.430872+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-18T06:34:40.430872+00:00","source":"crossref"},{"observed_at":"2026-08-18T06:34:34.496301+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T11:32:22.406933Z","title":"Directly Comparing the Listening Strategies of Humans and Machines,","venue":null,"work_id":"c7d957d5-6b7e-495e-91da-89ff59d1d144","year":2021},"citing_paper":{"arxiv_id":"2506.02181","last_updated":"2025-06-02T19:11:16Z","snapshot_observed_at":"2026-08-08T05:50:50.489478Z","submitted_at":"2025-06-02T19:11:16Z","title":"Echoes of Phonetics: Unveiling Relevant Acoustic Cues for ASR via Feature Attribution","version":1},"reference_index":22,"source":"pdf_text","source_observed_at":"2026-08-07T11:32:15.360537Z"},"links":{"citing_paper":"/paper/2506.02181"},"observation_digest":"sha256:1ea663cdc027fb35467de56351a207b20f57e7db7f027502d58962d14d71a82f","observation_id":"34779c97-1dc5-4c41-89fe-784ecac7585e","resolution":{"observed_at":"2026-08-07T11:32:22.489366Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-18T06:34:40.430872+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-18T06:34:40.430872+00:00","source":"crossref"},{"observed_at":"2026-08-18T06:34:34.496301+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2411.01710","last_updated":"2025-03-14T22:20:26Z","snapshot_observed_at":"2026-08-16T13:03:21.601569Z","submitted_at":"2024-11-03T23:02:30Z","title":"SPES: Spectrogram Perturbation for Explainable Speech-to-Text Generation","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2411.01710","snapshot_observed_at":"2026-08-07T11:32:15.428570Z","title":"SPES: Spectrogram Perturbation for Explainable Speech-to-Text Generation,","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2506.02181","last_updated":"2025-06-02T19:11:16Z","snapshot_observed_at":"2026-08-08T05:50:50.489478Z","submitted_at":"2025-06-02T19:11:16Z","title":"Echoes of Phonetics: Unveiling Relevant Acoustic Cues for ASR via Feature Attribution","version":1},"reference_index":23,"source":"pdf_text","source_observed_at":"2026-08-07T11:32:15.428570Z"},"links":{"cited_paper":"/paper/2411.01710","citing_paper":"/paper/2506.02181"},"observation_digest":"sha256:ea9982d373457d785b2324b54515512a66b550874295d6819462f698094b875b","observation_id":"c0ff9443-cf9f-4f20-b0a5-8365f0174d6d","resolution":{"observed_at":"2026-08-07T11:32:15.428570Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T11:32:25.311401Z","title":null,"venue":null,"work_id":"db5aefe5-8bcd-4feb-8234-70254a98ea7e","year":null},"citing_paper":{"arxiv_id":"2506.02181","last_updated":"2025-06-02T19:11:16Z","snapshot_observed_at":"2026-08-08T05:50:50.489478Z","submitted_at":"2025-06-02T19:11:16Z","title":"Echoes of Phonetics: Unveiling Relevant Acoustic Cues for ASR via Feature Attribution","version":1},"reference_index":24,"source":"pdf_text","source_observed_at":"2026-08-07T11:32:13.923390Z"},"links":{"citing_paper":"/paper/2506.02181"},"observation_digest":"sha256:be3f24f93bbd321a10fd72753dd56e9f3a8a61429f18d5b8ace66ebea35d9f8b","observation_id":"b03f39f8-479e-4e1c-85ab-9dd34483a66b","resolution":{"observed_at":"2026-08-07T11:32:25.364852Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-18T06:34:40.430872+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-18T06:34:40.430872+00:00","source":"crossref"},{"observed_at":"2026-08-18T06:34:34.496301+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T11:32:22.041855Z","title":"Burst and Transition Cues to V oicing Perception for Spo- ken Initial Stops by Impaired- and Normal-Hearing Listeners,","venue":null,"work_id":"6546530a-41c0-47db-8d0d-aeb2d1553d78","year":1987},"citing_paper":{"arxiv_id":"2506.02181","last_updated":"2025-06-02T19:11:16Z","snapshot_observed_at":"2026-08-08T05:50:50.489478Z","submitted_at":"2025-06-02T19:11:16Z","title":"Echoes of Phonetics: Unveiling Relevant Acoustic Cues for ASR via Feature Attribution","version":1},"reference_index":25,"source":"pdf_text","source_observed_at":"2026-08-07T11:32:15.594193Z"},"links":{"citing_paper":"/paper/2506.02181"},"observation_digest":"sha256:2d42a68435f5705f220459f31286c2d764d63f552b8b035996b7b0040f084550","observation_id":"9a5d5fc3-9bff-4c7e-9af2-d9bccee1abff","resolution":{"observed_at":"2026-08-07T11:32:22.124476Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-18T06:34:40.430872+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-18T06:34:40.430872+00:00","source":"crossref"},{"observed_at":"2026-08-18T06:34:34.496301+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T11:32:21.809818Z","title":"Acoustic cues of voiced and voiceless plosives for determining place of articulation,","venue":null,"work_id":"b02fc274-08d7-4f3d-8d22-e02db4799edc","year":2001},"citing_paper":{"arxiv_id":"2506.02181","last_updated":"2025-06-02T19:11:16Z","snapshot_observed_at":"2026-08-08T05:50:50.489478Z","submitted_at":"2025-06-02T19:11:16Z","title":"Echoes of Phonetics: Unveiling Relevant Acoustic Cues for ASR via Feature Attribution","version":1},"reference_index":26,"source":"pdf_text","source_observed_at":"2026-08-07T11:32:15.692957Z"},"links":{"citing_paper":"/paper/2506.02181"},"observation_digest":"sha256:0d9fbf9d0a23fc0f9947a8021db66029fc67d79daa0b86b70ff302842ad923b0","observation_id":"4c5491e7-8eea-45f3-95e2-2a5e777da75e","resolution":{"observed_at":"2026-08-07T11:32:21.905743Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-18T06:34:40.430872+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-18T06:34:40.430872+00:00","source":"crossref"},{"observed_at":"2026-08-18T06:34:34.496301+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T11:32:21.599194Z","title":"Acoustic characteristics of English fricatives,","venue":null,"work_id":"bdfe37b0-0ed8-44c0-b723-560b500ae84d","year":2000},"citing_paper":{"arxiv_id":"2506.02181","last_updated":"2025-06-02T19:11:16Z","snapshot_observed_at":"2026-08-08T05:50:50.489478Z","submitted_at":"2025-06-02T19:11:16Z","title":"Echoes of Phonetics: Unveiling Relevant Acoustic Cues for ASR via Feature Attribution","version":1},"reference_index":27,"source":"pdf_text","source_observed_at":"2026-08-07T11:32:15.779334Z"},"links":{"citing_paper":"/paper/2506.02181"},"observation_digest":"sha256:80543be495fa3e9c466a15f415efd5e7cb88360b7c48d6b3a36929c8d6587958","observation_id":"1fe81541-41d6-48f7-95b1-c105c5a05685","resolution":{"observed_at":"2026-08-07T11:32:21.687005Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-18T06:34:40.430872+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-18T06:34:40.430872+00:00","source":"crossref"},{"observed_at":"2026-08-18T06:34:34.496301+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T11:32:21.419865Z","title":"Acoustic characteristics of clearly spoken English fricatives,","venue":null,"work_id":"43a55625-354e-4664-b7b0-1efb44b2fcdc","year":2009},"citing_paper":{"arxiv_id":"2506.02181","last_updated":"2025-06-02T19:11:16Z","snapshot_observed_at":"2026-08-08T05:50:50.489478Z","submitted_at":"2025-06-02T19:11:16Z","title":"Echoes of Phonetics: Unveiling Relevant Acoustic Cues for ASR via Feature Attribution","version":1},"reference_index":28,"source":"pdf_text","source_observed_at":"2026-08-07T11:32:15.889219Z"},"links":{"citing_paper":"/paper/2506.02181"},"observation_digest":"sha256:84f40688b840fd0c7fe20abea44e3a4cc52664c1f15aa5e76d736d2ab96fda52","observation_id":"24f587b8-0313-400c-866b-b6335394924b","resolution":{"observed_at":"2026-08-07T11:32:21.511061Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-18T06:34:40.430872+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-18T06:34:40.430872+00:00","source":"crossref"},{"observed_at":"2026-08-18T06:34:34.496301+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T11:32:21.143897Z","title":"Conformer: Convolution-augmented Transformer for Speech Recognition,","venue":null,"work_id":"b41f86fb-10d4-4998-a958-7116ff4a4d0d","year":2020},"citing_paper":{"arxiv_id":"2506.02181","last_updated":"2025-06-02T19:11:16Z","snapshot_observed_at":"2026-08-08T05:50:50.489478Z","submitted_at":"2025-06-02T19:11:16Z","title":"Echoes of Phonetics: Unveiling Relevant Acoustic Cues for ASR via Feature Attribution","version":1},"reference_index":29,"source":"pdf_text","source_observed_at":"2026-08-07T11:32:15.968308Z"},"links":{"citing_paper":"/paper/2506.02181"},"observation_digest":"sha256:7d87403fd981474f21a6c8180048fa07abce03beebd9dcaa0377dc88c3794bb4","observation_id":"a2da6397-2f44-4d45-ad16-4085f91b2d35","resolution":{"observed_at":"2026-08-07T11:32:21.258248Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-18T06:34:40.430872+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-18T06:34:40.430872+00:00","source":"crossref"},{"observed_at":"2026-08-18T06:34:34.496301+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T11:32:16.075519Z","title":"Introducing Parsel- mouth: A Python interface to Praat,","venue":null,"work_id":null,"year":2018},"citing_paper":{"arxiv_id":"2506.02181","last_updated":"2025-06-02T19:11:16Z","snapshot_observed_at":"2026-08-08T05:50:50.489478Z","submitted_at":"2025-06-02T19:11:16Z","title":"Echoes of Phonetics: Unveiling Relevant Acoustic Cues for ASR via Feature Attribution","version":1},"reference_index":30,"source":"pdf_text","source_observed_at":"2026-08-07T11:32:16.075519Z"},"links":{"citing_paper":"/paper/2506.02181"},"observation_digest":"sha256:99ba3ebcd727512b072c565532ebe85235ae90e0cf90c3215189901a38dd53c9","observation_id":"8a09b4e1-738f-41ad-b7aa-3375bf1f9000","resolution":{"observed_at":"2026-08-07T11:32:16.075519Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T11:32:20.908533Z","title":"Pykaldi: A Python Wrapper for Kaldi,","venue":null,"work_id":"095c063e-e364-4923-8350-d2a7cfa9d216","year":2018},"citing_paper":{"arxiv_id":"2506.02181","last_updated":"2025-06-02T19:11:16Z","snapshot_observed_at":"2026-08-08T05:50:50.489478Z","submitted_at":"2025-06-02T19:11:16Z","title":"Echoes of Phonetics: Unveiling Relevant Acoustic Cues for ASR via Feature Attribution","version":1},"reference_index":31,"source":"pdf_text","source_observed_at":"2026-08-07T11:32:16.143157Z"},"links":{"citing_paper":"/paper/2506.02181"},"observation_digest":"sha256:88c0712189aee60419280c96f57c35903212068357e44e51c5854e3759328d0f","observation_id":"af314bcb-d75a-494b-bde2-ee3b78eecc9f","resolution":{"observed_at":"2026-08-07T11:32:21.020089Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-18T06:34:40.430872+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-18T06:34:40.430872+00:00","source":"crossref"},{"observed_at":"2026-08-18T06:34:34.496301+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T11:32:20.679750Z","title":"Neural Machine Transla- tion of Rare Words with Subword Units,","venue":null,"work_id":"7a0538be-2b33-476e-9014-293523d345a7","year":2016},"citing_paper":{"arxiv_id":"2506.02181","last_updated":"2025-06-02T19:11:16Z","snapshot_observed_at":"2026-08-08T05:50:50.489478Z","submitted_at":"2025-06-02T19:11:16Z","title":"Echoes of Phonetics: Unveiling Relevant Acoustic Cues for ASR via Feature Attribution","version":1},"reference_index":32,"source":"pdf_text","source_observed_at":"2026-08-07T11:32:16.218548Z"},"links":{"citing_paper":"/paper/2506.02181"},"observation_digest":"sha256:11996b5da00cee0fd92da1d7b3babb5396ee9d667834e48062ee29c4054381b0","observation_id":"3effd402-33aa-4488-80dd-0fb1fe60c79c","resolution":{"observed_at":"2026-08-07T11:32:20.786942Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-18T06:34:40.430872+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-18T06:34:40.430872+00:00","source":"crossref"},{"observed_at":"2026-08-18T06:34:34.496301+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T11:32:20.412515Z","title":"SentencePiece: A simple and lan- guage independent subword tokenizer and detokenizer for Neural Text Processing,","venue":null,"work_id":"74419b70-7fc9-4b41-a7d7-1c345fff2832","year":2018},"citing_paper":{"arxiv_id":"2506.02181","last_updated":"2025-06-02T19:11:16Z","snapshot_observed_at":"2026-08-08T05:50:50.489478Z","submitted_at":"2025-06-02T19:11:16Z","title":"Echoes of Phonetics: Unveiling Relevant Acoustic Cues for ASR via Feature Attribution","version":1},"reference_index":33,"source":"pdf_text","source_observed_at":"2026-08-07T11:32:16.297689Z"},"links":{"citing_paper":"/paper/2506.02181"},"observation_digest":"sha256:34cb56aecde2dcc90c123e43a13ecf7bac8841672c7e2f0100371506aa2c4774","observation_id":"592131ef-154a-4488-bb6f-b1564daa0dd9","resolution":{"observed_at":"2026-08-07T11:32:20.523379Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-18T06:34:40.430872+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-18T06:34:40.430872+00:00","source":"crossref"},{"observed_at":"2026-08-18T06:34:34.496301+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T11:32:20.199628Z","title":"Attention is All you Need,","venue":null,"work_id":"0edc4597-63f1-421e-994b-50a7f1d9cef8","year":2017},"citing_paper":{"arxiv_id":"2506.02181","last_updated":"2025-06-02T19:11:16Z","snapshot_observed_at":"2026-08-08T05:50:50.489478Z","submitted_at":"2025-06-02T19:11:16Z","title":"Echoes of Phonetics: Unveiling Relevant Acoustic Cues for ASR via Feature Attribution","version":1},"reference_index":34,"source":"pdf_text","source_observed_at":"2026-08-07T11:32:16.359174Z"},"links":{"citing_paper":"/paper/2506.02181"},"observation_digest":"sha256:48d46653d1c0495ef23e80b44e648fea8013667cecbb94ef6bdfaa2c6566eb7a","observation_id":"3f931957-e4ad-4cb7-95f1-c7a131c9a156","resolution":{"observed_at":"2026-08-07T11:32:20.287706Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-18T06:34:40.430872+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-18T06:34:40.430872+00:00","source":"crossref"},{"observed_at":"2026-08-18T06:34:34.496301+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T11:32:16.455994Z","title":"Fairseq S2T: Fast Speech-to-Text Modeling with Fairseq,","venue":null,"work_id":null,"year":2020},"citing_paper":{"arxiv_id":"2506.02181","last_updated":"2025-06-02T19:11:16Z","snapshot_observed_at":"2026-08-08T05:50:50.489478Z","submitted_at":"2025-06-02T19:11:16Z","title":"Echoes of Phonetics: Unveiling Relevant Acoustic Cues for ASR via Feature Attribution","version":1},"reference_index":35,"source":"pdf_text","source_observed_at":"2026-08-07T11:32:16.455994Z"},"links":{"citing_paper":"/paper/2506.02181"},"observation_digest":"sha256:b5b9cb22694590103d1de769ab697bcb87276f743296736175a587ccbe824f78","observation_id":"fe1b260f-6a28-4646-9634-947f5a0e244d","resolution":{"observed_at":"2026-08-07T11:32:16.455994Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T11:32:19.927921Z","title":"Com- mon V oice: A Massively-Multilingual Speech Corpus,","venue":null,"work_id":"74f6e01b-281a-4014-9e25-037734b375d4","year":2020},"citing_paper":{"arxiv_id":"2506.02181","last_updated":"2025-06-02T19:11:16Z","snapshot_observed_at":"2026-08-08T05:50:50.489478Z","submitted_at":"2025-06-02T19:11:16Z","title":"Echoes of Phonetics: Unveiling Relevant Acoustic Cues for ASR via Feature Attribution","version":1},"reference_index":36,"source":"pdf_text","source_observed_at":"2026-08-07T11:32:16.536895Z"},"links":{"citing_paper":"/paper/2506.02181"},"observation_digest":"sha256:47929d3447ed93f4c3c782c8687ee76b4e0cc88d8ea9e4e88e180347080c2a94","observation_id":"6d166b02-a9eb-4aad-b7b8-5a8562bf286a","resolution":{"observed_at":"2026-08-07T11:32:20.058654Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-18T06:34:40.430872+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-18T06:34:40.430872+00:00","source":"crossref"},{"observed_at":"2026-08-18T06:34:34.496301+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T11:32:16.626379Z","title":"Lib- rispeech: An ASR corpus based on public domain audio books,","venue":null,"work_id":null,"year":2015},"citing_paper":{"arxiv_id":"2506.02181","last_updated":"2025-06-02T19:11:16Z","snapshot_observed_at":"2026-08-08T05:50:50.489478Z","submitted_at":"2025-06-02T19:11:16Z","title":"Echoes of Phonetics: Unveiling Relevant Acoustic Cues for ASR via Feature Attribution","version":1},"reference_index":37,"source":"pdf_text","source_observed_at":"2026-08-07T11:32:16.626379Z"},"links":{"citing_paper":"/paper/2506.02181"},"observation_digest":"sha256:c8cff63ac2a6c048ccec8801592a5a938799d21b68db4fc3326cc56bb9f9c21c","observation_id":"5b9d54ab-ea59-49dd-a540-7eca7aa6a668","resolution":{"observed_at":"2026-08-07T11:32:16.626379Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T11:32:19.661217Z","title":"TED-LIUM 3: Twice as Much Data and Corpus Repartition for Experiments on Speaker Adaptation,","venue":null,"work_id":"5c535cd8-5966-48e6-867c-102b0ae68833","year":2018},"citing_paper":{"arxiv_id":"2506.02181","last_updated":"2025-06-02T19:11:16Z","snapshot_observed_at":"2026-08-08T05:50:50.489478Z","submitted_at":"2025-06-02T19:11:16Z","title":"Echoes of Phonetics: Unveiling Relevant Acoustic Cues for ASR via Feature Attribution","version":1},"reference_index":38,"source":"pdf_text","source_observed_at":"2026-08-07T11:32:16.724921Z"},"links":{"citing_paper":"/paper/2506.02181"},"observation_digest":"sha256:117254a10acae45313959f0e2d56bc152c26b7e6087bd79745b8195fdf5a707e","observation_id":"db3fa82a-0329-4f53-950d-1afe1a917983","resolution":{"observed_at":"2026-08-07T11:32:19.771580Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-18T06:34:40.430872+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-18T06:34:40.430872+00:00","source":"crossref"},{"observed_at":"2026-08-18T06:34:34.496301+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T11:32:19.478238Z","title":"V oxPopuli: A Large- Scale Multilingual Speech Corpus for Representation Learning, Semi-Supervised Learning and Interpretation,","venue":null,"work_id":"3e8e3a38-6175-4a61-a3b1-134ce3952d9e","year":2021},"citing_paper":{"arxiv_id":"2506.02181","last_updated":"2025-06-02T19:11:16Z","snapshot_observed_at":"2026-08-08T05:50:50.489478Z","submitted_at":"2025-06-02T19:11:16Z","title":"Echoes of Phonetics: Unveiling Relevant Acoustic Cues for ASR via Feature Attribution","version":1},"reference_index":39,"source":"pdf_text","source_observed_at":"2026-08-07T11:32:16.814577Z"},"links":{"citing_paper":"/paper/2506.02181"},"observation_digest":"sha256:e181edfe44b5090841eed162268318fff6a1494a5bef5796b295ead182854e2d","observation_id":"824f9cd2-f374-4c14-b28d-3409165a5104","resolution":{"observed_at":"2026-08-07T11:32:19.558003Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-18T06:34:40.430872+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-18T06:34:40.430872+00:00","source":"crossref"},{"observed_at":"2026-08-18T06:34:34.496301+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T11:32:19.294807Z","title":"Rethinking the Inception Architecture for Computer Vision,","venue":null,"work_id":"b84d7c60-74d6-4652-85f6-4a79909824da","year":2016},"citing_paper":{"arxiv_id":"2506.02181","last_updated":"2025-06-02T19:11:16Z","snapshot_observed_at":"2026-08-08T05:50:50.489478Z","submitted_at":"2025-06-02T19:11:16Z","title":"Echoes of Phonetics: Unveiling Relevant Acoustic Cues for ASR via Feature Attribution","version":1},"reference_index":40,"source":"pdf_text","source_observed_at":"2026-08-07T11:32:16.949565Z"},"links":{"citing_paper":"/paper/2506.02181"},"observation_digest":"sha256:8f02b8f71743182ed8974103922da40847ff825870c4fe98c8d1d11ee8468383","observation_id":"26f4e42a-4f62-481e-9f70-c3a43f6a9d06","resolution":{"observed_at":"2026-08-07T11:32:19.372863Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-18T06:34:40.430872+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-18T06:34:40.430872+00:00","source":"crossref"},{"observed_at":"2026-08-18T06:34:34.496301+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T11:32:19.096398Z","title":"Con- nectionist temporal classification: labelling unsegmented se- quence data with recurrent neural networks,","venue":null,"work_id":"33bfad77-731c-4ff5-a11b-df98ba562b9b","year":2006},"citing_paper":{"arxiv_id":"2506.02181","last_updated":"2025-06-02T19:11:16Z","snapshot_observed_at":"2026-08-08T05:50:50.489478Z","submitted_at":"2025-06-02T19:11:16Z","title":"Echoes of Phonetics: Unveiling Relevant Acoustic Cues for ASR via Feature Attribution","version":1},"reference_index":41,"source":"pdf_text","source_observed_at":"2026-08-07T11:32:17.048361Z"},"links":{"citing_paper":"/paper/2506.02181"},"observation_digest":"sha256:237ccca8702d8744de0a7be4bdc015fa4e11d948ec0eb1ce57032ef252735940","observation_id":"800f6bae-2e69-4255-b17b-ccd2ea389988","resolution":{"observed_at":"2026-08-07T11:32:19.203475Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-18T06:34:40.430872+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-18T06:34:40.430872+00:00","source":"crossref"},{"observed_at":"2026-08-18T06:34:34.496301+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T11:32:17.167561Z","title":"Adam: A Method for Stochastic Opti- mization,","venue":null,"work_id":null,"year":2015},"citing_paper":{"arxiv_id":"2506.02181","last_updated":"2025-06-02T19:11:16Z","snapshot_observed_at":"2026-08-08T05:50:50.489478Z","submitted_at":"2025-06-02T19:11:16Z","title":"Echoes of Phonetics: Unveiling Relevant Acoustic Cues for ASR via Feature Attribution","version":1},"reference_index":42,"source":"pdf_text","source_observed_at":"2026-08-07T11:32:17.167561Z"},"links":{"citing_paper":"/paper/2506.02181"},"observation_digest":"sha256:182296812eb433f25bf1c288149a97466eedadfe5a1f9d31a3c8957a44d71fa5","observation_id":"b5794d0f-42b7-49d1-aad3-55220a5f8f8a","resolution":{"observed_at":"2026-08-07T11:32:17.167561Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T11:32:18.801656Z","title":"SpecAugment: A Simple Data Aug- mentation Method for Automatic Speech Recognition,","venue":null,"work_id":"8a3c5360-8f0c-428d-a7b3-e23071204c7e","year":2019},"citing_paper":{"arxiv_id":"2506.02181","last_updated":"2025-06-02T19:11:16Z","snapshot_observed_at":"2026-08-08T05:50:50.489478Z","submitted_at":"2025-06-02T19:11:16Z","title":"Echoes of Phonetics: Unveiling Relevant Acoustic Cues for ASR via Feature Attribution","version":1},"reference_index":43,"source":"pdf_text","source_observed_at":"2026-08-07T11:32:17.259642Z"},"links":{"citing_paper":"/paper/2506.02181"},"observation_digest":"sha256:b4b6754c7145b8fb99a878857db86b39fb5cb3bbf1a793c15db14731bef28ae9","observation_id":"62b93cdb-19e3-43f4-b3ce-ada0bfe6a740","resolution":{"observed_at":"2026-08-07T11:32:18.943826Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-18T06:34:40.430872+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-18T06:34:40.430872+00:00","source":"crossref"},{"observed_at":"2026-08-18T06:34:34.496301+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T11:32:18.537308Z","title":"Darpa timit acoustic-phonetic continuous speech corpus cd-rom TIMIT,","venue":null,"work_id":"29f180a3-2833-4ab0-aad7-27098039c339","year":1993},"citing_paper":{"arxiv_id":"2506.02181","last_updated":"2025-06-02T19:11:16Z","snapshot_observed_at":"2026-08-08T05:50:50.489478Z","submitted_at":"2025-06-02T19:11:16Z","title":"Echoes of Phonetics: Unveiling Relevant Acoustic Cues for ASR via Feature Attribution","version":1},"reference_index":44,"source":"pdf_text","source_observed_at":"2026-08-07T11:32:17.396416Z"},"links":{"citing_paper":"/paper/2506.02181"},"observation_digest":"sha256:ac30e91bafebe0a9a248007643f8de3fab837051feea7e72035a88590bc1d30c","observation_id":"57088e64-c670-4bf6-a24c-a8f1cf64f16e","resolution":{"observed_at":"2026-08-07T11:32:18.661884Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-18T06:34:40.430872+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-18T06:34:40.430872+00:00","source":"crossref"},{"observed_at":"2026-08-18T06:34:34.496301+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T11:32:18.343714Z","title":"Twists, humps, and pebbles: Multilingual speech recognition models exhibit gen- der performance gaps,","venue":null,"work_id":"6411c528-8609-450c-9774-35a2d1a54c26","year":2024},"citing_paper":{"arxiv_id":"2506.02181","last_updated":"2025-06-02T19:11:16Z","snapshot_observed_at":"2026-08-08T05:50:50.489478Z","submitted_at":"2025-06-02T19:11:16Z","title":"Echoes of Phonetics: Unveiling Relevant Acoustic Cues for ASR via Feature Attribution","version":1},"reference_index":45,"source":"pdf_text","source_observed_at":"2026-08-07T11:32:17.502417Z"},"links":{"citing_paper":"/paper/2506.02181"},"observation_digest":"sha256:3336b013a48d916274e355ecc05d8818c7a046da9a84168161444b918a67f082","observation_id":"53152cd0-8d1c-437f-9961-44133f0028fe","resolution":{"observed_at":"2026-08-07T11:32:18.424129Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-18T06:34:40.430872+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-18T06:34:40.430872+00:00","source":"crossref"},{"observed_at":"2026-08-18T06:34:34.496301+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T11:32:18.184609Z","title":"Burst spectrum as a cue for the stop voicing contrast in American English,","venue":null,"work_id":"4404596a-cc3b-49de-8985-678e63195389","year":2014},"citing_paper":{"arxiv_id":"2506.02181","last_updated":"2025-06-02T19:11:16Z","snapshot_observed_at":"2026-08-08T05:50:50.489478Z","submitted_at":"2025-06-02T19:11:16Z","title":"Echoes of Phonetics: Unveiling Relevant Acoustic Cues for ASR via Feature Attribution","version":1},"reference_index":46,"source":"pdf_text","source_observed_at":"2026-08-07T11:32:17.610645Z"},"links":{"citing_paper":"/paper/2506.02181"},"observation_digest":"sha256:55dd18895b3dacf903ac7fe7be5504e281f83ad476981eed23931b361788dbe7","observation_id":"a73cfbae-9702-461d-be97-e71fab6edf86","resolution":{"observed_at":"2026-08-07T11:32:18.247636Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-18T06:34:40.430872+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-18T06:34:40.430872+00:00","source":"crossref"},{"observed_at":"2026-08-18T06:34:34.496301+00:00","source":"retraction_watch"}],"state":"measured"}}],"paper":{"arxiv_id":"2506.02181","last_updated":"2025-06-02T19:11:16Z","latest_version":1,"primary_category":"cs.CL","snapshot_observed_at":"2026-08-08T05:50:50.489478Z","submitted_at":"2025-06-02T19:11:16Z","title":"Echoes of Phonetics: Unveiling Relevant Acoustic Cues for ASR via Feature Attribution"},"reference_resolution":{"displayed":46,"state_counts":{"malformed_identifier":0,"metadata_mismatch":1,"parse_uncertain":0,"unresolved":8,"verified_exact":1,"verified_fuzzy":36},"total_outbound_references":46},"refusal":"A citation records a reference. It does not transfer a finding from one paper to another.","schema":"pith.paper-citation-record.v1","standing_sources":[{"observed_at":"2026-08-18T06:34:40.430872+00:00","source":"crossref"},{"observed_at":"2026-08-18T06:34:34.496301+00:00","source":"retraction_watch"}],"thesis":"As of 18 August 2026, this Paper Citation Record lists 46 of 46 outbound references and 1 inbound Pith citation observation for arXiv:2506.02181."}