{"as_of":"2026-08-16T20:56:00Z","caps":{"database_statements":6,"inbound":100,"outbound":100},"context_digest":"sha256:6f84e364cda9b422f520c77c8c098f4eec550b7f79c1582eae098c284b2f406b","coverage":[{"denominator":30,"lane":"reference_resolution","note":"Typed states for the displayed outbound observations.","records_observed":30,"source":"paper_references, paper_reference_links","source_observed_at":"2026-08-07T05:06:37.433193Z","state":"measured"},{"denominator":30,"lane":"standing_notices","note":"One-hop event checks from named stored sources.","records_observed":30,"source":"scholarly_work_events, retraction_status_cache","source_observed_at":"2026-08-16T06:30:59.297886+00:00","state":"measured"},{"denominator":0,"lane":"inbound_itemization","note":"Pith citing papers itemized under the disclosed page cap.","records_observed":0,"source":"paper_references, paper_reference_links","source_observed_at":null,"state":"measured"},{"denominator":1,"lane":"external_citation_measurements","note":"A source-named dated measurement, never combined with another source.","records_observed":0,"source":"cited_works","source_observed_at":null,"state":"measured"}],"external_citation_measurements":[],"inbound":[],"links":{"evidence":"/evidence","html":"/paper/2506.08836/citation-record","integrity":"/paper/2506.08836/integrity","json":"/paper/2506.08836/citation-record.json","paper":"/paper/2506.08836"},"outbound":[{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T05:06:37.319054Z","title":"ss\" instead of","venue":null,"work_id":null,"year":2022},"citing_paper":{"arxiv_id":"2506.08836","last_updated":"2025-06-10T14:22:48Z","snapshot_observed_at":"2026-08-15T14:37:25.055261Z","submitted_at":"2025-06-10T14:22:48Z","title":"Advancing STT for Low-Resource Real-World Speech","version":1},"reference_index":1,"source":"pdf_text","source_observed_at":"2026-08-07T05:06:37.319054Z"},"links":{"citing_paper":"/paper/2506.08836"},"observation_digest":"sha256:1cb80e885bc69965021c046c540ecb73ec80f8128d3d9f235559c525710ce903","observation_id":"6b5c2e59-d554-4e47-bc31-1350d214273e","resolution":{"observed_at":"2026-08-07T05:06:37.319054Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T05:06:38.114492Z","title":"O’Reilly Media Inc","venue":null,"work_id":"98377288-0e08-404b-a41e-6787c9df3259","year":2009},"citing_paper":{"arxiv_id":"2506.08836","last_updated":"2025-06-10T14:22:48Z","snapshot_observed_at":"2026-08-15T14:37:25.055261Z","submitted_at":"2025-06-10T14:22:48Z","title":"Advancing STT for Low-Resource Real-World Speech","version":1},"reference_index":2,"source":"pdf_text","source_observed_at":"2026-08-07T05:06:37.324007Z"},"links":{"citing_paper":"/paper/2506.08836"},"observation_digest":"sha256:14dff2c0e12c7bd08d67d7959f1088e4783481820f9e868f4ee88e106cbbbae2","observation_id":"1d8532df-3e00-45d9-a91a-fd544949a730","resolution":{"observed_at":"2026-08-07T05:06:38.118282Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-16T06:30:59.297886+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-16T06:30:59.297886+00:00","source":"crossref"},{"observed_at":"2026-08-16T06:30:54.164669+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"1604.06174","last_updated":"2016-04-22T19:21:36Z","snapshot_observed_at":"2026-08-14T14:33:36.303679Z","submitted_at":"2016-04-21T04:15:27Z","title":"Training Deep Nets with Sublinear Memory Cost","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"1604.06174","snapshot_observed_at":"2026-08-07T05:06:37.328092Z","title":"https://doi.org/10.48550/arXiv.1604.06174, http: //arxiv.org/abs/1604.06174, arXiv:1604.06174 [cs] version: 2","venue":null,"work_id":null,"year":2016},"citing_paper":{"arxiv_id":"2506.08836","last_updated":"2025-06-10T14:22:48Z","snapshot_observed_at":"2026-08-15T14:37:25.055261Z","submitted_at":"2025-06-10T14:22:48Z","title":"Advancing STT for Low-Resource Real-World Speech","version":1},"reference_index":3,"source":"pdf_text","source_observed_at":"2026-08-07T05:06:37.328092Z"},"links":{"cited_paper":"/paper/1604.06174","citing_paper":"/paper/2506.08836"},"observation_digest":"sha256:e2b83c9e460f5ac46e6406b06e063704085c948550db70f7a8ef7eb4edc45185","observation_id":"9aa1ceef-7b5c-4708-a268-ad9995cafd0a","resolution":{"observed_at":"2026-08-07T05:06:37.328092Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T05:06:38.102118Z","title":null,"venue":null,"work_id":"8d1018f5-fcb7-4105-a235-788f562f26ec","year":2024},"citing_paper":{"arxiv_id":"2506.08836","last_updated":"2025-06-10T14:22:48Z","snapshot_observed_at":"2026-08-15T14:37:25.055261Z","submitted_at":"2025-06-10T14:22:48Z","title":"Advancing STT for Low-Resource Real-World Speech","version":1},"reference_index":4,"source":"pdf_text","source_observed_at":"2026-08-07T05:06:37.332473Z"},"links":{"citing_paper":"/paper/2506.08836"},"observation_digest":"sha256:69d81349649c2b200a6ac423dacf694b9949f1183858a7c44293f01dd8fcc632","observation_id":"867c8dcf-1392-4925-b852-b04b1019e69e","resolution":{"observed_at":"2026-08-07T05:06:38.105913Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-16T06:30:59.297886+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-16T06:30:59.297886+00:00","source":"crossref"},{"observed_at":"2026-08-16T06:30:54.164669+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2309.10299","last_updated":"2023-09-19T04:04:14Z","snapshot_observed_at":"2026-08-16T14:59:34.541879Z","submitted_at":"2023-09-19T04:04:14Z","title":"Using fine-tuning and min lookahead beam search to improve Whisper","version":1},"cited_work":{"arxiv_id":"2309.10299","doi":null,"metadata_source":"pith","pith_arxiv_id":"2309.10299","snapshot_observed_at":"2026-08-07T05:06:37.982734Z","title":"Using fine-tuning and min lookahead beam search to improve Whisper","venue":"eess.AS","work_id":"e606724e-a4e8-4e81-8cf3-947ef596731e","year":2023},"citing_paper":{"arxiv_id":"2506.08836","last_updated":"2025-06-10T14:22:48Z","snapshot_observed_at":"2026-08-15T14:37:25.055261Z","submitted_at":"2025-06-10T14:22:48Z","title":"Advancing STT for Low-Resource Real-World Speech","version":1},"reference_index":5,"source":"pdf_text","source_observed_at":"2026-08-07T05:06:37.336417Z"},"links":{"cited_paper":"/paper/2309.10299","citing_paper":"/paper/2506.08836"},"observation_digest":"sha256:7e59291a0f4cd605bd979d763954ab5a826f7820b6cf405172c86968ade9cabb","observation_id":"ae460027-1a19-42de-9574-d34a26ae44db","resolution":{"observed_at":"2026-08-07T05:06:37.987225Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-16T06:30:59.297886+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-16T06:30:59.297886+00:00","source":"crossref"},{"observed_at":"2026-08-16T06:30:54.164669+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2103.11401","last_updated":"2021-03-21T14:00:09Z","snapshot_observed_at":"2026-08-16T18:37:29.334505Z","submitted_at":"2021-03-21T14:00:09Z","title":"SwissDial: Parallel Multidialectal Corpus of Spoken Swiss German","version":1},"cited_work":{"arxiv_id":"2103.11401","doi":null,"metadata_source":"pith","pith_arxiv_id":"2103.11401","snapshot_observed_at":"2026-08-07T05:06:37.963006Z","title":"SwissDial: Parallel Multidialectal Corpus of Spoken Swiss German","venue":"cs.CL","work_id":"5d8ff65d-0d69-4a14-8e3f-279e664c1577","year":2021},"citing_paper":{"arxiv_id":"2506.08836","last_updated":"2025-06-10T14:22:48Z","snapshot_observed_at":"2026-08-15T14:37:25.055261Z","submitted_at":"2025-06-10T14:22:48Z","title":"Advancing STT for Low-Resource Real-World Speech","version":1},"reference_index":6,"source":"pdf_text","source_observed_at":"2026-08-07T05:06:37.340733Z"},"links":{"cited_paper":"/paper/2103.11401","citing_paper":"/paper/2506.08836"},"observation_digest":"sha256:3c5c356dcadb1e881cf894433497cf10413109b81e2b170a9340703bad5d67e2","observation_id":"eaaf9627-300f-4f8c-84f7-75cd3496724c","resolution":{"observed_at":"2026-08-07T05:06:37.968157Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-16T06:30:59.297886+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-16T06:30:59.297886+00:00","source":"crossref"},{"observed_at":"2026-08-16T06:30:54.164669+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T05:06:37.345312Z","title":"In: Scherrer, Y., Jauhiainen, T., Ljubešić, N., Zampieri, M., Nakov, P., Tiedemann, J","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2506.08836","last_updated":"2025-06-10T14:22:48Z","snapshot_observed_at":"2026-08-15T14:37:25.055261Z","submitted_at":"2025-06-10T14:22:48Z","title":"Advancing STT for Low-Resource Real-World Speech","version":1},"reference_index":7,"source":"pdf_text","source_observed_at":"2026-08-07T05:06:37.345312Z"},"links":{"citing_paper":"/paper/2506.08836"},"observation_digest":"sha256:783608918758d3dc29947b115d859b80aafa38b753139ab00afd916bcb2a325a","observation_id":"fbc4e6e2-6eef-43eb-8b3c-0dfafc04ca60","resolution":{"observed_at":"2026-08-07T05:06:37.345312Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T05:06:37.349048Z","title":"In: ICASSP 2024 - 2024 IEEE International Conference on Acoustics, Speech and Signal Processing (ICASSP)","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2506.08836","last_updated":"2025-06-10T14:22:48Z","snapshot_observed_at":"2026-08-15T14:37:25.055261Z","submitted_at":"2025-06-10T14:22:48Z","title":"Advancing STT for Low-Resource Real-World Speech","version":1},"reference_index":8,"source":"pdf_text","source_observed_at":"2026-08-07T05:06:37.349048Z"},"links":{"citing_paper":"/paper/2506.08836"},"observation_digest":"sha256:380b12b140f6a7abcaba9a1c931efc727132570aa2c65c5bffc981363fb962f0","observation_id":"a283b7b9-f3cf-4152-9e6c-0cf4b0c498d4","resolution":{"observed_at":"2026-08-07T05:06:37.349048Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T05:06:38.089074Z","title":null,"venue":null,"work_id":"3e4c7e0a-1dbb-457d-98db-e07961fc83db","year":2022},"citing_paper":{"arxiv_id":"2506.08836","last_updated":"2025-06-10T14:22:48Z","snapshot_observed_at":"2026-08-15T14:37:25.055261Z","submitted_at":"2025-06-10T14:22:48Z","title":"Advancing STT for Low-Resource Real-World Speech","version":1},"reference_index":9,"source":"pdf_text","source_observed_at":"2026-08-07T05:06:37.352822Z"},"links":{"citing_paper":"/paper/2506.08836"},"observation_digest":"sha256:44a77bd40ffe26543ff21bab119037c6f7f2db17bf27804a0521b89b50e23ee3","observation_id":"6b580da3-d246-4983-ba21-cb2d20fdd5ed","resolution":{"observed_at":"2026-08-07T05:06:38.093299Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-16T06:30:59.297886+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-16T06:30:59.297886+00:00","source":"crossref"},{"observed_at":"2026-08-16T06:30:54.164669+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T05:06:37.356410Z","title":"In: Interspeech 2020","venue":null,"work_id":null,"year":2020},"citing_paper":{"arxiv_id":"2506.08836","last_updated":"2025-06-10T14:22:48Z","snapshot_observed_at":"2026-08-15T14:37:25.055261Z","submitted_at":"2025-06-10T14:22:48Z","title":"Advancing STT for Low-Resource Real-World Speech","version":1},"reference_index":10,"source":"pdf_text","source_observed_at":"2026-08-07T05:06:37.356410Z"},"links":{"citing_paper":"/paper/2506.08836"},"observation_digest":"sha256:eb893527b67f48cbd8241f3880cade6df1289faca3f78beffb78eca403fa957f","observation_id":"37b6b99c-72ca-4d93-875c-9971a476ee8a","resolution":{"observed_at":"2026-08-07T05:06:37.356410Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2409.10429","last_updated":"2025-06-14T16:46:01Z","snapshot_observed_at":"2026-08-16T13:18:07.192568Z","submitted_at":"2024-09-16T16:04:16Z","title":"SMILE: Speech Meta In-Context Learning for Low-Resource Language Automatic Speech Recognition","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2409.10429","snapshot_observed_at":"2026-08-07T05:06:37.360152Z","title":"https://doi.org/10.48550/arXiv","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2506.08836","last_updated":"2025-06-10T14:22:48Z","snapshot_observed_at":"2026-08-15T14:37:25.055261Z","submitted_at":"2025-06-10T14:22:48Z","title":"Advancing STT for Low-Resource Real-World Speech","version":1},"reference_index":11,"source":"pdf_text","source_observed_at":"2026-08-07T05:06:37.360152Z"},"links":{"cited_paper":"/paper/2409.10429","citing_paper":"/paper/2506.08836"},"observation_digest":"sha256:089482f26a60404de9e1d6352d37c1cd94c3caa17dd0b25f02268c2a57117233","observation_id":"6ab666c8-9a4a-44cc-a65d-22212e507764","resolution":{"observed_at":"2026-08-07T05:06:37.360152Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T05:06:37.364239Z","title":"EURASIP Journal on Audio, Speech, and Music Process- ing 2024(1), 29 (Jun 2024)","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2506.08836","last_updated":"2025-06-10T14:22:48Z","snapshot_observed_at":"2026-08-15T14:37:25.055261Z","submitted_at":"2025-06-10T14:22:48Z","title":"Advancing STT for Low-Resource Real-World Speech","version":1},"reference_index":12,"source":"pdf_text","source_observed_at":"2026-08-07T05:06:37.364239Z"},"links":{"citing_paper":"/paper/2506.08836"},"observation_digest":"sha256:b3e2ddd488268f2f1f04aa801fa3bb3ef7a2d62cca971fefb6ab45dd9c2f4b3a","observation_id":"8ea9578c-aecc-4f2f-ad7a-815e594633a2","resolution":{"observed_at":"2026-08-07T05:06:37.364239Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T05:06:38.075349Z","title":null,"venue":null,"work_id":"36ca8b22-2430-46d2-9ec3-cd5aec053ed0","year":2023},"citing_paper":{"arxiv_id":"2506.08836","last_updated":"2025-06-10T14:22:48Z","snapshot_observed_at":"2026-08-15T14:37:25.055261Z","submitted_at":"2025-06-10T14:22:48Z","title":"Advancing STT for Low-Resource Real-World Speech","version":1},"reference_index":13,"source":"pdf_text","source_observed_at":"2026-08-07T05:06:37.367955Z"},"links":{"citing_paper":"/paper/2506.08836"},"observation_digest":"sha256:66946b6c3b8b8a2d85b6d20232be5516978e927635058dcb1b2b89e2dc940740","observation_id":"e5f2f6c4-ebb4-4807-bc86-097330cfa496","resolution":{"observed_at":"2026-08-07T05:06:38.079769Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-16T06:30:59.297886+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-16T06:30:59.297886+00:00","source":"crossref"},{"observed_at":"2026-08-16T06:30:54.164669+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T05:06:38.061351Z","title":"In: 7th Inter- national Conference on Learning Representations, ICLR 2019, New Orleans, LA, USA, May 6-9, 2019","venue":null,"work_id":"4676d0f1-4ebf-4fa1-a17c-24a5ce11963c","year":2019},"citing_paper":{"arxiv_id":"2506.08836","last_updated":"2025-06-10T14:22:48Z","snapshot_observed_at":"2026-08-15T14:37:25.055261Z","submitted_at":"2025-06-10T14:22:48Z","title":"Advancing STT for Low-Resource Real-World Speech","version":1},"reference_index":14,"source":"pdf_text","source_observed_at":"2026-08-07T05:06:37.371450Z"},"links":{"citing_paper":"/paper/2506.08836"},"observation_digest":"sha256:66c991ecd3b6729e8cb65a6f72cc3ed7161996e0a9e333fd0fec0b9a11425fc3","observation_id":"53d79b48-2038-4aa2-8a40-f4952e7b80f2","resolution":{"observed_at":"2026-08-07T05:06:38.065513Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-16T06:30:59.297886+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-16T06:30:59.297886+00:00","source":"crossref"},{"observed_at":"2026-08-16T06:30:54.164669+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":"10.21437/slate.2023-20","metadata_source":"openalex","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-10T05:30:23.456663Z","title":"In: 9th Workshop on Speech and Language Tech- nology in Education (SLaTE)","venue":null,"work_id":"ab2d12d3-da34-472c-bf4e-cccf82ea1a95","year":2023},"citing_paper":{"arxiv_id":"2506.08836","last_updated":"2025-06-10T14:22:48Z","snapshot_observed_at":"2026-08-15T14:37:25.055261Z","submitted_at":"2025-06-10T14:22:48Z","title":"Advancing STT for Low-Resource Real-World Speech","version":1},"reference_index":15,"source":"pdf_text","source_observed_at":"2026-08-07T05:06:37.374904Z"},"links":{"citing_paper":"/paper/2506.08836"},"observation_digest":"sha256:96832ef9591a56568b91e6bc8ad1145ee9d37bd6c39be4a73d0a3d868f34fa3b","observation_id":"136106d7-1175-4976-86f6-5497f7947566","resolution":{"observed_at":"2026-08-07T05:06:37.555540Z","resolver_source":"doi","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-16T06:30:59.297886+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-16T06:30:59.297886+00:00","source":"crossref"},{"observed_at":"2026-08-16T06:30:54.164669+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":"10.18653/v1/2023.findings-emnlp.1018","metadata_source":"openalex","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-10T05:30:23.456663Z","title":"In: Bouamor, H., Pino, J., Bali, K","venue":null,"work_id":"b4eaa086-6b93-4e70-b704-a138aae132c0","year":2023},"citing_paper":{"arxiv_id":"2506.08836","last_updated":"2025-06-10T14:22:48Z","snapshot_observed_at":"2026-08-15T14:37:25.055261Z","submitted_at":"2025-06-10T14:22:48Z","title":"Advancing STT for Low-Resource Real-World Speech","version":1},"reference_index":16,"source":"pdf_text","source_observed_at":"2026-08-07T05:06:37.378461Z"},"links":{"citing_paper":"/paper/2506.08836"},"observation_digest":"sha256:60b030eb89e6b5daf9331aa370f358d29983b46b97ebf6a0b46ab24a430467ca","observation_id":"3cd49efa-28ac-4cd1-985b-198c398a2c55","resolution":{"observed_at":"2026-08-07T05:06:37.542081Z","resolver_source":"doi","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-16T06:30:59.297886+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-16T06:30:59.297886+00:00","source":"crossref"},{"observed_at":"2026-08-16T06:30:54.164669+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T05:06:37.382189Z","title":"In: Isabelle, P., Charniak, E., Lin, D","venue":null,"work_id":null,"year":2002},"citing_paper":{"arxiv_id":"2506.08836","last_updated":"2025-06-10T14:22:48Z","snapshot_observed_at":"2026-08-15T14:37:25.055261Z","submitted_at":"2025-06-10T14:22:48Z","title":"Advancing STT for Low-Resource Real-World Speech","version":1},"reference_index":17,"source":"pdf_text","source_observed_at":"2026-08-07T05:06:37.382189Z"},"links":{"citing_paper":"/paper/2506.08836"},"observation_digest":"sha256:03fd131d6523c113cf2d30d0786c3fe82a722e01d4cebfc8a500a8743885e85c","observation_id":"f5a0ff9b-dc99-417d-920e-991dbdda1a9d","resolution":{"observed_at":"2026-08-07T05:06:37.382189Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2411.04573","last_updated":"2024-11-07T09:57:57Z","snapshot_observed_at":"2026-08-16T13:02:09.845837Z","submitted_at":"2024-11-07T09:57:57Z","title":"Multistage Fine-tuning Strategies for Automatic Speech Recognition in Low-resource Languages","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2411.04573","snapshot_observed_at":"2026-08-07T05:06:37.385874Z","title":"https://doi.org/10.48550/arXiv.2411.04573, http://arxiv.org/abs/ 2411.04573, arXiv:2411.04573 [cs]","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2506.08836","last_updated":"2025-06-10T14:22:48Z","snapshot_observed_at":"2026-08-15T14:37:25.055261Z","submitted_at":"2025-06-10T14:22:48Z","title":"Advancing STT for Low-Resource Real-World Speech","version":1},"reference_index":18,"source":"pdf_text","source_observed_at":"2026-08-07T05:06:37.385874Z"},"links":{"cited_paper":"/paper/2411.04573","citing_paper":"/paper/2506.08836"},"observation_digest":"sha256:db0bb05f631c676a0e3f3a949e6d3983fd65abad7b80750296761422255a2986","observation_id":"04d042c8-45a9-47f2-b339-274358dae5e0","resolution":{"observed_at":"2026-08-07T05:06:37.385874Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":"10.21437/interspeech.2024-734","metadata_source":"openalex","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-10T05:30:23.456663Z","title":"In: Interspeech 2024","venue":null,"work_id":"3bf67a5f-a42b-4e12-8eab-ed5f9b43283f","year":2024},"citing_paper":{"arxiv_id":"2506.08836","last_updated":"2025-06-10T14:22:48Z","snapshot_observed_at":"2026-08-15T14:37:25.055261Z","submitted_at":"2025-06-10T14:22:48Z","title":"Advancing STT for Low-Resource Real-World Speech","version":1},"reference_index":19,"source":"pdf_text","source_observed_at":"2026-08-07T05:06:37.389869Z"},"links":{"citing_paper":"/paper/2506.08836"},"observation_digest":"sha256:1813aad80c471f4c514f5cd037d38b4d87e789f9c0b8313316489d23e3b10704","observation_id":"3e3958dc-469f-4e10-8d31-1021ef23885b","resolution":{"observed_at":"2026-08-07T05:06:37.515017Z","resolver_source":"doi","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-16T06:30:59.297886+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-16T06:30:59.297886+00:00","source":"crossref"},{"observed_at":"2026-08-16T06:30:54.164669+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":"10.18653/v1/2023.a","metadata_source":"doi_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T05:06:37.496470Z","title":"In: Rogers, A., Boyd-Graber, J., Okazaki, N","venue":null,"work_id":"fdfcab14-7297-40d1-a669-71b8c8d27df1","year":2023},"citing_paper":{"arxiv_id":"2506.08836","last_updated":"2025-06-10T14:22:48Z","snapshot_observed_at":"2026-08-15T14:37:25.055261Z","submitted_at":"2025-06-10T14:22:48Z","title":"Advancing STT for Low-Resource Real-World Speech","version":1},"reference_index":20,"source":"pdf_text","source_observed_at":"2026-08-07T05:06:37.393485Z"},"links":{"citing_paper":"/paper/2506.08836"},"observation_digest":"sha256:fb728bf9ba8234bfb8780b1decff1c0d76354f0cab784e1e71224e97ad9b04a1","observation_id":"8f9a8d39-d8c7-4485-a10f-79acd3ad1df1","resolution":{"observed_at":"2026-08-07T05:06:37.501868Z","resolver_source":"doi","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-16T06:30:59.297886+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-16T06:30:59.297886+00:00","source":"crossref"},{"observed_at":"2026-08-16T06:30:54.164669+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T05:06:38.047209Z","title":"In: Calzolari, N., Béchet, F., Blache, P., Choukri, K., Cieri, C., Declerck, T., Goggi, S., Isahara, H., Maegaard, B., Mariani, J., Mazo, H., Odijk, J., Piperidis, S","venue":null,"work_id":"35bb95eb-db47-4c66-8a5a-9de8508a75ac","year":2022},"citing_paper":{"arxiv_id":"2506.08836","last_updated":"2025-06-10T14:22:48Z","snapshot_observed_at":"2026-08-15T14:37:25.055261Z","submitted_at":"2025-06-10T14:22:48Z","title":"Advancing STT for Low-Resource Real-World Speech","version":1},"reference_index":21,"source":"pdf_text","source_observed_at":"2026-08-07T05:06:37.397644Z"},"links":{"citing_paper":"/paper/2506.08836"},"observation_digest":"sha256:039f3855d8b878eab49e12859bcab7592ebdfe261533f2d8046e639e1ab43b92","observation_id":"86e2eb03-195b-417b-a424-ba7d03c4439e","resolution":{"observed_at":"2026-08-07T05:06:38.051401Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-16T06:30:59.297886+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-16T06:30:59.297886+00:00","source":"crossref"},{"observed_at":"2026-08-16T06:30:54.164669+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T05:06:38.033832Z","title":"In: Proceedings of the Swiss Text Analytics Conference 2021","venue":null,"work_id":"c1c23743-b37f-4291-8f83-81c2c59208e2","year":2021},"citing_paper":{"arxiv_id":"2506.08836","last_updated":"2025-06-10T14:22:48Z","snapshot_observed_at":"2026-08-15T14:37:25.055261Z","submitted_at":"2025-06-10T14:22:48Z","title":"Advancing STT for Low-Resource Real-World Speech","version":1},"reference_index":22,"source":"pdf_text","source_observed_at":"2026-08-07T05:06:37.401819Z"},"links":{"citing_paper":"/paper/2506.08836"},"observation_digest":"sha256:2617159a7f60b69aaebf3e5e4be581f83c6402cd784999d5e22f5fbfd886e22a","observation_id":"10733f40-0781-4d46-80bd-6e61a9ac3530","resolution":{"observed_at":"2026-08-07T05:06:38.037955Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-16T06:30:59.297886+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-16T06:30:59.297886+00:00","source":"crossref"},{"observed_at":"2026-08-16T06:30:54.164669+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T05:06:37.405885Z","title":"In: Interspeech 2024","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2506.08836","last_updated":"2025-06-10T14:22:48Z","snapshot_observed_at":"2026-08-15T14:37:25.055261Z","submitted_at":"2025-06-10T14:22:48Z","title":"Advancing STT for Low-Resource Real-World Speech","version":1},"reference_index":23,"source":"pdf_text","source_observed_at":"2026-08-07T05:06:37.405885Z"},"links":{"citing_paper":"/paper/2506.08836"},"observation_digest":"sha256:2f486714ee9b3085f6434df12d296963049c33b1b278be622dc31dd2b7591367","observation_id":"ed2de395-598e-49b1-85ad-c6c732b404a4","resolution":{"observed_at":"2026-08-07T05:06:37.405885Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T05:06:37.410060Z","title":"In: Proceedings of the 40th International Conference on Machine Learning","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2506.08836","last_updated":"2025-06-10T14:22:48Z","snapshot_observed_at":"2026-08-15T14:37:25.055261Z","submitted_at":"2025-06-10T14:22:48Z","title":"Advancing STT for Low-Resource Real-World Speech","version":1},"reference_index":24,"source":"pdf_text","source_observed_at":"2026-08-07T05:06:37.410060Z"},"links":{"citing_paper":"/paper/2506.08836"},"observation_digest":"sha256:3d0965f4fbba02ddd2bb219534c12bad55a497c5526980aa81e84d397f601d1f","observation_id":"cd581552-c1b9-4d91-9ae2-5d0a3203d0b2","resolution":{"observed_at":"2026-08-07T05:06:37.410060Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2207.00412","last_updated":"2022-11-14T10:35:45Z","snapshot_observed_at":"2026-08-16T16:48:57.998644Z","submitted_at":"2022-07-01T13:43:06Z","title":"Swiss German Speech to Text system evaluation","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2207.00412","snapshot_observed_at":"2026-08-07T05:06:37.413778Z","title":null,"venue":null,"work_id":null,"year":2022},"citing_paper":{"arxiv_id":"2506.08836","last_updated":"2025-06-10T14:22:48Z","snapshot_observed_at":"2026-08-15T14:37:25.055261Z","submitted_at":"2025-06-10T14:22:48Z","title":"Advancing STT for Low-Resource Real-World Speech","version":1},"reference_index":25,"source":"pdf_text","source_observed_at":"2026-08-07T05:06:37.413778Z"},"links":{"cited_paper":"/paper/2207.00412","citing_paper":"/paper/2506.08836"},"observation_digest":"sha256:abb9cadc4d1ae1ebf35c7dc9dd0496354c096aa0d7788b180be17d43cfc1da37","observation_id":"7a64d7fc-545e-4c06-a1ae-ce2a4d4bad54","resolution":{"observed_at":"2026-08-07T05:06:37.413778Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T05:06:38.020430Z","title":"In: Ghorbel, H., Sokhn, M., Cieliebak, M., Hür- limann, M., de Salis, E., Guerne, J","venue":null,"work_id":"9ab1a0ee-53ae-4d97-a2bf-aa5937959633","year":2023},"citing_paper":{"arxiv_id":"2506.08836","last_updated":"2025-06-10T14:22:48Z","snapshot_observed_at":"2026-08-15T14:37:25.055261Z","submitted_at":"2025-06-10T14:22:48Z","title":"Advancing STT for Low-Resource Real-World Speech","version":1},"reference_index":26,"source":"pdf_text","source_observed_at":"2026-08-07T05:06:37.418138Z"},"links":{"citing_paper":"/paper/2506.08836"},"observation_digest":"sha256:1adcf116247578cd2579cb1e44a2f55d6d1f88f717a4a4261d8deeb5da936aac","observation_id":"1f520525-bc8b-4dca-8b04-3691dff1830a","resolution":{"observed_at":"2026-08-07T05:06:38.024910Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-16T06:30:59.297886+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-16T06:30:59.297886+00:00","source":"crossref"},{"observed_at":"2026-08-16T06:30:54.164669+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T05:06:38.006941Z","title":null,"venue":null,"work_id":"e7dc1afc-19ac-4a56-a144-e5154ab85cfc","year":2023},"citing_paper":{"arxiv_id":"2506.08836","last_updated":"2025-06-10T14:22:48Z","snapshot_observed_at":"2026-08-15T14:37:25.055261Z","submitted_at":"2025-06-10T14:22:48Z","title":"Advancing STT for Low-Resource Real-World Speech","version":1},"reference_index":27,"source":"pdf_text","source_observed_at":"2026-08-07T05:06:37.421916Z"},"links":{"citing_paper":"/paper/2506.08836"},"observation_digest":"sha256:7ada1914afeacbf711a2c03eac61144bdd1f13475fa21a3f155fea09b8d1224c","observation_id":"31a1a4ff-47f5-40c7-be05-36745d931d07","resolution":{"observed_at":"2026-08-07T05:06:38.011178Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-16T06:30:59.297886+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-16T06:30:59.297886+00:00","source":"crossref"},{"observed_at":"2026-08-16T06:30:54.164669+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2412.15726","last_updated":"2025-04-22T12:09:39Z","snapshot_observed_at":"2026-08-15T14:36:47.527646Z","submitted_at":"2024-12-20T09:49:02Z","title":"Fine-tuning Whisper on Low-Resource Languages for Real-World Applications","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2412.15726","snapshot_observed_at":"2026-08-07T05:06:37.425701Z","title":"https://doi.org/10.48550/arXiv.2412.15726, http://arxiv.org/abs/ 2412.15726, arXiv:2412.15726 [cs]","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2506.08836","last_updated":"2025-06-10T14:22:48Z","snapshot_observed_at":"2026-08-15T14:37:25.055261Z","submitted_at":"2025-06-10T14:22:48Z","title":"Advancing STT for Low-Resource Real-World Speech","version":1},"reference_index":28,"source":"pdf_text","source_observed_at":"2026-08-07T05:06:37.425701Z"},"links":{"cited_paper":"/paper/2412.15726","citing_paper":"/paper/2506.08836"},"observation_digest":"sha256:ad0af4d649c3ddfd33ecd5d616e93a37de6ca9725b3012e136379ba046002c36","observation_id":"c243bc8e-1f80-44e3-a30b-0f469af91b36","resolution":{"observed_at":"2026-08-07T05:06:37.425701Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T05:06:37.429598Z","title":"In: Advances in Neural In- formation Processing Systems","venue":null,"work_id":null,"year":2017},"citing_paper":{"arxiv_id":"2506.08836","last_updated":"2025-06-10T14:22:48Z","snapshot_observed_at":"2026-08-15T14:37:25.055261Z","submitted_at":"2025-06-10T14:22:48Z","title":"Advancing STT for Low-Resource Real-World Speech","version":1},"reference_index":29,"source":"pdf_text","source_observed_at":"2026-08-07T05:06:37.429598Z"},"links":{"citing_paper":"/paper/2506.08836"},"observation_digest":"sha256:109d7f7e5a59e99a772d47190eddc576986578ba24754c848338bd7d3a5beeb8","observation_id":"8490111f-f9a8-43c2-829f-c5ec3883ba90","resolution":{"observed_at":"2026-08-07T05:06:37.429598Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T05:06:37.433193Z","title":"In: Liu, Q., Schlangen, D","venue":null,"work_id":null,"year":2020},"citing_paper":{"arxiv_id":"2506.08836","last_updated":"2025-06-10T14:22:48Z","snapshot_observed_at":"2026-08-15T14:37:25.055261Z","submitted_at":"2025-06-10T14:22:48Z","title":"Advancing STT for Low-Resource Real-World Speech","version":1},"reference_index":30,"source":"pdf_text","source_observed_at":"2026-08-07T05:06:37.433193Z"},"links":{"citing_paper":"/paper/2506.08836"},"observation_digest":"sha256:d83568a962d63fc3fbd72aca2aa7a80378275b4e97be523a8d3196e7339a6b4a","observation_id":"eafbbb55-e5ca-4f16-9b79-d0df68d0e743","resolution":{"observed_at":"2026-08-07T05:06:37.433193Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}}],"paper":{"arxiv_id":"2506.08836","last_updated":"2025-06-10T14:22:48Z","latest_version":1,"primary_category":"cs.CL","snapshot_observed_at":"2026-08-15T14:37:25.055261Z","submitted_at":"2025-06-10T14:22:48Z","title":"Advancing STT for Low-Resource Real-World Speech"},"reference_resolution":{"displayed":30,"state_counts":{"malformed_identifier":0,"metadata_mismatch":0,"parse_uncertain":0,"unresolved":19,"verified_exact":6,"verified_fuzzy":5},"total_outbound_references":30},"refusal":"A citation records a reference. It does not transfer a finding from one paper to another.","schema":"pith.paper-citation-record.v1","standing_sources":[{"observed_at":"2026-08-16T06:30:59.297886+00:00","source":"crossref"},{"observed_at":"2026-08-16T06:30:54.164669+00:00","source":"retraction_watch"}],"thesis":"As of 16 August 2026, this Paper Citation Record lists 30 of 30 outbound references and 0 inbound Pith citation observations for arXiv:2506.08836."}