{"as_of":"2026-08-09T21:55:00Z","caps":{"database_statements":6,"inbound":100,"outbound":100},"context_digest":"sha256:7a211b0234441c0d7c26ba5ae4fa7f98d95af9cae4ae2f032445e7b32cee803f","coverage":[{"denominator":31,"lane":"reference_resolution","note":"Typed states for the displayed outbound observations.","records_observed":31,"source":"paper_references, paper_reference_links","source_observed_at":"2026-08-08T17:12:07.484406Z","state":"measured"},{"denominator":31,"lane":"standing_notices","note":"One-hop event checks from named stored sources.","records_observed":31,"source":"scholarly_work_events, retraction_status_cache","source_observed_at":"2026-08-09T06:31:02.800959+00:00","state":"measured"},{"denominator":0,"lane":"inbound_itemization","note":"Pith citing papers itemized under the disclosed page cap.","records_observed":0,"source":"paper_references, paper_reference_links","source_observed_at":null,"state":"measured"},{"denominator":1,"lane":"external_citation_measurements","note":"A source-named dated measurement, never combined with another source.","records_observed":0,"source":"cited_works","source_observed_at":null,"state":"measured"}],"external_citation_measurements":[],"inbound":[],"links":{"evidence":"/evidence","html":"/paper/2502.05980/citation-record","integrity":"/paper/2502.05980/integrity","json":"/paper/2502.05980/citation-record.json","paper":"/paper/2502.05980"},"outbound":[{"citation":{"cited_paper":{"arxiv_id":"1904.06037","last_updated":"2019-06-25T21:34:10Z","snapshot_observed_at":"2026-07-06T07:45:36.636186Z","submitted_at":"2019-04-12T05:15:31Z","title":"Direct speech-to-speech translation with a sequence-to-sequence model","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"1904.06037","snapshot_observed_at":"2026-08-08T17:12:07.362105Z","title":"Direct speech-to-speech translation with a sequence-to-sequence model","venue":null,"work_id":null,"year":1904},"citing_paper":{"arxiv_id":"2502.05980","last_updated":"2025-02-19T21:39:35Z","snapshot_observed_at":"2026-08-09T05:22:54.297592Z","submitted_at":"2025-02-09T18:15:00Z","title":"Speech to Speech Translation with Translatotron: A State of the Art Review","version":2},"reference_index":1,"source":"pdf_text","source_observed_at":"2026-08-08T17:12:07.362105Z"},"links":{"cited_paper":"/paper/1904.06037","citing_paper":"/paper/2502.05980"},"observation_digest":"sha256:09c24cd51c5935dcac876fb88be58c5198b6b3b5dbbcb1f6e0a014884e2ee919","observation_id":"8381addb-f247-4dde-af1f-eb81c71232b9","resolution":{"observed_at":"2026-08-08T17:12:07.362105Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-08T17:12:07.896310Z","title":"Translatotron 2: High-quality di- rect speech-to-speech translation with voice preservation","venue":null,"work_id":"7311ebfa-d385-40dc-8e83-c0cc071e8a31","year":2022},"citing_paper":{"arxiv_id":"2502.05980","last_updated":"2025-02-19T21:39:35Z","snapshot_observed_at":"2026-08-09T05:22:54.297592Z","submitted_at":"2025-02-09T18:15:00Z","title":"Speech to Speech Translation with Translatotron: A State of the Art Review","version":2},"reference_index":2,"source":"pdf_text","source_observed_at":"2026-08-08T17:12:07.366961Z"},"links":{"citing_paper":"/paper/2502.05980"},"observation_digest":"sha256:82c492518d26c767ea8a1766bd403f94b6b6ba61b4dd3e438c714fe3e275116b","observation_id":"79e58410-27fc-4cc3-b6db-42de0158263c","resolution":{"observed_at":"2026-08-08T17:12:07.901737Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-08T17:12:07.884422Z","title":"Translatotron 3: Speech-to-speech translation with monolingual data","venue":null,"work_id":"93d0c0cd-c787-4664-a148-5add5253325d","year":2024},"citing_paper":{"arxiv_id":"2502.05980","last_updated":"2025-02-19T21:39:35Z","snapshot_observed_at":"2026-08-09T05:22:54.297592Z","submitted_at":"2025-02-09T18:15:00Z","title":"Speech to Speech Translation with Translatotron: A State of the Art Review","version":2},"reference_index":3,"source":"pdf_text","source_observed_at":"2026-08-08T17:12:07.370831Z"},"links":{"citing_paper":"/paper/2502.05980"},"observation_digest":"sha256:5a70060d32807f7930d8dbad80e5824abc8d78e5d27727c2458516e4b23eb360","observation_id":"118791ae-5d7a-49f4-b5fa-223ccb0e9ec6","resolution":{"observed_at":"2026-08-08T17:12:07.888622Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2211.06474","last_updated":"2022-11-11T20:21:38Z","snapshot_observed_at":"2026-08-03T02:00:26.609861Z","submitted_at":"2022-11-11T20:21:38Z","title":"Speech-to-Speech Translation For A Real-world Unwritten Language","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2211.06474","snapshot_observed_at":"2026-08-08T17:12:07.375203Z","title":"Speech-to-speech translation for a real-world un- written language","venue":null,"work_id":null,"year":2022},"citing_paper":{"arxiv_id":"2502.05980","last_updated":"2025-02-19T21:39:35Z","snapshot_observed_at":"2026-08-09T05:22:54.297592Z","submitted_at":"2025-02-09T18:15:00Z","title":"Speech to Speech Translation with Translatotron: A State of the Art Review","version":2},"reference_index":4,"source":"pdf_text","source_observed_at":"2026-08-08T17:12:07.375203Z"},"links":{"cited_paper":"/paper/2211.06474","citing_paper":"/paper/2502.05980"},"observation_digest":"sha256:28842bddd40b87a5a4004d367651fcad046c02daf7783c17fdc157484920f8b3","observation_id":"f5169b67-b03b-4e8d-8d81-39fed7a7b571","resolution":{"observed_at":"2026-08-08T17:12:07.375203Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2205.12523","last_updated":"2023-03-02T09:17:01Z","snapshot_observed_at":"2026-08-09T19:12:03.939391Z","submitted_at":"2022-05-25T06:34:14Z","title":"TranSpeech: Speech-to-Speech Translation With Bilateral Perturbation","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2205.12523","snapshot_observed_at":"2026-08-08T17:12:07.380308Z","title":"Transpeech: Speech-to-speech translation with bilateral perturbation","venue":null,"work_id":null,"year":2022},"citing_paper":{"arxiv_id":"2502.05980","last_updated":"2025-02-19T21:39:35Z","snapshot_observed_at":"2026-08-09T05:22:54.297592Z","submitted_at":"2025-02-09T18:15:00Z","title":"Speech to Speech Translation with Translatotron: A State of the Art Review","version":2},"reference_index":5,"source":"pdf_text","source_observed_at":"2026-08-08T17:12:07.380308Z"},"links":{"cited_paper":"/paper/2205.12523","citing_paper":"/paper/2502.05980"},"observation_digest":"sha256:b36a60ab0bb64c53cfe37cf6110acda09c088d2561b7f1f7c6b92bca020b3ffc","observation_id":"c8ea78ef-3bce-4218-ba9a-4ee8686bdebb","resolution":{"observed_at":"2026-08-08T17:12:07.380308Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-08T17:12:07.872321Z","title":"Finite-state speech-to-speech translation","venue":null,"work_id":"8948acd5-0423-4ee5-ae5b-315752f5159b","year":1997},"citing_paper":{"arxiv_id":"2502.05980","last_updated":"2025-02-19T21:39:35Z","snapshot_observed_at":"2026-08-09T05:22:54.297592Z","submitted_at":"2025-02-09T18:15:00Z","title":"Speech to Speech Translation with Translatotron: A State of the Art Review","version":2},"reference_index":6,"source":"pdf_text","source_observed_at":"2026-08-08T17:12:07.384431Z"},"links":{"citing_paper":"/paper/2502.05980"},"observation_digest":"sha256:e831da5b04f8319742eae3e28f75b5546cdba0fe7d718809e81cf912b3673089","observation_id":"bac74e50-4b60-4929-a135-8201e712de0c","resolution":{"observed_at":"2026-08-08T17:12:07.876665Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-08T17:12:07.860131Z","title":"The ATR multilingual speech-to-speech trans- lation system","venue":null,"work_id":"fb0ed592-b75b-474d-a65e-b6318ebe1b4c","year":2006},"citing_paper":{"arxiv_id":"2502.05980","last_updated":"2025-02-19T21:39:35Z","snapshot_observed_at":"2026-08-09T05:22:54.297592Z","submitted_at":"2025-02-09T18:15:00Z","title":"Speech to Speech Translation with Translatotron: A State of the Art Review","version":2},"reference_index":7,"source":"pdf_text","source_observed_at":"2026-08-08T17:12:07.388844Z"},"links":{"citing_paper":"/paper/2502.05980"},"observation_digest":"sha256:cf89f0639208e83422cfe96bccf5d7f3e92896dc80085664cd0878bd91469788","observation_id":"74a3a4ba-dbfe-4648-bf71-4997f43ce0c4","resolution":{"observed_at":"2026-08-08T17:12:07.864670Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2107.05604","last_updated":"2022-03-21T20:00:14Z","snapshot_observed_at":"2026-08-06T15:09:20.249960Z","submitted_at":"2021-07-12T17:40:43Z","title":"Direct speech-to-speech translation with discrete units","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2107.05604","snapshot_observed_at":"2026-08-08T17:12:07.392609Z","title":"Direct speech-to-speech translation with discrete units","venue":null,"work_id":null,"year":2021},"citing_paper":{"arxiv_id":"2502.05980","last_updated":"2025-02-19T21:39:35Z","snapshot_observed_at":"2026-08-09T05:22:54.297592Z","submitted_at":"2025-02-09T18:15:00Z","title":"Speech to Speech Translation with Translatotron: A State of the Art Review","version":2},"reference_index":8,"source":"pdf_text","source_observed_at":"2026-08-08T17:12:07.392609Z"},"links":{"cited_paper":"/paper/2107.05604","citing_paper":"/paper/2502.05980"},"observation_digest":"sha256:47e4657a71f7efe15829181c6b699f317d6760070afd97b7d97792fc9560c4cd","observation_id":"b19e3aa1-19bc-44ed-97e5-15545b1308ee","resolution":{"observed_at":"2026-08-08T17:12:07.392609Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"1710.04087","last_updated":"2018-01-30T14:41:51Z","snapshot_observed_at":"2026-07-06T06:03:43.625152Z","submitted_at":"2017-10-11T14:24:28Z","title":"Word Translation Without Parallel Data","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"1710.04087","snapshot_observed_at":"2026-08-08T17:12:07.396651Z","title":"Word translation without parallel data","venue":null,"work_id":null,"year":2017},"citing_paper":{"arxiv_id":"2502.05980","last_updated":"2025-02-19T21:39:35Z","snapshot_observed_at":"2026-08-09T05:22:54.297592Z","submitted_at":"2025-02-09T18:15:00Z","title":"Speech to Speech Translation with Translatotron: A State of the Art Review","version":2},"reference_index":9,"source":"pdf_text","source_observed_at":"2026-08-08T17:12:07.396651Z"},"links":{"cited_paper":"/paper/1710.04087","citing_paper":"/paper/2502.05980"},"observation_digest":"sha256:9973df2a47d1247b62044548a04879bf6a9234d42590d59b08d0a7cac3ff78b6","observation_id":"2787ee62-27f4-4a6c-99f1-b3fa7b2789a5","resolution":{"observed_at":"2026-08-08T17:12:07.396651Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-08T17:12:07.848129Z","title":"Normalized word embedding and orthogonal trans- form for bilingual word translation","venue":null,"work_id":"36129b3e-bca7-42b1-bdd0-5eca1fe14e6e","year":2015},"citing_paper":{"arxiv_id":"2502.05980","last_updated":"2025-02-19T21:39:35Z","snapshot_observed_at":"2026-08-09T05:22:54.297592Z","submitted_at":"2025-02-09T18:15:00Z","title":"Speech to Speech Translation with Translatotron: A State of the Art Review","version":2},"reference_index":10,"source":"pdf_text","source_observed_at":"2026-08-08T17:12:07.400773Z"},"links":{"citing_paper":"/paper/2502.05980"},"observation_digest":"sha256:cdf4d631baf7679be712df8503209fd9305a8599ae01de6fed921b9bb858e20d","observation_id":"38e69443-4f43-4220-96c9-55c7a72ca5d7","resolution":{"observed_at":"2026-08-08T17:12:07.852158Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-08T17:12:07.835423Z","title":"Verbmobil: foundations of speech-to-speech translation","venue":null,"work_id":"ec020f43-eb13-48da-ad6b-1fd4d20af004","year":2013},"citing_paper":{"arxiv_id":"2502.05980","last_updated":"2025-02-19T21:39:35Z","snapshot_observed_at":"2026-08-09T05:22:54.297592Z","submitted_at":"2025-02-09T18:15:00Z","title":"Speech to Speech Translation with Translatotron: A State of the Art Review","version":2},"reference_index":11,"source":"pdf_text","source_observed_at":"2026-08-08T17:12:07.404505Z"},"links":{"citing_paper":"/paper/2502.05980"},"observation_digest":"sha256:ee11bf65481b2c9e7b830d9d838d6a7871ad1e612a5e2f38083e7f76d19ad9d8","observation_id":"3134ab2c-d514-498a-adca-52cbfec822ad","resolution":{"observed_at":"2026-08-08T17:12:07.839387Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-08T17:12:07.822280Z","title":"Ethnologue","venue":null,"work_id":"b4db7dd6-b81a-4b2d-a5c3-dd3bce631307","year":2010},"citing_paper":{"arxiv_id":"2502.05980","last_updated":"2025-02-19T21:39:35Z","snapshot_observed_at":"2026-08-09T05:22:54.297592Z","submitted_at":"2025-02-09T18:15:00Z","title":"Speech to Speech Translation with Translatotron: A State of the Art Review","version":2},"reference_index":12,"source":"pdf_text","source_observed_at":"2026-08-08T17:12:07.408634Z"},"links":{"citing_paper":"/paper/2502.05980"},"observation_digest":"sha256:993d6a8d80c4416c8fb4558610288369e05a238e86e0ca1cef348c5e47deb19b","observation_id":"80763036-0f63-42b3-ba1e-82c27a3f3d10","resolution":{"observed_at":"2026-08-08T17:12:07.826772Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2106.01045","last_updated":"2021-06-02T09:37:37Z","snapshot_observed_at":"2026-08-09T05:54:43.349264Z","submitted_at":"2021-06-02T09:37:37Z","title":"Cascade versus Direct Speech Translation: Do the Differences Still Make a Difference?","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2106.01045","snapshot_observed_at":"2026-08-08T17:12:07.413119Z","title":"Cascade versus direct speech translation: Do the differences still make a difference?","venue":null,"work_id":null,"year":2021},"citing_paper":{"arxiv_id":"2502.05980","last_updated":"2025-02-19T21:39:35Z","snapshot_observed_at":"2026-08-09T05:22:54.297592Z","submitted_at":"2025-02-09T18:15:00Z","title":"Speech to Speech Translation with Translatotron: A State of the Art Review","version":2},"reference_index":13,"source":"pdf_text","source_observed_at":"2026-08-08T17:12:07.413119Z"},"links":{"cited_paper":"/paper/2106.01045","citing_paper":"/paper/2502.05980"},"observation_digest":"sha256:65e7ca7e273dce7a2ae26e587ee4b50b5c75737e2af5a009f220e8869e270777","observation_id":"90c36234-49b0-4685-b503-a9111a6766a6","resolution":{"observed_at":"2026-08-08T17:12:07.413119Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-08T17:12:07.808550Z","title":"Automatic speech recognition","venue":null,"work_id":"4b7d655c-bee8-4fa2-9d51-09695e8629ae","year":2016},"citing_paper":{"arxiv_id":"2502.05980","last_updated":"2025-02-19T21:39:35Z","snapshot_observed_at":"2026-08-09T05:22:54.297592Z","submitted_at":"2025-02-09T18:15:00Z","title":"Speech to Speech Translation with Translatotron: A State of the Art Review","version":2},"reference_index":14,"source":"pdf_text","source_observed_at":"2026-08-08T17:12:07.417019Z"},"links":{"citing_paper":"/paper/2502.05980"},"observation_digest":"sha256:55b006603df2165d1081dcafb9a765729e3fc67241d1bc105c28e9294db25250","observation_id":"ac2740cd-f3aa-4169-9d83-d19dddcc9ed1","resolution":{"observed_at":"2026-08-08T17:12:07.813018Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-08T17:12:07.794898Z","title":"Statistical machine translation","venue":null,"work_id":"ca026046-35ba-4278-a6d1-deb7c5880937","year":2009},"citing_paper":{"arxiv_id":"2502.05980","last_updated":"2025-02-19T21:39:35Z","snapshot_observed_at":"2026-08-09T05:22:54.297592Z","submitted_at":"2025-02-09T18:15:00Z","title":"Speech to Speech Translation with Translatotron: A State of the Art Review","version":2},"reference_index":15,"source":"pdf_text","source_observed_at":"2026-08-08T17:12:07.420802Z"},"links":{"citing_paper":"/paper/2502.05980"},"observation_digest":"sha256:ab2b787a1c9255c5a2bc84c65ad6f87f17ea18d28793d1e86bc3084a0c424739","observation_id":"685f6479-e28f-47ae-a305-78efd86b91f4","resolution":{"observed_at":"2026-08-08T17:12:07.799835Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-08T17:12:07.784234Z","title":"An introduction to text-to-speech synthesis","venue":null,"work_id":"c3188576-2925-4335-b14a-613cd4d02915","year":1997},"citing_paper":{"arxiv_id":"2502.05980","last_updated":"2025-02-19T21:39:35Z","snapshot_observed_at":"2026-08-09T05:22:54.297592Z","submitted_at":"2025-02-09T18:15:00Z","title":"Speech to Speech Translation with Translatotron: A State of the Art Review","version":2},"reference_index":16,"source":"pdf_text","source_observed_at":"2026-08-08T17:12:07.424544Z"},"links":{"citing_paper":"/paper/2502.05980"},"observation_digest":"sha256:586b425d5cd9c102cd10846e4e1f381d04276d7017bf9ea527457e1d59f2d40a","observation_id":"3964122a-60f1-4e89-a9f5-b0b2a97c0b87","resolution":{"observed_at":"2026-08-08T17:12:07.788023Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-08T17:12:07.771350Z","title":"InIberSPEECH 2018 Nov 21 (pp","venue":null,"work_id":"f6276666-d8e2-435b-9f72-d1a7704d1c52","year":2018},"citing_paper":{"arxiv_id":"2502.05980","last_updated":"2025-02-19T21:39:35Z","snapshot_observed_at":"2026-08-09T05:22:54.297592Z","submitted_at":"2025-02-09T18:15:00Z","title":"Speech to Speech Translation with Translatotron: A State of the Art Review","version":2},"reference_index":17,"source":"pdf_text","source_observed_at":"2026-08-08T17:12:07.428441Z"},"links":{"citing_paper":"/paper/2502.05980"},"observation_digest":"sha256:55f1c4ccdce464d2eb35c98feb20db7a13f2564bc5c0442fdb829340b1c027a5","observation_id":"8c57deab-8536-4ac7-855e-89f295a902c0","resolution":{"observed_at":"2026-08-08T17:12:07.776112Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-08T17:12:07.759622Z","title":"Neural machine translation: A review","venue":null,"work_id":"927fd16f-7d22-4259-a4b9-717acfb9ec50","year":2020},"citing_paper":{"arxiv_id":"2502.05980","last_updated":"2025-02-19T21:39:35Z","snapshot_observed_at":"2026-08-09T05:22:54.297592Z","submitted_at":"2025-02-09T18:15:00Z","title":"Speech to Speech Translation with Translatotron: A State of the Art Review","version":2},"reference_index":18,"source":"pdf_text","source_observed_at":"2026-08-08T17:12:07.432322Z"},"links":{"citing_paper":"/paper/2502.05980"},"observation_digest":"sha256:827499ad0f8b08e8492759645a84bca90fb08a476bccc52d19da4158b6bf731f","observation_id":"01c29fad-cfcc-4e3b-af76-2728645007ad","resolution":{"observed_at":"2026-08-08T17:12:07.763610Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-08T17:12:07.745925Z","title":"Real-time translation of Indian sign language using LSTM","venue":null,"work_id":"84ab47d9-58d8-4e4c-8520-35f0ffb19afe","year":2019},"citing_paper":{"arxiv_id":"2502.05980","last_updated":"2025-02-19T21:39:35Z","snapshot_observed_at":"2026-08-09T05:22:54.297592Z","submitted_at":"2025-02-09T18:15:00Z","title":"Speech to Speech Translation with Translatotron: A State of the Art Review","version":2},"reference_index":19,"source":"pdf_text","source_observed_at":"2026-08-08T17:12:07.437603Z"},"links":{"citing_paper":"/paper/2502.05980"},"observation_digest":"sha256:263cd7d6d4b8b363f0802e124c735b2a233961b6e6e3fa7f433fb48dce2183ba","observation_id":"4b82adb2-5334-42c2-a2a8-0e51da8a479c","resolution":{"observed_at":"2026-08-08T17:12:07.750851Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-08T17:12:07.733360Z","title":"The multimodal approach in audiovisual translation","venue":null,"work_id":"045ef554-4dc3-401e-b065-ccad2b30191b","year":2016},"citing_paper":{"arxiv_id":"2502.05980","last_updated":"2025-02-19T21:39:35Z","snapshot_observed_at":"2026-08-09T05:22:54.297592Z","submitted_at":"2025-02-09T18:15:00Z","title":"Speech to Speech Translation with Translatotron: A State of the Art Review","version":2},"reference_index":20,"source":"pdf_text","source_observed_at":"2026-08-08T17:12:07.441667Z"},"links":{"citing_paper":"/paper/2502.05980"},"observation_digest":"sha256:20076cc2fea01b6c28432812df84ca084e1328e4c5dbb041df42414f1b4c3cbb","observation_id":"94586f86-118b-45f9-9140-a0dd36162769","resolution":{"observed_at":"2026-08-08T17:12:07.737552Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-08T17:12:07.720820Z","title":"Selection criteria for low resource lan- guage programs","venue":null,"work_id":"b9871cd9-7193-496b-958c-e67b46e00091","year":2016},"citing_paper":{"arxiv_id":"2502.05980","last_updated":"2025-02-19T21:39:35Z","snapshot_observed_at":"2026-08-09T05:22:54.297592Z","submitted_at":"2025-02-09T18:15:00Z","title":"Speech to Speech Translation with Translatotron: A State of the Art Review","version":2},"reference_index":21,"source":"pdf_text","source_observed_at":"2026-08-08T17:12:07.445712Z"},"links":{"citing_paper":"/paper/2502.05980"},"observation_digest":"sha256:bd0e12b7853fd5c306fae7d6d87f02731d4d3198b8845779715b198483fdeb15","observation_id":"b78c2b5e-61b7-4970-b5bf-353c22b73a03","resolution":{"observed_at":"2026-08-08T17:12:07.725396Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-08T17:12:07.707179Z","title":"Google’s multilingual neural machine trans- lation system: Enabling zero-shot translation","venue":null,"work_id":"f4d57778-5065-4890-925d-5515230ba9ba","year":2017},"citing_paper":{"arxiv_id":"2502.05980","last_updated":"2025-02-19T21:39:35Z","snapshot_observed_at":"2026-08-09T05:22:54.297592Z","submitted_at":"2025-02-09T18:15:00Z","title":"Speech to Speech Translation with Translatotron: A State of the Art Review","version":2},"reference_index":22,"source":"pdf_text","source_observed_at":"2026-08-08T17:12:07.449509Z"},"links":{"citing_paper":"/paper/2502.05980"},"observation_digest":"sha256:ddae53ce52dd81381d554a85ae94a95e180a0c980cf1bc7cb00537d6d6831f9c","observation_id":"f5233d0c-f9b5-422d-9ef7-1f4e5cefff79","resolution":{"observed_at":"2026-08-08T17:12:07.712013Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-08T17:12:07.694489Z","title":"Systems of prosodic and paralinguistic features in English","venue":null,"work_id":"30786947-cfb8-4c3d-bdd0-2ebfd28bd400","year":2021},"citing_paper":{"arxiv_id":"2502.05980","last_updated":"2025-02-19T21:39:35Z","snapshot_observed_at":"2026-08-09T05:22:54.297592Z","submitted_at":"2025-02-09T18:15:00Z","title":"Speech to Speech Translation with Translatotron: A State of the Art Review","version":2},"reference_index":23,"source":"pdf_text","source_observed_at":"2026-08-08T17:12:07.453391Z"},"links":{"citing_paper":"/paper/2502.05980"},"observation_digest":"sha256:3eb1c5fc7dbfbe7885007252478adc94ac4122a9f7603a2e2a02b10d9b1e448a","observation_id":"a0edfd7e-4f0f-4f6a-b16d-3d816ea5be5d","resolution":{"observed_at":"2026-08-08T17:12:07.698824Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-08T17:12:07.681126Z","title":"Translation performance from the user’s perspective of large language models and neural machine translation systems","venue":null,"work_id":"541fa022-d6c9-4c49-8578-d1e9f2f8ede8","year":2023},"citing_paper":{"arxiv_id":"2502.05980","last_updated":"2025-02-19T21:39:35Z","snapshot_observed_at":"2026-08-09T05:22:54.297592Z","submitted_at":"2025-02-09T18:15:00Z","title":"Speech to Speech Translation with Translatotron: A State of the Art Review","version":2},"reference_index":24,"source":"pdf_text","source_observed_at":"2026-08-08T17:12:07.456907Z"},"links":{"citing_paper":"/paper/2502.05980"},"observation_digest":"sha256:1a842d65f8f11613b4131ca84859a16edb1e6b914c3dda109258a5a899955f27","observation_id":"e7e3726e-af5e-474e-a71b-092fd1475cfb","resolution":{"observed_at":"2026-08-08T17:12:07.685765Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"1703.08581","last_updated":"2017-06-12T13:54:12Z","snapshot_observed_at":"2026-07-06T05:35:07.974673Z","submitted_at":"2017-03-24T19:45:24Z","title":"Sequence-to-Sequence Models Can Directly Translate Foreign Speech","version":2},"cited_work":{"arxiv_id":"1703.08581","doi":null,"metadata_source":"pith","pith_arxiv_id":"1703.08581","snapshot_observed_at":"2026-08-08T17:12:07.542897Z","title":"Sequence-to-Sequence Models Can Directly Translate Foreign Speech","venue":"cs.CL","work_id":"b6cfa66b-bb3d-45b3-8d50-48e93ca61449","year":2017},"citing_paper":{"arxiv_id":"2502.05980","last_updated":"2025-02-19T21:39:35Z","snapshot_observed_at":"2026-08-09T05:22:54.297592Z","submitted_at":"2025-02-09T18:15:00Z","title":"Speech to Speech Translation with Translatotron: A State of the Art Review","version":2},"reference_index":25,"source":"pdf_text","source_observed_at":"2026-08-08T17:12:07.460697Z"},"links":{"cited_paper":"/paper/1703.08581","citing_paper":"/paper/2502.05980"},"observation_digest":"sha256:7608931b904f20cbc9051342d4ddf5ddc7d6f2a4da1356cabb0bb7b1f3f83a92","observation_id":"8a871752-3ec1-46f6-a16f-d70386d80dbc","resolution":{"observed_at":"2026-08-08T17:12:07.548978Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-08T17:12:07.668760Z","title":"An Application for Build- ing a Polish Telephone Speech Corpus","venue":null,"work_id":"618a78a8-2f5a-4ca2-b54f-9a296738ef25","year":2018},"citing_paper":{"arxiv_id":"2502.05980","last_updated":"2025-02-19T21:39:35Z","snapshot_observed_at":"2026-08-09T05:22:54.297592Z","submitted_at":"2025-02-09T18:15:00Z","title":"Speech to Speech Translation with Translatotron: A State of the Art Review","version":2},"reference_index":26,"source":"pdf_text","source_observed_at":"2026-08-08T17:12:07.464474Z"},"links":{"citing_paper":"/paper/2502.05980"},"observation_digest":"sha256:41fbd1a65d9075e895f9b98f0806a3a1db31451f6576148b26889edd70302c35","observation_id":"e88f21ac-7ace-4d9e-91cd-28198b7a08d0","resolution":{"observed_at":"2026-08-08T17:12:07.673064Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2101.00390","last_updated":"2021-07-27T04:04:57Z","snapshot_observed_at":"2026-08-09T19:09:23.880235Z","submitted_at":"2021-01-02T07:24:21Z","title":"VoxPopuli: A Large-Scale Multilingual Speech Corpus for Representation Learning, Semi-Supervised Learning and Interpretation","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2101.00390","snapshot_observed_at":"2026-08-08T17:12:07.468458Z","title":"VoxPopuli: A large-scale multilingual speech corpus for rep- resentation learning, semi-supervised learning and interpretation","venue":null,"work_id":null,"year":2021},"citing_paper":{"arxiv_id":"2502.05980","last_updated":"2025-02-19T21:39:35Z","snapshot_observed_at":"2026-08-09T05:22:54.297592Z","submitted_at":"2025-02-09T18:15:00Z","title":"Speech to Speech Translation with Translatotron: A State of the Art Review","version":2},"reference_index":27,"source":"pdf_text","source_observed_at":"2026-08-08T17:12:07.468458Z"},"links":{"cited_paper":"/paper/2101.00390","citing_paper":"/paper/2502.05980"},"observation_digest":"sha256:ab633c0435fbae89dd78881726f1651b87d74b5054a5a7176fab03f6d41e219a","observation_id":"1c453f19-e2a0-4c06-8853-114f6e5ab342","resolution":{"observed_at":"2026-08-08T17:12:07.468458Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-08T17:12:07.656210Z","title":"ÌròyìnSpeech: A multi-purpose Yorùbá Speech Corpus","venue":null,"work_id":"0065a1e4-5f62-44f4-9e60-3d1baf1f2fb6","year":null},"citing_paper":{"arxiv_id":"2502.05980","last_updated":"2025-02-19T21:39:35Z","snapshot_observed_at":"2026-08-09T05:22:54.297592Z","submitted_at":"2025-02-09T18:15:00Z","title":"Speech to Speech Translation with Translatotron: A State of the Art Review","version":2},"reference_index":28,"source":"pdf_text","source_observed_at":"2026-08-08T17:12:07.472380Z"},"links":{"citing_paper":"/paper/2502.05980"},"observation_digest":"sha256:f45c0b35636449adb8ff892d337ff6271f7aeeb95dad6c0411837f056f981d9d","observation_id":"caa28949-7473-4404-a56b-da300154f2f3","resolution":{"observed_at":"2026-08-08T17:12:07.660674Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-08T17:12:07.643293Z","title":"EnglishtoYorubashortmessageservicespeechandtexttranslatorforandroidphones","venue":null,"work_id":"53ba6ca9-577e-4b78-93d9-abdff91a1d33","year":2021},"citing_paper":{"arxiv_id":"2502.05980","last_updated":"2025-02-19T21:39:35Z","snapshot_observed_at":"2026-08-09T05:22:54.297592Z","submitted_at":"2025-02-09T18:15:00Z","title":"Speech to Speech Translation with Translatotron: A State of the Art Review","version":2},"reference_index":29,"source":"pdf_text","source_observed_at":"2026-08-08T17:12:07.476229Z"},"links":{"citing_paper":"/paper/2502.05980"},"observation_digest":"sha256:0476f7c04b7cc6f2422c078bad24d4202f0ad1fe6a445ef2287e504df2952509","observation_id":"e603c8dd-fbd6-4a32-8dd4-ddfa904a2fca","resolution":{"observed_at":"2026-08-08T17:12:07.647918Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2406.19564","last_updated":"2024-06-27T22:38:04Z","snapshot_observed_at":"2026-08-07T05:01:53.776740Z","submitted_at":"2024-06-27T22:38:04Z","title":"Voices Unheard: NLP Resources and Models for Yor\\`ub\\'a Regional Dialects","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2406.19564","snapshot_observed_at":"2026-08-08T17:12:07.480162Z","title":"Voices Unheard: NLP Resources and Models for Yorub’a Regional Di- alects","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2502.05980","last_updated":"2025-02-19T21:39:35Z","snapshot_observed_at":"2026-08-09T05:22:54.297592Z","submitted_at":"2025-02-09T18:15:00Z","title":"Speech to Speech Translation with Translatotron: A State of the Art Review","version":2},"reference_index":30,"source":"pdf_text","source_observed_at":"2026-08-08T17:12:07.480162Z"},"links":{"cited_paper":"/paper/2406.19564","citing_paper":"/paper/2502.05980"},"observation_digest":"sha256:2d8f1334d37576b9174fecfeaf838763e8740e981b5dea9dd5906c541f271329","observation_id":"4ec75216-b521-416b-b6e2-7c621772f8fb","resolution":{"observed_at":"2026-08-08T17:12:07.480162Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-08T17:12:07.630756Z","title":"Developing an open-source corpus of yoruba speech","venue":null,"work_id":"4f81c35b-d248-4ac4-8db3-81f01c67de15","year":null},"citing_paper":{"arxiv_id":"2502.05980","last_updated":"2025-02-19T21:39:35Z","snapshot_observed_at":"2026-08-09T05:22:54.297592Z","submitted_at":"2025-02-09T18:15:00Z","title":"Speech to Speech Translation with Translatotron: A State of the Art Review","version":2},"reference_index":31,"source":"pdf_text","source_observed_at":"2026-08-08T17:12:07.484406Z"},"links":{"citing_paper":"/paper/2502.05980"},"observation_digest":"sha256:077431b4f782938888f699995a690e5dbd0819c699283e89c7f01854ec156bd2","observation_id":"f1ee0de4-ab25-42ad-aeee-bfd09b944e08","resolution":{"observed_at":"2026-08-08T17:12:07.634992Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}}],"paper":{"arxiv_id":"2502.05980","last_updated":"2025-02-19T21:39:35Z","latest_version":2,"primary_category":"cs.CL","snapshot_observed_at":"2026-08-09T05:22:54.297592Z","submitted_at":"2025-02-09T18:15:00Z","title":"Speech to Speech Translation with Translatotron: A State of the Art Review"},"reference_resolution":{"displayed":31,"state_counts":{"malformed_identifier":0,"metadata_mismatch":0,"parse_uncertain":0,"unresolved":8,"verified_exact":1,"verified_fuzzy":22},"total_outbound_references":31},"refusal":"A citation records a reference. It does not transfer a finding from one paper to another.","schema":"pith.paper-citation-record.v1","standing_sources":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"thesis":"As of 9 August 2026, this Paper Citation Record lists 31 of 31 outbound references and 0 inbound Pith citation observations for arXiv:2502.05980."}