{"as_of":"2026-08-20T14:50:00Z","caps":{"database_statements":6,"inbound":100,"outbound":100},"context_digest":"sha256:253bfa3a28954960b02e1827c0de28cf075b54af4f8eb338c0928f99a2c22207","coverage":[{"denominator":31,"lane":"reference_resolution","note":"Typed states for the displayed outbound observations.","records_observed":31,"source":"paper_references, paper_reference_links","source_observed_at":"2026-08-06T01:05:46.447560Z","state":"measured"},{"denominator":32,"lane":"standing_notices","note":"One-hop event checks from named stored sources.","records_observed":32,"source":"scholarly_work_events, retraction_status_cache","source_observed_at":"2026-08-20T06:33:59.587034+00:00","state":"measured"},{"denominator":1,"lane":"inbound_itemization","note":"Pith citing papers itemized under the disclosed page cap.","records_observed":1,"source":"paper_references, paper_reference_links","source_observed_at":"2026-08-06T01:05:43.913842Z","state":"measured"},{"denominator":1,"lane":"external_citation_measurements","note":"A source-named dated measurement, never combined with another source.","records_observed":0,"source":"pith","source_observed_at":"2026-08-06T01:05:46.601121Z","state":"measured"}],"external_citation_measurements":[],"inbound":[{"citation":{"cited_paper":{"arxiv_id":"2508.10009","last_updated":"2025-08-05T23:56:11Z","snapshot_observed_at":"2026-08-14T03:48:48.293793Z","submitted_at":"2025-08-05T23:56:11Z","title":"Beyond Hard Sharing: Efficient Multi-Task Speech-to-Text Modeling with Supervised Mixture of Experts","version":1},"cited_work":{"arxiv_id":"2508.10009","doi":null,"metadata_source":"pith","pith_arxiv_id":"2508.10009","snapshot_observed_at":"2026-08-06T01:05:46.601121Z","title":"Beyond Hard Sharing: Efficient Multi-Task Speech-to-Text Modeling with Supervised Mixture of Experts","venue":"cs.CL","work_id":"9135e8a7-caac-45ee-a4bd-79d8430a1be0","year":2025},"citing_paper":{"arxiv_id":"2508.10009","last_updated":"2025-08-05T23:56:11Z","snapshot_observed_at":"2026-08-14T03:48:48.293793Z","submitted_at":"2025-08-05T23:56:11Z","title":"Beyond Hard Sharing: Efficient Multi-Task Speech-to-Text Modeling with Supervised Mixture of Experts","version":1},"reference_index":1,"source":"pdf_text","source_observed_at":"2026-08-06T01:05:43.913842Z"},"links":{"cited_paper":"/paper/2508.10009","citing_paper":"/paper/2508.10009"},"observation_digest":"sha256:99c26a2cf56b8c772703cb57a9ff94d6b017675fb316b19fd78bc2d9d3ba34e4","observation_id":"586c0fd9-9f91-47e6-b24d-e0a6ec76ebec","resolution":{"observed_at":"2026-08-06T01:05:46.690100Z","resolver_source":"local_arxiv","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-20T06:33:59.587034+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-20T06:33:59.587034+00:00","source":"crossref"},{"observed_at":"2026-08-20T06:33:54.927442+00:00","source":"retraction_watch"}],"state":"measured"}}],"links":{"evidence":"/evidence","html":"/paper/2508.10009/citation-record","integrity":"/paper/2508.10009/integrity","json":"/paper/2508.10009/citation-record.json","paper":"/paper/2508.10009"},"outbound":[{"citation":{"cited_paper":{"arxiv_id":"2508.10009","last_updated":"2025-08-05T23:56:11Z","snapshot_observed_at":"2026-08-14T03:48:48.293793Z","submitted_at":"2025-08-05T23:56:11Z","title":"Beyond Hard Sharing: Efficient Multi-Task Speech-to-Text Modeling with Supervised Mixture of Experts","version":1},"cited_work":{"arxiv_id":"2508.10009","doi":null,"metadata_source":"pith","pith_arxiv_id":"2508.10009","snapshot_observed_at":"2026-08-06T01:05:46.601121Z","title":"Beyond Hard Sharing: Efficient Multi-Task Speech-to-Text Modeling with Supervised Mixture of Experts","venue":"cs.CL","work_id":"9135e8a7-caac-45ee-a4bd-79d8430a1be0","year":2025},"citing_paper":{"arxiv_id":"2508.10009","last_updated":"2025-08-05T23:56:11Z","snapshot_observed_at":"2026-08-14T03:48:48.293793Z","submitted_at":"2025-08-05T23:56:11Z","title":"Beyond Hard Sharing: Efficient Multi-Task Speech-to-Text Modeling with Supervised Mixture of Experts","version":1},"reference_index":1,"source":"pdf_text","source_observed_at":"2026-08-06T01:05:43.913842Z"},"links":{"cited_paper":"/paper/2508.10009","citing_paper":"/paper/2508.10009"},"observation_digest":"sha256:99c26a2cf56b8c772703cb57a9ff94d6b017675fb316b19fd78bc2d9d3ba34e4","observation_id":"586c0fd9-9f91-47e6-b24d-e0a6ec76ebec","resolution":{"observed_at":"2026-08-06T01:05:46.690100Z","resolver_source":"local_arxiv","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-20T06:33:59.587034+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-20T06:33:59.587034+00:00","source":"crossref"},{"observed_at":"2026-08-20T06:33:54.927442+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T01:05:50.363980Z","title":"Further details are presented in the following subsections","venue":null,"work_id":"a69746be-20c0-4230-a2c3-6df1d36b08f5","year":null},"citing_paper":{"arxiv_id":"2508.10009","last_updated":"2025-08-05T23:56:11Z","snapshot_observed_at":"2026-08-14T03:48:48.293793Z","submitted_at":"2025-08-05T23:56:11Z","title":"Beyond Hard Sharing: Efficient Multi-Task Speech-to-Text Modeling with Supervised Mixture of Experts","version":1},"reference_index":2,"source":"pdf_text","source_observed_at":"2026-08-06T01:05:43.958842Z"},"links":{"citing_paper":"/paper/2508.10009"},"observation_digest":"sha256:beec6d0c22adb7b38cf1ab059ced59fca50e1dafd03ebdb29d4d96e4ea84e855","observation_id":"fbb2e7e8-46e7-41d7-b8d2-f8a1c289cea2","resolution":{"observed_at":"2026-08-06T01:05:50.422775Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-20T06:33:59.587034+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-20T06:33:59.587034+00:00","source":"crossref"},{"observed_at":"2026-08-20T06:33:54.927442+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T01:05:50.238410Z","title":"Datasets 3.1.1","venue":null,"work_id":"731f3427-3851-4b83-8ad3-1ea514fa81c9","year":null},"citing_paper":{"arxiv_id":"2508.10009","last_updated":"2025-08-05T23:56:11Z","snapshot_observed_at":"2026-08-14T03:48:48.293793Z","submitted_at":"2025-08-05T23:56:11Z","title":"Beyond Hard Sharing: Efficient Multi-Task Speech-to-Text Modeling with Supervised Mixture of Experts","version":1},"reference_index":3,"source":"pdf_text","source_observed_at":"2026-08-06T01:05:44.019470Z"},"links":{"citing_paper":"/paper/2508.10009"},"observation_digest":"sha256:e9bc4ef04c7aa543750e8a1cc90296f65ac8e65f8ea5471fbdc47e3f11b70696","observation_id":"e6ec2d64-78eb-446e-a424-3b027d711362","resolution":{"observed_at":"2026-08-06T01:05:50.305824Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-20T06:33:59.587034+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-20T06:33:59.587034+00:00","source":"crossref"},{"observed_at":"2026-08-20T06:33:54.927442+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T01:05:49.946843Z","title":"Decoder S-MoE To validate the effectiveness of S-MoE applied to the decoder blocks, we conducted experiments on two tasks: Korean-to- English (ko2en) ST and Korean ASR (ko-ASR)","venue":null,"work_id":"edfbb1db-54c7-4a08-9b63-34cb297f5a05","year":null},"citing_paper":{"arxiv_id":"2508.10009","last_updated":"2025-08-05T23:56:11Z","snapshot_observed_at":"2026-08-14T03:48:48.293793Z","submitted_at":"2025-08-05T23:56:11Z","title":"Beyond Hard Sharing: Efficient Multi-Task Speech-to-Text Modeling with Supervised Mixture of Experts","version":1},"reference_index":4,"source":"pdf_text","source_observed_at":"2026-08-06T01:05:44.179893Z"},"links":{"citing_paper":"/paper/2508.10009"},"observation_digest":"sha256:52ad224628c6a408a1d80ec2fbdf92f2b464865a6e5057ed50788e0cd98de0a4","observation_id":"0aeff68f-136c-4610-8c23-4e739c4276f6","resolution":{"observed_at":"2026-08-06T01:05:50.012876Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-20T06:33:59.587034+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-20T06:33:59.587034+00:00","source":"crossref"},{"observed_at":"2026-08-20T06:33:54.927442+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T01:05:49.589231Z","title":"By us- ing guiding tokens instead of dynamic gating functions, S-MoE ensures efficient training and inference while improving per- formance across various tasks","venue":null,"work_id":"e89906b2-e7aa-4cd8-a5f2-26c24d58a9a7","year":null},"citing_paper":{"arxiv_id":"2508.10009","last_updated":"2025-08-05T23:56:11Z","snapshot_observed_at":"2026-08-14T03:48:48.293793Z","submitted_at":"2025-08-05T23:56:11Z","title":"Beyond Hard Sharing: Efficient Multi-Task Speech-to-Text Modeling with Supervised Mixture of Experts","version":1},"reference_index":5,"source":"pdf_text","source_observed_at":"2026-08-06T01:05:44.304450Z"},"links":{"citing_paper":"/paper/2508.10009"},"observation_digest":"sha256:721b839de15553c6115fdf880bc4ba7e4f5457b728af5a28936911099f9bbb02","observation_id":"8aa07d1d-aefc-43ca-bda8-c453b431af92","resolution":{"observed_at":"2026-08-06T01:05:49.696543Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-20T06:33:59.587034+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-20T06:33:59.587034+00:00","source":"crossref"},{"observed_at":"2026-08-20T06:33:54.927442+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T01:05:49.271138Z","title":"Sources of degradation of speech recognition in the telephone network,","venue":null,"work_id":"7258e59f-e74c-4680-a010-cf3adc6a76b2","year":1994},"citing_paper":{"arxiv_id":"2508.10009","last_updated":"2025-08-05T23:56:11Z","snapshot_observed_at":"2026-08-14T03:48:48.293793Z","submitted_at":"2025-08-05T23:56:11Z","title":"Beyond Hard Sharing: Efficient Multi-Task Speech-to-Text Modeling with Supervised Mixture of Experts","version":1},"reference_index":6,"source":"pdf_text","source_observed_at":"2026-08-06T01:05:44.870745Z"},"links":{"citing_paper":"/paper/2508.10009"},"observation_digest":"sha256:4acc4b56bafee4ca9bd9f58378091ad6dd411e7ab5afd6e783b4849e79024728","observation_id":"d9e30dc1-fdae-407f-aa2d-6e8400a5ed5a","resolution":{"observed_at":"2026-08-06T01:05:49.372762Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-20T06:33:59.587034+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-20T06:33:59.587034+00:00","source":"crossref"},{"observed_at":"2026-08-20T06:33:54.927442+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T01:05:49.140370Z","title":"Training wideband acoustic mod- els using mixed-bandwidth training data for speech recognition,","venue":null,"work_id":"e57bbf2a-9f1d-4d6f-9c43-2895d81c6498","year":2006},"citing_paper":{"arxiv_id":"2508.10009","last_updated":"2025-08-05T23:56:11Z","snapshot_observed_at":"2026-08-14T03:48:48.293793Z","submitted_at":"2025-08-05T23:56:11Z","title":"Beyond Hard Sharing: Efficient Multi-Task Speech-to-Text Modeling with Supervised Mixture of Experts","version":1},"reference_index":7,"source":"pdf_text","source_observed_at":"2026-08-06T01:05:44.974020Z"},"links":{"citing_paper":"/paper/2508.10009"},"observation_digest":"sha256:872a2df26a6f95ca4ed85990be658b8577832cf63bffd90ba2c0f73b5397d45e","observation_id":"52a8f881-d60d-4db8-ad6b-d52e9d955011","resolution":{"observed_at":"2026-08-06T01:05:49.194971Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-20T06:33:59.587034+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-20T06:33:59.587034+00:00","source":"crossref"},{"observed_at":"2026-08-20T06:33:54.927442+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2009.09796","last_updated":"2020-09-10T19:31:04Z","snapshot_observed_at":"2026-08-09T08:30:49.737024Z","submitted_at":"2020-09-10T19:31:04Z","title":"Multi-Task Learning with Deep Neural Networks: A Survey","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2009.09796","snapshot_observed_at":"2026-08-06T01:05:44.392830Z","title":"Multi-task learning with deep neural networks: A survey,","venue":null,"work_id":null,"year":2009},"citing_paper":{"arxiv_id":"2508.10009","last_updated":"2025-08-05T23:56:11Z","snapshot_observed_at":"2026-08-14T03:48:48.293793Z","submitted_at":"2025-08-05T23:56:11Z","title":"Beyond Hard Sharing: Efficient Multi-Task Speech-to-Text Modeling with Supervised Mixture of Experts","version":1},"reference_index":8,"source":"pdf_text","source_observed_at":"2026-08-06T01:05:44.392830Z"},"links":{"cited_paper":"/paper/2009.09796","citing_paper":"/paper/2508.10009"},"observation_digest":"sha256:002c01a62988cf8017aa4716a5ee277a4ea69e2283465ff3a4205506392b00c8","observation_id":"ad7de39c-ef70-43ce-85c2-34c2e3cdf340","resolution":{"observed_at":"2026-08-06T01:05:44.392830Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"1706.05098","last_updated":"2017-06-15T21:38:12Z","snapshot_observed_at":"2026-08-13T21:53:02.384757Z","submitted_at":"2017-06-15T21:38:12Z","title":"An Overview of Multi-Task Learning in Deep Neural Networks","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"1706.05098","snapshot_observed_at":"2026-08-06T01:05:44.487656Z","title":"An overview of multi-task learning in deep neural net- works,","venue":null,"work_id":null,"year":2017},"citing_paper":{"arxiv_id":"2508.10009","last_updated":"2025-08-05T23:56:11Z","snapshot_observed_at":"2026-08-14T03:48:48.293793Z","submitted_at":"2025-08-05T23:56:11Z","title":"Beyond Hard Sharing: Efficient Multi-Task Speech-to-Text Modeling with Supervised Mixture of Experts","version":1},"reference_index":9,"source":"pdf_text","source_observed_at":"2026-08-06T01:05:44.487656Z"},"links":{"cited_paper":"/paper/1706.05098","citing_paper":"/paper/2508.10009"},"observation_digest":"sha256:2b5a1384e469949e5e69433a246f854ad83a337695f9150ded3261a586366d19","observation_id":"8b4dfa92-13ca-4b43-b161-c2d33b5e5357","resolution":{"observed_at":"2026-08-06T01:05:44.487656Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T01:05:49.445786Z","title":"A survey on multi-task learning ieee transactions on knowledge and data engineering,","venue":null,"work_id":"02802804-b2a3-4431-b20e-84709a1ee8ef","year":2021},"citing_paper":{"arxiv_id":"2508.10009","last_updated":"2025-08-05T23:56:11Z","snapshot_observed_at":"2026-08-14T03:48:48.293793Z","submitted_at":"2025-08-05T23:56:11Z","title":"Beyond Hard Sharing: Efficient Multi-Task Speech-to-Text Modeling with Supervised Mixture of Experts","version":1},"reference_index":10,"source":"pdf_text","source_observed_at":"2026-08-06T01:05:44.579694Z"},"links":{"citing_paper":"/paper/2508.10009"},"observation_digest":"sha256:80802cbbfced2030ef0c7e71895890abdb0ff7cb0fe16ec0014e804c7ecb741b","observation_id":"3df0ff5c-f5c7-4ec8-9e47-2ad73d12ec50","resolution":{"observed_at":"2026-08-06T01:05:49.508010Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-20T06:33:59.587034+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-20T06:33:59.587034+00:00","source":"crossref"},{"observed_at":"2026-08-20T06:33:54.927442+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2204.07689","last_updated":"2022-04-16T00:56:12Z","snapshot_observed_at":"2026-08-16T17:06:37.440640Z","submitted_at":"2022-04-16T00:56:12Z","title":"Sparsely Activated Mixture-of-Experts are Robust Multi-Task Learners","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2204.07689","snapshot_observed_at":"2026-08-06T01:05:44.679323Z","title":"Sparsely activated mixture-of-experts are robust multi-task learners,","venue":null,"work_id":null,"year":2022},"citing_paper":{"arxiv_id":"2508.10009","last_updated":"2025-08-05T23:56:11Z","snapshot_observed_at":"2026-08-14T03:48:48.293793Z","submitted_at":"2025-08-05T23:56:11Z","title":"Beyond Hard Sharing: Efficient Multi-Task Speech-to-Text Modeling with Supervised Mixture of Experts","version":1},"reference_index":11,"source":"pdf_text","source_observed_at":"2026-08-06T01:05:44.679323Z"},"links":{"cited_paper":"/paper/2204.07689","citing_paper":"/paper/2508.10009"},"observation_digest":"sha256:53e6551387feb712bec69588a18e9245f4ef4d9a35ac3898c4dba8dbee8e3571","observation_id":"13ecb848-b814-4af7-853b-f9a0e5895d14","resolution":{"observed_at":"2026-08-06T01:05:44.679323Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"1701.06538","last_updated":"2017-01-23T18:10:00Z","snapshot_observed_at":"2026-08-19T12:32:00.517027Z","submitted_at":"2017-01-23T18:10:00Z","title":"Outrageously Large Neural Networks: The Sparsely-Gated Mixture-of-Experts Layer","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"1701.06538","snapshot_observed_at":"2026-08-06T01:05:44.775645Z","title":"Outrageously large neural networks: The sparsely-gated mixture-of-experts layer,","venue":null,"work_id":null,"year":2017},"citing_paper":{"arxiv_id":"2508.10009","last_updated":"2025-08-05T23:56:11Z","snapshot_observed_at":"2026-08-14T03:48:48.293793Z","submitted_at":"2025-08-05T23:56:11Z","title":"Beyond Hard Sharing: Efficient Multi-Task Speech-to-Text Modeling with Supervised Mixture of Experts","version":1},"reference_index":12,"source":"pdf_text","source_observed_at":"2026-08-06T01:05:44.775645Z"},"links":{"cited_paper":"/paper/1701.06538","citing_paper":"/paper/2508.10009"},"observation_digest":"sha256:08b5ccb95216c986433c15a89f83a8ae946849746be4a5264c2477a8f44b0625","observation_id":"b14f26aa-863c-4f9f-8fb7-d577a9f976d5","resolution":{"observed_at":"2026-08-06T01:05:44.775645Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T01:05:47.948496Z","title":"[Online]","venue":null,"work_id":"a8f7a88b-d274-4100-a97d-bc8559e93aae","year":2020},"citing_paper":{"arxiv_id":"2508.10009","last_updated":"2025-08-05T23:56:11Z","snapshot_observed_at":"2026-08-14T03:48:48.293793Z","submitted_at":"2025-08-05T23:56:11Z","title":"Beyond Hard Sharing: Efficient Multi-Task Speech-to-Text Modeling with Supervised Mixture of Experts","version":1},"reference_index":13,"source":"pdf_text","source_observed_at":"2026-08-06T01:05:45.460528Z"},"links":{"citing_paper":"/paper/2508.10009"},"observation_digest":"sha256:a1cc78b5bcab53e8d17af9e30077d4d078a498c072400394d3a6800e9c590825","observation_id":"acd40c53-6442-45f1-b63d-2d5296943386","resolution":{"observed_at":"2026-08-06T01:05:48.012664Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-20T06:33:59.587034+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-20T06:33:59.587034+00:00","source":"crossref"},{"observed_at":"2026-08-20T06:33:54.927442+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T01:05:47.778205Z","title":"[Online]","venue":null,"work_id":"f436f8d3-89cf-4582-bc19-58d06cd278eb","year":2022},"citing_paper":{"arxiv_id":"2508.10009","last_updated":"2025-08-05T23:56:11Z","snapshot_observed_at":"2026-08-14T03:48:48.293793Z","submitted_at":"2025-08-05T23:56:11Z","title":"Beyond Hard Sharing: Efficient Multi-Task Speech-to-Text Modeling with Supervised Mixture of Experts","version":1},"reference_index":14,"source":"pdf_text","source_observed_at":"2026-08-06T01:05:45.556818Z"},"links":{"citing_paper":"/paper/2508.10009"},"observation_digest":"sha256:d794a78bd3fb414b2854d30f3285de1d55a89465d92d04dd31aa18549dc28c00","observation_id":"1ffa1fea-0693-407b-934a-1c9022565ecd","resolution":{"observed_at":"2026-08-06T01:05:47.877416Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-20T06:33:59.587034+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-20T06:33:59.587034+00:00","source":"crossref"},{"observed_at":"2026-08-20T06:33:54.927442+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T01:05:48.974625Z","title":"[Online]","venue":null,"work_id":"25273d15-5a7b-42fe-a7bc-63888f0af054","year":2022},"citing_paper":{"arxiv_id":"2508.10009","last_updated":"2025-08-05T23:56:11Z","snapshot_observed_at":"2026-08-14T03:48:48.293793Z","submitted_at":"2025-08-05T23:56:11Z","title":"Beyond Hard Sharing: Efficient Multi-Task Speech-to-Text Modeling with Supervised Mixture of Experts","version":1},"reference_index":15,"source":"pdf_text","source_observed_at":"2026-08-06T01:05:45.073418Z"},"links":{"citing_paper":"/paper/2508.10009"},"observation_digest":"sha256:aefcc54ddbc1a267cb8aca77bbc1d8ff5e550cee540a0688ff5c357e2b241d9b","observation_id":"eaeb04fd-c4b1-4509-9e5f-27669a752e2d","resolution":{"observed_at":"2026-08-06T01:05:49.060196Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-20T06:33:59.587034+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-20T06:33:59.587034+00:00","source":"crossref"},{"observed_at":"2026-08-20T06:33:54.927442+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T01:05:48.757364Z","title":"[Online]","venue":null,"work_id":"c98c873b-665e-44c0-a686-021e72a894eb","year":2021},"citing_paper":{"arxiv_id":"2508.10009","last_updated":"2025-08-05T23:56:11Z","snapshot_observed_at":"2026-08-14T03:48:48.293793Z","submitted_at":"2025-08-05T23:56:11Z","title":"Beyond Hard Sharing: Efficient Multi-Task Speech-to-Text Modeling with Supervised Mixture of Experts","version":1},"reference_index":16,"source":"pdf_text","source_observed_at":"2026-08-06T01:05:45.170383Z"},"links":{"citing_paper":"/paper/2508.10009"},"observation_digest":"sha256:add2d90bf1c160bad221af62242488f7953b6d61441f66b965589f30139cbe5f","observation_id":"bed821b7-7dc2-48d8-8567-32de69b8f85a","resolution":{"observed_at":"2026-08-06T01:05:48.848835Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-20T06:33:59.587034+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-20T06:33:59.587034+00:00","source":"crossref"},{"observed_at":"2026-08-20T06:33:54.927442+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T01:05:48.518044Z","title":"[Online]","venue":null,"work_id":"4a1520de-e4d7-4ede-89b9-16ee2da0da01","year":2018},"citing_paper":{"arxiv_id":"2508.10009","last_updated":"2025-08-05T23:56:11Z","snapshot_observed_at":"2026-08-14T03:48:48.293793Z","submitted_at":"2025-08-05T23:56:11Z","title":"Beyond Hard Sharing: Efficient Multi-Task Speech-to-Text Modeling with Supervised Mixture of Experts","version":1},"reference_index":17,"source":"pdf_text","source_observed_at":"2026-08-06T01:05:45.264475Z"},"links":{"citing_paper":"/paper/2508.10009"},"observation_digest":"sha256:b8aa5f82533d48ee2622fc8adbc2597ea5bb8794845901807b74d14d116fb0bf","observation_id":"f46f5ba2-6f3d-4384-b665-a628fc6e8ef8","resolution":{"observed_at":"2026-08-06T01:05:48.642320Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-20T06:33:59.587034+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-20T06:33:59.587034+00:00","source":"crossref"},{"observed_at":"2026-08-20T06:33:54.927442+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T01:05:48.304732Z","title":"[Online]","venue":null,"work_id":"4c23a9f6-00a7-4ed7-be3e-17a9fe45980a","year":2021},"citing_paper":{"arxiv_id":"2508.10009","last_updated":"2025-08-05T23:56:11Z","snapshot_observed_at":"2026-08-14T03:48:48.293793Z","submitted_at":"2025-08-05T23:56:11Z","title":"Beyond Hard Sharing: Efficient Multi-Task Speech-to-Text Modeling with Supervised Mixture of Experts","version":1},"reference_index":18,"source":"pdf_text","source_observed_at":"2026-08-06T01:05:45.321574Z"},"links":{"citing_paper":"/paper/2508.10009"},"observation_digest":"sha256:40592bdbaeaa95d8b25a381507852f93e0add4980246cb7106853bcb13c45ef6","observation_id":"d5367c58-a785-4998-84df-386c5f577384","resolution":{"observed_at":"2026-08-06T01:05:48.440198Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-20T06:33:59.587034+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-20T06:33:59.587034+00:00","source":"crossref"},{"observed_at":"2026-08-20T06:33:54.927442+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T01:05:48.062483Z","title":"[Online]","venue":null,"work_id":"e096be50-f116-40f0-aada-0476fcf338c0","year":2022},"citing_paper":{"arxiv_id":"2508.10009","last_updated":"2025-08-05T23:56:11Z","snapshot_observed_at":"2026-08-14T03:48:48.293793Z","submitted_at":"2025-08-05T23:56:11Z","title":"Beyond Hard Sharing: Efficient Multi-Task Speech-to-Text Modeling with Supervised Mixture of Experts","version":1},"reference_index":19,"source":"pdf_text","source_observed_at":"2026-08-06T01:05:45.394744Z"},"links":{"citing_paper":"/paper/2508.10009"},"observation_digest":"sha256:62e99a2dc9370e93c2bdfcc2fc47b9280011a6841043c5f1ca127fe1fc8b6ad0","observation_id":"db6ba20b-e15a-44fb-8a18-ec018b8b582e","resolution":{"observed_at":"2026-08-06T01:05:48.164816Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-20T06:33:59.587034+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-20T06:33:59.587034+00:00","source":"crossref"},{"observed_at":"2026-08-20T06:33:54.927442+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T01:05:46.944251Z","title":"Pulse code modulation (pcm) of voice fre- quencies,","venue":null,"work_id":"be243061-2235-4d22-8619-e3b38370decf","year":1988},"citing_paper":{"arxiv_id":"2508.10009","last_updated":"2025-08-05T23:56:11Z","snapshot_observed_at":"2026-08-14T03:48:48.293793Z","submitted_at":"2025-08-05T23:56:11Z","title":"Beyond Hard Sharing: Efficient Multi-Task Speech-to-Text Modeling with Supervised Mixture of Experts","version":1},"reference_index":20,"source":"pdf_text","source_observed_at":"2026-08-06T01:05:46.071817Z"},"links":{"citing_paper":"/paper/2508.10009"},"observation_digest":"sha256:2101f2f79c7286944e79ca56a045b4b49204e45ae78a0fb3d5d4f3dd3b7ca6e7","observation_id":"072708d6-4e41-4654-a0ed-d3ab503c465a","resolution":{"observed_at":"2026-08-06T01:05:46.982149Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-20T06:33:59.587034+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-20T06:33:59.587034+00:00","source":"crossref"},{"observed_at":"2026-08-20T06:33:54.927442+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T01:05:50.097219Z","title":"Our in-house test sets consist of 1,000 samples of male and female speech from daily conversations, with refer- ence translations curated by professional translators","venue":null,"work_id":"b464c3d1-db46-4410-969f-c0ea84be9b3e","year":2048},"citing_paper":{"arxiv_id":"2508.10009","last_updated":"2025-08-05T23:56:11Z","snapshot_observed_at":"2026-08-14T03:48:48.293793Z","submitted_at":"2025-08-05T23:56:11Z","title":"Beyond Hard Sharing: Efficient Multi-Task Speech-to-Text Modeling with Supervised Mixture of Experts","version":1},"reference_index":21,"source":"pdf_text","source_observed_at":"2026-08-06T01:05:44.122596Z"},"links":{"citing_paper":"/paper/2508.10009"},"observation_digest":"sha256:b51dc374e4c3a140235b3f305b82c8d20bd82ae42ffeefe114dbf66685149903","observation_id":"33a791ee-bc8a-4d32-8f6d-292134b25192","resolution":{"observed_at":"2026-08-06T01:05:50.162275Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-20T06:33:59.587034+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-20T06:33:59.587034+00:00","source":"crossref"},{"observed_at":"2026-08-20T06:33:54.927442+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T01:05:47.546438Z","title":"[Online]","venue":null,"work_id":"200ee500-2b10-4ca3-ac97-142d7dd5123a","year":2022},"citing_paper":{"arxiv_id":"2508.10009","last_updated":"2025-08-05T23:56:11Z","snapshot_observed_at":"2026-08-14T03:48:48.293793Z","submitted_at":"2025-08-05T23:56:11Z","title":"Beyond Hard Sharing: Efficient Multi-Task Speech-to-Text Modeling with Supervised Mixture of Experts","version":1},"reference_index":22,"source":"pdf_text","source_observed_at":"2026-08-06T01:05:45.618663Z"},"links":{"citing_paper":"/paper/2508.10009"},"observation_digest":"sha256:5be11b138bc939b26dd2f016a0340769aca8b6a98a87d4bfb269e80a65839fa0","observation_id":"aa6474cd-aa62-4796-975f-820f642c40e4","resolution":{"observed_at":"2026-08-06T01:05:47.654187Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-20T06:33:59.587034+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-20T06:33:59.587034+00:00","source":"crossref"},{"observed_at":"2026-08-20T06:33:54.927442+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T01:05:47.391104Z","title":"[Online]","venue":null,"work_id":"0c77fc28-fbe4-4050-9a23-abbfef53948a","year":2021},"citing_paper":{"arxiv_id":"2508.10009","last_updated":"2025-08-05T23:56:11Z","snapshot_observed_at":"2026-08-14T03:48:48.293793Z","submitted_at":"2025-08-05T23:56:11Z","title":"Beyond Hard Sharing: Efficient Multi-Task Speech-to-Text Modeling with Supervised Mixture of Experts","version":1},"reference_index":23,"source":"pdf_text","source_observed_at":"2026-08-06T01:05:45.683072Z"},"links":{"citing_paper":"/paper/2508.10009"},"observation_digest":"sha256:1c445593e8fec0f492670884cc439ccfbb1ac390d675ffeb35a97c791d0f2645","observation_id":"a1f6a03c-ea88-490c-be24-e1e607f20541","resolution":{"observed_at":"2026-08-06T01:05:47.474698Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-20T06:33:59.587034+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-20T06:33:59.587034+00:00","source":"crossref"},{"observed_at":"2026-08-20T06:33:54.927442+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T01:05:49.781782Z","title":"Whisper is a widely used mul- tilingual speech model trained with diverse language pairs","venue":null,"work_id":"7842dfd8-8dc9-4574-a468-59bd0a51509d","year":null},"citing_paper":{"arxiv_id":"2508.10009","last_updated":"2025-08-05T23:56:11Z","snapshot_observed_at":"2026-08-14T03:48:48.293793Z","submitted_at":"2025-08-05T23:56:11Z","title":"Beyond Hard Sharing: Efficient Multi-Task Speech-to-Text Modeling with Supervised Mixture of Experts","version":1},"reference_index":24,"source":"pdf_text","source_observed_at":"2026-08-06T01:05:44.239321Z"},"links":{"citing_paper":"/paper/2508.10009"},"observation_digest":"sha256:0517dcd937216358730d8c288b72fc0af5a9715ad06b32ca245b248aff3d9408","observation_id":"38a3220f-17f0-4c27-8d6f-3a9caec67139","resolution":{"observed_at":"2026-08-06T01:05:49.880492Z","resolver_source":"raw_fallback","status":"malformed_identifier"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-20T06:33:59.587034+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-20T06:33:59.587034+00:00","source":"crossref"},{"observed_at":"2026-08-20T06:33:54.927442+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T01:05:47.198996Z","title":"[Online]","venue":null,"work_id":"2446631d-1459-4344-8b40-4d6b85b1735c","year":2021},"citing_paper":{"arxiv_id":"2508.10009","last_updated":"2025-08-05T23:56:11Z","snapshot_observed_at":"2026-08-14T03:48:48.293793Z","submitted_at":"2025-08-05T23:56:11Z","title":"Beyond Hard Sharing: Efficient Multi-Task Speech-to-Text Modeling with Supervised Mixture of Experts","version":1},"reference_index":25,"source":"pdf_text","source_observed_at":"2026-08-06T01:05:45.783547Z"},"links":{"citing_paper":"/paper/2508.10009"},"observation_digest":"sha256:811b88610ec2999a0c66396c1cb5ce699297227ff51f639442009f11ed6537d1","observation_id":"49206302-ec76-4c0c-8fc2-e0053d1bb54b","resolution":{"observed_at":"2026-08-06T01:05:47.280919Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-20T06:33:59.587034+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-20T06:33:59.587034+00:00","source":"crossref"},{"observed_at":"2026-08-20T06:33:54.927442+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T01:05:47.067287Z","title":"[Online]","venue":null,"work_id":"9c4375f1-b783-46b7-8a2a-20e39da04106","year":2022},"citing_paper":{"arxiv_id":"2508.10009","last_updated":"2025-08-05T23:56:11Z","snapshot_observed_at":"2026-08-14T03:48:48.293793Z","submitted_at":"2025-08-05T23:56:11Z","title":"Beyond Hard Sharing: Efficient Multi-Task Speech-to-Text Modeling with Supervised Mixture of Experts","version":1},"reference_index":26,"source":"pdf_text","source_observed_at":"2026-08-06T01:05:45.863666Z"},"links":{"citing_paper":"/paper/2508.10009"},"observation_digest":"sha256:783e57a016e5707d76d893fe54da80b5e62abdc254c4f5f8c3012aecec16d14c","observation_id":"c4cad7bc-fb68-44ec-99b9-fde756a834a2","resolution":{"observed_at":"2026-08-06T01:05:47.116778Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-20T06:33:59.587034+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-20T06:33:59.587034+00:00","source":"crossref"},{"observed_at":"2026-08-20T06:33:54.927442+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T01:05:45.952971Z","title":"The adaptive multirate wideband speech codec (amr-wb),","venue":null,"work_id":null,"year":2002},"citing_paper":{"arxiv_id":"2508.10009","last_updated":"2025-08-05T23:56:11Z","snapshot_observed_at":"2026-08-14T03:48:48.293793Z","submitted_at":"2025-08-05T23:56:11Z","title":"Beyond Hard Sharing: Efficient Multi-Task Speech-to-Text Modeling with Supervised Mixture of Experts","version":1},"reference_index":27,"source":"pdf_text","source_observed_at":"2026-08-06T01:05:45.952971Z"},"links":{"citing_paper":"/paper/2508.10009"},"observation_digest":"sha256:ebb0c73fc0e1c8c21e1b84a11dd344a739bc8454743fc6381f120166d84a5918","observation_id":"e2fefd12-a43e-47c1-9e49-927204004dd2","resolution":{"observed_at":"2026-08-06T01:05:45.952971Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T01:05:46.175496Z","title":"Fleurs: Few-shot learning evaluation of universal representations of speech,","venue":null,"work_id":null,"year":2022},"citing_paper":{"arxiv_id":"2508.10009","last_updated":"2025-08-05T23:56:11Z","snapshot_observed_at":"2026-08-14T03:48:48.293793Z","submitted_at":"2025-08-05T23:56:11Z","title":"Beyond Hard Sharing: Efficient Multi-Task Speech-to-Text Modeling with Supervised Mixture of Experts","version":1},"reference_index":28,"source":"pdf_text","source_observed_at":"2026-08-06T01:05:46.175496Z"},"links":{"citing_paper":"/paper/2508.10009"},"observation_digest":"sha256:38b75aa6f8300a0c2da6aa579a357684a8130f0a613a9d1098d527916345fbae","observation_id":"ad761f03-d506-4a96-bd8f-686fc5e454ef","resolution":{"observed_at":"2026-08-06T01:05:46.175496Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T01:05:46.789354Z","title":"Attention is all you need,","venue":null,"work_id":"1f3d24f6-c76b-486f-a973-bb7aa3188ec5","year":2017},"citing_paper":{"arxiv_id":"2508.10009","last_updated":"2025-08-05T23:56:11Z","snapshot_observed_at":"2026-08-14T03:48:48.293793Z","submitted_at":"2025-08-05T23:56:11Z","title":"Beyond Hard Sharing: Efficient Multi-Task Speech-to-Text Modeling with Supervised Mixture of Experts","version":1},"reference_index":29,"source":"pdf_text","source_observed_at":"2026-08-06T01:05:46.281915Z"},"links":{"citing_paper":"/paper/2508.10009"},"observation_digest":"sha256:63ef78fda9600aea1d1a3c5ae687092c11d4bf8098d095b0b248f917e6595c3e","observation_id":"29b34952-b47b-4dee-9584-9a971a3bf21d","resolution":{"observed_at":"2026-08-06T01:05:46.876936Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-20T06:33:59.587034+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-20T06:33:59.587034+00:00","source":"crossref"},{"observed_at":"2026-08-20T06:33:54.927442+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T01:05:46.369512Z","title":"Bleu: a method for automatic evaluation of machine translation,","venue":null,"work_id":null,"year":2002},"citing_paper":{"arxiv_id":"2508.10009","last_updated":"2025-08-05T23:56:11Z","snapshot_observed_at":"2026-08-14T03:48:48.293793Z","submitted_at":"2025-08-05T23:56:11Z","title":"Beyond Hard Sharing: Efficient Multi-Task Speech-to-Text Modeling with Supervised Mixture of Experts","version":1},"reference_index":30,"source":"pdf_text","source_observed_at":"2026-08-06T01:05:46.369512Z"},"links":{"citing_paper":"/paper/2508.10009"},"observation_digest":"sha256:b92639dea274046f5422404691e063e46967426b62a83f2d927d938c66e099f0","observation_id":"bd80abc3-b9aa-49ce-9548-76222386d705","resolution":{"observed_at":"2026-08-06T01:05:46.369512Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T01:05:46.447560Z","title":"Robust speech recognition via large-scale weak supervision,","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2508.10009","last_updated":"2025-08-05T23:56:11Z","snapshot_observed_at":"2026-08-14T03:48:48.293793Z","submitted_at":"2025-08-05T23:56:11Z","title":"Beyond Hard Sharing: Efficient Multi-Task Speech-to-Text Modeling with Supervised Mixture of Experts","version":1},"reference_index":31,"source":"pdf_text","source_observed_at":"2026-08-06T01:05:46.447560Z"},"links":{"citing_paper":"/paper/2508.10009"},"observation_digest":"sha256:1f3ea01ea58d16e2c082087bf3197e46cde338f6e21c11f641dd6e3e4269c949","observation_id":"2be105f0-24c8-4d44-b516-90f8e82b548b","resolution":{"observed_at":"2026-08-06T01:05:46.447560Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}}],"paper":{"arxiv_id":"2508.10009","last_updated":"2025-08-05T23:56:11Z","latest_version":1,"primary_category":"cs.CL","snapshot_observed_at":"2026-08-14T03:48:48.293793Z","submitted_at":"2025-08-05T23:56:11Z","title":"Beyond Hard Sharing: Efficient Multi-Task Speech-to-Text Modeling with Supervised Mixture of Experts"},"reference_resolution":{"displayed":31,"state_counts":{"malformed_identifier":1,"metadata_mismatch":1,"parse_uncertain":0,"unresolved":8,"verified_exact":0,"verified_fuzzy":21},"total_outbound_references":31},"refusal":"A citation records a reference. It does not transfer a finding from one paper to another.","schema":"pith.paper-citation-record.v1","standing_sources":[{"observed_at":"2026-08-20T06:33:59.587034+00:00","source":"crossref"},{"observed_at":"2026-08-20T06:33:54.927442+00:00","source":"retraction_watch"}],"thesis":"As of 20 August 2026, this Paper Citation Record lists 31 of 31 outbound references and 1 inbound Pith citation observation for arXiv:2508.10009."}