{"as_of":"2026-08-10T01:18:00Z","caps":{"database_statements":6,"inbound":100,"outbound":100},"context_digest":"sha256:4b5ccbdfbbaf330b8b4026db119f98a6d51d255bbf6c48df45b245c34803cf1b","coverage":[{"denominator":18,"lane":"reference_resolution","note":"Typed states for the displayed outbound observations.","records_observed":18,"source":"paper_references, paper_reference_links","source_observed_at":"2026-08-08T12:19:08.006769Z","state":"measured"},{"denominator":19,"lane":"standing_notices","note":"One-hop event checks from named stored sources.","records_observed":19,"source":"scholarly_work_events, retraction_status_cache","source_observed_at":"2026-08-09T06:31:02.800959+00:00","state":"measured"},{"denominator":1,"lane":"inbound_itemization","note":"Pith citing papers itemized under the disclosed page cap.","records_observed":1,"source":"paper_references, paper_reference_links","source_observed_at":"2026-08-06T23:10:57.314825Z","state":"measured"},{"denominator":1,"lane":"external_citation_measurements","note":"A source-named dated measurement, never combined with another source.","records_observed":0,"source":"pith","source_observed_at":"2026-08-06T23:10:57.510789Z","state":"measured"}],"external_citation_measurements":[],"inbound":[{"citation":{"cited_paper":{"arxiv_id":"2502.07575","last_updated":"2025-02-21T04:10:54Z","snapshot_observed_at":"2026-08-08T12:13:15.265253Z","submitted_at":"2025-02-11T14:17:29Z","title":"Towards Efficient and Multifaceted Computer-assisted Pronunciation Training Leveraging Hierarchical Selective State Space Model and Decoupled Cross-entropy Loss","version":2},"cited_work":{"arxiv_id":"2502.07575","doi":null,"metadata_source":"pith","pith_arxiv_id":"2502.07575","snapshot_observed_at":"2026-08-06T23:10:57.510789Z","title":"Towards Efficient and Multifaceted Computer-assisted Pronunciation Training Leveraging Hierarchical Selective State Space Model and Decoupled Cross-entropy Loss","venue":"eess.AS","work_id":"9871365c-8156-4032-9418-ee2202a7c25b","year":2025},"citing_paper":{"arxiv_id":"2506.19315","last_updated":"2025-07-25T09:26:59Z","snapshot_observed_at":"2026-08-10T00:19:16.139069Z","submitted_at":"2025-06-24T05:12:32Z","title":"JCAPT: A Joint Modeling Approach for CAPT","version":2},"reference_index":16,"source":"pdf_text","source_observed_at":"2026-08-06T23:10:57.314825Z"},"links":{"cited_paper":"/paper/2502.07575","citing_paper":"/paper/2506.19315"},"observation_digest":"sha256:a47875d696f41a57fe0543bb1758b0a58f613b927e1e300029fa159c675675a2","observation_id":"a2f6d2ec-6c8c-4498-9198-851eefda73ac","resolution":{"observed_at":"2026-08-06T23:10:57.516689Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}}],"links":{"evidence":"/evidence","html":"/paper/2502.07575/citation-record","integrity":"/paper/2502.07575/integrity","json":"/paper/2502.07575/citation-record.json","paper":"/paper/2502.07575"},"outbound":[{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-08T12:19:08.387728Z","title":"However, it is generally expected that a full-fledged CAPT system should perform both functionalities simultaneously and efficiently","venue":null,"work_id":"782a2c82-de48-4ed1-b069-a98d3ed70718","year":2009},"citing_paper":{"arxiv_id":"2502.07575","last_updated":"2025-02-21T04:10:54Z","snapshot_observed_at":"2026-08-08T12:13:15.265253Z","submitted_at":"2025-02-11T14:17:29Z","title":"Towards Efficient and Multifaceted Computer-assisted Pronunciation Training Leveraging Hierarchical Selective State Space Model and Decoupled Cross-entropy Loss","version":2},"reference_index":1,"source":"pdf_text","source_observed_at":"2026-08-08T12:19:07.921568Z"},"links":{"citing_paper":"/paper/2502.07575"},"observation_digest":"sha256:84de8deac3fee1cbeacf4c6b4a357d053dcfaf3df250f68e7905e0604ac44ccf","observation_id":"f2126658-2576-4a79-b045-59896f80d958","resolution":{"observed_at":"2026-08-08T12:19:08.392519Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-08T12:19:08.342023Z","title":"These modules collectively generate the corresponding aspect score sequence 𝐬𝑔 for each linguistic granularity 𝑔, as well as the phonetic error states 𝐞 and diagnosis 𝐲","venue":null,"work_id":"4c502535-919c-4507-98b4-e2d03b412b02","year":2000},"citing_paper":{"arxiv_id":"2502.07575","last_updated":"2025-02-21T04:10:54Z","snapshot_observed_at":"2026-08-08T12:13:15.265253Z","submitted_at":"2025-02-11T14:17:29Z","title":"Towards Efficient and Multifaceted Computer-assisted Pronunciation Training Leveraging Hierarchical Selective State Space Model and Decoupled Cross-entropy Loss","version":2},"reference_index":2,"source":"pdf_text","source_observed_at":"2026-08-08T12:19:07.937362Z"},"links":{"citing_paper":"/paper/2502.07575"},"observation_digest":"sha256:336d32eb0991415b9ce564301ae10a390d4ce36a50fdd04e775a6c592efe8d3e","observation_id":"5f041df5-8e2a-46a8-b284-765efafd6e7c","resolution":{"observed_at":"2026-08-08T12:19:08.347525Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-08T12:19:08.356577Z","title":"These errors usually have clear-cut distinctions between correct and incorrect ones, and can be easily quantified through deletions, substitutions, and insertions","venue":null,"work_id":"f1abde2a-9597-43e9-b6fb-7d122ed9be68","year":2021},"citing_paper":{"arxiv_id":"2502.07575","last_updated":"2025-02-21T04:10:54Z","snapshot_observed_at":"2026-08-08T12:13:15.265253Z","submitted_at":"2025-02-11T14:17:29Z","title":"Towards Efficient and Multifaceted Computer-assisted Pronunciation Training Leveraging Hierarchical Selective State Space Model and Decoupled Cross-entropy Loss","version":2},"reference_index":3,"source":"pdf_text","source_observed_at":"2026-08-08T12:19:07.932079Z"},"links":{"citing_paper":"/paper/2502.07575"},"observation_digest":"sha256:e99d829095e2826fd1981c708f2b7bd3eaf0cc09c18e48f1c31d6b052722839f","observation_id":"d43c3eba-9f01-4768-8695-0434ed7c6a80","resolution":{"observed_at":"2026-08-08T12:19:08.362769Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-08T12:19:08.328478Z","title":"Notably, there are several studies investigating the bidirectional processing of Mamba (Liang et al., 2024; Zhang et al., 2024; Jiang et al., 2024)","venue":null,"work_id":"65714285-5b0a-44b6-aec0-eecd7378664c","year":2024},"citing_paper":{"arxiv_id":"2502.07575","last_updated":"2025-02-21T04:10:54Z","snapshot_observed_at":"2026-08-08T12:13:15.265253Z","submitted_at":"2025-02-11T14:17:29Z","title":"Towards Efficient and Multifaceted Computer-assisted Pronunciation Training Leveraging Hierarchical Selective State Space Model and Decoupled Cross-entropy Loss","version":2},"reference_index":5,"source":"pdf_text","source_observed_at":"2026-08-08T12:19:07.942760Z"},"links":{"citing_paper":"/paper/2502.07575"},"observation_digest":"sha256:14d7a096807d9e725ebcecbe57c2416252641bc8a6229ad7da22af12ae099106","observation_id":"d44e73af-d10a-4229-8d33-c02c717ac438","resolution":{"observed_at":"2026-08-08T12:19:08.332810Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-08T12:19:08.298517Z","title":"The APA module contains one regressor that aims to predict the phone-level aspect score 𝑠0𝑔𝑝ℎ𝑛(accuracy)","venue":null,"work_id":"3f29bedf-e568-4276-8568-a446b79aa06d","year":2022},"citing_paper":{"arxiv_id":"2502.07575","last_updated":"2025-02-21T04:10:54Z","snapshot_observed_at":"2026-08-08T12:13:15.265253Z","submitted_at":"2025-02-11T14:17:29Z","title":"Towards Efficient and Multifaceted Computer-assisted Pronunciation Training Leveraging Hierarchical Selective State Space Model and Decoupled Cross-entropy Loss","version":2},"reference_index":7,"source":"pdf_text","source_observed_at":"2026-08-08T12:19:07.953380Z"},"links":{"citing_paper":"/paper/2502.07575"},"observation_digest":"sha256:ab507efe515e0ed983eb8fd0a445256ebb7d367c70512046733865d1aa402396","observation_id":"7a149287-4e5f-4743-8832-f9df68db3bc1","resolution":{"observed_at":"2026-08-08T12:19:08.303935Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2312.00752","last_updated":"2024-05-31T17:55:27Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2023-12-01T18:01:34Z","title":"Mamba: Linear-Time Sequence Modeling with Selective State Spaces","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2312.00752","snapshot_observed_at":"2026-08-08T12:19:07.976059Z","title":"arXiv preprint arXiv:2312.00752","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2502.07575","last_updated":"2025-02-21T04:10:54Z","snapshot_observed_at":"2026-08-08T12:13:15.265253Z","submitted_at":"2025-02-11T14:17:29Z","title":"Towards Efficient and Multifaceted Computer-assisted Pronunciation Training Leveraging Hierarchical Selective State Space Model and Decoupled Cross-entropy Loss","version":2},"reference_index":12,"source":"pdf_text","source_observed_at":"2026-08-08T12:19:07.976059Z"},"links":{"cited_paper":"/paper/2312.00752","citing_paper":"/paper/2502.07575"},"observation_digest":"sha256:c4f0f6c6cf156987bfcd537e91527bc990cf52dd1e3e9907e9853b04a7dc8f83","observation_id":"d33575be-4913-41ca-b1c2-f55852cc5c33","resolution":{"observed_at":"2026-08-08T12:19:07.976059Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2404.15772","last_updated":"2024-06-27T03:31:25Z","snapshot_observed_at":"2026-08-06T08:19:23.367190Z","submitted_at":"2024-04-24T09:45:48Z","title":"Bi-Mamba+: Bidirectional Mamba for Time Series Forecasting","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2404.15772","snapshot_observed_at":"2026-08-08T12:19:07.985835Z","title":"arXiv preprint arXiv:2404.15772","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2502.07575","last_updated":"2025-02-21T04:10:54Z","snapshot_observed_at":"2026-08-08T12:13:15.265253Z","submitted_at":"2025-02-11T14:17:29Z","title":"Towards Efficient and Multifaceted Computer-assisted Pronunciation Training Leveraging Hierarchical Selective State Space Model and Decoupled Cross-entropy Loss","version":2},"reference_index":14,"source":"pdf_text","source_observed_at":"2026-08-08T12:19:07.985835Z"},"links":{"cited_paper":"/paper/2404.15772","citing_paper":"/paper/2502.07575"},"observation_digest":"sha256:04a8736fcc3713bbbdeef41db45a29ef9cd5038551f6d1f85cc2461b93ee6900","observation_id":"d03425cd-6d26-4c28-b4e7-88107123665a","resolution":{"observed_at":"2026-08-08T12:19:07.985835Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2401.09417","last_updated":"2024-11-14T02:00:33Z","snapshot_observed_at":"2026-07-06T17:16:59.193820Z","submitted_at":"2024-01-17T18:56:18Z","title":"Vision Mamba: Efficient Visual Representation Learning with Bidirectional State Space Model","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2401.09417","snapshot_observed_at":"2026-08-08T12:19:07.990806Z","title":"arXiv preprint arXiv:2401.09417","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2502.07575","last_updated":"2025-02-21T04:10:54Z","snapshot_observed_at":"2026-08-08T12:13:15.265253Z","submitted_at":"2025-02-11T14:17:29Z","title":"Towards Efficient and Multifaceted Computer-assisted Pronunciation Training Leveraging Hierarchical Selective State Space Model and Decoupled Cross-entropy Loss","version":2},"reference_index":15,"source":"pdf_text","source_observed_at":"2026-08-08T12:19:07.990806Z"},"links":{"cited_paper":"/paper/2401.09417","citing_paper":"/paper/2502.07575"},"observation_digest":"sha256:5816137297431a9841e874e32c1d04b98bc540da9db26674c5de2b9c1330461f","observation_id":"0d299bed-cb26-457e-b369-8990e051a3ef","resolution":{"observed_at":"2026-08-08T12:19:07.990806Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2405.12609","last_updated":"2025-04-27T05:17:15Z","snapshot_observed_at":"2026-07-06T18:17:12.404441Z","submitted_at":"2024-05-21T09:04:48Z","title":"Mamba in Speech: Towards an Alternative to Self-Attention","version":6},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2405.12609","snapshot_observed_at":"2026-08-08T12:19:07.995703Z","title":"arXiv preprint arXiv:2405.12609","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2502.07575","last_updated":"2025-02-21T04:10:54Z","snapshot_observed_at":"2026-08-08T12:13:15.265253Z","submitted_at":"2025-02-11T14:17:29Z","title":"Towards Efficient and Multifaceted Computer-assisted Pronunciation Training Leveraging Hierarchical Selective State Space Model and Decoupled Cross-entropy Loss","version":2},"reference_index":16,"source":"pdf_text","source_observed_at":"2026-08-08T12:19:07.995703Z"},"links":{"cited_paper":"/paper/2405.12609","citing_paper":"/paper/2502.07575"},"observation_digest":"sha256:a1ba8760fbd7bf15323676496ecbdc36f79423f1ea8b0d3a9825753c868f310e","observation_id":"c6c79692-9b67-4d9b-80da-5f9a76193c3f","resolution":{"observed_at":"2026-08-08T12:19:07.995703Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-08T12:19:08.236550Z","title":"The combining weights 𝜔𝑔 for APA loss are uniformly set to 1.0 for each granularity level 𝑔","venue":null,"work_id":"3a3b4c92-0a06-4084-990a-a979fe9c6bec","year":2022},"citing_paper":{"arxiv_id":"2502.07575","last_updated":"2025-02-21T04:10:54Z","snapshot_observed_at":"2026-08-08T12:13:15.265253Z","submitted_at":"2025-02-11T14:17:29Z","title":"Towards Efficient and Multifaceted Computer-assisted Pronunciation Training Leveraging Hierarchical Selective State Space Model and Decoupled Cross-entropy Loss","version":2},"reference_index":17,"source":"pdf_text","source_observed_at":"2026-08-08T12:19:08.001187Z"},"links":{"citing_paper":"/paper/2502.07575"},"observation_digest":"sha256:68a2654a8fc8e0c9a45e371625137e20e8ff728316d46e3a7be17e9606f8b1de","observation_id":"fc9f1182-ca35-4b55-82ad-d3eb0038fcfd","resolution":{"observed_at":"2026-08-08T12:19:08.241622Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-08T12:19:08.220710Z","title":null,"venue":null,"work_id":"61c813a3-d803-49ce-ab5b-564942607289","year":2024},"citing_paper":{"arxiv_id":"2502.07575","last_updated":"2025-02-21T04:10:54Z","snapshot_observed_at":"2026-08-08T12:13:15.265253Z","submitted_at":"2025-02-11T14:17:29Z","title":"Towards Efficient and Multifaceted Computer-assisted Pronunciation Training Leveraging Hierarchical Selective State Space Model and Decoupled Cross-entropy Loss","version":2},"reference_index":18,"source":"pdf_text","source_observed_at":"2026-08-08T12:19:08.006769Z"},"links":{"citing_paper":"/paper/2502.07575"},"observation_digest":"sha256:e3f9ca458ccad3bf1349ebab86141ec9dce662b1a02bde28d82d094111ceb4cc","observation_id":"ad280e6f-533a-4f2b-b9a7-04ed58894fec","resolution":{"observed_at":"2026-08-08T12:19:08.226328Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-08T12:19:08.372995Z","title":"reading-aloud","venue":null,"work_id":"f1714f11-7ebe-4b2b-b960-e98493866cf3","year":2023},"citing_paper":{"arxiv_id":"2502.07575","last_updated":"2025-02-21T04:10:54Z","snapshot_observed_at":"2026-08-08T12:13:15.265253Z","submitted_at":"2025-02-11T14:17:29Z","title":"Towards Efficient and Multifaceted Computer-assisted Pronunciation Training Leveraging Hierarchical Selective State Space Model and Decoupled Cross-entropy Loss","version":2},"reference_index":2009,"source":"pdf_text","source_observed_at":"2026-08-08T12:19:07.927281Z"},"links":{"citing_paper":"/paper/2502.07575"},"observation_digest":"sha256:197ad6498ea4b74af3caf40f8acd66ecc0fdd6cde0c6f0a1d8022648a083a373","observation_id":"0113c96c-b848-4660-8370-04299622ba88","resolution":{"observed_at":"2026-08-08T12:19:08.377896Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-08T12:19:08.251371Z","title":"ETS Research Report Series 2015(1):1–11","venue":null,"work_id":"b0c95b30-90cb-4174-831b-e671b34f2abe","year":2015},"citing_paper":{"arxiv_id":"2502.07575","last_updated":"2025-02-21T04:10:54Z","snapshot_observed_at":"2026-08-08T12:13:15.265253Z","submitted_at":"2025-02-11T14:17:29Z","title":"Towards Efficient and Multifaceted Computer-assisted Pronunciation Training Leveraging Hierarchical Selective State Space Model and Decoupled Cross-entropy Loss","version":2},"reference_index":2015,"source":"pdf_text","source_observed_at":"2026-08-08T12:19:07.967573Z"},"links":{"citing_paper":"/paper/2502.07575"},"observation_digest":"sha256:8a4115363420dd10cd9ef7a41e04f52814fc577c7f7844cd9cbe253a8f6bf9ae","observation_id":"79022266-41fe-4f0e-8d26-41c3171136ea","resolution":{"observed_at":"2026-08-08T12:19:08.256312Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-08T12:19:08.282750Z","title":null,"venue":null,"work_id":"0e702a40-2000-47c3-835b-d9efa7eb0b40","year":2021},"citing_paper":{"arxiv_id":"2502.07575","last_updated":"2025-02-21T04:10:54Z","snapshot_observed_at":"2026-08-08T12:13:15.265253Z","submitted_at":"2025-02-11T14:17:29Z","title":"Towards Efficient and Multifaceted Computer-assisted Pronunciation Training Leveraging Hierarchical Selective State Space Model and Decoupled Cross-entropy Loss","version":2},"reference_index":2017,"source":"pdf_text","source_observed_at":"2026-08-08T12:19:07.958489Z"},"links":{"citing_paper":"/paper/2502.07575"},"observation_digest":"sha256:d0eb56c66c117e596d8c1a074bae420352d80bdbbbaaa524a089ca03ba3e0700","observation_id":"e3ff5edd-324f-47b2-b4ca-24e708ea7001","resolution":{"observed_at":"2026-08-08T12:19:08.287512Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2104.08428","last_updated":"2021-04-17T03:11:41Z","snapshot_observed_at":"2026-07-06T11:00:50.937337Z","submitted_at":"2021-04-17T03:11:41Z","title":"A Full Text-Dependent End to End Mispronunciation Detection and Diagnosis with Easy Data Augmentation Techniques","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2104.08428","snapshot_observed_at":"2026-08-08T12:19:07.971301Z","title":"arXiv preprint arXiv:2104.08428","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2502.07575","last_updated":"2025-02-21T04:10:54Z","snapshot_observed_at":"2026-08-08T12:13:15.265253Z","submitted_at":"2025-02-11T14:17:29Z","title":"Towards Efficient and Multifaceted Computer-assisted Pronunciation Training Leveraging Hierarchical Selective State Space Model and Decoupled Cross-entropy Loss","version":2},"reference_index":2021,"source":"pdf_text","source_observed_at":"2026-08-08T12:19:07.971301Z"},"links":{"cited_paper":"/paper/2104.08428","citing_paper":"/paper/2502.07575"},"observation_digest":"sha256:a573d08114118f847d79f38ec335c00cfe508022f6e23b85d70a9dd7d8715cf2","observation_id":"a861861f-8133-43c8-b8a7-bc5772cc47c0","resolution":{"observed_at":"2026-08-08T12:19:07.971301Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-08T12:19:08.267038Z","title":null,"venue":null,"work_id":"1764a459-5ac7-42a3-9b5f-8840a5fdc173","year":2023},"citing_paper":{"arxiv_id":"2502.07575","last_updated":"2025-02-21T04:10:54Z","snapshot_observed_at":"2026-08-08T12:13:15.265253Z","submitted_at":"2025-02-11T14:17:29Z","title":"Towards Efficient and Multifaceted Computer-assisted Pronunciation Training Leveraging Hierarchical Selective State Space Model and Decoupled Cross-entropy Loss","version":2},"reference_index":2022,"source":"pdf_text","source_observed_at":"2026-08-08T12:19:07.963458Z"},"links":{"citing_paper":"/paper/2502.07575"},"observation_digest":"sha256:ffcab167564c05d2dc56f1dad22f45d97f6413abc0b7c4b45e4e30803e3e704b","observation_id":"ef42ce8a-df5a-4f20-9090-d87136f429f4","resolution":{"observed_at":"2026-08-08T12:19:08.272001Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-08T12:19:08.314585Z","title":null,"venue":null,"work_id":"46a8e6c4-df5a-4153-bc7a-4e055e14b641","year":2022},"citing_paper":{"arxiv_id":"2502.07575","last_updated":"2025-02-21T04:10:54Z","snapshot_observed_at":"2026-08-08T12:13:15.265253Z","submitted_at":"2025-02-11T14:17:29Z","title":"Towards Efficient and Multifaceted Computer-assisted Pronunciation Training Leveraging Hierarchical Selective State Space Model and Decoupled Cross-entropy Loss","version":2},"reference_index":2023,"source":"pdf_text","source_observed_at":"2026-08-08T12:19:07.948120Z"},"links":{"citing_paper":"/paper/2502.07575"},"observation_digest":"sha256:bb67a6b22b8396d627a3fabf44a068c460d36422477783e1cd60478ad025f321","observation_id":"399c6f3b-d3df-4dd3-b4a5-4f9650e4de89","resolution":{"observed_at":"2026-08-08T12:19:08.319060Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2403.18257","last_updated":"2024-05-01T00:36:13Z","snapshot_observed_at":"2026-08-05T13:52:45.580279Z","submitted_at":"2024-03-27T05:00:08Z","title":"Dual-path Mamba: Short and Long-term Bidirectional Selective Structured State Space Models for Speech Separation","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2403.18257","snapshot_observed_at":"2026-08-08T12:19:07.980987Z","title":"Yassine Kheir, Ahmed Ali, and Shammur Chowdhury","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2502.07575","last_updated":"2025-02-21T04:10:54Z","snapshot_observed_at":"2026-08-08T12:13:15.265253Z","submitted_at":"2025-02-11T14:17:29Z","title":"Towards Efficient and Multifaceted Computer-assisted Pronunciation Training Leveraging Hierarchical Selective State Space Model and Decoupled Cross-entropy Loss","version":2},"reference_index":2024,"source":"pdf_text","source_observed_at":"2026-08-08T12:19:07.980987Z"},"links":{"cited_paper":"/paper/2403.18257","citing_paper":"/paper/2502.07575"},"observation_digest":"sha256:58d8cbedeeead4ce5ab124ee8ebc82565b0f6b89d54fb8634d003d9f4dc07076","observation_id":"f86411e5-577b-4b4b-ad97-12f63edb952f","resolution":{"observed_at":"2026-08-08T12:19:07.980987Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}}],"paper":{"arxiv_id":"2502.07575","last_updated":"2025-02-21T04:10:54Z","latest_version":2,"primary_category":"eess.AS","snapshot_observed_at":"2026-08-08T12:13:15.265253Z","submitted_at":"2025-02-11T14:17:29Z","title":"Towards Efficient and Multifaceted Computer-assisted Pronunciation Training Leveraging Hierarchical Selective State Space Model and Decoupled Cross-entropy Loss"},"reference_resolution":{"displayed":18,"state_counts":{"malformed_identifier":0,"metadata_mismatch":0,"parse_uncertain":0,"unresolved":10,"verified_exact":0,"verified_fuzzy":8},"total_outbound_references":18},"refusal":"A citation records a reference. It does not transfer a finding from one paper to another.","schema":"pith.paper-citation-record.v1","standing_sources":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"thesis":"As of 10 August 2026, this Paper Citation Record lists 18 of 18 outbound references and 1 inbound Pith citation observation for arXiv:2502.07575."}