{"as_of":"2026-08-11T15:07:00Z","caps":{"database_statements":6,"inbound":100,"outbound":100},"context_digest":"sha256:856c3ff766b3972c197adb52632db27de914adf3839cb48fa94a0fbb3e49a034","coverage":[{"denominator":37,"lane":"reference_resolution","note":"Typed states for the displayed outbound observations.","records_observed":37,"source":"paper_references, paper_reference_links","source_observed_at":"2026-08-07T04:31:33.099103Z","state":"measured"},{"denominator":37,"lane":"standing_notices","note":"One-hop event checks from named stored sources.","records_observed":37,"source":"scholarly_work_events, retraction_status_cache","source_observed_at":"2026-08-11T06:34:44.6726+00:00","state":"measured"},{"denominator":0,"lane":"inbound_itemization","note":"Pith citing papers itemized under the disclosed page cap.","records_observed":0,"source":"paper_references, paper_reference_links","source_observed_at":null,"state":"measured"},{"denominator":1,"lane":"external_citation_measurements","note":"A source-named dated measurement, never combined with another source.","records_observed":0,"source":"cited_works","source_observed_at":null,"state":"measured"}],"external_citation_measurements":[],"inbound":[],"links":{"evidence":"/evidence","html":"/paper/2506.21563/citation-record","integrity":"/paper/2506.21563/integrity","json":"/paper/2506.21563/citation-record.json","paper":"/paper/2506.21563"},"outbound":[{"citation":{"cited_paper":{"arxiv_id":"2406.03368","last_updated":"2025-01-23T17:57:28Z","snapshot_observed_at":"2026-08-08T14:58:02.421377Z","submitted_at":"2024-06-05T15:23:08Z","title":"IrokoBench: A New Benchmark for African Languages in the Age of Large Language Models","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2406.03368","snapshot_observed_at":"2026-08-07T04:31:30.363100Z","title":null,"venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2506.21563","last_updated":"2025-06-12T07:02:28Z","snapshot_observed_at":"2026-08-08T09:20:55.054492Z","submitted_at":"2025-06-12T07:02:28Z","title":"FormosanBench: Benchmarking Low-Resource Austronesian Languages in the Era of Large Language Models","version":1},"reference_index":1,"source":"arxiv_source","source_observed_at":"2026-08-07T04:31:30.363100Z"},"links":{"cited_paper":"/paper/2406.03368","citing_paper":"/paper/2506.21563"},"observation_digest":"sha256:9c79a6e58d7b0d19f40e4ce097a5281d1c2c95d1d097d3d383ae36b8fb2916a8","observation_id":"a5b30d12-c67c-455e-9a11-52cd85156143","resolution":{"observed_at":"2026-08-07T04:31:30.363100Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T04:31:30.432519Z","title":null,"venue":null,"work_id":null,"year":2019},"citing_paper":{"arxiv_id":"2506.21563","last_updated":"2025-06-12T07:02:28Z","snapshot_observed_at":"2026-08-08T09:20:55.054492Z","submitted_at":"2025-06-12T07:02:28Z","title":"FormosanBench: Benchmarking Low-Resource Austronesian Languages in the Era of Large Language Models","version":1},"reference_index":2,"source":"arxiv_source","source_observed_at":"2026-08-07T04:31:30.432519Z"},"links":{"citing_paper":"/paper/2506.21563"},"observation_digest":"sha256:f9bd25ff92720d89c543aad31cd4cf6166cfb9c80edcb9ee8119fa631b58a892","observation_id":"df8dd6df-1e94-465c-a327-3b8df2b2c947","resolution":{"observed_at":"2026-08-07T04:31:30.432519Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2406.19564","last_updated":"2024-06-27T22:38:04Z","snapshot_observed_at":"2026-08-07T05:01:53.776740Z","submitted_at":"2024-06-27T22:38:04Z","title":"Voices Unheard: NLP Resources and Models for Yor\\`ub\\'a Regional Dialects","version":1},"cited_work":{"arxiv_id":"2406.19564","doi":null,"metadata_source":"pith","pith_arxiv_id":"2406.19564","snapshot_observed_at":"2026-08-07T04:31:34.422100Z","title":"Voices Unheard: NLP Resources and Models for Yor\\`ub\\'a Regional Dialects","venue":"cs.CL","work_id":"20cead16-ce14-4584-9093-1b795c8ac58b","year":2024},"citing_paper":{"arxiv_id":"2506.21563","last_updated":"2025-06-12T07:02:28Z","snapshot_observed_at":"2026-08-08T09:20:55.054492Z","submitted_at":"2025-06-12T07:02:28Z","title":"FormosanBench: Benchmarking Low-Resource Austronesian Languages in the Era of Large Language Models","version":1},"reference_index":3,"source":"arxiv_source","source_observed_at":"2026-08-07T04:31:30.472828Z"},"links":{"cited_paper":"/paper/2406.19564","citing_paper":"/paper/2506.21563"},"observation_digest":"sha256:309e1d9f893f603498b865210dd917e82607a2988da2e64b2738d4c2b2840134","observation_id":"6957bc18-f3ce-4acf-8397-84450df4dd81","resolution":{"observed_at":"2026-08-07T04:31:34.473716Z","resolver_source":"local_arxiv","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-11T06:34:44.6726+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-11T06:34:44.6726+00:00","source":"crossref"},{"observed_at":"2026-08-11T06:34:36.301508+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T04:31:35.469031Z","title":null,"venue":null,"work_id":"cc57aecf-7fd6-4949-9c0d-683a47e8f487","year":2024},"citing_paper":{"arxiv_id":"2506.21563","last_updated":"2025-06-12T07:02:28Z","snapshot_observed_at":"2026-08-08T09:20:55.054492Z","submitted_at":"2025-06-12T07:02:28Z","title":"FormosanBench: Benchmarking Low-Resource Austronesian Languages in the Era of Large Language Models","version":1},"reference_index":4,"source":"arxiv_source","source_observed_at":"2026-08-07T04:31:30.506855Z"},"links":{"citing_paper":"/paper/2506.21563"},"observation_digest":"sha256:c347f78a1b50940bb9bc505565575116ddb3578ef71ad6555fc08f0e7de9953c","observation_id":"8efe6793-4637-43c6-9994-c769c5df5d90","resolution":{"observed_at":"2026-08-07T04:31:35.518130Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-11T06:34:44.6726+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-11T06:34:44.6726+00:00","source":"crossref"},{"observed_at":"2026-08-11T06:34:36.301508+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2006.11477","last_updated":"2020-10-22T06:09:10Z","snapshot_observed_at":"2026-08-07T05:13:16.889179Z","submitted_at":"2020-06-20T02:35:02Z","title":"wav2vec 2.0: A Framework for Self-Supervised Learning of Speech Representations","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2006.11477","snapshot_observed_at":"2026-08-07T04:31:30.566703Z","title":null,"venue":null,"work_id":null,"year":2020},"citing_paper":{"arxiv_id":"2506.21563","last_updated":"2025-06-12T07:02:28Z","snapshot_observed_at":"2026-08-08T09:20:55.054492Z","submitted_at":"2025-06-12T07:02:28Z","title":"FormosanBench: Benchmarking Low-Resource Austronesian Languages in the Era of Large Language Models","version":1},"reference_index":5,"source":"arxiv_source","source_observed_at":"2026-08-07T04:31:30.566703Z"},"links":{"cited_paper":"/paper/2006.11477","citing_paper":"/paper/2506.21563"},"observation_digest":"sha256:6779fecbeb15c6f3fd859e55b25a54a05c40523c20f64145b5f6107da62e8333","observation_id":"1f91c863-cfea-40e8-b772-b06a76bc67e4","resolution":{"observed_at":"2026-08-07T04:31:30.566703Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":"stable/4292810","doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T04:31:34.295915Z","title":null,"venue":null,"work_id":"bf27afae-e4c3-4f12-98b1-b253fe2171a0","year":1984},"citing_paper":{"arxiv_id":"2506.21563","last_updated":"2025-06-12T07:02:28Z","snapshot_observed_at":"2026-08-08T09:20:55.054492Z","submitted_at":"2025-06-12T07:02:28Z","title":"FormosanBench: Benchmarking Low-Resource Austronesian Languages in the Era of Large Language Models","version":1},"reference_index":7,"source":"arxiv_source","source_observed_at":"2026-08-07T04:31:30.734618Z"},"links":{"citing_paper":"/paper/2506.21563"},"observation_digest":"sha256:f4890548b1f9f5c95c263915e3e46f84d1ada5e6b7bb97cfe84ed33c349cfaf3","observation_id":"7b83a6b9-0234-4ab4-98aa-5e109ef48fc8","resolution":{"observed_at":"2026-08-07T04:31:34.371982Z","resolver_source":"raw_fallback","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-11T06:34:44.6726+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-11T06:34:44.6726+00:00","source":"crossref"},{"observed_at":"2026-08-11T06:34:36.301508+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":"10.1146/annurev-linguistics-011718-012440","metadata_source":"openalex","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-10T05:30:23.456663Z","title":null,"venue":"Annual Review of Linguistics","work_id":"6ebd30c4-f8ea-47d8-94db-9e6e66eb3aad","year":2019},"citing_paper":{"arxiv_id":"2506.21563","last_updated":"2025-06-12T07:02:28Z","snapshot_observed_at":"2026-08-08T09:20:55.054492Z","submitted_at":"2025-06-12T07:02:28Z","title":"FormosanBench: Benchmarking Low-Resource Austronesian Languages in the Era of Large Language Models","version":1},"reference_index":8,"source":"arxiv_source","source_observed_at":"2026-08-07T04:31:30.775246Z"},"links":{"citing_paper":"/paper/2506.21563"},"observation_digest":"sha256:ec6e53dbc5335145e3cfad81bddc967d22c9a3e664f6e4fa9bf6de3a5bb5eba5","observation_id":"cacb66be-c3aa-40ff-9982-6a33105bcf1c","resolution":{"observed_at":"2026-08-07T04:31:33.450980Z","resolver_source":"doi","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-11T06:34:44.6726+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-11T06:34:44.6726+00:00","source":"crossref"},{"observed_at":"2026-08-11T06:34:36.301508+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T04:31:35.377008Z","title":null,"venue":null,"work_id":"45bcde4c-c7ae-4bdd-8961-0ffa2771f834","year":2018},"citing_paper":{"arxiv_id":"2506.21563","last_updated":"2025-06-12T07:02:28Z","snapshot_observed_at":"2026-08-08T09:20:55.054492Z","submitted_at":"2025-06-12T07:02:28Z","title":"FormosanBench: Benchmarking Low-Resource Austronesian Languages in the Era of Large Language Models","version":1},"reference_index":9,"source":"arxiv_source","source_observed_at":"2026-08-07T04:31:30.881962Z"},"links":{"citing_paper":"/paper/2506.21563"},"observation_digest":"sha256:82339db7d8f172866ced524e5d4b71af624f68e90fb7d66e48c5dbebe0ecffcc","observation_id":"ebbc6500-378d-4de3-ba55-f2e3094fdbd1","resolution":{"observed_at":"2026-08-07T04:31:35.431448Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-11T06:34:44.6726+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-11T06:34:44.6726+00:00","source":"crossref"},{"observed_at":"2026-08-11T06:34:36.301508+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2308.11596","last_updated":"2023-10-25T03:52:07Z","snapshot_observed_at":"2026-07-06T16:09:09.018523Z","submitted_at":"2023-08-22T17:44:18Z","title":"SeamlessM4T: Massively Multilingual & Multimodal Machine Translation","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2308.11596","snapshot_observed_at":"2026-08-07T04:31:30.949740Z","title":null,"venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2506.21563","last_updated":"2025-06-12T07:02:28Z","snapshot_observed_at":"2026-08-08T09:20:55.054492Z","submitted_at":"2025-06-12T07:02:28Z","title":"FormosanBench: Benchmarking Low-Resource Austronesian Languages in the Era of Large Language Models","version":1},"reference_index":10,"source":"arxiv_source","source_observed_at":"2026-08-07T04:31:30.949740Z"},"links":{"cited_paper":"/paper/2308.11596","citing_paper":"/paper/2506.21563"},"observation_digest":"sha256:da692263e283b33a6225fb911308e52c53c08bf06b8675dc3a19b0dd751d1d03","observation_id":"c480a112-5d40-4e21-b8ee-6ebf9fb6cc97","resolution":{"observed_at":"2026-08-07T04:31:30.949740Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T04:31:31.021524Z","title":null,"venue":null,"work_id":null,"year":2020},"citing_paper":{"arxiv_id":"2506.21563","last_updated":"2025-06-12T07:02:28Z","snapshot_observed_at":"2026-08-08T09:20:55.054492Z","submitted_at":"2025-06-12T07:02:28Z","title":"FormosanBench: Benchmarking Low-Resource Austronesian Languages in the Era of Large Language Models","version":1},"reference_index":11,"source":"arxiv_source","source_observed_at":"2026-08-07T04:31:31.021524Z"},"links":{"citing_paper":"/paper/2506.21563"},"observation_digest":"sha256:8e41646cfff4bc4da2090668a07515c65074df0f2c2fe30ae5e5383a11e08051","observation_id":"c6c0092b-d33a-44a7-8d6f-a23f1ed76edc","resolution":{"observed_at":"2026-08-07T04:31:31.021524Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T04:31:35.273949Z","title":null,"venue":null,"work_id":"95e77c7b-9701-46a3-ac1b-251f5e225d86","year":2024},"citing_paper":{"arxiv_id":"2506.21563","last_updated":"2025-06-12T07:02:28Z","snapshot_observed_at":"2026-08-08T09:20:55.054492Z","submitted_at":"2025-06-12T07:02:28Z","title":"FormosanBench: Benchmarking Low-Resource Austronesian Languages in the Era of Large Language Models","version":1},"reference_index":12,"source":"arxiv_source","source_observed_at":"2026-08-07T04:31:31.130520Z"},"links":{"citing_paper":"/paper/2506.21563"},"observation_digest":"sha256:e782fda7b6988f77a3f83582c2defdc6754a0149653d32592cf7127af2cbdae1","observation_id":"d32c9d44-7f86-4ad6-bbd0-f5ac505e5e7b","resolution":{"observed_at":"2026-08-07T04:31:35.290629Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-11T06:34:44.6726+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-11T06:34:44.6726+00:00","source":"crossref"},{"observed_at":"2026-08-11T06:34:36.301508+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2010.12309","last_updated":"2021-04-09T13:48:02Z","snapshot_observed_at":"2026-07-06T10:07:35.145338Z","submitted_at":"2020-10-23T11:22:01Z","title":"A Survey on Recent Approaches for Natural Language Processing in Low-Resource Scenarios","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2010.12309","snapshot_observed_at":"2026-08-07T04:31:31.238324Z","title":"Hedderich, Lukas Lange, Heike Adel, Jannik Strötgen, and Dietrich Klakow","venue":null,"work_id":null,"year":2021},"citing_paper":{"arxiv_id":"2506.21563","last_updated":"2025-06-12T07:02:28Z","snapshot_observed_at":"2026-08-08T09:20:55.054492Z","submitted_at":"2025-06-12T07:02:28Z","title":"FormosanBench: Benchmarking Low-Resource Austronesian Languages in the Era of Large Language Models","version":1},"reference_index":13,"source":"arxiv_source","source_observed_at":"2026-08-07T04:31:31.238324Z"},"links":{"cited_paper":"/paper/2010.12309","citing_paper":"/paper/2506.21563"},"observation_digest":"sha256:cfa1c7ac51921f2ca3c305ad4582716ebc66be81ed4305f96714d479eff63c4d","observation_id":"0e96d213-d578-41f7-972d-2519e9e8c88d","resolution":{"observed_at":"2026-08-07T04:31:31.238324Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2009.03300","last_updated":"2021-01-12T18:57:11Z","snapshot_observed_at":"2026-08-10T12:35:09.020030Z","submitted_at":"2020-09-07T17:59:25Z","title":"Measuring Massive Multitask Language Understanding","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2009.03300","snapshot_observed_at":"2026-08-07T04:31:31.282994Z","title":null,"venue":null,"work_id":null,"year":2020},"citing_paper":{"arxiv_id":"2506.21563","last_updated":"2025-06-12T07:02:28Z","snapshot_observed_at":"2026-08-08T09:20:55.054492Z","submitted_at":"2025-06-12T07:02:28Z","title":"FormosanBench: Benchmarking Low-Resource Austronesian Languages in the Era of Large Language Models","version":1},"reference_index":14,"source":"arxiv_source","source_observed_at":"2026-08-07T04:31:31.282994Z"},"links":{"cited_paper":"/paper/2009.03300","citing_paper":"/paper/2506.21563"},"observation_digest":"sha256:a34679134c87bd6ea643a23360130cc9934935d70684922634d7233005fc96f5","observation_id":"260faef9-0701-4723-8f10-ed2ee41b2ca4","resolution":{"observed_at":"2026-08-07T04:31:31.282994Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2310.06825","last_updated":"2023-10-10T17:54:58Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2023-10-10T17:54:58Z","title":"Mistral 7B","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2310.06825","snapshot_observed_at":"2026-08-07T04:31:31.321024Z","title":null,"venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2506.21563","last_updated":"2025-06-12T07:02:28Z","snapshot_observed_at":"2026-08-08T09:20:55.054492Z","submitted_at":"2025-06-12T07:02:28Z","title":"FormosanBench: Benchmarking Low-Resource Austronesian Languages in the Era of Large Language Models","version":1},"reference_index":15,"source":"arxiv_source","source_observed_at":"2026-08-07T04:31:31.321024Z"},"links":{"cited_paper":"/paper/2310.06825","citing_paper":"/paper/2506.21563"},"observation_digest":"sha256:21a91ddb2d418786f668466073961387d9498726df4c904a6416fc5c38c8b595","observation_id":"bcf45092-c979-42b2-a045-40ccbed26508","resolution":{"observed_at":"2026-08-07T04:31:31.321024Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T04:31:31.459291Z","title":null,"venue":null,"work_id":null,"year":2020},"citing_paper":{"arxiv_id":"2506.21563","last_updated":"2025-06-12T07:02:28Z","snapshot_observed_at":"2026-08-08T09:20:55.054492Z","submitted_at":"2025-06-12T07:02:28Z","title":"FormosanBench: Benchmarking Low-Resource Austronesian Languages in the Era of Large Language Models","version":1},"reference_index":16,"source":"arxiv_source","source_observed_at":"2026-08-07T04:31:31.459291Z"},"links":{"citing_paper":"/paper/2506.21563"},"observation_digest":"sha256:c922f65eb0e44707b63a2573bcfa4b1b1c33b7a43e12fa539ec155af577d91da","observation_id":"0cc83782-2970-48f0-b8b4-656f6399542f","resolution":{"observed_at":"2026-08-07T04:31:31.459291Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2410.14815","last_updated":"2025-04-21T05:29:01Z","snapshot_observed_at":"2026-08-08T14:58:04.830848Z","submitted_at":"2024-10-18T18:35:19Z","title":"Adapting Multilingual LLMs to Low-Resource Languages using Continued Pre-training and Synthetic Corpus","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2410.14815","snapshot_observed_at":"2026-08-07T04:31:31.516463Z","title":null,"venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2506.21563","last_updated":"2025-06-12T07:02:28Z","snapshot_observed_at":"2026-08-08T09:20:55.054492Z","submitted_at":"2025-06-12T07:02:28Z","title":"FormosanBench: Benchmarking Low-Resource Austronesian Languages in the Era of Large Language Models","version":1},"reference_index":17,"source":"arxiv_source","source_observed_at":"2026-08-07T04:31:31.516463Z"},"links":{"cited_paper":"/paper/2410.14815","citing_paper":"/paper/2506.21563"},"observation_digest":"sha256:d1bd90a225fdc1ffee724d678f45cd0f95d7d555ec9812a55bbc9ac3a9b9d5f7","observation_id":"61758689-51d3-4526-85a8-a1000599c8d8","resolution":{"observed_at":"2026-08-07T04:31:31.516463Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T04:31:35.147073Z","title":null,"venue":null,"work_id":"57743ccb-14d1-497b-bc1c-0c72f02df59c","year":2023},"citing_paper":{"arxiv_id":"2506.21563","last_updated":"2025-06-12T07:02:28Z","snapshot_observed_at":"2026-08-08T09:20:55.054492Z","submitted_at":"2025-06-12T07:02:28Z","title":"FormosanBench: Benchmarking Low-Resource Austronesian Languages in the Era of Large Language Models","version":1},"reference_index":18,"source":"arxiv_source","source_observed_at":"2026-08-07T04:31:31.571240Z"},"links":{"citing_paper":"/paper/2506.21563"},"observation_digest":"sha256:fb0f76b3e7189abeac4bad407c3f24e0df4b0050a971b1aa1dbbe04878426c7f","observation_id":"ba92b765-d2f2-4182-9df7-30e0d2fbe2d5","resolution":{"observed_at":"2026-08-07T04:31:35.204583Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-11T06:34:44.6726+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-11T06:34:44.6726+00:00","source":"crossref"},{"observed_at":"2026-08-11T06:34:36.301508+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T04:31:31.675173Z","title":null,"venue":null,"work_id":null,"year":2004},"citing_paper":{"arxiv_id":"2506.21563","last_updated":"2025-06-12T07:02:28Z","snapshot_observed_at":"2026-08-08T09:20:55.054492Z","submitted_at":"2025-06-12T07:02:28Z","title":"FormosanBench: Benchmarking Low-Resource Austronesian Languages in the Era of Large Language Models","version":1},"reference_index":19,"source":"arxiv_source","source_observed_at":"2026-08-07T04:31:31.675173Z"},"links":{"citing_paper":"/paper/2506.21563"},"observation_digest":"sha256:25223138cfccfa2e87cdb1212dba2f95313c6ad3fe1b9f9a2fe4b9e73756a2a8","observation_id":"178af65d-fa52-4424-98a0-ffa85ae19f9b","resolution":{"observed_at":"2026-08-07T04:31:31.675173Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"1801.10198","last_updated":"2018-01-30T20:07:01Z","snapshot_observed_at":"2026-07-06T06:21:01.297755Z","submitted_at":"2018-01-30T20:07:01Z","title":"Generating Wikipedia by Summarizing Long Sequences","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"1801.10198","snapshot_observed_at":"2026-08-07T04:31:31.751157Z","title":"Liu, Mohammad Saleh, Etienne Pot, Ben Goodrich, Ryan Sepassi, Lukasz Kaiser, and Noam Shazeer","venue":null,"work_id":null,"year":2018},"citing_paper":{"arxiv_id":"2506.21563","last_updated":"2025-06-12T07:02:28Z","snapshot_observed_at":"2026-08-08T09:20:55.054492Z","submitted_at":"2025-06-12T07:02:28Z","title":"FormosanBench: Benchmarking Low-Resource Austronesian Languages in the Era of Large Language Models","version":1},"reference_index":20,"source":"arxiv_source","source_observed_at":"2026-08-07T04:31:31.751157Z"},"links":{"cited_paper":"/paper/1801.10198","citing_paper":"/paper/2506.21563"},"observation_digest":"sha256:9f8a37d5e018d6f51133710871de0e5017c50dc5e072bbcfb594571c2b640b56","observation_id":"d0f9db94-70d7-44d6-be3c-6b1ae4637b34","resolution":{"observed_at":"2026-08-07T04:31:31.751157Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2006.07264","last_updated":"2020-06-12T15:21:57Z","snapshot_observed_at":"2026-08-08T14:56:31.646029Z","submitted_at":"2020-06-12T15:21:57Z","title":"Low-resource Languages: A Review of Past Work and Future Challenges","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2006.07264","snapshot_observed_at":"2026-08-07T04:31:31.878832Z","title":null,"venue":null,"work_id":null,"year":2020},"citing_paper":{"arxiv_id":"2506.21563","last_updated":"2025-06-12T07:02:28Z","snapshot_observed_at":"2026-08-08T09:20:55.054492Z","submitted_at":"2025-06-12T07:02:28Z","title":"FormosanBench: Benchmarking Low-Resource Austronesian Languages in the Era of Large Language Models","version":1},"reference_index":21,"source":"arxiv_source","source_observed_at":"2026-08-07T04:31:31.878832Z"},"links":{"cited_paper":"/paper/2006.07264","citing_paper":"/paper/2506.21563"},"observation_digest":"sha256:9ace8d5f4120bc605512ce7f629dfff76a37c3b7f3e162ab8c5394950cdfb3e4","observation_id":"b9e9c457-362f-462c-affa-df624cd02927","resolution":{"observed_at":"2026-08-07T04:31:31.878832Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T04:31:35.059739Z","title":null,"venue":null,"work_id":"e50f3dc7-7dc7-4973-8f2d-19064d983fee","year":2025},"citing_paper":{"arxiv_id":"2506.21563","last_updated":"2025-06-12T07:02:28Z","snapshot_observed_at":"2026-08-08T09:20:55.054492Z","submitted_at":"2025-06-12T07:02:28Z","title":"FormosanBench: Benchmarking Low-Resource Austronesian Languages in the Era of Large Language Models","version":1},"reference_index":22,"source":"arxiv_source","source_observed_at":"2026-08-07T04:31:31.935731Z"},"links":{"citing_paper":"/paper/2506.21563"},"observation_digest":"sha256:7f179824e50e594769e66c27ddc1b387f4e0752a9b4ee2812ff4fe45c481daf2","observation_id":"bff62830-dbad-48d2-b27f-9c509931fdec","resolution":{"observed_at":"2026-08-07T04:31:35.094606Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-11T06:34:44.6726+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-11T06:34:44.6726+00:00","source":"crossref"},{"observed_at":"2026-08-11T06:34:36.301508+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T04:31:34.928478Z","title":null,"venue":null,"work_id":"59a65fef-cf1a-40da-b6e5-ecc2befcc1c3","year":2024},"citing_paper":{"arxiv_id":"2506.21563","last_updated":"2025-06-12T07:02:28Z","snapshot_observed_at":"2026-08-08T09:20:55.054492Z","submitted_at":"2025-06-12T07:02:28Z","title":"FormosanBench: Benchmarking Low-Resource Austronesian Languages in the Era of Large Language Models","version":1},"reference_index":23,"source":"arxiv_source","source_observed_at":"2026-08-07T04:31:31.978172Z"},"links":{"citing_paper":"/paper/2506.21563"},"observation_digest":"sha256:9f3851b21dedf121e89da68af027cad98317eb9ed9bf25eaf5cb48a55489f1e4","observation_id":"65d21a59-4a1e-4fe5-aac0-546fbc4e1806","resolution":{"observed_at":"2026-08-07T04:31:34.978925Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-11T06:34:44.6726+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-11T06:34:44.6726+00:00","source":"crossref"},{"observed_at":"2026-08-11T06:34:36.301508+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T04:31:32.061531Z","title":null,"venue":null,"work_id":null,"year":2002},"citing_paper":{"arxiv_id":"2506.21563","last_updated":"2025-06-12T07:02:28Z","snapshot_observed_at":"2026-08-08T09:20:55.054492Z","submitted_at":"2025-06-12T07:02:28Z","title":"FormosanBench: Benchmarking Low-Resource Austronesian Languages in the Era of Large Language Models","version":1},"reference_index":24,"source":"arxiv_source","source_observed_at":"2026-08-07T04:31:32.061531Z"},"links":{"citing_paper":"/paper/2506.21563"},"observation_digest":"sha256:5cecd7b76ebd3d2c2309a942e9838707850d7f2035553905156c4b0cf73f3ba8","observation_id":"5a7035ce-feaf-465a-a360-cf8e610a2f89","resolution":{"observed_at":"2026-08-07T04:31:32.061531Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T04:31:34.804584Z","title":null,"venue":null,"work_id":"106d5742-4936-4eaa-9f5d-14317c8122d7","year":2023},"citing_paper":{"arxiv_id":"2506.21563","last_updated":"2025-06-12T07:02:28Z","snapshot_observed_at":"2026-08-08T09:20:55.054492Z","submitted_at":"2025-06-12T07:02:28Z","title":"FormosanBench: Benchmarking Low-Resource Austronesian Languages in the Era of Large Language Models","version":1},"reference_index":25,"source":"arxiv_source","source_observed_at":"2026-08-07T04:31:32.174020Z"},"links":{"citing_paper":"/paper/2506.21563"},"observation_digest":"sha256:87c3a8481d246137c24ea10d22528785ae07818dd378f2abc7a6c9fa7acc3d78","observation_id":"86fdd33e-2d4f-4696-aa82-abde46c30dd6","resolution":{"observed_at":"2026-08-07T04:31:34.886231Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-11T06:34:44.6726+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-11T06:34:44.6726+00:00","source":"crossref"},{"observed_at":"2026-08-11T06:34:36.301508+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2212.04356","last_updated":"2022-12-06T18:46:04Z","snapshot_observed_at":"2026-07-06T14:28:21.844826Z","submitted_at":"2022-12-06T18:46:04Z","title":"Robust Speech Recognition via Large-Scale Weak Supervision","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2212.04356","snapshot_observed_at":"2026-08-07T04:31:32.218638Z","title":null,"venue":null,"work_id":null,"year":2022},"citing_paper":{"arxiv_id":"2506.21563","last_updated":"2025-06-12T07:02:28Z","snapshot_observed_at":"2026-08-08T09:20:55.054492Z","submitted_at":"2025-06-12T07:02:28Z","title":"FormosanBench: Benchmarking Low-Resource Austronesian Languages in the Era of Large Language Models","version":1},"reference_index":26,"source":"arxiv_source","source_observed_at":"2026-08-07T04:31:32.218638Z"},"links":{"cited_paper":"/paper/2212.04356","citing_paper":"/paper/2506.21563"},"observation_digest":"sha256:af5c447406c2b3b42fcce7f37e72e55d065b6a8ccefaf11387c2e1a3628793fb","observation_id":"30b1a339-1ffb-4fff-b9a5-23f5fe2b993d","resolution":{"observed_at":"2026-08-07T04:31:32.218638Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2206.09755","last_updated":"2022-06-20T13:04:23Z","snapshot_observed_at":"2026-08-09T13:16:52.138390Z","submitted_at":"2022-06-20T13:04:23Z","title":"Square One Bias in NLP: Towards a Multi-Dimensional Exploration of the Research Manifold","version":1},"cited_work":{"arxiv_id":"2206.09755","doi":null,"metadata_source":"pith","pith_arxiv_id":"2206.09755","snapshot_observed_at":"2026-08-07T04:31:33.970320Z","title":"Square One Bias in NLP: Towards a Multi-Dimensional Exploration of the Research Manifold","venue":"cs.CL","work_id":"ae07fdc6-8646-4f50-98f7-4866b99c9146","year":2022},"citing_paper":{"arxiv_id":"2506.21563","last_updated":"2025-06-12T07:02:28Z","snapshot_observed_at":"2026-08-08T09:20:55.054492Z","submitted_at":"2025-06-12T07:02:28Z","title":"FormosanBench: Benchmarking Low-Resource Austronesian Languages in the Era of Large Language Models","version":1},"reference_index":27,"source":"arxiv_source","source_observed_at":"2026-08-07T04:31:32.302736Z"},"links":{"cited_paper":"/paper/2206.09755","citing_paper":"/paper/2506.21563"},"observation_digest":"sha256:c4ec79c8be7483ea5ea453195d3537371ef3a8be264f33c245a0a93d07085ab6","observation_id":"4c9618a1-1672-42de-83dd-09a57f6b02ee","resolution":{"observed_at":"2026-08-07T04:31:34.034800Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-11T06:34:44.6726+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-11T06:34:44.6726+00:00","source":"crossref"},{"observed_at":"2026-08-11T06:34:36.301508+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2409.17912","last_updated":"2024-11-11T22:14:04Z","snapshot_observed_at":"2026-07-06T19:22:50.028253Z","submitted_at":"2024-09-26T14:56:38Z","title":"Atlas-Chat: Adapting Large Language Models for Low-Resource Moroccan Arabic Dialect","version":2},"cited_work":{"arxiv_id":"2409.17912","doi":null,"metadata_source":"pith","pith_arxiv_id":"2409.17912","snapshot_observed_at":"2026-08-07T04:31:33.870380Z","title":"Atlas-Chat: Adapting Large Language Models for Low-Resource Moroccan Arabic Dialect","venue":"cs.CL","work_id":"0c53073f-8e67-4d68-9fae-50d03c9c5f56","year":2024},"citing_paper":{"arxiv_id":"2506.21563","last_updated":"2025-06-12T07:02:28Z","snapshot_observed_at":"2026-08-08T09:20:55.054492Z","submitted_at":"2025-06-12T07:02:28Z","title":"FormosanBench: Benchmarking Low-Resource Austronesian Languages in the Era of Large Language Models","version":1},"reference_index":28,"source":"arxiv_source","source_observed_at":"2026-08-07T04:31:32.351721Z"},"links":{"cited_paper":"/paper/2409.17912","citing_paper":"/paper/2506.21563"},"observation_digest":"sha256:df79fc12c077e3f9eb7eb483068111af586753c37ab5a590767de287dd394b93","observation_id":"59ace668-8ea7-4501-8084-1c68df9541a7","resolution":{"observed_at":"2026-08-07T04:31:33.891421Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-11T06:34:44.6726+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-11T06:34:44.6726+00:00","source":"crossref"},{"observed_at":"2026-08-11T06:34:36.301508+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2209.13101","last_updated":"2022-09-27T01:28:02Z","snapshot_observed_at":"2026-08-10T16:26:49.292748Z","submitted_at":"2022-09-27T01:28:02Z","title":"WikiDes: A Wikipedia-Based Dataset for Generating Short Descriptions from Paragraphs","version":1},"cited_work":{"arxiv_id":"2209.13101","doi":null,"metadata_source":"pith","pith_arxiv_id":"2209.13101","snapshot_observed_at":"2026-08-07T04:31:33.756701Z","title":"WikiDes: A Wikipedia-Based Dataset for Generating Short Descriptions from Paragraphs","venue":"cs.CL","work_id":"d35e2440-36bb-4388-a70a-dd4d131c25f1","year":2022},"citing_paper":{"arxiv_id":"2506.21563","last_updated":"2025-06-12T07:02:28Z","snapshot_observed_at":"2026-08-08T09:20:55.054492Z","submitted_at":"2025-06-12T07:02:28Z","title":"FormosanBench: Benchmarking Low-Resource Austronesian Languages in the Era of Large Language Models","version":1},"reference_index":29,"source":"arxiv_source","source_observed_at":"2026-08-07T04:31:32.397777Z"},"links":{"cited_paper":"/paper/2209.13101","citing_paper":"/paper/2506.21563"},"observation_digest":"sha256:6c731a2ddbaa7a2b3af9e11852f7645bf0e79984e28d07c38607b897a193c27c","observation_id":"e4166388-5f91-45f3-9524-a7aa0dc6a4b3","resolution":{"observed_at":"2026-08-07T04:31:33.812297Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-11T06:34:44.6726+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-11T06:34:44.6726+00:00","source":"crossref"},{"observed_at":"2026-08-11T06:34:36.301508+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2403.08295","last_updated":"2024-04-16T12:52:47Z","snapshot_observed_at":"2026-08-03T03:29:01.959523Z","submitted_at":"2024-03-13T06:59:16Z","title":"Gemma: Open Models Based on Gemini Research and Technology","version":4},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2403.08295","snapshot_observed_at":"2026-08-07T04:31:32.448951Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2506.21563","last_updated":"2025-06-12T07:02:28Z","snapshot_observed_at":"2026-08-08T09:20:55.054492Z","submitted_at":"2025-06-12T07:02:28Z","title":"FormosanBench: Benchmarking Low-Resource Austronesian Languages in the Era of Large Language Models","version":1},"reference_index":30,"source":"arxiv_source","source_observed_at":"2026-08-07T04:31:32.448951Z"},"links":{"cited_paper":"/paper/2403.08295","citing_paper":"/paper/2506.21563"},"observation_digest":"sha256:5eddf38548ac0d1953486161bcea93869df44979959297c1f54374e915a75413","observation_id":"3d36df8b-fbd3-4be2-9879-372423d4f663","resolution":{"observed_at":"2026-08-07T04:31:32.448951Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2207.04672","last_updated":"2022-08-25T17:10:53Z","snapshot_observed_at":"2026-07-06T13:29:47.927628Z","submitted_at":"2022-07-11T07:33:36Z","title":"No Language Left Behind: Scaling Human-Centered Machine Translation","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2207.04672","snapshot_observed_at":"2026-08-07T04:31:32.577746Z","title":null,"venue":null,"work_id":null,"year":2022},"citing_paper":{"arxiv_id":"2506.21563","last_updated":"2025-06-12T07:02:28Z","snapshot_observed_at":"2026-08-08T09:20:55.054492Z","submitted_at":"2025-06-12T07:02:28Z","title":"FormosanBench: Benchmarking Low-Resource Austronesian Languages in the Era of Large Language Models","version":1},"reference_index":31,"source":"arxiv_source","source_observed_at":"2026-08-07T04:31:32.577746Z"},"links":{"cited_paper":"/paper/2207.04672","citing_paper":"/paper/2506.21563"},"observation_digest":"sha256:405919c5f33aaddda8e2a35941c7f401152af9f06fcd76785d0e17e37c58bfe7","observation_id":"96af276d-395f-45a2-b384-923f0e6faf4d","resolution":{"observed_at":"2026-08-07T04:31:32.577746Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T04:31:34.691355Z","title":null,"venue":null,"work_id":"10bc16f8-023f-4c8f-bc82-0b1f8f86c5a3","year":2020},"citing_paper":{"arxiv_id":"2506.21563","last_updated":"2025-06-12T07:02:28Z","snapshot_observed_at":"2026-08-08T09:20:55.054492Z","submitted_at":"2025-06-12T07:02:28Z","title":"FormosanBench: Benchmarking Low-Resource Austronesian Languages in the Era of Large Language Models","version":1},"reference_index":32,"source":"arxiv_source","source_observed_at":"2026-08-07T04:31:32.631866Z"},"links":{"citing_paper":"/paper/2506.21563"},"observation_digest":"sha256:25890ca73adb61809c5e8a44abf9d2b05c619975bc5b1ea1bf4707d5dac06a07","observation_id":"e2b1444c-cb88-440e-a0da-628e846a8dbe","resolution":{"observed_at":"2026-08-07T04:31:34.733079Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-11T06:34:44.6726+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-11T06:34:44.6726+00:00","source":"crossref"},{"observed_at":"2026-08-11T06:34:36.301508+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2410.15687","last_updated":"2024-10-21T06:55:35Z","snapshot_observed_at":"2026-07-06T19:36:49.949466Z","submitted_at":"2024-10-21T06:55:35Z","title":"DomainSum: A Hierarchical Benchmark for Fine-Grained Domain Shift in Abstractive Text Summarization","version":1},"cited_work":{"arxiv_id":"2410.15687","doi":null,"metadata_source":"pith","pith_arxiv_id":"2410.15687","snapshot_observed_at":"2026-08-07T04:31:33.641290Z","title":"DomainSum: A Hierarchical Benchmark for Fine-Grained Domain Shift in Abstractive Text Summarization","venue":"cs.CL","work_id":"c171b263-5dd5-4ec9-92e0-a181c2bd299a","year":2024},"citing_paper":{"arxiv_id":"2506.21563","last_updated":"2025-06-12T07:02:28Z","snapshot_observed_at":"2026-08-08T09:20:55.054492Z","submitted_at":"2025-06-12T07:02:28Z","title":"FormosanBench: Benchmarking Low-Resource Austronesian Languages in the Era of Large Language Models","version":1},"reference_index":33,"source":"arxiv_source","source_observed_at":"2026-08-07T04:31:32.718308Z"},"links":{"cited_paper":"/paper/2410.15687","citing_paper":"/paper/2506.21563"},"observation_digest":"sha256:61bfd32cd4052bd7108e04cbf148161b22aec4a4796c92da97353e3477da28d1","observation_id":"701be416-8541-471c-8858-094eab2f4c8e","resolution":{"observed_at":"2026-08-07T04:31:33.683192Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-11T06:34:44.6726+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-11T06:34:44.6726+00:00","source":"crossref"},{"observed_at":"2026-08-11T06:34:36.301508+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2207.02263","last_updated":"2022-07-05T18:52:19Z","snapshot_observed_at":"2026-08-09T16:04:13.738635Z","submitted_at":"2022-07-05T18:52:19Z","title":"Improving the Faithfulness of Abstractive Summarization via Entity Coverage Control","version":1},"cited_work":{"arxiv_id":"2207.02263","doi":null,"metadata_source":"pith","pith_arxiv_id":"2207.02263","snapshot_observed_at":"2026-08-07T04:31:33.501683Z","title":"Improving the Faithfulness of Abstractive Summarization via Entity Coverage Control","venue":"cs.CL","work_id":"c5e85ec0-09d5-47f4-bfa0-1a4607aea605","year":2022},"citing_paper":{"arxiv_id":"2506.21563","last_updated":"2025-06-12T07:02:28Z","snapshot_observed_at":"2026-08-08T09:20:55.054492Z","submitted_at":"2025-06-12T07:02:28Z","title":"FormosanBench: Benchmarking Low-Resource Austronesian Languages in the Era of Large Language Models","version":1},"reference_index":34,"source":"arxiv_source","source_observed_at":"2026-08-07T04:31:32.808261Z"},"links":{"cited_paper":"/paper/2207.02263","citing_paper":"/paper/2506.21563"},"observation_digest":"sha256:1ddcbf4a0a3696b26a19edfa6bd1e9c0ebc878d61af195e3f4dd108a6b1aff3a","observation_id":"557bd216-e470-4afc-b7e7-1f9dc14a112c","resolution":{"observed_at":"2026-08-07T04:31:33.557432Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-11T06:34:44.6726+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-11T06:34:44.6726+00:00","source":"crossref"},{"observed_at":"2026-08-11T06:34:36.301508+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T04:31:34.561665Z","title":null,"venue":null,"work_id":"794e44b9-0d9a-47cf-ad12-0a440f4716bc","year":2024},"citing_paper":{"arxiv_id":"2506.21563","last_updated":"2025-06-12T07:02:28Z","snapshot_observed_at":"2026-08-08T09:20:55.054492Z","submitted_at":"2025-06-12T07:02:28Z","title":"FormosanBench: Benchmarking Low-Resource Austronesian Languages in the Era of Large Language Models","version":1},"reference_index":35,"source":"arxiv_source","source_observed_at":"2026-08-07T04:31:32.850555Z"},"links":{"citing_paper":"/paper/2506.21563"},"observation_digest":"sha256:1a1107a1f79e29c37af0d4553674bbe0bba14f7e5cc31a731444bf98899b1dd4","observation_id":"a2ebfee1-8714-4f0b-9ffc-534a28e09ffb","resolution":{"observed_at":"2026-08-07T04:31:34.621348Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-11T06:34:44.6726+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-11T06:34:44.6726+00:00","source":"crossref"},{"observed_at":"2026-08-11T06:34:36.301508+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":"10.18653/v1/2024.findings-acl.670","metadata_source":"openalex","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-10T05:30:23.456663Z","title":null,"venue":null,"work_id":"9d790a51-4b0f-423b-8f05-8abf0741cad0","year":2024},"citing_paper":{"arxiv_id":"2506.21563","last_updated":"2025-06-12T07:02:28Z","snapshot_observed_at":"2026-08-08T09:20:55.054492Z","submitted_at":"2025-06-12T07:02:28Z","title":"FormosanBench: Benchmarking Low-Resource Austronesian Languages in the Era of Large Language Models","version":1},"reference_index":36,"source":"arxiv_source","source_observed_at":"2026-08-07T04:31:32.995797Z"},"links":{"citing_paper":"/paper/2506.21563"},"observation_digest":"sha256:ff8e0fa9db0f76cb4b8e83f9718a69d2d433d63945923dcfe26f90c1724ad52b","observation_id":"588a99b2-7339-48aa-a86f-7e162a7d0f47","resolution":{"observed_at":"2026-08-07T04:31:33.278902Z","resolver_source":"doi","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-11T06:34:44.6726+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-11T06:34:44.6726+00:00","source":"crossref"},{"observed_at":"2026-08-11T06:34:36.301508+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T04:31:33.051143Z","title":"online\" 'onlinestring :=","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2506.21563","last_updated":"2025-06-12T07:02:28Z","snapshot_observed_at":"2026-08-08T09:20:55.054492Z","submitted_at":"2025-06-12T07:02:28Z","title":"FormosanBench: Benchmarking Low-Resource Austronesian Languages in the Era of Large Language Models","version":1},"reference_index":37,"source":"arxiv_source","source_observed_at":"2026-08-07T04:31:33.051143Z"},"links":{"citing_paper":"/paper/2506.21563"},"observation_digest":"sha256:c9c842a9a4768060b1912f30fe5c1ff29a95298c334225757a33c92fbffe9576","observation_id":"97f5be0a-62b4-4625-a7cd-b2cb78bc6265","resolution":{"observed_at":"2026-08-07T04:31:33.051143Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T04:31:33.099103Z","title":"write newline","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2506.21563","last_updated":"2025-06-12T07:02:28Z","snapshot_observed_at":"2026-08-08T09:20:55.054492Z","submitted_at":"2025-06-12T07:02:28Z","title":"FormosanBench: Benchmarking Low-Resource Austronesian Languages in the Era of Large Language Models","version":1},"reference_index":38,"source":"arxiv_source","source_observed_at":"2026-08-07T04:31:33.099103Z"},"links":{"citing_paper":"/paper/2506.21563"},"observation_digest":"sha256:b2562972881050fdf078d0608d00d90d225092b9b8208882c881f78ba15256d4","observation_id":"eaae9465-ab14-4a3e-b57a-f400572bdb69","resolution":{"observed_at":"2026-08-07T04:31:33.099103Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}}],"paper":{"arxiv_id":"2506.21563","last_updated":"2025-06-12T07:02:28Z","latest_version":1,"primary_category":"cs.CL","snapshot_observed_at":"2026-08-08T09:20:55.054492Z","submitted_at":"2025-06-12T07:02:28Z","title":"FormosanBench: Benchmarking Low-Resource Austronesian Languages in the Era of Large Language Models"},"reference_resolution":{"displayed":37,"state_counts":{"malformed_identifier":0,"metadata_mismatch":1,"parse_uncertain":0,"unresolved":28,"verified_exact":8,"verified_fuzzy":0},"total_outbound_references":37},"refusal":"A citation records a reference. It does not transfer a finding from one paper to another.","schema":"pith.paper-citation-record.v1","standing_sources":[{"observed_at":"2026-08-11T06:34:44.6726+00:00","source":"crossref"},{"observed_at":"2026-08-11T06:34:36.301508+00:00","source":"retraction_watch"}],"thesis":"As of 11 August 2026, this Paper Citation Record lists 37 of 37 outbound references and 0 inbound Pith citation observations for arXiv:2506.21563."}