{"as_of":"2026-08-10T13:31:00Z","caps":{"database_statements":6,"inbound":100,"outbound":100},"context_digest":"sha256:a91874de7eb657263b6db7b1e910986f9015eab7d061b4dd818169fbc0a65ceb","coverage":[{"denominator":35,"lane":"reference_resolution","note":"Typed states for the displayed outbound observations.","records_observed":35,"source":"paper_references, paper_reference_links","source_observed_at":"2026-08-07T11:37:06.635382Z","state":"measured"},{"denominator":35,"lane":"standing_notices","note":"One-hop event checks from named stored sources.","records_observed":35,"source":"scholarly_work_events, retraction_status_cache","source_observed_at":"2026-08-10T06:31:04.303077+00:00","state":"measured"},{"denominator":0,"lane":"inbound_itemization","note":"Pith citing papers itemized under the disclosed page cap.","records_observed":0,"source":"paper_references, paper_reference_links","source_observed_at":null,"state":"measured"},{"denominator":1,"lane":"external_citation_measurements","note":"A source-named dated measurement, never combined with another source.","records_observed":0,"source":"cited_works","source_observed_at":null,"state":"measured"}],"external_citation_measurements":[],"inbound":[],"links":{"evidence":"/evidence","html":"/paper/2506.01808/citation-record","integrity":"/paper/2506.01808/integrity","json":"/paper/2506.01808/citation-record.json","paper":"/paper/2506.01808"},"outbound":[{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T11:37:07.071330Z","title":"McCrae, Salima Mdhaffar, Yasmin Moslem, Kenton Murray, Satoshi Nakamura, Matteo Negri, Jan Niehues, Atul Kr","venue":null,"work_id":"17071f52-35dc-467d-b212-6c7298eeaaf0","year":2025},"citing_paper":{"arxiv_id":"2506.01808","last_updated":"2025-06-02T15:52:57Z","snapshot_observed_at":"2026-08-09T16:51:07.160206Z","submitted_at":"2025-06-02T15:52:57Z","title":"NAVER LABS Europe Submission to the Instruction-following Track","version":1},"reference_index":1,"source":"pdf_text","source_observed_at":"2026-08-07T11:37:06.499009Z"},"links":{"citing_paper":"/paper/2506.01808"},"observation_digest":"sha256:cdafa35c34eb61595fdde6f62bbb791fee590939232acf13bd9b147a0a5103be","observation_id":"7c3fec06-d499-4612-b211-78e096641959","resolution":{"observed_at":"2026-08-07T11:37:07.074974Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2303.08774","last_updated":"2024-03-04T06:01:33Z","snapshot_observed_at":"2026-08-07T07:30:12.213965Z","submitted_at":"2023-03-15T17:15:04Z","title":"GPT-4 Technical Report","version":6},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2303.08774","snapshot_observed_at":"2026-08-07T11:37:06.503289Z","title":"Gpt-4 technical report.arXiv preprint arXiv:2303.08774, 2023","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2506.01808","last_updated":"2025-06-02T15:52:57Z","snapshot_observed_at":"2026-08-09T16:51:07.160206Z","submitted_at":"2025-06-02T15:52:57Z","title":"NAVER LABS Europe Submission to the Instruction-following Track","version":1},"reference_index":2,"source":"pdf_text","source_observed_at":"2026-08-07T11:37:06.503289Z"},"links":{"cited_paper":"/paper/2303.08774","citing_paper":"/paper/2506.01808"},"observation_digest":"sha256:0bf89186447809aed5636dfad281b09d7c4769fb2adeac49dc2d395b00a413ea","observation_id":"69e96998-aabf-4311-952d-110a99aacc46","resolution":{"observed_at":"2026-08-07T11:37:06.503289Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T11:37:07.060591Z","title":null,"venue":null,"work_id":"6852bd0f-fa03-4991-aa47-d484fafbdc01","year":2024},"citing_paper":{"arxiv_id":"2506.01808","last_updated":"2025-06-02T15:52:57Z","snapshot_observed_at":"2026-08-09T16:51:07.160206Z","submitted_at":"2025-06-02T15:52:57Z","title":"NAVER LABS Europe Submission to the Instruction-following Track","version":1},"reference_index":3,"source":"pdf_text","source_observed_at":"2026-08-07T11:37:06.506761Z"},"links":{"citing_paper":"/paper/2506.01808"},"observation_digest":"sha256:84ed238feda1b9250b6eb43b202e4a4d0c01d14d750ec54cba22c230d575d6dc","observation_id":"713f3730-07bf-4cb9-b4ff-ef6cf8a553d3","resolution":{"observed_at":"2026-08-07T11:37:07.063918Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T11:37:06.510426Z","title":"From tower to spire: Adding the speech modality to a text-only llm.arXiv preprint arXiv:2503.10620, 2025","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2506.01808","last_updated":"2025-06-02T15:52:57Z","snapshot_observed_at":"2026-08-09T16:51:07.160206Z","submitted_at":"2025-06-02T15:52:57Z","title":"NAVER LABS Europe Submission to the Instruction-following Track","version":1},"reference_index":4,"source":"pdf_text","source_observed_at":"2026-08-07T11:37:06.510426Z"},"links":{"citing_paper":"/paper/2506.01808"},"observation_digest":"sha256:2b4f080421ce8bf304f1ead74f7b46e708963750eaa859d0ab58f559af116a58","observation_id":"a48b94cd-a300-44c4-a7d9-51d03d7d982d","resolution":{"observed_at":"2026-08-07T11:37:06.510426Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2308.11596","last_updated":"2023-10-25T03:52:07Z","snapshot_observed_at":"2026-07-06T16:09:09.018523Z","submitted_at":"2023-08-22T17:44:18Z","title":"SeamlessM4T: Massively Multilingual & Multimodal Machine Translation","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2308.11596","snapshot_observed_at":"2026-08-07T11:37:06.514410Z","title":"Seamlessm4t: Massively multilingual & multimodal machine translation","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2506.01808","last_updated":"2025-06-02T15:52:57Z","snapshot_observed_at":"2026-08-09T16:51:07.160206Z","submitted_at":"2025-06-02T15:52:57Z","title":"NAVER LABS Europe Submission to the Instruction-following Track","version":1},"reference_index":5,"source":"pdf_text","source_observed_at":"2026-08-07T11:37:06.514410Z"},"links":{"cited_paper":"/paper/2308.11596","citing_paper":"/paper/2506.01808"},"observation_digest":"sha256:cf384d4fff18bd303650fc72aa5c9e0f3e7a4f6af87ae4364059ea365aa43eaa","observation_id":"29dbeb19-39f9-4ba6-ace1-20d89998f4fe","resolution":{"observed_at":"2026-08-07T11:37:06.514410Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2410.00037","last_updated":"2024-10-02T09:11:45Z","snapshot_observed_at":"2026-07-30T10:21:14.474746Z","submitted_at":"2024-09-17T17:55:39Z","title":"Moshi: a speech-text foundation model for real-time dialogue","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2410.00037","snapshot_observed_at":"2026-08-07T11:37:06.518264Z","title":"Moshi: a speech-text foun- dation model for real-time dialogue","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2506.01808","last_updated":"2025-06-02T15:52:57Z","snapshot_observed_at":"2026-08-09T16:51:07.160206Z","submitted_at":"2025-06-02T15:52:57Z","title":"NAVER LABS Europe Submission to the Instruction-following Track","version":1},"reference_index":6,"source":"pdf_text","source_observed_at":"2026-08-07T11:37:06.518264Z"},"links":{"cited_paper":"/paper/2410.00037","citing_paper":"/paper/2506.01808"},"observation_digest":"sha256:05637fcc573ed1c3f1a9f6786c6f38f5aa93a70a34e2c32afb075a007168416f","observation_id":"9e235be1-43f9-4d9c-a77c-18d8c2621ad0","resolution":{"observed_at":"2026-08-07T11:37:06.518264Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2303.03378","last_updated":"2023-03-06T18:58:06Z","snapshot_observed_at":"2026-08-08T22:04:13.117781Z","submitted_at":"2023-03-06T18:58:06Z","title":"PaLM-E: An Embodied Multimodal Language Model","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2303.03378","snapshot_observed_at":"2026-08-07T11:37:06.522198Z","title":"Palm- e: An embodied multimodal language model.arXiv preprint arXiv:2303.03378, 2023","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2506.01808","last_updated":"2025-06-02T15:52:57Z","snapshot_observed_at":"2026-08-09T16:51:07.160206Z","submitted_at":"2025-06-02T15:52:57Z","title":"NAVER LABS Europe Submission to the Instruction-following Track","version":1},"reference_index":7,"source":"pdf_text","source_observed_at":"2026-08-07T11:37:06.522198Z"},"links":{"cited_paper":"/paper/2303.03378","citing_paper":"/paper/2506.01808"},"observation_digest":"sha256:4f578a89a7c77a83e3c1d01f4b090fb7ccc5d3a9bc5c15b9ba65c84877fd99d6","observation_id":"6c8ffb9c-fd9f-4064-bc0c-1f614b804218","resolution":{"observed_at":"2026-08-07T11:37:06.522198Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2407.21783","last_updated":"2024-11-23T23:27:33Z","snapshot_observed_at":"2026-07-06T18:55:11.576666Z","submitted_at":"2024-07-31T17:54:27Z","title":"The Llama 3 Herd of Models","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2407.21783","snapshot_observed_at":"2026-08-07T11:37:06.526560Z","title":"The llama 3 herd of models","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2506.01808","last_updated":"2025-06-02T15:52:57Z","snapshot_observed_at":"2026-08-09T16:51:07.160206Z","submitted_at":"2025-06-02T15:52:57Z","title":"NAVER LABS Europe Submission to the Instruction-following Track","version":1},"reference_index":8,"source":"pdf_text","source_observed_at":"2026-08-07T11:37:06.526560Z"},"links":{"cited_paper":"/paper/2407.21783","citing_paper":"/paper/2506.01808"},"observation_digest":"sha256:a74375099dac55eb549368ffc3fe0c252099c2606986d1f65e2c254a21090faa","observation_id":"78473648-fb55-408c-becd-bcd5199fc2f3","resolution":{"observed_at":"2026-08-07T11:37:06.526560Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T11:37:07.049298Z","title":"LoRA: Low-rank adaptation of large language 7 NAVER LABS Europe Submission to the Instruction-following Track models","venue":null,"work_id":"9ad09866-a0dd-48d5-9529-596e889cd8a5","year":2022},"citing_paper":{"arxiv_id":"2506.01808","last_updated":"2025-06-02T15:52:57Z","snapshot_observed_at":"2026-08-09T16:51:07.160206Z","submitted_at":"2025-06-02T15:52:57Z","title":"NAVER LABS Europe Submission to the Instruction-following Track","version":1},"reference_index":9,"source":"pdf_text","source_observed_at":"2026-08-07T11:37:06.535763Z"},"links":{"citing_paper":"/paper/2506.01808"},"observation_digest":"sha256:edb8576ea60b119c0ec8b1b264806eaca37ad5b51bdbfc3daab46359175fa962","observation_id":"0527e5fb-6d4e-4c48-a7bc-12898a03eb3a","resolution":{"observed_at":"2026-08-07T11:37:07.053020Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2404.00656","last_updated":"2024-09-21T15:27:30Z","snapshot_observed_at":"2026-08-09T20:04:27.847882Z","submitted_at":"2024-03-31T12:01:32Z","title":"WavLLM: Towards Robust and Adaptive Speech Large Language Model","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2404.00656","snapshot_observed_at":"2026-08-07T11:37:06.539254Z","title":"Wavllm: Towards ro- bust and adaptive speech large language model.arXiv preprint arXiv:2404.00656, 2024","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2506.01808","last_updated":"2025-06-02T15:52:57Z","snapshot_observed_at":"2026-08-09T16:51:07.160206Z","submitted_at":"2025-06-02T15:52:57Z","title":"NAVER LABS Europe Submission to the Instruction-following Track","version":1},"reference_index":10,"source":"pdf_text","source_observed_at":"2026-08-07T11:37:06.539254Z"},"links":{"cited_paper":"/paper/2404.00656","citing_paper":"/paper/2506.01808"},"observation_digest":"sha256:dae3a2cc011d48a090481eaefee59f65983814ed736529ace42eb90e607c3f5d","observation_id":"e44347e4-7050-40ad-adc0-08546cffa511","resolution":{"observed_at":"2026-08-07T11:37:06.539254Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T11:37:06.543185Z","title":"Audiogpt: Understand- ing and generating speech, music, sound, and talking head","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2506.01808","last_updated":"2025-06-02T15:52:57Z","snapshot_observed_at":"2026-08-09T16:51:07.160206Z","submitted_at":"2025-06-02T15:52:57Z","title":"NAVER LABS Europe Submission to the Instruction-following Track","version":1},"reference_index":11,"source":"pdf_text","source_observed_at":"2026-08-07T11:37:06.543185Z"},"links":{"citing_paper":"/paper/2506.01808"},"observation_digest":"sha256:3014912c3c13e6ac2c7d802e18ad0e3aeee38f504058d51a57483452c72b181d","observation_id":"8eda5c9f-a019-4a83-abf5-9999fef582d8","resolution":{"observed_at":"2026-08-07T11:37:06.543185Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T11:37:06.546654Z","title":"Iranzo-Sánchez, J","venue":null,"work_id":null,"year":2020},"citing_paper":{"arxiv_id":"2506.01808","last_updated":"2025-06-02T15:52:57Z","snapshot_observed_at":"2026-08-09T16:51:07.160206Z","submitted_at":"2025-06-02T15:52:57Z","title":"NAVER LABS Europe Submission to the Instruction-following Track","version":1},"reference_index":12,"source":"pdf_text","source_observed_at":"2026-08-07T11:37:06.546654Z"},"links":{"citing_paper":"/paper/2506.01808"},"observation_digest":"sha256:086552340605099b46d9af05f91d03a09aef0d0ade45fdee8d1385ac22582791","observation_id":"268ac7ca-3cac-4d7d-aa18-08e6c5e59813","resolution":{"observed_at":"2026-08-07T11:37:06.546654Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2401.04088","last_updated":"2024-01-08T18:47:34Z","snapshot_observed_at":"2026-08-08T06:16:25.839566Z","submitted_at":"2024-01-08T18:47:34Z","title":"Mixtral of Experts","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2401.04088","snapshot_observed_at":"2026-08-07T11:37:06.550805Z","title":"Mixtral of experts.arXiv preprint arXiv:2401.04088, 2024","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2506.01808","last_updated":"2025-06-02T15:52:57Z","snapshot_observed_at":"2026-08-09T16:51:07.160206Z","submitted_at":"2025-06-02T15:52:57Z","title":"NAVER LABS Europe Submission to the Instruction-following Track","version":1},"reference_index":13,"source":"pdf_text","source_observed_at":"2026-08-07T11:37:06.550805Z"},"links":{"cited_paper":"/paper/2401.04088","citing_paper":"/paper/2506.01808"},"observation_digest":"sha256:879cf5c304550ce9c3747b2e0a258cf60f9b94599912e487a0c7a73edac3bf02","observation_id":"a2d57ed4-61f6-49a2-b936-66c799a19564","resolution":{"observed_at":"2026-08-07T11:37:06.550805Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2405.02246","last_updated":"2024-05-03T17:00:00Z","snapshot_observed_at":"2026-08-08T17:33:57.727194Z","submitted_at":"2024-05-03T17:00:00Z","title":"What matters when building vision-language models?","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2405.02246","snapshot_observed_at":"2026-08-07T11:37:06.554978Z","title":"What matters when building vision- language models?, 2024","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2506.01808","last_updated":"2025-06-02T15:52:57Z","snapshot_observed_at":"2026-08-09T16:51:07.160206Z","submitted_at":"2025-06-02T15:52:57Z","title":"NAVER LABS Europe Submission to the Instruction-following Track","version":1},"reference_index":14,"source":"pdf_text","source_observed_at":"2026-08-07T11:37:06.554978Z"},"links":{"cited_paper":"/paper/2405.02246","citing_paper":"/paper/2506.01808"},"observation_digest":"sha256:cdaca72e83190a16f3396c16a62e3480ff16faea54c4ef8f01218a5afbb4ae6a","observation_id":"7b786262-dace-4649-815d-49deec82777d","resolution":{"observed_at":"2026-08-07T11:37:06.554978Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T11:37:07.025794Z","title":"Spoken squad: A study of mitigating the impact ofspeechrecognitionerrorsonlisteningcomprehension","venue":null,"work_id":"64e5c969-bffb-4b44-b62b-f024a907f210","year":2018},"citing_paper":{"arxiv_id":"2506.01808","last_updated":"2025-06-02T15:52:57Z","snapshot_observed_at":"2026-08-09T16:51:07.160206Z","submitted_at":"2025-06-02T15:52:57Z","title":"NAVER LABS Europe Submission to the Instruction-following Track","version":1},"reference_index":15,"source":"pdf_text","source_observed_at":"2026-08-07T11:37:06.558652Z"},"links":{"citing_paper":"/paper/2506.01808"},"observation_digest":"sha256:1e46bbf061688f39293840d1828155cf5934aea542597721c925195df1ce06ab","observation_id":"90ff5998-a95c-4657-b95c-5e87a9c4e4a7","resolution":{"observed_at":"2026-08-07T11:37:07.029748Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T11:37:07.014292Z","title":"ROUGE: A package for automatic eval- uation of summaries","venue":null,"work_id":"20b7cf8d-348c-4f06-9556-b8240c344b73","year":2004},"citing_paper":{"arxiv_id":"2506.01808","last_updated":"2025-06-02T15:52:57Z","snapshot_observed_at":"2026-08-09T16:51:07.160206Z","submitted_at":"2025-06-02T15:52:57Z","title":"NAVER LABS Europe Submission to the Instruction-following Track","version":1},"reference_index":16,"source":"pdf_text","source_observed_at":"2026-08-07T11:37:06.562104Z"},"links":{"citing_paper":"/paper/2506.01808"},"observation_digest":"sha256:ea3ce727cccdf4b22dc61789c939ad91e6d09ac0480cbf6e01444d488e79e318","observation_id":"67ea7aa6-3736-46d4-9236-bc298fdd2d9d","resolution":{"observed_at":"2026-08-07T11:37:07.018393Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2304.08485","last_updated":"2023-12-11T17:46:14Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2023-04-17T17:59:25Z","title":"Visual Instruction Tuning","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2304.08485","snapshot_observed_at":"2026-08-07T11:37:06.566706Z","title":"Visual Instruction Tuning (LLaVA), 2023","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2506.01808","last_updated":"2025-06-02T15:52:57Z","snapshot_observed_at":"2026-08-09T16:51:07.160206Z","submitted_at":"2025-06-02T15:52:57Z","title":"NAVER LABS Europe Submission to the Instruction-following Track","version":1},"reference_index":17,"source":"pdf_text","source_observed_at":"2026-08-07T11:37:06.566706Z"},"links":{"cited_paper":"/paper/2304.08485","citing_paper":"/paper/2506.01808"},"observation_digest":"sha256:8cf3aed6b8692b479677714970f6b32d5595eddab2df31bd30bd31ba4f3f3834","observation_id":"529dc7ec-0d63-4f5c-b0c1-57b19886e295","resolution":{"observed_at":"2026-08-07T11:37:06.566706Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T11:37:07.002983Z","title":"Guerreiro, Ricardo Rei, Duarte M","venue":null,"work_id":"dee13529-43a1-4b82-b511-0bb75bb7613b","year":2024},"citing_paper":{"arxiv_id":"2506.01808","last_updated":"2025-06-02T15:52:57Z","snapshot_observed_at":"2026-08-09T16:51:07.160206Z","submitted_at":"2025-06-02T15:52:57Z","title":"NAVER LABS Europe Submission to the Instruction-following Track","version":1},"reference_index":18,"source":"pdf_text","source_observed_at":"2026-08-07T11:37:06.571065Z"},"links":{"citing_paper":"/paper/2506.01808"},"observation_digest":"sha256:f24213d6e5a9dee81c1cbe803da3f4788b1ec848a16eed95e37557980fdcfd29","observation_id":"48b3f722-9daa-4496-a77a-3203c0b1c53c","resolution":{"observed_at":"2026-08-07T11:37:07.007019Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2409.16235","last_updated":"2024-09-24T16:51:36Z","snapshot_observed_at":"2026-07-06T19:21:20.826877Z","submitted_at":"2024-09-24T16:51:36Z","title":"EuroLLM: Multilingual Language Models for Europe","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2409.16235","snapshot_observed_at":"2026-08-07T11:37:06.575250Z","title":"Eurollm: Multilingual language mod- els for europe.arXiv preprint arXiv:2409.16235, 2024","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2506.01808","last_updated":"2025-06-02T15:52:57Z","snapshot_observed_at":"2026-08-09T16:51:07.160206Z","submitted_at":"2025-06-02T15:52:57Z","title":"NAVER LABS Europe Submission to the Instruction-following Track","version":1},"reference_index":19,"source":"pdf_text","source_observed_at":"2026-08-07T11:37:06.575250Z"},"links":{"cited_paper":"/paper/2409.16235","citing_paper":"/paper/2506.01808"},"observation_digest":"sha256:1ee044cee0e0477a1db79eeb8ec8e32f3c8800510bed3bbc12ae8c64d04f366f","observation_id":"2920b360-4f54-485c-a034-4d39e88c725b","resolution":{"observed_at":"2026-08-07T11:37:06.575250Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T11:37:06.579223Z","title":"Spirit-lm: Interleaved spoken and written language model.Transactions of the Association for Computational Linguistics, 13:30–52, 2025","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2506.01808","last_updated":"2025-06-02T15:52:57Z","snapshot_observed_at":"2026-08-09T16:51:07.160206Z","submitted_at":"2025-06-02T15:52:57Z","title":"NAVER LABS Europe Submission to the Instruction-following Track","version":1},"reference_index":20,"source":"pdf_text","source_observed_at":"2026-08-07T11:37:06.579223Z"},"links":{"citing_paper":"/paper/2506.01808"},"observation_digest":"sha256:dd51f21f1a90d783ba11160dc1dfa4845ffdd0dea694a409c4391d466f94b1dd","observation_id":"5fdc741b-b046-466f-a913-11c68e3f2221","resolution":{"observed_at":"2026-08-07T11:37:06.579223Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2503.22577","last_updated":"2025-05-20T10:29:41Z","snapshot_observed_at":"2026-08-07T16:29:48.727258Z","submitted_at":"2025-03-28T16:26:52Z","title":"Breaking Language Barriers in Visual Language Models via Multilingual Textual Regularization","version":2},"cited_work":{"arxiv_id":"2503.22577","doi":null,"metadata_source":"pith","pith_arxiv_id":"2503.22577","snapshot_observed_at":"2026-08-07T11:37:06.706367Z","title":"Breaking Language Barriers in Visual Language Models via Multilingual Textual Regularization","venue":"cs.CV","work_id":"c19ef716-546f-4228-9f3e-6c8ed4748b25","year":2025},"citing_paper":{"arxiv_id":"2506.01808","last_updated":"2025-06-02T15:52:57Z","snapshot_observed_at":"2026-08-09T16:51:07.160206Z","submitted_at":"2025-06-02T15:52:57Z","title":"NAVER LABS Europe Submission to the Instruction-following Track","version":1},"reference_index":21,"source":"pdf_text","source_observed_at":"2026-08-07T11:37:06.582912Z"},"links":{"cited_paper":"/paper/2503.22577","citing_paper":"/paper/2506.01808"},"observation_digest":"sha256:bd46978f58869fa803c818426d8db707fe5f746258dd6e3ef6e3c696a35e4b95","observation_id":"8517dfc8-5bcd-484f-aed7-718695d290c6","resolution":{"observed_at":"2026-08-07T11:37:06.712079Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T11:37:06.985809Z","title":"A call for clarity in reporting BLEU scores","venue":null,"work_id":"22eafb73-4382-4b01-a374-521068d3beb7","year":null},"citing_paper":{"arxiv_id":"2506.01808","last_updated":"2025-06-02T15:52:57Z","snapshot_observed_at":"2026-08-09T16:51:07.160206Z","submitted_at":"2025-06-02T15:52:57Z","title":"NAVER LABS Europe Submission to the Instruction-following Track","version":1},"reference_index":22,"source":"pdf_text","source_observed_at":"2026-08-07T11:37:06.586670Z"},"links":{"citing_paper":"/paper/2506.01808"},"observation_digest":"sha256:d5c5a5756d800b8c93b401cadc802ad208dabfc2884cba450c92b63c8c0cb6cd","observation_id":"8f305645-c550-446f-abe4-20f706ecd273","resolution":{"observed_at":"2026-08-07T11:37:06.989451Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T11:37:06.594248Z","title":"Scaling speech technology to 1,000+ languages.Jour- nal of Machine Learning Research, 25(97):1–52, 2024","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2506.01808","last_updated":"2025-06-02T15:52:57Z","snapshot_observed_at":"2026-08-09T16:51:07.160206Z","submitted_at":"2025-06-02T15:52:57Z","title":"NAVER LABS Europe Submission to the Instruction-following Track","version":1},"reference_index":23,"source":"pdf_text","source_observed_at":"2026-08-07T11:37:06.594248Z"},"links":{"citing_paper":"/paper/2506.01808"},"observation_digest":"sha256:f81535f9a60915a62194215515b2ada334afad96ac162d4d4a5eb1410a855761","observation_id":"8e793ba6-d5c3-4bef-b36d-73488094783e","resolution":{"observed_at":"2026-08-07T11:37:06.594248Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T11:37:06.955202Z","title":"BERGEN: A benchmarking library for retrieval-augmented generation","venue":null,"work_id":"decb62ae-3ad9-46ea-9b4b-2419e4ee1de0","year":2024},"citing_paper":{"arxiv_id":"2506.01808","last_updated":"2025-06-02T15:52:57Z","snapshot_observed_at":"2026-08-09T16:51:07.160206Z","submitted_at":"2025-06-02T15:52:57Z","title":"NAVER LABS Europe Submission to the Instruction-following Track","version":1},"reference_index":24,"source":"pdf_text","source_observed_at":"2026-08-07T11:37:06.598180Z"},"links":{"citing_paper":"/paper/2506.01808"},"observation_digest":"sha256:5afa0914dc9d91fde0550508aa3335624ab3194c54d8a86482ee0f68a0a0001d","observation_id":"28db40f3-5550-4fa2-bf91-c9a8a521bed4","resolution":{"observed_at":"2026-08-07T11:37:06.959433Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T11:37:06.943700Z","title":null,"venue":null,"work_id":"29e7a67f-9cc6-4b04-97b7-8390235da712","year":2022},"citing_paper":{"arxiv_id":"2506.01808","last_updated":"2025-06-02T15:52:57Z","snapshot_observed_at":"2026-08-09T16:51:07.160206Z","submitted_at":"2025-06-02T15:52:57Z","title":"NAVER LABS Europe Submission to the Instruction-following Track","version":1},"reference_index":25,"source":"pdf_text","source_observed_at":"2026-08-07T11:37:06.601596Z"},"links":{"citing_paper":"/paper/2506.01808"},"observation_digest":"sha256:d712e1a33441c459ee1c6244faeb833fb2e009f145ff922b7f9504c5c3cce39a","observation_id":"fca08bbe-71e6-472f-954f-5c7bae3edea8","resolution":{"observed_at":"2026-08-07T11:37:06.947205Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2306.12925","last_updated":"2023-06-22T14:37:54Z","snapshot_observed_at":"2026-08-07T01:18:02.068918Z","submitted_at":"2023-06-22T14:37:54Z","title":"AudioPaLM: A Large Language Model That Can Speak and Listen","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2306.12925","snapshot_observed_at":"2026-08-07T11:37:06.604936Z","title":"Audiopalm: A large language model that can speak and listen.arXiv preprint arXiv:2306.12925, 2023","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2506.01808","last_updated":"2025-06-02T15:52:57Z","snapshot_observed_at":"2026-08-09T16:51:07.160206Z","submitted_at":"2025-06-02T15:52:57Z","title":"NAVER LABS Europe Submission to the Instruction-following Track","version":1},"reference_index":26,"source":"pdf_text","source_observed_at":"2026-08-07T11:37:06.604936Z"},"links":{"cited_paper":"/paper/2306.12925","citing_paper":"/paper/2506.01808"},"observation_digest":"sha256:8ec57622e83aa68d17d6909728e2528514c190876a62cc658d5ab2939ad0b02d","observation_id":"f79628f3-5985-4d3b-a747-e3bf10535588","resolution":{"observed_at":"2026-08-07T11:37:06.604936Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T11:37:06.932681Z","title":"Evaluating multilingual speech translation under realistic condi- tions with resegmentation and terminology","venue":null,"work_id":"2979ad94-6df9-48e1-b264-846c046d57fb","year":2023},"citing_paper":{"arxiv_id":"2506.01808","last_updated":"2025-06-02T15:52:57Z","snapshot_observed_at":"2026-08-09T16:51:07.160206Z","submitted_at":"2025-06-02T15:52:57Z","title":"NAVER LABS Europe Submission to the Instruction-following Track","version":1},"reference_index":27,"source":"pdf_text","source_observed_at":"2026-08-07T11:37:06.608488Z"},"links":{"citing_paper":"/paper/2506.01808"},"observation_digest":"sha256:f99321d8b4f60161b4afbe817b47ae7a907270a512da399abed6561e2138b01a","observation_id":"0b0423d7-3a1f-4ae8-870d-1fdb098dbd99","resolution":{"observed_at":"2026-08-07T11:37:06.936469Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2310.13289","last_updated":"2024-04-08T06:12:52Z","snapshot_observed_at":"2026-08-02T21:51:36.809095Z","submitted_at":"2023-10-20T05:41:57Z","title":"SALMONN: Towards Generic Hearing Abilities for Large Language Models","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2310.13289","snapshot_observed_at":"2026-08-07T11:37:06.612173Z","title":"Salmonn: Towardsgenerichearingabilitiesforlargelan- guage models.arXiv preprint arXiv:2310.13289, 2023","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2506.01808","last_updated":"2025-06-02T15:52:57Z","snapshot_observed_at":"2026-08-09T16:51:07.160206Z","submitted_at":"2025-06-02T15:52:57Z","title":"NAVER LABS Europe Submission to the Instruction-following Track","version":1},"reference_index":28,"source":"pdf_text","source_observed_at":"2026-08-07T11:37:06.612173Z"},"links":{"cited_paper":"/paper/2310.13289","citing_paper":"/paper/2506.01808"},"observation_digest":"sha256:13646131542bcfc5d14321dd988b38ddc86625c28d81e500725ca8a1546ecdab","observation_id":"ed556425-ebb2-49de-852d-dd04fbaa7627","resolution":{"observed_at":"2026-08-07T11:37:06.612173Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2312.11805","last_updated":"2025-05-09T21:04:06Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2023-12-19T02:39:27Z","title":"Gemini: A Family of Highly Capable Multimodal Models","version":5},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2312.11805","snapshot_observed_at":"2026-08-07T11:37:06.616463Z","title":"Gemini: a family of highly capable multimodal models","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2506.01808","last_updated":"2025-06-02T15:52:57Z","snapshot_observed_at":"2026-08-09T16:51:07.160206Z","submitted_at":"2025-06-02T15:52:57Z","title":"NAVER LABS Europe Submission to the Instruction-following Track","version":1},"reference_index":29,"source":"pdf_text","source_observed_at":"2026-08-07T11:37:06.616463Z"},"links":{"cited_paper":"/paper/2312.11805","citing_paper":"/paper/2506.01808"},"observation_digest":"sha256:56274009b97ac9d89c2d2084905aff121875ae1b8dd415fee177b25f3cba0ef8","observation_id":"e494bf00-d5f7-4d07-9158-8e15878564ef","resolution":{"observed_at":"2026-08-07T11:37:06.616463Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T11:37:06.921289Z","title":null,"venue":null,"work_id":"435d6ca1-a044-4c57-85cb-8568dcda180d","year":2025},"citing_paper":{"arxiv_id":"2506.01808","last_updated":"2025-06-02T15:52:57Z","snapshot_observed_at":"2026-08-09T16:51:07.160206Z","submitted_at":"2025-06-02T15:52:57Z","title":"NAVER LABS Europe Submission to the Instruction-following Track","version":1},"reference_index":30,"source":"pdf_text","source_observed_at":"2026-08-07T11:37:06.620780Z"},"links":{"citing_paper":"/paper/2506.01808"},"observation_digest":"sha256:f9b095988de49064e67824f7b142ccb78a42edc9257f339f68d5a4d0d92cc22c","observation_id":"b934a0aa-6c4f-4856-b6ff-6ae52e3985c8","resolution":{"observed_at":"2026-08-07T11:37:06.925827Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T11:37:06.910153Z","title":"torchtune: Py- torch’s finetuning library, 2024","venue":null,"work_id":"b3400eb8-11f3-4929-8f89-59a9179e1ed2","year":2024},"citing_paper":{"arxiv_id":"2506.01808","last_updated":"2025-06-02T15:52:57Z","snapshot_observed_at":"2026-08-09T16:51:07.160206Z","submitted_at":"2025-06-02T15:52:57Z","title":"NAVER LABS Europe Submission to the Instruction-following Track","version":1},"reference_index":31,"source":"pdf_text","source_observed_at":"2026-08-07T11:37:06.625535Z"},"links":{"citing_paper":"/paper/2506.01808"},"observation_digest":"sha256:64a2517a45b12718aa387460ee90dae26717b2dd4321ac3515c41fa10a991212","observation_id":"59dcf11b-3005-4317-8278-092b9be9ec2e","resolution":{"observed_at":"2026-08-07T11:37:06.914231Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T11:37:06.899057Z","title":"Llama 2: Open foundation and fine-tuned chat models, 2023","venue":null,"work_id":"cda4a8dc-7603-43e4-b243-9aed4e89d58a","year":2023},"citing_paper":{"arxiv_id":"2506.01808","last_updated":"2025-06-02T15:52:57Z","snapshot_observed_at":"2026-08-09T16:51:07.160206Z","submitted_at":"2025-06-02T15:52:57Z","title":"NAVER LABS Europe Submission to the Instruction-following Track","version":1},"reference_index":32,"source":"pdf_text","source_observed_at":"2026-08-07T11:37:06.629098Z"},"links":{"citing_paper":"/paper/2506.01808"},"observation_digest":"sha256:0855b4d4c28c4d73b71003625bdfcd1031f5c5ab0133a5b559f4283311842b67","observation_id":"7205e1a5-7ee9-43a0-ace1-77133abc0661","resolution":{"observed_at":"2026-08-07T11:37:06.903300Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T11:37:06.632229Z","title":"Covost 2: A massively multilingual speech-to-text translation corpus, 2020","venue":null,"work_id":null,"year":2020},"citing_paper":{"arxiv_id":"2506.01808","last_updated":"2025-06-02T15:52:57Z","snapshot_observed_at":"2026-08-09T16:51:07.160206Z","submitted_at":"2025-06-02T15:52:57Z","title":"NAVER LABS Europe Submission to the Instruction-following Track","version":1},"reference_index":33,"source":"pdf_text","source_observed_at":"2026-08-07T11:37:06.632229Z"},"links":{"citing_paper":"/paper/2506.01808"},"observation_digest":"sha256:f4123bdbd3e94cdd82d5210b618788b51b2f4fd93f4e43fcf46919d110017b74","observation_id":"40ce9703-3feb-44d9-82d6-6a15240c47ca","resolution":{"observed_at":"2026-08-07T11:37:06.632229Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2412.15115","last_updated":"2025-01-03T02:18:21Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2024-12-19T17:56:09Z","title":"Qwen2.5 Technical Report","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2412.15115","snapshot_observed_at":"2026-08-07T11:37:06.635382Z","title":"1”, instead of “0","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2506.01808","last_updated":"2025-06-02T15:52:57Z","snapshot_observed_at":"2026-08-09T16:51:07.160206Z","submitted_at":"2025-06-02T15:52:57Z","title":"NAVER LABS Europe Submission to the Instruction-following Track","version":1},"reference_index":34,"source":"pdf_text","source_observed_at":"2026-08-07T11:37:06.635382Z"},"links":{"cited_paper":"/paper/2412.15115","citing_paper":"/paper/2506.01808"},"observation_digest":"sha256:61be4f4e87a2cb6ba8ed461fd86fd79c71675e43f688c8f68d7d1e66fb2bdac9","observation_id":"ad24d2c1-314d-4162-b1b1-2e27ab122559","resolution":{"observed_at":"2026-08-07T11:37:06.635382Z","resolver_source":null,"status":"malformed_identifier"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T11:37:06.974138Z","title":null,"venue":null,"work_id":"7cf6ac76-0edd-42ee-81a7-790925a30048","year":null},"citing_paper":{"arxiv_id":"2506.01808","last_updated":"2025-06-02T15:52:57Z","snapshot_observed_at":"2026-08-09T16:51:07.160206Z","submitted_at":"2025-06-02T15:52:57Z","title":"NAVER LABS Europe Submission to the Instruction-following Track","version":1},"reference_index":2018,"source":"pdf_text","source_observed_at":"2026-08-07T11:37:06.590475Z"},"links":{"citing_paper":"/paper/2506.01808"},"observation_digest":"sha256:affce8c8db903bc3ad9618f1ee90a8bb146f3a472c718a7917fab37e8069324e","observation_id":"da1a30ef-2e35-40d7-bb0a-2e32bdb6cd22","resolution":{"observed_at":"2026-08-07T11:37:06.978146Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}}],"paper":{"arxiv_id":"2506.01808","last_updated":"2025-06-02T15:52:57Z","latest_version":1,"primary_category":"cs.CL","snapshot_observed_at":"2026-08-09T16:51:07.160206Z","submitted_at":"2025-06-02T15:52:57Z","title":"NAVER LABS Europe Submission to the Instruction-following Track"},"reference_resolution":{"displayed":35,"state_counts":{"malformed_identifier":1,"metadata_mismatch":0,"parse_uncertain":0,"unresolved":23,"verified_exact":1,"verified_fuzzy":10},"total_outbound_references":35},"refusal":"A citation records a reference. It does not transfer a finding from one paper to another.","schema":"pith.paper-citation-record.v1","standing_sources":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"thesis":"As of 10 August 2026, this Paper Citation Record lists 35 of 35 outbound references and 0 inbound Pith citation observations for arXiv:2506.01808."}