{"as_of":"2026-08-10T03:00:00Z","caps":{"database_statements":6,"inbound":100,"outbound":100},"context_digest":"sha256:5068bd35524765d9525af4c04be9d59175087d329c584cb3e7386b447fcf6d4d","coverage":[{"denominator":0,"lane":"reference_resolution","note":"Typed states for the displayed outbound observations.","records_observed":0,"source":"paper_references, paper_reference_links","source_observed_at":null,"state":"measured"},{"denominator":44,"lane":"standing_notices","note":"One-hop event checks from named stored sources.","records_observed":44,"source":"scholarly_work_events, retraction_status_cache","source_observed_at":"2026-08-09T06:31:02.800959+00:00","state":"measured"},{"denominator":44,"lane":"inbound_itemization","note":"Pith citing papers itemized under the disclosed page cap.","records_observed":44,"source":"paper_references, paper_reference_links","source_observed_at":"2026-08-08T20:08:13.602909Z","state":"measured"},{"denominator":1,"lane":"external_citation_measurements","note":"A source-named dated measurement, never combined with another source.","records_observed":1,"source":"arxiv_reference","source_observed_at":"2026-08-05T02:28:24.338817Z","state":"measured"}],"external_citation_measurements":[{"count":50,"observed_at":"2026-08-05T02:28:24.338817Z","source":"arxiv_reference"}],"inbound":[{"citation":{"cited_paper":{"arxiv_id":"2306.04528","last_updated":"2024-07-16T07:29:49Z","snapshot_observed_at":"2026-08-07T20:45:44.253504Z","submitted_at":"2023-06-07T15:37:00Z","title":"PromptRobust: Towards Evaluating the Robustness of Large Language Models on Adversarial Prompts","version":5},"cited_work":{"arxiv_id":"2306.04528","doi":"10.48550/arxiv.2306.04528","metadata_source":"arxiv_reference","pith_arxiv_id":"2306.04528","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Promptbench: Towards evaluating the robustness of large language models on adversarial prompts","venue":"arXiv (Cornell University)","work_id":"e65e5731-8f09-44e1-94d5-1674f9e08988","year":2024},"citing_paper":{"arxiv_id":"2307.15043","last_updated":"2023-12-20T20:48:57Z","snapshot_observed_at":"2026-07-06T15:59:23.019044Z","submitted_at":"2023-07-27T17:49:12Z","title":"Universal and Transferable Adversarial Attacks on Aligned Language Models","version":2},"reference_index":29,"source":"pdf_text","source_observed_at":"2026-05-24T07:42:09.112946Z"},"links":{"cited_paper":"/paper/2306.04528","citing_paper":"/paper/2307.15043"},"observation_digest":"sha256:02490e0f7a549b2763317cbe1dcbba7250c2ddcafb2ccb203223bf315c552a70","observation_id":"2b2e6116-13c7-4016-8e76-f738ee64014d","resolution":{"observed_at":"2026-05-24T07:44:08.483210Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2306.04528","last_updated":"2024-07-16T07:29:49Z","snapshot_observed_at":"2026-08-07T20:45:44.253504Z","submitted_at":"2023-06-07T15:37:00Z","title":"PromptRobust: Towards Evaluating the Robustness of Large Language Models on Adversarial Prompts","version":5},"cited_work":{"arxiv_id":"2306.04528","doi":"10.48550/arxiv.2306.04528","metadata_source":"arxiv_reference","pith_arxiv_id":"2306.04528","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Promptbench: Towards evaluating the robustness of large language models on adversarial prompts","venue":"arXiv (Cornell University)","work_id":"e65e5731-8f09-44e1-94d5-1674f9e08988","year":2024},"citing_paper":{"arxiv_id":"2309.00614","last_updated":"2023-09-04T17:47:36Z","snapshot_observed_at":"2026-07-06T16:13:23.343694Z","submitted_at":"2023-09-01T17:59:44Z","title":"Baseline Defenses for Adversarial Attacks Against Aligned Language Models","version":2},"reference_index":65,"source":"arxiv_source","source_observed_at":"2026-05-13T23:24:39.835347Z"},"links":{"cited_paper":"/paper/2306.04528","citing_paper":"/paper/2309.00614"},"observation_digest":"sha256:971b22ec3d03c53a2a6957ee038195599f9f20a16db0569bcf626d9953c7924f","observation_id":"d4a95b09-3832-4a2f-bc31-9fac787648ee","resolution":{"observed_at":"2026-05-13T23:24:40.082946Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2306.04528","last_updated":"2024-07-16T07:29:49Z","snapshot_observed_at":"2026-08-07T20:45:44.253504Z","submitted_at":"2023-06-07T15:37:00Z","title":"PromptRobust: Towards Evaluating the Robustness of Large Language Models on Adversarial Prompts","version":5},"cited_work":{"arxiv_id":"2306.04528","doi":"10.48550/arxiv.2306.04528","metadata_source":"arxiv_reference","pith_arxiv_id":"2306.04528","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Promptbench: Towards evaluating the robustness of large language models on adversarial prompts","venue":"arXiv (Cornell University)","work_id":"e65e5731-8f09-44e1-94d5-1674f9e08988","year":2024},"citing_paper":{"arxiv_id":"2309.01219","last_updated":"2025-09-14T09:34:46Z","snapshot_observed_at":"2026-07-06T16:13:46.112815Z","submitted_at":"2023-09-03T16:56:48Z","title":"Siren's Song in the AI Ocean: A Survey on Hallucination in Large Language Models","version":3},"reference_index":30,"source":"pdf_text","source_observed_at":"2026-05-12T14:21:16.453610Z"},"links":{"cited_paper":"/paper/2306.04528","citing_paper":"/paper/2309.01219"},"observation_digest":"sha256:3d058e7d96381d5f995e3201c0b8cadd61fc5ea3694e1f312c0e2c95eb45aa45","observation_id":"c781a3b8-e332-4632-8ca8-f7ec654ec293","resolution":{"observed_at":"2026-05-12T14:21:16.477948Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2306.04528","last_updated":"2024-07-16T07:29:49Z","snapshot_observed_at":"2026-08-07T20:45:44.253504Z","submitted_at":"2023-06-07T15:37:00Z","title":"PromptRobust: Towards Evaluating the Robustness of Large Language Models on Adversarial Prompts","version":5},"cited_work":{"arxiv_id":"2306.04528","doi":"10.48550/arxiv.2306.04528","metadata_source":"arxiv_reference","pith_arxiv_id":"2306.04528","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Promptbench: Towards evaluating the robustness of large language models on adversarial prompts","venue":"arXiv (Cornell University)","work_id":"e65e5731-8f09-44e1-94d5-1674f9e08988","year":2024},"citing_paper":{"arxiv_id":"2309.08532","last_updated":"2025-05-01T11:56:52Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2023-09-15T16:50:09Z","title":"EvoPrompt: Connecting LLMs with Evolutionary Algorithms Yields Powerful Prompt Optimizers","version":3},"reference_index":137,"source":"arxiv_source","source_observed_at":"2026-05-16T06:11:49.475825Z"},"links":{"cited_paper":"/paper/2306.04528","citing_paper":"/paper/2309.08532"},"observation_digest":"sha256:a6ab9f972792bc8c1062ce5849b455abc50e3cdf9aeb787744d3456f08234622","observation_id":"f2ed5e7c-03b1-4f84-a0fd-c4e85a3b4c55","resolution":{"observed_at":"2026-05-16T06:11:49.585490Z","resolver_source":"arxiv_id","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2306.04528","last_updated":"2024-07-16T07:29:49Z","snapshot_observed_at":"2026-08-07T20:45:44.253504Z","submitted_at":"2023-06-07T15:37:00Z","title":"PromptRobust: Towards Evaluating the Robustness of Large Language Models on Adversarial Prompts","version":5},"cited_work":{"arxiv_id":"2306.04528","doi":"10.48550/arxiv.2306.04528","metadata_source":"arxiv_reference","pith_arxiv_id":"2306.04528","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Promptbench: Towards evaluating the robustness of large language models on adversarial prompts","venue":"arXiv (Cornell University)","work_id":"e65e5731-8f09-44e1-94d5-1674f9e08988","year":2024},"citing_paper":{"arxiv_id":"2401.05561","last_updated":"2024-09-30T10:17:12Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2024-01-10T22:07:21Z","title":"TrustLLM: Trustworthiness in Large Language Models","version":6},"reference_index":172,"source":"pdf_text","source_observed_at":"2026-05-18T11:17:08.108565Z"},"links":{"cited_paper":"/paper/2306.04528","citing_paper":"/paper/2401.05561"},"observation_digest":"sha256:45535631cb28e1a3e5e7a97be78b7eecf99d956a5f073e072d735bcb2cc1763a","observation_id":"6503384a-e3c4-4163-9717-5493a50b0b00","resolution":{"observed_at":"2026-05-18T11:17:08.543894Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2306.04528","last_updated":"2024-07-16T07:29:49Z","snapshot_observed_at":"2026-08-07T20:45:44.253504Z","submitted_at":"2023-06-07T15:37:00Z","title":"PromptRobust: Towards Evaluating the Robustness of Large Language Models on Adversarial Prompts","version":5},"cited_work":{"arxiv_id":"2306.04528","doi":"10.48550/arxiv.2306.04528","metadata_source":"arxiv_reference","pith_arxiv_id":"2306.04528","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Promptbench: Towards evaluating the robustness of large language models on adversarial prompts","venue":"arXiv (Cornell University)","work_id":"e65e5731-8f09-44e1-94d5-1674f9e08988","year":2024},"citing_paper":{"arxiv_id":"2402.06922","last_updated":"2026-04-21T14:06:34Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2024-02-10T11:07:24Z","title":"Whispers in the Machine: Confidentiality in Agentic Systems","version":5},"reference_index":30,"source":"pdf_text","source_observed_at":"2026-05-24T03:59:03.972043Z"},"links":{"cited_paper":"/paper/2306.04528","citing_paper":"/paper/2402.06922"},"observation_digest":"sha256:aa7ac97b46fe5494f2705d336bfa15f7d9f0f0d05fa912eccbe4f613b7e7aced","observation_id":"f5db751b-c80a-4b12-a3b9-929b9718e4c5","resolution":{"observed_at":"2026-05-24T04:03:53.865582Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2306.04528","last_updated":"2024-07-16T07:29:49Z","snapshot_observed_at":"2026-08-07T20:45:44.253504Z","submitted_at":"2023-06-07T15:37:00Z","title":"PromptRobust: Towards Evaluating the Robustness of Large Language Models on Adversarial Prompts","version":5},"cited_work":{"arxiv_id":"2306.04528","doi":"10.48550/arxiv.2306.04528","metadata_source":"arxiv_reference","pith_arxiv_id":"2306.04528","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Promptbench: Towards evaluating the robustness of large language models on adversarial prompts","venue":"arXiv (Cornell University)","work_id":"e65e5731-8f09-44e1-94d5-1674f9e08988","year":2024},"citing_paper":{"arxiv_id":"2406.04244","last_updated":"2024-06-06T16:41:39Z","snapshot_observed_at":"2026-07-30T15:43:06.151242Z","submitted_at":"2024-06-06T16:41:39Z","title":"Benchmark Data Contamination of Large Language Models: A Survey","version":1},"reference_index":191,"source":"pdf_text","source_observed_at":"2026-05-22T23:10:40.420241Z"},"links":{"cited_paper":"/paper/2306.04528","citing_paper":"/paper/2406.04244"},"observation_digest":"sha256:ca6babb3a7cb146dd32e1c127b0eb351d8cc8e48ebbc153d82c8196bf3376369","observation_id":"6f1e8eaf-774d-4692-a572-11f1d10e44ce","resolution":{"observed_at":"2026-05-22T23:10:41.169831Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2306.04528","last_updated":"2024-07-16T07:29:49Z","snapshot_observed_at":"2026-08-07T20:45:44.253504Z","submitted_at":"2023-06-07T15:37:00Z","title":"PromptRobust: Towards Evaluating the Robustness of Large Language Models on Adversarial Prompts","version":5},"cited_work":{"arxiv_id":"2306.04528","doi":"10.48550/arxiv.2306.04528","metadata_source":"arxiv_reference","pith_arxiv_id":"2306.04528","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Promptbench: Towards evaluating the robustness of large language models on adversarial prompts","venue":"arXiv (Cornell University)","work_id":"e65e5731-8f09-44e1-94d5-1674f9e08988","year":2024},"citing_paper":{"arxiv_id":"2409.10102","last_updated":"2026-05-16T07:15:07Z","snapshot_observed_at":"2026-08-04T19:12:49.508911Z","submitted_at":"2024-09-16T09:06:44Z","title":"Trustworthiness in Retrieval-Augmented Generation Systems: A Survey","version":2},"reference_index":92,"source":"pdf_text","source_observed_at":"2026-05-23T21:08:11.787013Z"},"links":{"cited_paper":"/paper/2306.04528","citing_paper":"/paper/2409.10102"},"observation_digest":"sha256:2c3ac5035d4a9d51ab4ddb0c9e994959953bf6b98a373151d40564e9d97da013","observation_id":"b7e0043e-1fed-4f3a-b51b-9971c85280c1","resolution":{"observed_at":"2026-05-23T21:08:25.916992Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2306.04528","last_updated":"2024-07-16T07:29:49Z","snapshot_observed_at":"2026-08-07T20:45:44.253504Z","submitted_at":"2023-06-07T15:37:00Z","title":"PromptRobust: Towards Evaluating the Robustness of Large Language Models on Adversarial Prompts","version":5},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2306.04528","snapshot_observed_at":"2026-08-08T20:08:13.602909Z","title":null,"venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2502.05150","last_updated":"2025-02-07T18:26:15Z","snapshot_observed_at":"2026-08-09T15:49:28.774276Z","submitted_at":"2025-02-07T18:26:15Z","title":"CodeSCM: Causal Analysis for Multi-Modal Code Generation","version":1},"reference_index":54,"source":"arxiv_source","source_observed_at":"2026-08-08T20:08:13.602909Z"},"links":{"cited_paper":"/paper/2306.04528","citing_paper":"/paper/2502.05150"},"observation_digest":"sha256:774e3a2a3679f4b75803bfe944823647083d3a672cd8a9f34491ebc1729ee86d","observation_id":"cfbb0258-1cdc-4033-91d5-b4190fddb289","resolution":{"observed_at":"2026-08-08T20:08:13.602909Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2306.04528","last_updated":"2024-07-16T07:29:49Z","snapshot_observed_at":"2026-08-07T20:45:44.253504Z","submitted_at":"2023-06-07T15:37:00Z","title":"PromptRobust: Towards Evaluating the Robustness of Large Language Models on Adversarial Prompts","version":5},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2306.04528","snapshot_observed_at":"2026-08-08T16:55:45.553837Z","title":"ArXivabs/2306.04528(2023)","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2502.06065","last_updated":"2025-02-09T23:01:03Z","snapshot_observed_at":"2026-08-08T16:50:09.624997Z","submitted_at":"2025-02-09T23:01:03Z","title":"Benchmarking Prompt Sensitivity in Large Language Models","version":1},"reference_index":40,"source":"pdf_text","source_observed_at":"2026-08-08T16:55:45.553837Z"},"links":{"cited_paper":"/paper/2306.04528","citing_paper":"/paper/2502.06065"},"observation_digest":"sha256:2fdc7f24b33ae4c502c9cb0b78c1c5195a93aa21bd28fc962b82a3c181e16ae1","observation_id":"b7267df2-c433-47f8-99f8-1efbc8a7c390","resolution":{"observed_at":"2026-08-08T16:55:45.553837Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2306.04528","last_updated":"2024-07-16T07:29:49Z","snapshot_observed_at":"2026-08-07T20:45:44.253504Z","submitted_at":"2023-06-07T15:37:00Z","title":"PromptRobust: Towards Evaluating the Robustness of Large Language Models on Adversarial Prompts","version":5},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2306.04528","snapshot_observed_at":"2026-08-08T10:23:06.309263Z","title":null,"venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2502.08142","last_updated":"2025-02-12T05:48:57Z","snapshot_observed_at":"2026-08-09T00:12:26.816440Z","submitted_at":"2025-02-12T05:48:57Z","title":"Bridging the Safety Gap: A Guardrail Pipeline for Trustworthy LLM Inferences","version":1},"reference_index":84,"source":"pdf_text","source_observed_at":"2026-08-08T10:23:06.309263Z"},"links":{"cited_paper":"/paper/2306.04528","citing_paper":"/paper/2502.08142"},"observation_digest":"sha256:235479e037625add07ecbe74a62607a3469e2e1a9650ce27b724415a9b571ba7","observation_id":"5ae89962-a314-4ef0-ba48-b8b0beaff677","resolution":{"observed_at":"2026-08-08T10:23:06.309263Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2306.04528","last_updated":"2024-07-16T07:29:49Z","snapshot_observed_at":"2026-08-07T20:45:44.253504Z","submitted_at":"2023-06-07T15:37:00Z","title":"PromptRobust: Towards Evaluating the Robustness of Large Language Models on Adversarial Prompts","version":5},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2306.04528","snapshot_observed_at":"2026-08-07T23:35:42.927960Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2502.09670","last_updated":"2025-02-12T22:55:43Z","snapshot_observed_at":"2026-08-09T08:58:06.881734Z","submitted_at":"2025-02-12T22:55:43Z","title":"The Science of Evaluating Foundation Models","version":1},"reference_index":98,"source":"pdf_text","source_observed_at":"2026-08-07T23:35:42.927960Z"},"links":{"cited_paper":"/paper/2306.04528","citing_paper":"/paper/2502.09670"},"observation_digest":"sha256:6aaf7b1930174847521cd243f52129ee025d5d8681bb1bc319324f64324dee9e","observation_id":"defb6ea1-6aed-4bd9-b8e7-c2bba422d85b","resolution":{"observed_at":"2026-08-07T23:35:42.927960Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2306.04528","last_updated":"2024-07-16T07:29:49Z","snapshot_observed_at":"2026-08-07T20:45:44.253504Z","submitted_at":"2023-06-07T15:37:00Z","title":"PromptRobust: Towards Evaluating the Robustness of Large Language Models on Adversarial Prompts","version":5},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2306.04528","snapshot_observed_at":"2026-08-07T13:44:04.449448Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2505.21224","last_updated":"2025-05-27T14:10:30Z","snapshot_observed_at":"2026-08-07T20:46:21.434684Z","submitted_at":"2025-05-27T14:10:30Z","title":"A Representation Level Analysis of NMT Model Robustness to Grammatical Errors","version":1},"reference_index":59,"source":"arxiv_source","source_observed_at":"2026-08-07T13:44:04.449448Z"},"links":{"cited_paper":"/paper/2306.04528","citing_paper":"/paper/2505.21224"},"observation_digest":"sha256:a4e15694e0f18025e5e1577047770c3895c9e2837d12cebf20512c85509826db","observation_id":"3abcd471-5878-46b6-b2a7-526041342659","resolution":{"observed_at":"2026-08-07T13:44:04.449448Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2306.04528","last_updated":"2024-07-16T07:29:49Z","snapshot_observed_at":"2026-08-07T20:45:44.253504Z","submitted_at":"2023-06-07T15:37:00Z","title":"PromptRobust: Towards Evaluating the Robustness of Large Language Models on Adversarial Prompts","version":5},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2306.04528","snapshot_observed_at":"2026-08-07T13:35:56.355578Z","title":"Promptbench: Towards evaluating the robustness of large language models on adversarial prompts.arXiv preprint arXiv:2306.04528, 2023","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2505.21494","last_updated":"2025-05-27T17:56:57Z","snapshot_observed_at":"2026-08-07T21:17:47.870634Z","submitted_at":"2025-05-27T17:56:57Z","title":"Adversarial Attacks against Closed-Source MLLMs via Feature Optimal Alignment","version":1},"reference_index":63,"source":"pdf_text","source_observed_at":"2026-08-07T13:35:56.355578Z"},"links":{"cited_paper":"/paper/2306.04528","citing_paper":"/paper/2505.21494"},"observation_digest":"sha256:bb50adfc97be9aac0d17a3c2a9c64816d4573b440013c175a2068b5fcd2306d5","observation_id":"02365128-d5a3-4427-853b-65671b5f5a03","resolution":{"observed_at":"2026-08-07T13:35:56.355578Z","resolver_source":null,"status":"malformed_identifier"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2306.04528","last_updated":"2024-07-16T07:29:49Z","snapshot_observed_at":"2026-08-07T20:45:44.253504Z","submitted_at":"2023-06-07T15:37:00Z","title":"PromptRobust: Towards Evaluating the Robustness of Large Language Models on Adversarial Prompts","version":5},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2306.04528","snapshot_observed_at":"2026-08-07T13:15:55.842329Z","title":"q_proj\",","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2505.22355","last_updated":"2025-05-28T13:35:12Z","snapshot_observed_at":"2026-08-09T09:37:33.386478Z","submitted_at":"2025-05-28T13:35:12Z","title":"Look Within or Look Beyond? A Theoretical Comparison Between Parameter-Efficient and Full Fine-Tuning","version":1},"reference_index":74,"source":"pdf_text","source_observed_at":"2026-08-07T13:15:55.842329Z"},"links":{"cited_paper":"/paper/2306.04528","citing_paper":"/paper/2505.22355"},"observation_digest":"sha256:e692e6d9905cd0e9097f1c5495446042071d18e71c9617eee2d266f16dd9497a","observation_id":"1f3d2782-3fef-4ef9-acbe-fb7db3bb5404","resolution":{"observed_at":"2026-08-07T13:15:55.842329Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2306.04528","last_updated":"2024-07-16T07:29:49Z","snapshot_observed_at":"2026-08-07T20:45:44.253504Z","submitted_at":"2023-06-07T15:37:00Z","title":"PromptRobust: Towards Evaluating the Robustness of Large Language Models on Adversarial Prompts","version":5},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2306.04528","snapshot_observed_at":"2026-08-07T05:31:52.719510Z","title":"Promptbench: Towards evaluating the robustness of large language models on adversarial prompts.arXiv preprint arXiv:2306.04528, 2023","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2506.07645","last_updated":"2025-06-09T11:09:39Z","snapshot_observed_at":"2026-08-08T17:01:03.200885Z","submitted_at":"2025-06-09T11:09:39Z","title":"Evaluating LLMs Robustness in Less Resourced Languages with Proxy Models","version":1},"reference_index":31,"source":"pdf_text","source_observed_at":"2026-08-07T05:31:52.719510Z"},"links":{"cited_paper":"/paper/2306.04528","citing_paper":"/paper/2506.07645"},"observation_digest":"sha256:0939e8064e79397476f76f719e0266d17d185e731dda9082036a5b27bb44dca8","observation_id":"a07116a8-6252-46a8-9d49-de7c94e931c5","resolution":{"observed_at":"2026-08-07T05:31:52.719510Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2306.04528","last_updated":"2024-07-16T07:29:49Z","snapshot_observed_at":"2026-08-07T20:45:44.253504Z","submitted_at":"2023-06-07T15:37:00Z","title":"PromptRobust: Towards Evaluating the Robustness of Large Language Models on Adversarial Prompts","version":5},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2306.04528","snapshot_observed_at":"2026-08-07T05:42:31.720393Z","title":null,"venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2506.11111","last_updated":"2025-07-09T06:18:33Z","snapshot_observed_at":"2026-08-09T12:39:14.309278Z","submitted_at":"2025-06-08T16:20:12Z","title":"Evaluating and Improving Robustness in Large Language Models: A Survey and Future Directions","version":2},"reference_index":245,"source":"pdf_text","source_observed_at":"2026-08-07T05:42:31.720393Z"},"links":{"cited_paper":"/paper/2306.04528","citing_paper":"/paper/2506.11111"},"observation_digest":"sha256:c6838f3561312137545cf7f84b34e8555cfd73f6f7333d03611da4e2ab583d0c","observation_id":"aa8fdce4-2789-41cc-a44b-cb4d9432a2dc","resolution":{"observed_at":"2026-08-07T05:42:31.720393Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2306.04528","last_updated":"2024-07-16T07:29:49Z","snapshot_observed_at":"2026-08-07T20:45:44.253504Z","submitted_at":"2023-06-07T15:37:00Z","title":"PromptRobust: Towards Evaluating the Robustness of Large Language Models on Adversarial Prompts","version":5},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2306.04528","snapshot_observed_at":"2026-08-06T18:55:55.041443Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2507.06956","last_updated":"2025-07-09T15:39:17Z","snapshot_observed_at":"2026-08-07T20:46:16.086010Z","submitted_at":"2025-07-09T15:39:17Z","title":"Investigating the Robustness of Retrieval-Augmented Generation at the Query Level","version":1},"reference_index":44,"source":"arxiv_source","source_observed_at":"2026-08-06T18:55:55.041443Z"},"links":{"cited_paper":"/paper/2306.04528","citing_paper":"/paper/2507.06956"},"observation_digest":"sha256:6e1468756f98820ba685467f9f72ec2fd37dcbf1947ad8b153d669a6408692d9","observation_id":"92ea6742-49b2-4302-9880-1e77754532cb","resolution":{"observed_at":"2026-08-06T18:55:55.041443Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2306.04528","last_updated":"2024-07-16T07:29:49Z","snapshot_observed_at":"2026-08-07T20:45:44.253504Z","submitted_at":"2023-06-07T15:37:00Z","title":"PromptRobust: Towards Evaluating the Robustness of Large Language Models on Adversarial Prompts","version":5},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2306.04528","snapshot_observed_at":"2026-08-06T17:45:58.463121Z","title":"Promptbench: Towards evaluating the robustness of large language models on adversarial prompts,","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2507.10054","last_updated":"2025-07-23T10:43:29Z","snapshot_observed_at":"2026-08-09T02:38:07.011465Z","submitted_at":"2025-07-14T08:36:26Z","title":"Explicit Vulnerability Generation with LLMs: An Investigation Beyond Adversarial Attacks","version":2},"reference_index":8,"source":"pdf_text","source_observed_at":"2026-08-06T17:45:58.463121Z"},"links":{"cited_paper":"/paper/2306.04528","citing_paper":"/paper/2507.10054"},"observation_digest":"sha256:33b9d40722854d9f2608e4d4c48c7e48ab2eedf51a6528bc0c21f7a685e4a014","observation_id":"47108bc9-e6b0-48df-a5cc-ffe314623e52","resolution":{"observed_at":"2026-08-06T17:45:58.463121Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2306.04528","last_updated":"2024-07-16T07:29:49Z","snapshot_observed_at":"2026-08-07T20:45:44.253504Z","submitted_at":"2023-06-07T15:37:00Z","title":"PromptRobust: Towards Evaluating the Robustness of Large Language Models on Adversarial Prompts","version":5},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2306.04528","snapshot_observed_at":"2026-08-06T14:57:06.882060Z","title":"PromptBench: Towards evaluating the robustness of Large Language Models on adversarial prompts.arXiv: 2306.04528 [cs.CL], June 2023","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2507.17257","last_updated":"2025-07-23T06:56:15Z","snapshot_observed_at":"2026-08-06T18:04:24.113657Z","submitted_at":"2025-07-23T06:56:15Z","title":"Agent Identity Evals: Measuring Agentic Identity","version":1},"reference_index":77,"source":"pdf_text","source_observed_at":"2026-08-06T14:57:06.882060Z"},"links":{"cited_paper":"/paper/2306.04528","citing_paper":"/paper/2507.17257"},"observation_digest":"sha256:ec77e9465ef42bfa413e3cc7886c7c1c545162dfaddef6cc8ab55545d36cefc7","observation_id":"15fb529e-3caf-4a08-bb0b-0efca3b5e99f","resolution":{"observed_at":"2026-08-06T14:57:06.882060Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2306.04528","last_updated":"2024-07-16T07:29:49Z","snapshot_observed_at":"2026-08-07T20:45:44.253504Z","submitted_at":"2023-06-07T15:37:00Z","title":"PromptRobust: Towards Evaluating the Robustness of Large Language Models on Adversarial Prompts","version":5},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2306.04528","snapshot_observed_at":"2026-08-05T05:59:28.243698Z","title":"Promptrobust: Towards evaluating the robustness of large language models on adversarial prompts, 2024","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2509.04615","last_updated":"2025-09-04T18:59:07Z","snapshot_observed_at":"2026-08-08T07:38:10.301418Z","submitted_at":"2025-09-04T18:59:07Z","title":"Breaking to Build: A Threat Model of Prompt-Based Attacks for Securing LLMs","version":1},"reference_index":6,"source":"pdf_text","source_observed_at":"2026-08-05T05:59:28.243698Z"},"links":{"cited_paper":"/paper/2306.04528","citing_paper":"/paper/2509.04615"},"observation_digest":"sha256:c6df740c0cdd695acee82566dfbdad217feb528a5c97f591af8b903f98fd9f9b","observation_id":"7b63df69-81fe-40e7-a062-bae3a724ff7f","resolution":{"observed_at":"2026-08-05T05:59:28.243698Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2306.04528","last_updated":"2024-07-16T07:29:49Z","snapshot_observed_at":"2026-08-07T20:45:44.253504Z","submitted_at":"2023-06-07T15:37:00Z","title":"PromptRobust: Towards Evaluating the Robustness of Large Language Models on Adversarial Prompts","version":5},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2306.04528","snapshot_observed_at":"2026-08-03T06:46:26.373904Z","title":"Promptbench: Towards eval- uating the robustness of large language models on adversarial prompts.arXiv preprint arXiv:2306.04528, 2023","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2601.22025","last_updated":"2026-06-09T23:57:32Z","snapshot_observed_at":"2026-08-07T21:40:38.906811Z","submitted_at":"2026-01-29T17:32:34Z","title":"When Generic Prompt Improvements Hurt: Evaluation-Driven Iteration for LLM Applications","version":2},"reference_index":28,"source":"pdf_text","source_observed_at":"2026-08-03T06:46:26.373904Z"},"links":{"cited_paper":"/paper/2306.04528","citing_paper":"/paper/2601.22025"},"observation_digest":"sha256:90856f506f6abb04497b45efc596d37e58650a6bbf7cef15574c609ec51121a8","observation_id":"0495afe7-d37d-4a4b-8fe3-b7e7c29526c7","resolution":{"observed_at":"2026-08-03T06:46:26.373904Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2306.04528","last_updated":"2024-07-16T07:29:49Z","snapshot_observed_at":"2026-08-07T20:45:44.253504Z","submitted_at":"2023-06-07T15:37:00Z","title":"PromptRobust: Towards Evaluating the Robustness of Large Language Models on Adversarial Prompts","version":5},"cited_work":{"arxiv_id":"2306.04528","doi":"10.48550/arxiv.2306.04528","metadata_source":"arxiv_reference","pith_arxiv_id":"2306.04528","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Promptbench: Towards evaluating the robustness of large language models on adversarial prompts","venue":"arXiv (Cornell University)","work_id":"e65e5731-8f09-44e1-94d5-1674f9e08988","year":2024},"citing_paper":{"arxiv_id":"2603.03332","last_updated":"2026-04-16T23:41:41Z","snapshot_observed_at":"2026-08-04T02:21:24.392691Z","submitted_at":"2026-02-11T03:11:30Z","title":"Fragile Thoughts: How Large Language Models Handle Chain-of-Thought Perturbations","version":3},"reference_index":12,"source":"pdf_text","source_observed_at":"2026-05-16T03:43:18.987241Z"},"links":{"cited_paper":"/paper/2306.04528","citing_paper":"/paper/2603.03332"},"observation_digest":"sha256:e80861463ecaa6d92e9772c9aec6eed8d0885a994a45a2efc52e1401c9a35dc7","observation_id":"3bf2096c-140c-4017-84eb-0dd8394f2a8b","resolution":{"observed_at":"2026-05-16T03:47:15.338790Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2306.04528","last_updated":"2024-07-16T07:29:49Z","snapshot_observed_at":"2026-08-07T20:45:44.253504Z","submitted_at":"2023-06-07T15:37:00Z","title":"PromptRobust: Towards Evaluating the Robustness of Large Language Models on Adversarial Prompts","version":5},"cited_work":{"arxiv_id":"2306.04528","doi":"10.48550/arxiv.2306.04528","metadata_source":"arxiv_reference","pith_arxiv_id":"2306.04528","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Promptbench: Towards evaluating the robustness of large language models on adversarial prompts","venue":"arXiv (Cornell University)","work_id":"e65e5731-8f09-44e1-94d5-1674f9e08988","year":2024},"citing_paper":{"arxiv_id":"2603.10477","last_updated":"2026-04-08T07:03:59Z","snapshot_observed_at":"2026-08-03T13:46:05.416163Z","submitted_at":"2026-03-11T07:00:59Z","title":"PEEM: Prompt Engineering Evaluation Metrics for Interpretable Joint Evaluation of Prompts and Responses","version":2},"reference_index":43,"source":"pdf_text","source_observed_at":"2026-05-15T13:57:41.428695Z"},"links":{"cited_paper":"/paper/2306.04528","citing_paper":"/paper/2603.10477"},"observation_digest":"sha256:aec806221ccaf333baa1718fc09ea751ae81da6c3c7e17a72f457711ffe388ea","observation_id":"abd144a7-b5e5-4d2b-8c87-04c8cdef2ffd","resolution":{"observed_at":"2026-05-15T14:00:03.073494Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2306.04528","last_updated":"2024-07-16T07:29:49Z","snapshot_observed_at":"2026-08-07T20:45:44.253504Z","submitted_at":"2023-06-07T15:37:00Z","title":"PromptRobust: Towards Evaluating the Robustness of Large Language Models on Adversarial Prompts","version":5},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2306.04528","snapshot_observed_at":"2026-07-13T14:38:58.879673Z","title":"Promptbench: Towards evaluating the robustness of large language models on adversarial prompts,","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2604.01039","last_updated":"2026-06-05T19:49:24Z","snapshot_observed_at":"2026-08-07T20:47:26.471296Z","submitted_at":"2026-04-01T15:45:56Z","title":"Automated Framework to Evaluate and Harden LLM System Instructions against Encoding Attacks","version":2},"reference_index":21,"source":"pdf_text","source_observed_at":"2026-07-13T14:38:58.879673Z"},"links":{"cited_paper":"/paper/2306.04528","citing_paper":"/paper/2604.01039"},"observation_digest":"sha256:b0226d7a45c50643f0474466b943243aecbc4901268a4a9ac12791dedbc50ad3","observation_id":"282e714c-1f2b-4b3a-ac98-aedba4179bcb","resolution":{"observed_at":"2026-07-13T14:38:58.879673Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2306.04528","last_updated":"2024-07-16T07:29:49Z","snapshot_observed_at":"2026-08-07T20:45:44.253504Z","submitted_at":"2023-06-07T15:37:00Z","title":"PromptRobust: Towards Evaluating the Robustness of Large Language Models on Adversarial Prompts","version":5},"cited_work":{"arxiv_id":"2306.04528","doi":"10.48550/arxiv.2306.04528","metadata_source":"arxiv_reference","pith_arxiv_id":"2306.04528","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Promptbench: Towards evaluating the robustness of large language models on adversarial prompts","venue":"arXiv (Cornell University)","work_id":"e65e5731-8f09-44e1-94d5-1674f9e08988","year":2024},"citing_paper":{"arxiv_id":"2604.16421","last_updated":"2026-04-03T11:36:49Z","snapshot_observed_at":"2026-07-06T23:03:43.854612Z","submitted_at":"2026-04-03T11:36:49Z","title":"Measuring Representation Robustness in Large Language Models for Geometry","version":1},"reference_index":40,"source":"pdf_text","source_observed_at":"2026-05-13T19:35:32.531660Z"},"links":{"cited_paper":"/paper/2306.04528","citing_paper":"/paper/2604.16421"},"observation_digest":"sha256:596f2c80e2c350a4035d36b97a7989960d52e1a72285d2cfb741b67384a6a0c4","observation_id":"ed3491bb-86d1-4f3e-aa00-134f2634b233","resolution":{"observed_at":"2026-05-13T19:38:10.484726Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2306.04528","last_updated":"2024-07-16T07:29:49Z","snapshot_observed_at":"2026-08-07T20:45:44.253504Z","submitted_at":"2023-06-07T15:37:00Z","title":"PromptRobust: Towards Evaluating the Robustness of Large Language Models on Adversarial Prompts","version":5},"cited_work":{"arxiv_id":"2306.04528","doi":"10.48550/arxiv.2306.04528","metadata_source":"arxiv_reference","pith_arxiv_id":"2306.04528","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Promptbench: Towards evaluating the robustness of large language models on adversarial prompts","venue":"arXiv (Cornell University)","work_id":"e65e5731-8f09-44e1-94d5-1674f9e08988","year":2024},"citing_paper":{"arxiv_id":"2604.23135","last_updated":"2026-05-17T15:51:14Z","snapshot_observed_at":"2026-07-06T23:09:23.591544Z","submitted_at":"2026-04-25T04:26:19Z","title":"Characterizing Paraphrase-Induced Failures in Lean 4 Autoformalization","version":1},"reference_index":11,"source":"pdf_text","source_observed_at":"2026-05-08T08:33:34.423179Z"},"links":{"cited_paper":"/paper/2306.04528","citing_paper":"/paper/2604.23135"},"observation_digest":"sha256:53f7791487b04b57148f8f65ddfa39695f172e0069a9f43c8477f13a1787f7ae","observation_id":"11f307d6-7263-46c5-8027-d796552a41b7","resolution":{"observed_at":"2026-05-11T20:36:08.464012Z","resolver_source":"arxiv_id","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2306.04528","last_updated":"2024-07-16T07:29:49Z","snapshot_observed_at":"2026-08-07T20:45:44.253504Z","submitted_at":"2023-06-07T15:37:00Z","title":"PromptRobust: Towards Evaluating the Robustness of Large Language Models on Adversarial Prompts","version":5},"cited_work":{"arxiv_id":"2306.04528","doi":"10.48550/arxiv.2306.04528","metadata_source":"arxiv_reference","pith_arxiv_id":"2306.04528","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Promptbench: Towards evaluating the robustness of large language models on adversarial prompts","venue":"arXiv (Cornell University)","work_id":"e65e5731-8f09-44e1-94d5-1674f9e08988","year":2024},"citing_paper":{"arxiv_id":"2604.23135","last_updated":"2026-05-17T15:51:14Z","snapshot_observed_at":"2026-07-06T23:09:23.591544Z","submitted_at":"2026-04-25T04:26:19Z","title":"Characterizing Paraphrase-Induced Failures in Lean 4 Autoformalization","version":2},"reference_index":11,"source":"pdf_text","source_observed_at":"2026-05-21T00:49:58.959331Z"},"links":{"cited_paper":"/paper/2306.04528","citing_paper":"/paper/2604.23135"},"observation_digest":"sha256:4224f23d39f657d855d38e8b2309f31aaa4407734ac2b43c8db2c70df92f7062","observation_id":"be275719-33b2-4a79-a3bc-77875de9734d","resolution":{"observed_at":"2026-05-21T00:53:53.714254Z","resolver_source":"arxiv_id","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2306.04528","last_updated":"2024-07-16T07:29:49Z","snapshot_observed_at":"2026-08-07T20:45:44.253504Z","submitted_at":"2023-06-07T15:37:00Z","title":"PromptRobust: Towards Evaluating the Robustness of Large Language Models on Adversarial Prompts","version":5},"cited_work":{"arxiv_id":"2306.04528","doi":"10.48550/arxiv.2306.04528","metadata_source":"arxiv_reference","pith_arxiv_id":"2306.04528","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Promptbench: Towards evaluating the robustness of large language models on adversarial prompts","venue":"arXiv (Cornell University)","work_id":"e65e5731-8f09-44e1-94d5-1674f9e08988","year":2024},"citing_paper":{"arxiv_id":"2604.23338","last_updated":"2026-05-06T17:17:02Z","snapshot_observed_at":"2026-08-06T19:46:39.219000Z","submitted_at":"2026-04-25T14:57:15Z","title":"A Systematic Survey of Security Threats and Defenses in LLM-Based AI Agents: A Layered Attack Surface Framework","version":2},"reference_index":77,"source":"pdf_text","source_observed_at":"2026-05-08T07:53:13.746141Z"},"links":{"cited_paper":"/paper/2306.04528","citing_paper":"/paper/2604.23338"},"observation_digest":"sha256:82b182117343043feaf69cc4f83b963399ede9b974b13070bc8176502eca51b7","observation_id":"615c03fa-9e15-4524-9310-144f87348904","resolution":{"observed_at":"2026-05-11T20:51:09.098195Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2306.04528","last_updated":"2024-07-16T07:29:49Z","snapshot_observed_at":"2026-08-07T20:45:44.253504Z","submitted_at":"2023-06-07T15:37:00Z","title":"PromptRobust: Towards Evaluating the Robustness of Large Language Models on Adversarial Prompts","version":5},"cited_work":{"arxiv_id":"2306.04528","doi":"10.48550/arxiv.2306.04528","metadata_source":"arxiv_reference","pith_arxiv_id":"2306.04528","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Promptbench: Towards evaluating the robustness of large language models on adversarial prompts","venue":"arXiv (Cornell University)","work_id":"e65e5731-8f09-44e1-94d5-1674f9e08988","year":2024},"citing_paper":{"arxiv_id":"2604.24712","last_updated":"2026-04-27T17:21:09Z","snapshot_observed_at":"2026-07-06T23:10:42.926677Z","submitted_at":"2026-04-27T17:21:09Z","title":"When Prompt Under-Specification Improves Code Correctness: An Exploratory Study of Prompt Wording and Structure Effects on LLM-Based Code Generation","version":1},"reference_index":46,"source":"pdf_text","source_observed_at":"2026-05-08T03:00:26.137401Z"},"links":{"cited_paper":"/paper/2306.04528","citing_paper":"/paper/2604.24712"},"observation_digest":"sha256:dd13889c876448c239a1d224f9c3a443c63d797566baea2f35b239574c40debc","observation_id":"91f507b1-67ba-48d8-bccc-2c084d245b36","resolution":{"observed_at":"2026-05-11T22:17:07.253922Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2306.04528","last_updated":"2024-07-16T07:29:49Z","snapshot_observed_at":"2026-08-07T20:45:44.253504Z","submitted_at":"2023-06-07T15:37:00Z","title":"PromptRobust: Towards Evaluating the Robustness of Large Language Models on Adversarial Prompts","version":5},"cited_work":{"arxiv_id":"2306.04528","doi":"10.48550/arxiv.2306.04528","metadata_source":"arxiv_reference","pith_arxiv_id":"2306.04528","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Promptbench: Towards evaluating the robustness of large language models on adversarial prompts","venue":"arXiv (Cornell University)","work_id":"e65e5731-8f09-44e1-94d5-1674f9e08988","year":2024},"citing_paper":{"arxiv_id":"2605.04665","last_updated":"2026-05-09T22:09:59Z","snapshot_observed_at":"2026-07-06T23:17:23.641984Z","submitted_at":"2026-05-06T09:11:10Z","title":"Paraphrase-Induced Output-Mode Collapse: When LLMs Break Character Under Semantically Equivalent Inputs","version":1},"reference_index":5,"source":"pdf_text","source_observed_at":"2026-05-08T16:26:54.068997Z"},"links":{"cited_paper":"/paper/2306.04528","citing_paper":"/paper/2605.04665"},"observation_digest":"sha256:47d6b164c6f8cd9950dcae71f8523a09f8d78b69205f95e9e4328a8d81fe8fac","observation_id":"b9a6b961-0c06-4f58-8018-8c947129bdfd","resolution":{"observed_at":"2026-05-11T18:16:07.328226Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2306.04528","last_updated":"2024-07-16T07:29:49Z","snapshot_observed_at":"2026-08-07T20:45:44.253504Z","submitted_at":"2023-06-07T15:37:00Z","title":"PromptRobust: Towards Evaluating the Robustness of Large Language Models on Adversarial Prompts","version":5},"cited_work":{"arxiv_id":"2306.04528","doi":"10.48550/arxiv.2306.04528","metadata_source":"arxiv_reference","pith_arxiv_id":"2306.04528","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Promptbench: Towards evaluating the robustness of large language models on adversarial prompts","venue":"arXiv (Cornell University)","work_id":"e65e5731-8f09-44e1-94d5-1674f9e08988","year":2024},"citing_paper":{"arxiv_id":"2605.04665","last_updated":"2026-05-09T22:09:59Z","snapshot_observed_at":"2026-07-06T23:17:23.641984Z","submitted_at":"2026-05-06T09:11:10Z","title":"Paraphrase-Induced Output-Mode Collapse: When LLMs Break Character Under Semantically Equivalent Inputs","version":2},"reference_index":5,"source":"pdf_text","source_observed_at":"2026-05-12T02:13:48.672711Z"},"links":{"cited_paper":"/paper/2306.04528","citing_paper":"/paper/2605.04665"},"observation_digest":"sha256:ad54a891d91c59cc5f86dffcf40189796ef635988f7249f38994da7414c8158f","observation_id":"0810c81a-95cc-4e59-8a2b-0761e6fcc17e","resolution":{"observed_at":"2026-05-12T02:16:16.124945Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2306.04528","last_updated":"2024-07-16T07:29:49Z","snapshot_observed_at":"2026-08-07T20:45:44.253504Z","submitted_at":"2023-06-07T15:37:00Z","title":"PromptRobust: Towards Evaluating the Robustness of Large Language Models on Adversarial Prompts","version":5},"cited_work":{"arxiv_id":"2306.04528","doi":"10.48550/arxiv.2306.04528","metadata_source":"arxiv_reference","pith_arxiv_id":"2306.04528","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Promptbench: Towards evaluating the robustness of large language models on adversarial prompts","venue":"arXiv (Cornell University)","work_id":"e65e5731-8f09-44e1-94d5-1674f9e08988","year":2024},"citing_paper":{"arxiv_id":"2605.09041","last_updated":"2026-05-09T16:26:49Z","snapshot_observed_at":"2026-08-03T06:54:18.779987Z","submitted_at":"2026-05-09T16:26:49Z","title":"BiAxisAudit: A Novel Framework to Evaluate LLM Bias Across Prompt Sensitivity and Response-Layer Divergence","version":1},"reference_index":34,"source":"pdf_text","source_observed_at":"2026-05-12T03:26:18.375974Z"},"links":{"cited_paper":"/paper/2306.04528","citing_paper":"/paper/2605.09041"},"observation_digest":"sha256:ba84c1010cb014564849aa087f2300b068c127e66b8b17d78e7202be13c46691","observation_id":"8f2c0e70-523f-4127-bf25-7274bc3202c0","resolution":{"observed_at":"2026-05-12T07:21:26.819643Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2306.04528","last_updated":"2024-07-16T07:29:49Z","snapshot_observed_at":"2026-08-07T20:45:44.253504Z","submitted_at":"2023-06-07T15:37:00Z","title":"PromptRobust: Towards Evaluating the Robustness of Large Language Models on Adversarial Prompts","version":5},"cited_work":{"arxiv_id":"2306.04528","doi":"10.48550/arxiv.2306.04528","metadata_source":"arxiv_reference","pith_arxiv_id":"2306.04528","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Promptbench: Towards evaluating the robustness of large language models on adversarial prompts","venue":"arXiv (Cornell University)","work_id":"e65e5731-8f09-44e1-94d5-1674f9e08988","year":2024},"citing_paper":{"arxiv_id":"2605.10516","last_updated":"2026-05-11T13:06:24Z","snapshot_observed_at":"2026-07-06T23:22:28.318343Z","submitted_at":"2026-05-11T13:06:24Z","title":"Consistency as a Testable Property: Statistical Methods to Evaluate AI Agent Reliability","version":1},"reference_index":35,"source":"pdf_text","source_observed_at":"2026-05-12T04:41:15.286881Z"},"links":{"cited_paper":"/paper/2306.04528","citing_paper":"/paper/2605.10516"},"observation_digest":"sha256:2c41064fa2b401c25513407c8b16d7c04a4bc5c4a13d82a5723cb9d898ce55f1","observation_id":"cb4f7e11-7e6d-44ef-a76b-25a5eccf2e66","resolution":{"observed_at":"2026-05-12T04:41:21.819932Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2306.04528","last_updated":"2024-07-16T07:29:49Z","snapshot_observed_at":"2026-08-07T20:45:44.253504Z","submitted_at":"2023-06-07T15:37:00Z","title":"PromptRobust: Towards Evaluating the Robustness of Large Language Models on Adversarial Prompts","version":5},"cited_work":{"arxiv_id":"2306.04528","doi":"10.48550/arxiv.2306.04528","metadata_source":"arxiv_reference","pith_arxiv_id":"2306.04528","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Promptbench: Towards evaluating the robustness of large language models on adversarial prompts","venue":"arXiv (Cornell University)","work_id":"e65e5731-8f09-44e1-94d5-1674f9e08988","year":2024},"citing_paper":{"arxiv_id":"2606.01210","last_updated":"2026-05-31T13:00:57Z","snapshot_observed_at":"2026-08-08T11:31:09.932062Z","submitted_at":"2026-05-31T13:00:57Z","title":"Can we trust LLM Self-Explanations for Entity Resolution?","version":1},"reference_index":73,"source":"pdf_text","source_observed_at":"2026-06-28T16:13:14.653064Z"},"links":{"cited_paper":"/paper/2306.04528","citing_paper":"/paper/2606.01210"},"observation_digest":"sha256:9deeba2d9b69c1e215e50f2e2959b434862ab2f8df133b53222b96bd4929438e","observation_id":"ee06c8c2-73dc-445b-95b2-7422e7f31eb1","resolution":{"observed_at":"2026-07-01T21:56:15.179442Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2306.04528","last_updated":"2024-07-16T07:29:49Z","snapshot_observed_at":"2026-08-07T20:45:44.253504Z","submitted_at":"2023-06-07T15:37:00Z","title":"PromptRobust: Towards Evaluating the Robustness of Large Language Models on Adversarial Prompts","version":5},"cited_work":{"arxiv_id":"2306.04528","doi":"10.48550/arxiv.2306.04528","metadata_source":"arxiv_reference","pith_arxiv_id":"2306.04528","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Promptbench: Towards evaluating the robustness of large language models on adversarial prompts","venue":"arXiv (Cornell University)","work_id":"e65e5731-8f09-44e1-94d5-1674f9e08988","year":2024},"citing_paper":{"arxiv_id":"2606.01441","last_updated":"2026-05-31T20:20:53Z","snapshot_observed_at":"2026-08-05T08:36:53.588712Z","submitted_at":"2026-05-31T20:20:53Z","title":"Dive into Ambiguity: A*-Inspired Multi-Agents Commonsense Obfuscation Attack on LLM Prompts","version":1},"reference_index":43,"source":"pdf_text","source_observed_at":"2026-06-28T16:54:12.354178Z"},"links":{"cited_paper":"/paper/2306.04528","citing_paper":"/paper/2606.01441"},"observation_digest":"sha256:caa9116864f0b1e4506f0674cee83768c67bdc765de2146f55513519d966435c","observation_id":"f4b554b0-5b1b-4c6d-bf9d-daca796d5c9e","resolution":{"observed_at":"2026-07-01T21:26:16.387346Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2306.04528","last_updated":"2024-07-16T07:29:49Z","snapshot_observed_at":"2026-08-07T20:45:44.253504Z","submitted_at":"2023-06-07T15:37:00Z","title":"PromptRobust: Towards Evaluating the Robustness of Large Language Models on Adversarial Prompts","version":5},"cited_work":{"arxiv_id":"2306.04528","doi":"10.48550/arxiv.2306.04528","metadata_source":"arxiv_reference","pith_arxiv_id":"2306.04528","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Promptbench: Towards evaluating the robustness of large language models on adversarial prompts","venue":"arXiv (Cornell University)","work_id":"e65e5731-8f09-44e1-94d5-1674f9e08988","year":2024},"citing_paper":{"arxiv_id":"2606.06924","last_updated":"2026-06-05T05:42:00Z","snapshot_observed_at":"2026-08-06T11:49:12.989321Z","submitted_at":"2026-06-05T05:42:00Z","title":"From Sampled Outcomes to Capability Distributions: Rethinking Supervision for LLM Routing","version":1},"reference_index":158,"source":"arxiv_source","source_observed_at":"2026-06-27T22:54:28.452796Z"},"links":{"cited_paper":"/paper/2306.04528","citing_paper":"/paper/2606.06924"},"observation_digest":"sha256:051f4fe787f48a8fab021e7100c271244e0c7b8d890198fbe3d6749fcbf27853","observation_id":"5e033c3d-c72d-4347-a826-3d07a3b969b3","resolution":{"observed_at":"2026-07-02T16:17:08.772986Z","resolver_source":"arxiv_id","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2306.04528","last_updated":"2024-07-16T07:29:49Z","snapshot_observed_at":"2026-08-07T20:45:44.253504Z","submitted_at":"2023-06-07T15:37:00Z","title":"PromptRobust: Towards Evaluating the Robustness of Large Language Models on Adversarial Prompts","version":5},"cited_work":{"arxiv_id":"2306.04528","doi":"10.48550/arxiv.2306.04528","metadata_source":"arxiv_reference","pith_arxiv_id":"2306.04528","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Promptbench: Towards evaluating the robustness of large language models on adversarial prompts","venue":"arXiv (Cornell University)","work_id":"e65e5731-8f09-44e1-94d5-1674f9e08988","year":2024},"citing_paper":{"arxiv_id":"2606.12978","last_updated":"2026-06-11T07:12:17Z","snapshot_observed_at":"2026-08-03T00:41:34.363368Z","submitted_at":"2026-06-11T07:12:17Z","title":"Trajectory-Level Redirection Attacks on Vision-Language-Action Models","version":1},"reference_index":43,"source":"pdf_text","source_observed_at":"2026-06-27T06:33:53.013076Z"},"links":{"cited_paper":"/paper/2306.04528","citing_paper":"/paper/2606.12978"},"observation_digest":"sha256:320c7160e072d150ad43a44db8899fb110d287ad0b51a7da0af33ece64c97918","observation_id":"2ec8c2be-f86e-4035-be37-84907a53bfb9","resolution":{"observed_at":"2026-07-03T15:18:33.763432Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2306.04528","last_updated":"2024-07-16T07:29:49Z","snapshot_observed_at":"2026-08-07T20:45:44.253504Z","submitted_at":"2023-06-07T15:37:00Z","title":"PromptRobust: Towards Evaluating the Robustness of Large Language Models on Adversarial Prompts","version":5},"cited_work":{"arxiv_id":"2306.04528","doi":"10.48550/arxiv.2306.04528","metadata_source":"arxiv_reference","pith_arxiv_id":"2306.04528","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Promptbench: Towards evaluating the robustness of large language models on adversarial prompts","venue":"arXiv (Cornell University)","work_id":"e65e5731-8f09-44e1-94d5-1674f9e08988","year":2024},"citing_paper":{"arxiv_id":"2606.23716","last_updated":"2026-06-16T14:19:04Z","snapshot_observed_at":"2026-08-03T11:45:50.371507Z","submitted_at":"2026-06-16T14:19:04Z","title":"Legal Reasoning Is Not Lawyering: Rethinking Legal Benchmarks for Pro Se Access to Justice","version":1},"reference_index":17,"source":"arxiv_source","source_observed_at":"2026-06-26T22:26:54.239505Z"},"links":{"cited_paper":"/paper/2306.04528","citing_paper":"/paper/2606.23716"},"observation_digest":"sha256:767ed540b49df49b95b405ad8c91f6fbd71fe7c7a139d49b931fb296b0ec2bdc","observation_id":"25f70ffd-2c35-4b17-b7af-c765deb85791","resolution":{"observed_at":"2026-07-03T23:19:03.950183Z","resolver_source":"arxiv_id","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2306.04528","last_updated":"2024-07-16T07:29:49Z","snapshot_observed_at":"2026-08-07T20:45:44.253504Z","submitted_at":"2023-06-07T15:37:00Z","title":"PromptRobust: Towards Evaluating the Robustness of Large Language Models on Adversarial Prompts","version":5},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2306.04528","snapshot_observed_at":"2026-07-14T19:17:45.652076Z","title":"Zhu, K., Wang, J., Zhou, J., Wang, Z., Chen, H., Wang, Y., Yang, L., Ye, W., Gong, N","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2607.09665","last_updated":"2026-05-02T22:51:41Z","snapshot_observed_at":"2026-08-04T00:20:04.568648Z","submitted_at":"2026-05-02T22:51:41Z","title":"Format Sensitivity Index: Token-Controlled Prompt Wrapper Robustness and Schema Compliance in LLM Benchmarking","version":1},"reference_index":16,"source":"pdf_text","source_observed_at":"2026-07-14T19:17:45.652076Z"},"links":{"cited_paper":"/paper/2306.04528","citing_paper":"/paper/2607.09665"},"observation_digest":"sha256:dc005f08f2d19da5c3d36e9c0615d3fe0c9b62bf6c6447aeb073313eb5a7b56f","observation_id":"4448a36d-8194-49fc-9258-dbf5724d335e","resolution":{"observed_at":"2026-07-14T19:17:45.652076Z","resolver_source":null,"status":"malformed_identifier"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2306.04528","last_updated":"2024-07-16T07:29:49Z","snapshot_observed_at":"2026-08-07T20:45:44.253504Z","submitted_at":"2023-06-07T15:37:00Z","title":"PromptRobust: Towards Evaluating the Robustness of Large Language Models on Adversarial Prompts","version":5},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2306.04528","snapshot_observed_at":"2026-07-14T19:17:45.652076Z","title":"https://arxiv.org/abs/2306.04528","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2607.09665","last_updated":"2026-05-02T22:51:41Z","snapshot_observed_at":"2026-08-04T00:20:04.568648Z","submitted_at":"2026-05-02T22:51:41Z","title":"Format Sensitivity Index: Token-Controlled Prompt Wrapper Robustness and Schema Compliance in LLM Benchmarking","version":1},"reference_index":17,"source":"pdf_text","source_observed_at":"2026-07-14T19:17:45.652076Z"},"links":{"cited_paper":"/paper/2306.04528","citing_paper":"/paper/2607.09665"},"observation_digest":"sha256:fe64c90e25c6567291a6ba554e3424e6ed4ceeb6b3ce7c3a1de4514525bb542d","observation_id":"3e67b377-fb5f-4839-9e62-3a2f3fa3b155","resolution":{"observed_at":"2026-07-14T19:17:45.652076Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2306.04528","last_updated":"2024-07-16T07:29:49Z","snapshot_observed_at":"2026-08-07T20:45:44.253504Z","submitted_at":"2023-06-07T15:37:00Z","title":"PromptRobust: Towards Evaluating the Robustness of Large Language Models on Adversarial Prompts","version":5},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2306.04528","snapshot_observed_at":"2026-08-01T14:35:01.351523Z","title":"Promptbench: Towards evaluating the robustness of large language models on adversarial prompts,","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2607.18725","last_updated":"2026-07-21T05:32:40Z","snapshot_observed_at":"2026-08-07T20:46:20.302049Z","submitted_at":"2026-07-21T05:32:40Z","title":"Find Before You Fine-Tune: A Diagnostic Study of Small LLMs for Cybersecurity QA","version":1},"reference_index":16,"source":"pdf_text","source_observed_at":"2026-08-01T14:35:01.351523Z"},"links":{"cited_paper":"/paper/2306.04528","citing_paper":"/paper/2607.18725"},"observation_digest":"sha256:d211c0f5dc533431c4b3c3389f9bc7a2f664eceeec07b6cb7fe7ecf7b315d31a","observation_id":"0fb91a91-a8a6-4c45-b562-b96172b8ce3a","resolution":{"observed_at":"2026-08-01T14:35:01.351523Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2306.04528","last_updated":"2024-07-16T07:29:49Z","snapshot_observed_at":"2026-08-07T20:45:44.253504Z","submitted_at":"2023-06-07T15:37:00Z","title":"PromptRobust: Towards Evaluating the Robustness of Large Language Models on Adversarial Prompts","version":5},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2306.04528","snapshot_observed_at":"2026-08-02T07:07:45.246508Z","title":null,"venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2607.22683","last_updated":"2026-07-13T05:02:43Z","snapshot_observed_at":"2026-08-07T06:32:39.192476Z","submitted_at":"2026-07-13T05:02:43Z","title":"Imprompt: A Language Framework for Prompt Programming","version":1},"reference_index":66,"source":"pdf_text","source_observed_at":"2026-08-02T07:07:45.246508Z"},"links":{"cited_paper":"/paper/2306.04528","citing_paper":"/paper/2607.22683"},"observation_digest":"sha256:6e712c8773d0beb017a4ae05ff2b4cfb7202a96c24a38adb82b6fb0d1edbcea2","observation_id":"d1c5e8b1-a155-4fd8-817c-65fea8b70be4","resolution":{"observed_at":"2026-08-02T07:07:45.246508Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2306.04528","last_updated":"2024-07-16T07:29:49Z","snapshot_observed_at":"2026-08-07T20:45:44.253504Z","submitted_at":"2023-06-07T15:37:00Z","title":"PromptRobust: Towards Evaluating the Robustness of Large Language Models on Adversarial Prompts","version":5},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2306.04528","snapshot_observed_at":"2026-08-07T13:57:11.549507Z","title":"PromptBench: Towards evaluating the robustness of large language models on adversarial prompts,","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2608.06154","last_updated":"2026-08-06T15:21:42Z","snapshot_observed_at":"2026-08-09T23:13:01.112489Z","submitted_at":"2026-08-06T15:21:42Z","title":"Visual Grounding in Zero-Shot Vision-Language Control","version":1},"reference_index":37,"source":"pdf_text","source_observed_at":"2026-08-07T13:57:11.549507Z"},"links":{"cited_paper":"/paper/2306.04528","citing_paper":"/paper/2608.06154"},"observation_digest":"sha256:e71f8728f03e8abcb3880464dd1a8d8a27b7cc24f3f0d867e70394cc2ea1eee9","observation_id":"d7353651-c3b0-43ac-b1a7-9db0d51c80dc","resolution":{"observed_at":"2026-08-07T13:57:11.549507Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}}],"links":{"evidence":"/evidence","html":"/paper/2306.04528/citation-record","integrity":"/paper/2306.04528/integrity","json":"/paper/2306.04528/citation-record.json","paper":"/paper/2306.04528"},"outbound":[],"paper":{"arxiv_id":"2306.04528","last_updated":"2024-07-16T07:29:49Z","latest_version":5,"primary_category":"cs.CL","snapshot_observed_at":"2026-08-07T20:45:44.253504Z","submitted_at":"2023-06-07T15:37:00Z","title":"PromptRobust: Towards Evaluating the Robustness of Large Language Models on Adversarial Prompts"},"reference_resolution":{"displayed":0,"state_counts":{"malformed_identifier":0,"metadata_mismatch":0,"parse_uncertain":0,"unresolved":0,"verified_exact":0,"verified_fuzzy":0},"total_outbound_references":0},"refusal":"A citation records a reference. It does not transfer a finding from one paper to another.","schema":"pith.paper-citation-record.v1","standing_sources":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"thesis":"As of 10 August 2026, this Paper Citation Record lists 0 of 0 outbound references and 44 inbound Pith citation observations for arXiv:2306.04528."}