{"as_of":"2026-08-19T08:29:00Z","caps":{"database_statements":6,"inbound":100,"outbound":100},"context_digest":"sha256:538e6288fa5e05f858d46c59d3a05d8306e7cceff8a2890c877ce9cd0ba5554e","coverage":[{"denominator":96,"lane":"reference_resolution","note":"Typed states for the displayed outbound observations.","records_observed":96,"source":"paper_references, paper_reference_links","source_observed_at":"2026-08-12T19:35:30.417149Z","state":"measured"},{"denominator":96,"lane":"standing_notices","note":"One-hop event checks from named stored sources.","records_observed":96,"source":"scholarly_work_events, retraction_status_cache","source_observed_at":"2026-08-19T06:32:44.657259+00:00","state":"measured"},{"denominator":0,"lane":"inbound_itemization","note":"Pith citing papers itemized under the disclosed page cap.","records_observed":0,"source":"paper_references, paper_reference_links","source_observed_at":null,"state":"measured"},{"denominator":1,"lane":"external_citation_measurements","note":"A source-named dated measurement, never combined with another source.","records_observed":0,"source":"cited_works","source_observed_at":null,"state":"measured"}],"external_citation_measurements":[],"inbound":[],"links":{"evidence":"/evidence","html":"/paper/2411.10606/citation-record","integrity":"/paper/2411.10606/integrity","json":"/paper/2411.10606/citation-record.json","paper":"/paper/2411.10606"},"outbound":[{"citation":{"cited_paper":{"arxiv_id":"2307.09288","last_updated":"2023-07-19T17:08:59Z","snapshot_observed_at":"2026-08-07T12:56:43.323460Z","submitted_at":"2023-07-18T14:31:57Z","title":"Llama 2: Open Foundation and Fine-Tuned Chat Models","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2307.09288","snapshot_observed_at":"2026-08-12T19:35:29.873235Z","title":"Llama 2: Open foundation and fine-tuned chat models","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2411.10606","last_updated":"2024-11-15T22:02:28Z","snapshot_observed_at":"2026-08-14T13:50:36.265901Z","submitted_at":"2024-11-15T22:02:28Z","title":"AmoebaLLM: Constructing Any-Shape Large Language Models for Efficient and Instant Deployment","version":1},"reference_index":1,"source":"pdf_text","source_observed_at":"2026-08-12T19:35:29.873235Z"},"links":{"cited_paper":"/paper/2307.09288","citing_paper":"/paper/2411.10606"},"observation_digest":"sha256:d7aee076915b07dc0f827c70b31c1a8bd89c3c70a761f0bf5990afdcf6565cb9","observation_id":"7ac2c22a-3db1-4391-b9e4-b2faa01fe961","resolution":{"observed_at":"2026-08-12T19:35:29.873235Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-12T19:35:29.878998Z","title":"Introducing Meta Llama 3: The most capable openly available LLM to date, 2024","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2411.10606","last_updated":"2024-11-15T22:02:28Z","snapshot_observed_at":"2026-08-14T13:50:36.265901Z","submitted_at":"2024-11-15T22:02:28Z","title":"AmoebaLLM: Constructing Any-Shape Large Language Models for Efficient and Instant Deployment","version":1},"reference_index":2,"source":"pdf_text","source_observed_at":"2026-08-12T19:35:29.878998Z"},"links":{"citing_paper":"/paper/2411.10606"},"observation_digest":"sha256:d58ce86e7f00b57927a9c0d126a90f6e66a67f11d577f210348fda36e520f7b6","observation_id":"3a050736-faf3-42dd-af31-04a86d07851f","resolution":{"observed_at":"2026-08-12T19:35:29.878998Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2408.00118","last_updated":"2024-10-02T15:22:49Z","snapshot_observed_at":"2026-08-02T16:20:09.773989Z","submitted_at":"2024-07-31T19:13:07Z","title":"Gemma 2: Improving Open Language Models at a Practical Size","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2408.00118","snapshot_observed_at":"2026-08-12T19:35:29.883496Z","title":"Gemma 2: Improving open language models at a practical size","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2411.10606","last_updated":"2024-11-15T22:02:28Z","snapshot_observed_at":"2026-08-14T13:50:36.265901Z","submitted_at":"2024-11-15T22:02:28Z","title":"AmoebaLLM: Constructing Any-Shape Large Language Models for Efficient and Instant Deployment","version":1},"reference_index":3,"source":"pdf_text","source_observed_at":"2026-08-12T19:35:29.883496Z"},"links":{"cited_paper":"/paper/2408.00118","citing_paper":"/paper/2411.10606"},"observation_digest":"sha256:01631c74e41f9f8b56726d3dd7ba6467bdb049892b9f05c6ecaeda1164fb3de3","observation_id":"4af90698-19e9-4c2d-b08c-d9bcc438bc60","resolution":{"observed_at":"2026-08-12T19:35:29.883496Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2303.08774","last_updated":"2024-03-04T06:01:33Z","snapshot_observed_at":"2026-08-17T09:58:46.058102Z","submitted_at":"2023-03-15T17:15:04Z","title":"GPT-4 Technical Report","version":6},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2303.08774","snapshot_observed_at":"2026-08-12T19:35:29.888698Z","title":"Gpt-4 technical report","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2411.10606","last_updated":"2024-11-15T22:02:28Z","snapshot_observed_at":"2026-08-14T13:50:36.265901Z","submitted_at":"2024-11-15T22:02:28Z","title":"AmoebaLLM: Constructing Any-Shape Large Language Models for Efficient and Instant Deployment","version":1},"reference_index":4,"source":"pdf_text","source_observed_at":"2026-08-12T19:35:29.888698Z"},"links":{"cited_paper":"/paper/2303.08774","citing_paper":"/paper/2411.10606"},"observation_digest":"sha256:6b2539cb0f2a170b7dd1fc11df783f225fa38c9c3d8114e740dd0129aa3c201c","observation_id":"33cffe53-7a20-4394-b8c2-81a4d4a040de","resolution":{"observed_at":"2026-08-12T19:35:29.888698Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-12T19:35:29.893809Z","title":"SparseGPT: Massive language models can be accurately pruned in one-shot, 2023","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2411.10606","last_updated":"2024-11-15T22:02:28Z","snapshot_observed_at":"2026-08-14T13:50:36.265901Z","submitted_at":"2024-11-15T22:02:28Z","title":"AmoebaLLM: Constructing Any-Shape Large Language Models for Efficient and Instant Deployment","version":1},"reference_index":5,"source":"pdf_text","source_observed_at":"2026-08-12T19:35:29.893809Z"},"links":{"citing_paper":"/paper/2411.10606"},"observation_digest":"sha256:a75c9efa0b4f725d1484018d350ac895520ff95cf308ee6c90e37423c6797d8f","observation_id":"e39b3ff4-9a02-45c0-aeb8-91d111890281","resolution":{"observed_at":"2026-08-12T19:35:29.893809Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-12T19:35:29.898632Z","title":"A simple and effective pruning approach for large language models, 2023","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2411.10606","last_updated":"2024-11-15T22:02:28Z","snapshot_observed_at":"2026-08-14T13:50:36.265901Z","submitted_at":"2024-11-15T22:02:28Z","title":"AmoebaLLM: Constructing Any-Shape Large Language Models for Efficient and Instant Deployment","version":1},"reference_index":6,"source":"pdf_text","source_observed_at":"2026-08-12T19:35:29.898632Z"},"links":{"citing_paper":"/paper/2411.10606"},"observation_digest":"sha256:7479a796d492979ca52d13d44fcb363a86f02579923e7671655450b894f2bd46","observation_id":"074d9850-6d45-47ac-afa7-87d5fa80a25c","resolution":{"observed_at":"2026-08-12T19:35:29.898632Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-12T19:35:29.904206Z","title":"Llm-pruner: On the structural pruning of large language models","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2411.10606","last_updated":"2024-11-15T22:02:28Z","snapshot_observed_at":"2026-08-14T13:50:36.265901Z","submitted_at":"2024-11-15T22:02:28Z","title":"AmoebaLLM: Constructing Any-Shape Large Language Models for Efficient and Instant Deployment","version":1},"reference_index":7,"source":"pdf_text","source_observed_at":"2026-08-12T19:35:29.904206Z"},"links":{"citing_paper":"/paper/2411.10606"},"observation_digest":"sha256:d6b07e2724cf46e22ec29e43feed7d3e2e5e86eb9351caa2c1398bba2e46abb8","observation_id":"ad6d53fb-7e62-4ac8-b6a7-51c0daa62415","resolution":{"observed_at":"2026-08-12T19:35:29.904206Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-12T19:35:29.908874Z","title":"Fluctuation-based adaptive structured pruning for large language models","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2411.10606","last_updated":"2024-11-15T22:02:28Z","snapshot_observed_at":"2026-08-14T13:50:36.265901Z","submitted_at":"2024-11-15T22:02:28Z","title":"AmoebaLLM: Constructing Any-Shape Large Language Models for Efficient and Instant Deployment","version":1},"reference_index":8,"source":"pdf_text","source_observed_at":"2026-08-12T19:35:29.908874Z"},"links":{"citing_paper":"/paper/2411.10606"},"observation_digest":"sha256:fc805b98a9443fa7e8e77310f2e2d7270e5d3530ab7cc4014f067cba4788baf9","observation_id":"6e89a609-dcc2-491c-8720-5ca2b7947cae","resolution":{"observed_at":"2026-08-12T19:35:29.908874Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2402.02834","last_updated":"2024-06-23T08:45:33Z","snapshot_observed_at":"2026-08-18T12:31:28.811633Z","submitted_at":"2024-02-05T09:44:49Z","title":"Shortened LLaMA: Depth Pruning for Large Language Models with Comparison of Retraining Methods","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2402.02834","snapshot_observed_at":"2026-08-12T19:35:29.913456Z","title":"Shortened llama: A simple depth pruning for large language models","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2411.10606","last_updated":"2024-11-15T22:02:28Z","snapshot_observed_at":"2026-08-14T13:50:36.265901Z","submitted_at":"2024-11-15T22:02:28Z","title":"AmoebaLLM: Constructing Any-Shape Large Language Models for Efficient and Instant Deployment","version":1},"reference_index":9,"source":"pdf_text","source_observed_at":"2026-08-12T19:35:29.913456Z"},"links":{"cited_paper":"/paper/2402.02834","citing_paper":"/paper/2411.10606"},"observation_digest":"sha256:4dcae2de96fffb249461745d9f9ac3693f5aafdab83e773bdac087c727bd71c6","observation_id":"49e3ff13-fe7c-44d7-8c7f-f67d75b71927","resolution":{"observed_at":"2026-08-12T19:35:29.913456Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2310.06694","last_updated":"2024-04-11T01:18:06Z","snapshot_observed_at":"2026-08-16T14:53:24.619970Z","submitted_at":"2023-10-10T15:13:30Z","title":"Sheared LLaMA: Accelerating Language Model Pre-training via Structured Pruning","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2310.06694","snapshot_observed_at":"2026-08-12T19:35:30.016955Z","title":"Sheared llama: Accelerating language model pre-training via structured pruning","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2411.10606","last_updated":"2024-11-15T22:02:28Z","snapshot_observed_at":"2026-08-14T13:50:36.265901Z","submitted_at":"2024-11-15T22:02:28Z","title":"AmoebaLLM: Constructing Any-Shape Large Language Models for Efficient and Instant Deployment","version":1},"reference_index":10,"source":"pdf_text","source_observed_at":"2026-08-12T19:35:30.016955Z"},"links":{"cited_paper":"/paper/2310.06694","citing_paper":"/paper/2411.10606"},"observation_digest":"sha256:55ed1602760b5d6af0cc57faa2dc897513b349496495cc25ddc51b487a5b4fd8","observation_id":"912d186c-7927-4131-97bc-9acd9026720a","resolution":{"observed_at":"2026-08-12T19:35:30.016955Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-12T19:35:30.022979Z","title":"Bignas: Scaling up neural architecture search with big single-stage models","venue":null,"work_id":null,"year":2020},"citing_paper":{"arxiv_id":"2411.10606","last_updated":"2024-11-15T22:02:28Z","snapshot_observed_at":"2026-08-14T13:50:36.265901Z","submitted_at":"2024-11-15T22:02:28Z","title":"AmoebaLLM: Constructing Any-Shape Large Language Models for Efficient and Instant Deployment","version":1},"reference_index":11,"source":"pdf_text","source_observed_at":"2026-08-12T19:35:30.022979Z"},"links":{"citing_paper":"/paper/2411.10606"},"observation_digest":"sha256:4875a5a3b3e620d97960b7f4b6518be897ebc734a844b363993f208ca669b642","observation_id":"21254411-3deb-4f1a-897e-846df2ca4c09","resolution":{"observed_at":"2026-08-12T19:35:30.022979Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-12T19:35:30.028587Z","title":"Attentivenas: Improving neural architecture search via attentive sampling","venue":null,"work_id":null,"year":2021},"citing_paper":{"arxiv_id":"2411.10606","last_updated":"2024-11-15T22:02:28Z","snapshot_observed_at":"2026-08-14T13:50:36.265901Z","submitted_at":"2024-11-15T22:02:28Z","title":"AmoebaLLM: Constructing Any-Shape Large Language Models for Efficient and Instant Deployment","version":1},"reference_index":12,"source":"pdf_text","source_observed_at":"2026-08-12T19:35:30.028587Z"},"links":{"citing_paper":"/paper/2411.10606"},"observation_digest":"sha256:2bca124c3b11845dd8a016d69161073bb013751d59155452dd705b3020273535","observation_id":"918895aa-2039-461c-b132-04d9e7021f1a","resolution":{"observed_at":"2026-08-12T19:35:30.028587Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-12T19:35:30.033535Z","title":"Alphanet: Improved training of supernets with alpha-divergence","venue":null,"work_id":null,"year":2021},"citing_paper":{"arxiv_id":"2411.10606","last_updated":"2024-11-15T22:02:28Z","snapshot_observed_at":"2026-08-14T13:50:36.265901Z","submitted_at":"2024-11-15T22:02:28Z","title":"AmoebaLLM: Constructing Any-Shape Large Language Models for Efficient and Instant Deployment","version":1},"reference_index":13,"source":"pdf_text","source_observed_at":"2026-08-12T19:35:30.033535Z"},"links":{"citing_paper":"/paper/2411.10606"},"observation_digest":"sha256:725cdd1a9f31c6103b4116bf33b031f50419a69aaa509bf3d34d2fcbc1dcdb84","observation_id":"8ffcdd17-8bb3-4407-a790-b3f4b800c2c6","resolution":{"observed_at":"2026-08-12T19:35:30.033535Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-12T19:35:30.039214Z","title":"Nasvit: Neural architecture search for efficient vision transformers with gradient conflict-aware supernet training","venue":null,"work_id":null,"year":2022},"citing_paper":{"arxiv_id":"2411.10606","last_updated":"2024-11-15T22:02:28Z","snapshot_observed_at":"2026-08-14T13:50:36.265901Z","submitted_at":"2024-11-15T22:02:28Z","title":"AmoebaLLM: Constructing Any-Shape Large Language Models for Efficient and Instant Deployment","version":1},"reference_index":14,"source":"pdf_text","source_observed_at":"2026-08-12T19:35:30.039214Z"},"links":{"citing_paper":"/paper/2411.10606"},"observation_digest":"sha256:988c4f617d230446665c2b80d5841ddb4bbac9ed425a777b93ab00a1b79692b9","observation_id":"09e77131-0bb1-4ea7-a568-fee8226a24c7","resolution":{"observed_at":"2026-08-12T19:35:30.039214Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"1908.09791","last_updated":"2020-04-29T20:49:05Z","snapshot_observed_at":"2026-08-14T10:59:14.094474Z","submitted_at":"2019-08-26T16:46:23Z","title":"Once-for-All: Train One Network and Specialize it for Efficient Deployment","version":5},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"1908.09791","snapshot_observed_at":"2026-08-12T19:35:30.044626Z","title":"Once-for-all: Train one network and specialize it for efficient deployment","venue":null,"work_id":null,"year":1908},"citing_paper":{"arxiv_id":"2411.10606","last_updated":"2024-11-15T22:02:28Z","snapshot_observed_at":"2026-08-14T13:50:36.265901Z","submitted_at":"2024-11-15T22:02:28Z","title":"AmoebaLLM: Constructing Any-Shape Large Language Models for Efficient and Instant Deployment","version":1},"reference_index":15,"source":"pdf_text","source_observed_at":"2026-08-12T19:35:30.044626Z"},"links":{"cited_paper":"/paper/1908.09791","citing_paper":"/paper/2411.10606"},"observation_digest":"sha256:fa8d146e606835ebf4dd64cf5f4ba27cd8d21d0e3f550c35e0c1e01e502dcb44","observation_id":"e7171d6d-3dc5-4d10-b95b-6203468d0ef0","resolution":{"observed_at":"2026-08-12T19:35:30.044626Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-12T19:35:30.050418Z","title":"Gradient surgery for multi-task learning","venue":null,"work_id":null,"year":2020},"citing_paper":{"arxiv_id":"2411.10606","last_updated":"2024-11-15T22:02:28Z","snapshot_observed_at":"2026-08-14T13:50:36.265901Z","submitted_at":"2024-11-15T22:02:28Z","title":"AmoebaLLM: Constructing Any-Shape Large Language Models for Efficient and Instant Deployment","version":1},"reference_index":16,"source":"pdf_text","source_observed_at":"2026-08-12T19:35:30.050418Z"},"links":{"citing_paper":"/paper/2411.10606"},"observation_digest":"sha256:205764b2d8e169aea72d11bc46aac7ff8751919081593cb80d9f51612292fa77","observation_id":"99092a33-0125-4b28-b0c0-989611a220fd","resolution":{"observed_at":"2026-08-12T19:35:30.050418Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-12T19:35:30.055502Z","title":"Conflict-averse gradient descent for multi-task learning","venue":null,"work_id":null,"year":2021},"citing_paper":{"arxiv_id":"2411.10606","last_updated":"2024-11-15T22:02:28Z","snapshot_observed_at":"2026-08-14T13:50:36.265901Z","submitted_at":"2024-11-15T22:02:28Z","title":"AmoebaLLM: Constructing Any-Shape Large Language Models for Efficient and Instant Deployment","version":1},"reference_index":17,"source":"pdf_text","source_observed_at":"2026-08-12T19:35:30.055502Z"},"links":{"citing_paper":"/paper/2411.10606"},"observation_digest":"sha256:1d056482ce18708f15010146eae9ba9be6cbcc19e2a0619d8736fb7c387b19d7","observation_id":"23df3483-031f-4980-83b1-fa4f5f06e275","resolution":{"observed_at":"2026-08-12T19:35:30.055502Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2106.09685","last_updated":"2021-10-16T18:40:34Z","snapshot_observed_at":"2026-08-17T18:04:53.578114Z","submitted_at":"2021-06-17T17:37:18Z","title":"LoRA: Low-Rank Adaptation of Large Language Models","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2106.09685","snapshot_observed_at":"2026-08-12T19:35:30.061058Z","title":"Lora: Low-rank adaptation of large language models","venue":null,"work_id":null,"year":2021},"citing_paper":{"arxiv_id":"2411.10606","last_updated":"2024-11-15T22:02:28Z","snapshot_observed_at":"2026-08-14T13:50:36.265901Z","submitted_at":"2024-11-15T22:02:28Z","title":"AmoebaLLM: Constructing Any-Shape Large Language Models for Efficient and Instant Deployment","version":1},"reference_index":18,"source":"pdf_text","source_observed_at":"2026-08-12T19:35:30.061058Z"},"links":{"cited_paper":"/paper/2106.09685","citing_paper":"/paper/2411.10606"},"observation_digest":"sha256:25261af02dc13ca73ebb07a61dd39d84dfd2e89e24f019fbcb976eebc83af8f0","observation_id":"2a24467d-a2fa-4b9c-a123-16d08bb7ce07","resolution":{"observed_at":"2026-08-12T19:35:30.061058Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-12T19:35:30.066109Z","title":"Tensorrt-llm, 2024","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2411.10606","last_updated":"2024-11-15T22:02:28Z","snapshot_observed_at":"2026-08-14T13:50:36.265901Z","submitted_at":"2024-11-15T22:02:28Z","title":"AmoebaLLM: Constructing Any-Shape Large Language Models for Efficient and Instant Deployment","version":1},"reference_index":19,"source":"pdf_text","source_observed_at":"2026-08-12T19:35:30.066109Z"},"links":{"citing_paper":"/paper/2411.10606"},"observation_digest":"sha256:470f6973f60501641757104a94c75d0e384cdd3b61c178dc8ae6b12a1b8acfc6","observation_id":"6e2a80fe-defc-43da-81d3-5dc364a2e3a8","resolution":{"observed_at":"2026-08-12T19:35:30.066109Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-12T19:35:30.070491Z","title":"MLC-LLM, 2023","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2411.10606","last_updated":"2024-11-15T22:02:28Z","snapshot_observed_at":"2026-08-14T13:50:36.265901Z","submitted_at":"2024-11-15T22:02:28Z","title":"AmoebaLLM: Constructing Any-Shape Large Language Models for Efficient and Instant Deployment","version":1},"reference_index":20,"source":"pdf_text","source_observed_at":"2026-08-12T19:35:30.070491Z"},"links":{"citing_paper":"/paper/2411.10606"},"observation_digest":"sha256:7800941deb4b8360751d4f4575570bb6381942cf2c17fa5bbb206641020cd4e0","observation_id":"72be604c-366d-497e-919f-3598ab8f8f73","resolution":{"observed_at":"2026-08-12T19:35:30.070491Z","resolver_source":null,"status":"parse_uncertain"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-12T19:35:30.074742Z","title":"Pytorch: An imperative style, high-performance deep learning library","venue":null,"work_id":null,"year":2019},"citing_paper":{"arxiv_id":"2411.10606","last_updated":"2024-11-15T22:02:28Z","snapshot_observed_at":"2026-08-14T13:50:36.265901Z","submitted_at":"2024-11-15T22:02:28Z","title":"AmoebaLLM: Constructing Any-Shape Large Language Models for Efficient and Instant Deployment","version":1},"reference_index":21,"source":"pdf_text","source_observed_at":"2026-08-12T19:35:30.074742Z"},"links":{"citing_paper":"/paper/2411.10606"},"observation_digest":"sha256:639d32bf65d7894d2e097c4898b0638659dd9cf318173e1dcf86659aeb15eecb","observation_id":"f3395e26-8cd9-4141-9683-3d9f875dfe64","resolution":{"observed_at":"2026-08-12T19:35:30.074742Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-12T19:35:30.079145Z","title":"The theory of dynamic programming","venue":null,"work_id":null,"year":1954},"citing_paper":{"arxiv_id":"2411.10606","last_updated":"2024-11-15T22:02:28Z","snapshot_observed_at":"2026-08-14T13:50:36.265901Z","submitted_at":"2024-11-15T22:02:28Z","title":"AmoebaLLM: Constructing Any-Shape Large Language Models for Efficient and Instant Deployment","version":1},"reference_index":22,"source":"pdf_text","source_observed_at":"2026-08-12T19:35:30.079145Z"},"links":{"citing_paper":"/paper/2411.10606"},"observation_digest":"sha256:772c7d2dbc276888abd2d80a21abf7d07a8bd504a3c69d33671c2a06bc79e599","observation_id":"d3468837-168b-4cd9-b4c3-f073a01e5a3e","resolution":{"observed_at":"2026-08-12T19:35:30.079145Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2405.05904","last_updated":"2024-10-01T12:08:23Z","snapshot_observed_at":"2026-08-16T13:53:58.560866Z","submitted_at":"2024-05-09T17:00:22Z","title":"Does Fine-Tuning LLMs on New Knowledge Encourage Hallucinations?","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2405.05904","snapshot_observed_at":"2026-08-12T19:35:30.083461Z","title":"Does fine-tuning llms on new knowledge encourage hallucinations? arXiv preprint arXiv:2405.05904, 2024","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2411.10606","last_updated":"2024-11-15T22:02:28Z","snapshot_observed_at":"2026-08-14T13:50:36.265901Z","submitted_at":"2024-11-15T22:02:28Z","title":"AmoebaLLM: Constructing Any-Shape Large Language Models for Efficient and Instant Deployment","version":1},"reference_index":23,"source":"pdf_text","source_observed_at":"2026-08-12T19:35:30.083461Z"},"links":{"cited_paper":"/paper/2405.05904","citing_paper":"/paper/2411.10606"},"observation_digest":"sha256:9a066aabdb1dcc7c88aab2ee4900b5287adc3fc7c4e7d3aff71005b92cc3ae1a","observation_id":"31f05847-ef84-494d-971d-d9e1782e620e","resolution":{"observed_at":"2026-08-12T19:35:30.083461Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2401.01286","last_updated":"2024-11-17T06:50:44Z","snapshot_observed_at":"2026-08-17T20:32:47.955697Z","submitted_at":"2024-01-02T16:54:58Z","title":"A Comprehensive Study of Knowledge Editing for Large Language Models","version":5},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2401.01286","snapshot_observed_at":"2026-08-12T19:35:30.087787Z","title":"A comprehensive study of knowledge editing for large language models","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2411.10606","last_updated":"2024-11-15T22:02:28Z","snapshot_observed_at":"2026-08-14T13:50:36.265901Z","submitted_at":"2024-11-15T22:02:28Z","title":"AmoebaLLM: Constructing Any-Shape Large Language Models for Efficient and Instant Deployment","version":1},"reference_index":24,"source":"pdf_text","source_observed_at":"2026-08-12T19:35:30.087787Z"},"links":{"cited_paper":"/paper/2401.01286","citing_paper":"/paper/2411.10606"},"observation_digest":"sha256:232a3972122c21735ba111418797398f6c8294858b1926cacecc75d07d3990ed","observation_id":"2d961584-828c-441e-afe4-097cb45c8b44","resolution":{"observed_at":"2026-08-12T19:35:30.087787Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2104.08696","last_updated":"2022-03-10T02:28:59Z","snapshot_observed_at":"2026-08-18T21:13:11.433434Z","submitted_at":"2021-04-18T03:38:26Z","title":"Knowledge Neurons in Pretrained Transformers","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2104.08696","snapshot_observed_at":"2026-08-12T19:35:30.092157Z","title":"Knowledge neurons in pretrained transformers","venue":null,"work_id":null,"year":2021},"citing_paper":{"arxiv_id":"2411.10606","last_updated":"2024-11-15T22:02:28Z","snapshot_observed_at":"2026-08-14T13:50:36.265901Z","submitted_at":"2024-11-15T22:02:28Z","title":"AmoebaLLM: Constructing Any-Shape Large Language Models for Efficient and Instant Deployment","version":1},"reference_index":25,"source":"pdf_text","source_observed_at":"2026-08-12T19:35:30.092157Z"},"links":{"cited_paper":"/paper/2104.08696","citing_paper":"/paper/2411.10606"},"observation_digest":"sha256:754762de6b6e272b0c6200807f04dedfe076eacc1cf5c8070b93f1066c903641","observation_id":"15244494-8c89-43df-a223-6b0d52aace79","resolution":{"observed_at":"2026-08-12T19:35:30.092157Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2012.14913","last_updated":"2021-09-05T17:32:27Z","snapshot_observed_at":"2026-08-13T01:11:33.500647Z","submitted_at":"2020-12-29T19:12:05Z","title":"Transformer Feed-Forward Layers Are Key-Value Memories","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2012.14913","snapshot_observed_at":"2026-08-12T19:35:30.098960Z","title":"Transformer feed-forward layers are key-value memories","venue":null,"work_id":null,"year":2012},"citing_paper":{"arxiv_id":"2411.10606","last_updated":"2024-11-15T22:02:28Z","snapshot_observed_at":"2026-08-14T13:50:36.265901Z","submitted_at":"2024-11-15T22:02:28Z","title":"AmoebaLLM: Constructing Any-Shape Large Language Models for Efficient and Instant Deployment","version":1},"reference_index":26,"source":"pdf_text","source_observed_at":"2026-08-12T19:35:30.098960Z"},"links":{"cited_paper":"/paper/2012.14913","citing_paper":"/paper/2411.10606"},"observation_digest":"sha256:dd058249b651dd4efb967564693e2936500e27d936cce0f9a6d3fd82b98142a1","observation_id":"b0f2183a-70a2-4b89-90a6-1b841ee4eb84","resolution":{"observed_at":"2026-08-12T19:35:30.098960Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-12T19:35:30.103475Z","title":"What does bert learn about the structure of language? In ACL 2019-57th Annual Meeting of the Association for Computational Linguistics, 2019","venue":null,"work_id":null,"year":2019},"citing_paper":{"arxiv_id":"2411.10606","last_updated":"2024-11-15T22:02:28Z","snapshot_observed_at":"2026-08-14T13:50:36.265901Z","submitted_at":"2024-11-15T22:02:28Z","title":"AmoebaLLM: Constructing Any-Shape Large Language Models for Efficient and Instant Deployment","version":1},"reference_index":27,"source":"pdf_text","source_observed_at":"2026-08-12T19:35:30.103475Z"},"links":{"citing_paper":"/paper/2411.10606"},"observation_digest":"sha256:8bf7b3713533e4d2728f84bf6e141c6eb5e1f95e3610fb30a86526d855f2c109","observation_id":"b84bdce9-a71c-469c-a51e-6176cdee0a64","resolution":{"observed_at":"2026-08-12T19:35:30.103475Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-12T19:35:30.107989Z","title":"Locating and editing factual associations in gpt","venue":null,"work_id":null,"year":2022},"citing_paper":{"arxiv_id":"2411.10606","last_updated":"2024-11-15T22:02:28Z","snapshot_observed_at":"2026-08-14T13:50:36.265901Z","submitted_at":"2024-11-15T22:02:28Z","title":"AmoebaLLM: Constructing Any-Shape Large Language Models for Efficient and Instant Deployment","version":1},"reference_index":28,"source":"pdf_text","source_observed_at":"2026-08-12T19:35:30.107989Z"},"links":{"citing_paper":"/paper/2411.10606"},"observation_digest":"sha256:90f3d2b249fc660e7752f39e7f001c3ac9082dba3ed4ae546ecbfbaa63f07d35","observation_id":"bb0176d4-7c36-4e88-a59a-e045890fa948","resolution":{"observed_at":"2026-08-12T19:35:30.107989Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2402.16061","last_updated":"2024-03-04T13:37:48Z","snapshot_observed_at":"2026-08-18T08:30:13.765079Z","submitted_at":"2024-02-25T11:15:42Z","title":"How Large Language Models Encode Context Knowledge? A Layer-Wise Probing Study","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2402.16061","snapshot_observed_at":"2026-08-12T19:35:30.112667Z","title":"How large language models encode context knowledge? a layer-wise probing study","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2411.10606","last_updated":"2024-11-15T22:02:28Z","snapshot_observed_at":"2026-08-14T13:50:36.265901Z","submitted_at":"2024-11-15T22:02:28Z","title":"AmoebaLLM: Constructing Any-Shape Large Language Models for Efficient and Instant Deployment","version":1},"reference_index":29,"source":"pdf_text","source_observed_at":"2026-08-12T19:35:30.112667Z"},"links":{"cited_paper":"/paper/2402.16061","citing_paper":"/paper/2411.10606"},"observation_digest":"sha256:9ee2b45c17e4b95fe736fbea42a146bc3b31f126a8b4b48895f8803d41a50f36","observation_id":"ef8ef89b-5b45-4bb5-a05b-30d93c0b01aa","resolution":{"observed_at":"2026-08-12T19:35:30.112667Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-12T19:35:30.117360Z","title":"Deep residual learning for image recognition","venue":null,"work_id":null,"year":2016},"citing_paper":{"arxiv_id":"2411.10606","last_updated":"2024-11-15T22:02:28Z","snapshot_observed_at":"2026-08-14T13:50:36.265901Z","submitted_at":"2024-11-15T22:02:28Z","title":"AmoebaLLM: Constructing Any-Shape Large Language Models for Efficient and Instant Deployment","version":1},"reference_index":30,"source":"pdf_text","source_observed_at":"2026-08-12T19:35:30.117360Z"},"links":{"citing_paper":"/paper/2411.10606"},"observation_digest":"sha256:c541db78f7591aadf7084d7eaa99b66880713cdcd774e4891c871decdd574c8a","observation_id":"6672af4e-b6f4-4403-9c7c-661b58a03566","resolution":{"observed_at":"2026-08-12T19:35:30.117360Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2305.13172","last_updated":"2023-11-30T08:55:24Z","snapshot_observed_at":"2026-08-19T01:06:15.177630Z","submitted_at":"2023-05-22T16:00:00Z","title":"Editing Large Language Models: Problems, Methods, and Opportunities","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2305.13172","snapshot_observed_at":"2026-08-12T19:35:30.122172Z","title":"Editing large language models: Problems, methods, and opportunities","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2411.10606","last_updated":"2024-11-15T22:02:28Z","snapshot_observed_at":"2026-08-14T13:50:36.265901Z","submitted_at":"2024-11-15T22:02:28Z","title":"AmoebaLLM: Constructing Any-Shape Large Language Models for Efficient and Instant Deployment","version":1},"reference_index":31,"source":"pdf_text","source_observed_at":"2026-08-12T19:35:30.122172Z"},"links":{"cited_paper":"/paper/2305.13172","citing_paper":"/paper/2411.10606"},"observation_digest":"sha256:4df0ed20c30d811d2fab729fbade4d4ce020762c0bba7633552ee7af257e8bc4","observation_id":"c879663c-11fb-4568-bba5-ac0ec6260f3a","resolution":{"observed_at":"2026-08-12T19:35:30.122172Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2308.07269","last_updated":"2024-06-24T02:17:57Z","snapshot_observed_at":"2026-08-16T15:08:35.243070Z","submitted_at":"2023-08-14T16:52:42Z","title":"EasyEdit: An Easy-to-use Knowledge Editing Framework for Large Language Models","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2308.07269","snapshot_observed_at":"2026-08-12T19:35:30.126702Z","title":"Easyedit: An easy-to-use knowledge editing framework for large language models","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2411.10606","last_updated":"2024-11-15T22:02:28Z","snapshot_observed_at":"2026-08-14T13:50:36.265901Z","submitted_at":"2024-11-15T22:02:28Z","title":"AmoebaLLM: Constructing Any-Shape Large Language Models for Efficient and Instant Deployment","version":1},"reference_index":32,"source":"pdf_text","source_observed_at":"2026-08-12T19:35:30.126702Z"},"links":{"cited_paper":"/paper/2308.07269","citing_paper":"/paper/2411.10606"},"observation_digest":"sha256:b132f7bdc1dd74efca2b39ec9d26769b71491210f320c6269211d2cd852abe16","observation_id":"a1a04534-4b26-4923-9e5b-cabe5f33667b","resolution":{"observed_at":"2026-08-12T19:35:30.126702Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2403.17887","last_updated":"2025-03-03T17:02:05Z","snapshot_observed_at":"2026-08-16T14:06:10.494270Z","submitted_at":"2024-03-26T17:20:04Z","title":"The Unreasonable Ineffectiveness of the Deeper Layers","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2403.17887","snapshot_observed_at":"2026-08-12T19:35:30.131623Z","title":"The unreasonable ineffectiveness of the deeper layers","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2411.10606","last_updated":"2024-11-15T22:02:28Z","snapshot_observed_at":"2026-08-14T13:50:36.265901Z","submitted_at":"2024-11-15T22:02:28Z","title":"AmoebaLLM: Constructing Any-Shape Large Language Models for Efficient and Instant Deployment","version":1},"reference_index":33,"source":"pdf_text","source_observed_at":"2026-08-12T19:35:30.131623Z"},"links":{"cited_paper":"/paper/2403.17887","citing_paper":"/paper/2411.10606"},"observation_digest":"sha256:353032ccada8cf675d841b0441994ae29f9da2bb72fd9a872cf08660f0382fb2","observation_id":"392a327d-5116-4ffc-8d15-109b458dfbc8","resolution":{"observed_at":"2026-08-12T19:35:30.131623Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-12T19:35:31.704262Z","title":"Flexible group-level pruning of deep neural networks for on-device machine learning","venue":null,"work_id":"700a266d-5612-4278-9b16-ba951d71b8b6","year":2020},"citing_paper":{"arxiv_id":"2411.10606","last_updated":"2024-11-15T22:02:28Z","snapshot_observed_at":"2026-08-14T13:50:36.265901Z","submitted_at":"2024-11-15T22:02:28Z","title":"AmoebaLLM: Constructing Any-Shape Large Language Models for Efficient and Instant Deployment","version":1},"reference_index":34,"source":"pdf_text","source_observed_at":"2026-08-12T19:35:30.136613Z"},"links":{"citing_paper":"/paper/2411.10606"},"observation_digest":"sha256:32138444bce10e0bb5b1f0d0626a972296b4c6a5dc0773c180a3d462e9be1801","observation_id":"e61bbb7e-0b56-47b8-82c1-f037072c12f1","resolution":{"observed_at":"2026-08-12T19:35:31.709874Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-19T06:32:44.657259+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-19T06:32:44.657259+00:00","source":"crossref"},{"observed_at":"2026-08-19T06:32:39.956319+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-12T19:35:30.141013Z","title":"Efficient joint optimization of layer-adaptive weight pruning in deep neural networks","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2411.10606","last_updated":"2024-11-15T22:02:28Z","snapshot_observed_at":"2026-08-14T13:50:36.265901Z","submitted_at":"2024-11-15T22:02:28Z","title":"AmoebaLLM: Constructing Any-Shape Large Language Models for Efficient and Instant Deployment","version":1},"reference_index":35,"source":"pdf_text","source_observed_at":"2026-08-12T19:35:30.141013Z"},"links":{"citing_paper":"/paper/2411.10606"},"observation_digest":"sha256:7353e0b60f672ce2ce5ed30476dc3ec4fdf7d1df6fbe96467cb00f44e548a7c1","observation_id":"1d22120c-601d-4f95-9fb2-d78a6faa38a5","resolution":{"observed_at":"2026-08-12T19:35:30.141013Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"1701.06538","last_updated":"2017-01-23T18:10:00Z","snapshot_observed_at":"2026-08-13T11:35:07.866136Z","submitted_at":"2017-01-23T18:10:00Z","title":"Outrageously Large Neural Networks: The Sparsely-Gated Mixture-of-Experts Layer","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"1701.06538","snapshot_observed_at":"2026-08-12T19:35:30.145804Z","title":"Outrageously large neural networks: The sparsely-gated mixture-of-experts layer","venue":null,"work_id":null,"year":2017},"citing_paper":{"arxiv_id":"2411.10606","last_updated":"2024-11-15T22:02:28Z","snapshot_observed_at":"2026-08-14T13:50:36.265901Z","submitted_at":"2024-11-15T22:02:28Z","title":"AmoebaLLM: Constructing Any-Shape Large Language Models for Efficient and Instant Deployment","version":1},"reference_index":36,"source":"pdf_text","source_observed_at":"2026-08-12T19:35:30.145804Z"},"links":{"cited_paper":"/paper/1701.06538","citing_paper":"/paper/2411.10606"},"observation_digest":"sha256:906ad6856d0c16a615df330c14fe37ae074846c9b854d59afa41f761aab7b54c","observation_id":"0a017fc1-d8cb-4b73-a198-0bbf65be26f8","resolution":{"observed_at":"2026-08-12T19:35:30.145804Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-12T19:35:31.674167Z","title":"Mole: Mixture of lora experts","venue":null,"work_id":"66f18d68-031c-4d9f-bf90-e160c422bf02","year":2023},"citing_paper":{"arxiv_id":"2411.10606","last_updated":"2024-11-15T22:02:28Z","snapshot_observed_at":"2026-08-14T13:50:36.265901Z","submitted_at":"2024-11-15T22:02:28Z","title":"AmoebaLLM: Constructing Any-Shape Large Language Models for Efficient and Instant Deployment","version":1},"reference_index":37,"source":"pdf_text","source_observed_at":"2026-08-12T19:35:30.150557Z"},"links":{"citing_paper":"/paper/2411.10606"},"observation_digest":"sha256:2dda16dbff81ee6fbf54f38778f65a37617f6cbbf3ec7833634e2ccd5902a0be","observation_id":"75dc7c37-c307-4666-87fb-a6bb8813fc63","resolution":{"observed_at":"2026-08-12T19:35:31.679009Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-19T06:32:44.657259+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-19T06:32:44.657259+00:00","source":"crossref"},{"observed_at":"2026-08-19T06:32:39.956319+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2401.16160","last_updated":"2024-01-30T15:44:58Z","snapshot_observed_at":"2026-08-17T07:39:58.326895Z","submitted_at":"2024-01-29T13:48:36Z","title":"LLaVA-MoLE: Sparse Mixture of LoRA Experts for Mitigating Data Conflicts in Instruction Finetuning MLLMs","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2401.16160","snapshot_observed_at":"2026-08-12T19:35:30.155374Z","title":"Llava-mole: Sparse mixture of lora experts for mitigating data conflicts in instruction finetuning mllms","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2411.10606","last_updated":"2024-11-15T22:02:28Z","snapshot_observed_at":"2026-08-14T13:50:36.265901Z","submitted_at":"2024-11-15T22:02:28Z","title":"AmoebaLLM: Constructing Any-Shape Large Language Models for Efficient and Instant Deployment","version":1},"reference_index":38,"source":"pdf_text","source_observed_at":"2026-08-12T19:35:30.155374Z"},"links":{"cited_paper":"/paper/2401.16160","citing_paper":"/paper/2411.10606"},"observation_digest":"sha256:086e1e163affe581a3e2be6a69b26be01360d04693a376ae4d5b371168da06d1","observation_id":"4d0a2219-df45-403c-a336-546f01181b40","resolution":{"observed_at":"2026-08-12T19:35:30.155374Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2404.15159","last_updated":"2024-07-20T02:26:49Z","snapshot_observed_at":"2026-08-16T13:59:03.312702Z","submitted_at":"2024-04-22T02:15:52Z","title":"MixLoRA: Enhancing Large Language Models Fine-Tuning with LoRA-based Mixture of Experts","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2404.15159","snapshot_observed_at":"2026-08-12T19:35:30.160491Z","title":"Mixlora: Enhancing large language models fine-tuning with lora based mixture of experts","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2411.10606","last_updated":"2024-11-15T22:02:28Z","snapshot_observed_at":"2026-08-14T13:50:36.265901Z","submitted_at":"2024-11-15T22:02:28Z","title":"AmoebaLLM: Constructing Any-Shape Large Language Models for Efficient and Instant Deployment","version":1},"reference_index":39,"source":"pdf_text","source_observed_at":"2026-08-12T19:35:30.160491Z"},"links":{"cited_paper":"/paper/2404.15159","citing_paper":"/paper/2411.10606"},"observation_digest":"sha256:4f2b52f036e2ccaca81c2829a440239c237eec2dac13ea9c5cefc2695a18b0bb","observation_id":"50e26a20-33dc-4bd7-b0ca-8159b97d1588","resolution":{"observed_at":"2026-08-12T19:35:30.160491Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-12T19:35:31.657758Z","title":"Stanford alpaca: An instruction-following llama model","venue":null,"work_id":"16114f9b-3a29-4b0e-9694-8eab1e3f249a","year":2023},"citing_paper":{"arxiv_id":"2411.10606","last_updated":"2024-11-15T22:02:28Z","snapshot_observed_at":"2026-08-14T13:50:36.265901Z","submitted_at":"2024-11-15T22:02:28Z","title":"AmoebaLLM: Constructing Any-Shape Large Language Models for Efficient and Instant Deployment","version":1},"reference_index":40,"source":"pdf_text","source_observed_at":"2026-08-12T19:35:30.165671Z"},"links":{"citing_paper":"/paper/2411.10606"},"observation_digest":"sha256:c97120088d0c762ccb0b9deb3cd82667e32a429b05b69b28046f38fec928b33d","observation_id":"4389acc7-60c2-4e0f-97e4-d38bd9f159f2","resolution":{"observed_at":"2026-08-12T19:35:31.662834Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-19T06:32:44.657259+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-19T06:32:44.657259+00:00","source":"crossref"},{"observed_at":"2026-08-19T06:32:39.956319+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-12T19:35:31.640192Z","title":"Vicuna: An open-source chatbot impressing gpt-4 with 90%* chatgpt quality, 2023","venue":null,"work_id":"fc0daee4-20c6-40a5-acf4-c6f23a6e3299","year":2023},"citing_paper":{"arxiv_id":"2411.10606","last_updated":"2024-11-15T22:02:28Z","snapshot_observed_at":"2026-08-14T13:50:36.265901Z","submitted_at":"2024-11-15T22:02:28Z","title":"AmoebaLLM: Constructing Any-Shape Large Language Models for Efficient and Instant Deployment","version":1},"reference_index":41,"source":"pdf_text","source_observed_at":"2026-08-12T19:35:30.170468Z"},"links":{"citing_paper":"/paper/2411.10606"},"observation_digest":"sha256:f3be91b94bb822591b0c39555fee47f7b4990fe4713937b88b5e85a468fa273e","observation_id":"1ebad3c0-60a4-402e-bffb-1c94a342d1b7","resolution":{"observed_at":"2026-08-12T19:35:31.645889Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-19T06:32:44.657259+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-19T06:32:44.657259+00:00","source":"crossref"},{"observed_at":"2026-08-19T06:32:39.956319+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-12T19:35:31.622557Z","title":"Language model evaluation harness (package version caaf9ab)","venue":null,"work_id":"6e8c4821-00cf-4f1a-a015-f9d2df0152c8","year":2023},"citing_paper":{"arxiv_id":"2411.10606","last_updated":"2024-11-15T22:02:28Z","snapshot_observed_at":"2026-08-14T13:50:36.265901Z","submitted_at":"2024-11-15T22:02:28Z","title":"AmoebaLLM: Constructing Any-Shape Large Language Models for Efficient and Instant Deployment","version":1},"reference_index":42,"source":"pdf_text","source_observed_at":"2026-08-12T19:35:30.175250Z"},"links":{"citing_paper":"/paper/2411.10606"},"observation_digest":"sha256:b71e7154f7c6308a0dea1161a6b09a25aa618e9d57d2a617b0c63057095025a6","observation_id":"9c9b451a-0730-4739-8088-e61a5f3d8690","resolution":{"observed_at":"2026-08-12T19:35:31.627647Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-19T06:32:44.657259+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-19T06:32:44.657259+00:00","source":"crossref"},{"observed_at":"2026-08-19T06:32:39.956319+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-12T19:35:31.605155Z","title":"BoolQ: Exploring the surprising difficulty of natural yes/no questions","venue":null,"work_id":"29053399-f7fb-4004-bcdc-9849c465bce0","year":2019},"citing_paper":{"arxiv_id":"2411.10606","last_updated":"2024-11-15T22:02:28Z","snapshot_observed_at":"2026-08-14T13:50:36.265901Z","submitted_at":"2024-11-15T22:02:28Z","title":"AmoebaLLM: Constructing Any-Shape Large Language Models for Efficient and Instant Deployment","version":1},"reference_index":43,"source":"pdf_text","source_observed_at":"2026-08-12T19:35:30.180109Z"},"links":{"citing_paper":"/paper/2411.10606"},"observation_digest":"sha256:01b9accd6e528a74c164dfe19379f6b989cd7b02566e1ab87ba568c2cca31628","observation_id":"f4bca867-0d03-444a-9e1b-96ef646fadb3","resolution":{"observed_at":"2026-08-12T19:35:31.610413Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-19T06:32:44.657259+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-19T06:32:44.657259+00:00","source":"crossref"},{"observed_at":"2026-08-19T06:32:39.956319+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-12T19:35:30.184945Z","title":"Piqa: Reasoning about physical commonsense in natural language","venue":null,"work_id":null,"year":2020},"citing_paper":{"arxiv_id":"2411.10606","last_updated":"2024-11-15T22:02:28Z","snapshot_observed_at":"2026-08-14T13:50:36.265901Z","submitted_at":"2024-11-15T22:02:28Z","title":"AmoebaLLM: Constructing Any-Shape Large Language Models for Efficient and Instant Deployment","version":1},"reference_index":44,"source":"pdf_text","source_observed_at":"2026-08-12T19:35:30.184945Z"},"links":{"citing_paper":"/paper/2411.10606"},"observation_digest":"sha256:35105c8ab2a68e8e595b0b663fa1b21819c52971181bcc03f0fa4dd647d01b64","observation_id":"2fd0be4b-855b-4448-b225-18cc1d0fa9f4","resolution":{"observed_at":"2026-08-12T19:35:30.184945Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-12T19:35:31.572609Z","title":"Hellaswag: Can a machine really finish your sentence? In ACL, 2019","venue":null,"work_id":"89818540-ab7c-4fae-a4d5-931241b17eb8","year":2019},"citing_paper":{"arxiv_id":"2411.10606","last_updated":"2024-11-15T22:02:28Z","snapshot_observed_at":"2026-08-14T13:50:36.265901Z","submitted_at":"2024-11-15T22:02:28Z","title":"AmoebaLLM: Constructing Any-Shape Large Language Models for Efficient and Instant Deployment","version":1},"reference_index":45,"source":"pdf_text","source_observed_at":"2026-08-12T19:35:30.190059Z"},"links":{"citing_paper":"/paper/2411.10606"},"observation_digest":"sha256:a29453ffd8822c09480b0a7ab763d814cc01649bf26583a082beb6d454b2365a","observation_id":"5c7ae752-0f65-4a87-bb5a-7d46a4e11e92","resolution":{"observed_at":"2026-08-12T19:35:31.577408Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-19T06:32:44.657259+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-19T06:32:44.657259+00:00","source":"crossref"},{"observed_at":"2026-08-19T06:32:39.956319+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"1907.10641","last_updated":"2019-11-21T19:01:32Z","snapshot_observed_at":"2026-08-16T10:42:13.725165Z","submitted_at":"2019-07-24T18:11:59Z","title":"WinoGrande: An Adversarial Winograd Schema Challenge at Scale","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"1907.10641","snapshot_observed_at":"2026-08-12T19:35:30.194734Z","title":"Winogrande: An adversarial winograd schema challenge at scale","venue":null,"work_id":null,"year":1907},"citing_paper":{"arxiv_id":"2411.10606","last_updated":"2024-11-15T22:02:28Z","snapshot_observed_at":"2026-08-14T13:50:36.265901Z","submitted_at":"2024-11-15T22:02:28Z","title":"AmoebaLLM: Constructing Any-Shape Large Language Models for Efficient and Instant Deployment","version":1},"reference_index":46,"source":"pdf_text","source_observed_at":"2026-08-12T19:35:30.194734Z"},"links":{"cited_paper":"/paper/1907.10641","citing_paper":"/paper/2411.10606"},"observation_digest":"sha256:ae0821380283a5e01925607a1dddbddbbdb6e489b393f36950d77c178b279c57","observation_id":"4c388efe-3aac-439a-8f26-eed23ba10016","resolution":{"observed_at":"2026-08-12T19:35:30.194734Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"1803.05457","last_updated":"2018-03-14T18:04:21Z","snapshot_observed_at":"2026-08-14T19:36:07.505691Z","submitted_at":"2018-03-14T18:04:21Z","title":"Think you have Solved Question Answering? Try ARC, the AI2 Reasoning Challenge","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"1803.05457","snapshot_observed_at":"2026-08-12T19:35:30.199244Z","title":"Think you have solved question answering? try arc, the ai2 reasoning challenge","venue":null,"work_id":null,"year":2018},"citing_paper":{"arxiv_id":"2411.10606","last_updated":"2024-11-15T22:02:28Z","snapshot_observed_at":"2026-08-14T13:50:36.265901Z","submitted_at":"2024-11-15T22:02:28Z","title":"AmoebaLLM: Constructing Any-Shape Large Language Models for Efficient and Instant Deployment","version":1},"reference_index":47,"source":"pdf_text","source_observed_at":"2026-08-12T19:35:30.199244Z"},"links":{"cited_paper":"/paper/1803.05457","citing_paper":"/paper/2411.10606"},"observation_digest":"sha256:f46358156a51911f801cea8defa182334fe37bb29433c70eac57fef70a979833","observation_id":"b5373cc0-deb0-40d6-9b78-3119eacec35c","resolution":{"observed_at":"2026-08-12T19:35:30.199244Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-12T19:35:30.203624Z","title":"Can a suit of armor conduct electricity? a new dataset for open book question answering","venue":null,"work_id":null,"year":2018},"citing_paper":{"arxiv_id":"2411.10606","last_updated":"2024-11-15T22:02:28Z","snapshot_observed_at":"2026-08-14T13:50:36.265901Z","submitted_at":"2024-11-15T22:02:28Z","title":"AmoebaLLM: Constructing Any-Shape Large Language Models for Efficient and Instant Deployment","version":1},"reference_index":48,"source":"pdf_text","source_observed_at":"2026-08-12T19:35:30.203624Z"},"links":{"citing_paper":"/paper/2411.10606"},"observation_digest":"sha256:d4c2cce35f0505bc2e60d0ebcfa9df46c56226b94307aab6182ba2c22741430a","observation_id":"08a44a45-7036-4352-bbdf-04874b73c224","resolution":{"observed_at":"2026-08-12T19:35:30.203624Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2009.03300","last_updated":"2021-01-12T18:57:11Z","snapshot_observed_at":"2026-08-13T20:44:28.824685Z","submitted_at":"2020-09-07T17:59:25Z","title":"Measuring Massive Multitask Language Understanding","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2009.03300","snapshot_observed_at":"2026-08-12T19:35:30.207874Z","title":"Measuring massive multitask language understanding","venue":null,"work_id":null,"year":2009},"citing_paper":{"arxiv_id":"2411.10606","last_updated":"2024-11-15T22:02:28Z","snapshot_observed_at":"2026-08-14T13:50:36.265901Z","submitted_at":"2024-11-15T22:02:28Z","title":"AmoebaLLM: Constructing Any-Shape Large Language Models for Efficient and Instant Deployment","version":1},"reference_index":49,"source":"pdf_text","source_observed_at":"2026-08-12T19:35:30.207874Z"},"links":{"cited_paper":"/paper/2009.03300","citing_paper":"/paper/2411.10606"},"observation_digest":"sha256:c14eb2c2e8457dcb494b5e0433506601c04278bb962b2570784dd9e7902553b1","observation_id":"d4b33512-7f09-4ab1-a7d9-4e1cf8ab0af3","resolution":{"observed_at":"2026-08-12T19:35:30.207874Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-12T19:35:30.212614Z","title":"Pointer sentinel mixture models","venue":null,"work_id":null,"year":2017},"citing_paper":{"arxiv_id":"2411.10606","last_updated":"2024-11-15T22:02:28Z","snapshot_observed_at":"2026-08-14T13:50:36.265901Z","submitted_at":"2024-11-15T22:02:28Z","title":"AmoebaLLM: Constructing Any-Shape Large Language Models for Efficient and Instant Deployment","version":1},"reference_index":50,"source":"pdf_text","source_observed_at":"2026-08-12T19:35:30.212614Z"},"links":{"citing_paper":"/paper/2411.10606"},"observation_digest":"sha256:12b7f34ec4ac1780f1ca940c60baa18e31b778fcb0db90af8b7b5b810053c98f","observation_id":"be7616b8-9517-4e6d-b67a-9ed25d877c74","resolution":{"observed_at":"2026-08-12T19:35:30.212614Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-12T19:35:30.216816Z","title":"Aligning books and movies: Towards story-like visual explanations by watching movies and reading books","venue":null,"work_id":null,"year":2015},"citing_paper":{"arxiv_id":"2411.10606","last_updated":"2024-11-15T22:02:28Z","snapshot_observed_at":"2026-08-14T13:50:36.265901Z","submitted_at":"2024-11-15T22:02:28Z","title":"AmoebaLLM: Constructing Any-Shape Large Language Models for Efficient and Instant Deployment","version":1},"reference_index":51,"source":"pdf_text","source_observed_at":"2026-08-12T19:35:30.216816Z"},"links":{"citing_paper":"/paper/2411.10606"},"observation_digest":"sha256:1650f61cfdcb563292da82498e12c845060e1e9fc673030b8dd26a062edace49","observation_id":"83c901e8-269b-4f7e-8d6c-dd18bd80c43f","resolution":{"observed_at":"2026-08-12T19:35:30.216816Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-12T19:35:30.221053Z","title":"Attention is all you need","venue":null,"work_id":null,"year":2017},"citing_paper":{"arxiv_id":"2411.10606","last_updated":"2024-11-15T22:02:28Z","snapshot_observed_at":"2026-08-14T13:50:36.265901Z","submitted_at":"2024-11-15T22:02:28Z","title":"AmoebaLLM: Constructing Any-Shape Large Language Models for Efficient and Instant Deployment","version":1},"reference_index":52,"source":"pdf_text","source_observed_at":"2026-08-12T19:35:30.221053Z"},"links":{"citing_paper":"/paper/2411.10606"},"observation_digest":"sha256:6bf1100088890596d1ce4ac319b962212cd1fa25af6b1e054f24b5fdae9473d5","observation_id":"89bdc06e-e9b4-48e4-89e8-fe858dfb5870","resolution":{"observed_at":"2026-08-12T19:35:30.221053Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"1810.04805","last_updated":"2019-05-24T20:37:26Z","snapshot_observed_at":"2026-08-14T18:16:28.847993Z","submitted_at":"2018-10-11T00:50:01Z","title":"BERT: Pre-training of Deep Bidirectional Transformers for Language Understanding","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"1810.04805","snapshot_observed_at":"2026-08-12T19:35:30.225307Z","title":"Bert: Pre-training of deep bidirectional transformers for language understanding","venue":null,"work_id":null,"year":2018},"citing_paper":{"arxiv_id":"2411.10606","last_updated":"2024-11-15T22:02:28Z","snapshot_observed_at":"2026-08-14T13:50:36.265901Z","submitted_at":"2024-11-15T22:02:28Z","title":"AmoebaLLM: Constructing Any-Shape Large Language Models for Efficient and Instant Deployment","version":1},"reference_index":53,"source":"pdf_text","source_observed_at":"2026-08-12T19:35:30.225307Z"},"links":{"cited_paper":"/paper/1810.04805","citing_paper":"/paper/2411.10606"},"observation_digest":"sha256:166d3a88d7a83bee5a09d83c486c5fa876d76dbf4c25b7c42206bd83a02bc298","observation_id":"6d369cb1-5597-4b4d-84c0-08e19845b49f","resolution":{"observed_at":"2026-08-12T19:35:30.225307Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-12T19:35:30.229536Z","title":null,"venue":null,"work_id":null,"year":2020},"citing_paper":{"arxiv_id":"2411.10606","last_updated":"2024-11-15T22:02:28Z","snapshot_observed_at":"2026-08-14T13:50:36.265901Z","submitted_at":"2024-11-15T22:02:28Z","title":"AmoebaLLM: Constructing Any-Shape Large Language Models for Efficient and Instant Deployment","version":1},"reference_index":54,"source":"pdf_text","source_observed_at":"2026-08-12T19:35:30.229536Z"},"links":{"citing_paper":"/paper/2411.10606"},"observation_digest":"sha256:9e4d1d5a4952060efbb0a4ad176aa759a3f830c4dac0122ac0407fbec275eca4","observation_id":"77165d89-cad7-4935-a838-451e523d9a5a","resolution":{"observed_at":"2026-08-12T19:35:30.229536Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2203.17189","last_updated":"2022-03-31T17:12:13Z","snapshot_observed_at":"2026-08-16T17:10:11.301337Z","submitted_at":"2022-03-31T17:12:13Z","title":"Scaling Up Models and Data with $\\texttt{t5x}$ and $\\texttt{seqio}$","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2203.17189","snapshot_observed_at":"2026-08-12T19:35:30.233705Z","title":null,"venue":null,"work_id":null,"year":2022},"citing_paper":{"arxiv_id":"2411.10606","last_updated":"2024-11-15T22:02:28Z","snapshot_observed_at":"2026-08-14T13:50:36.265901Z","submitted_at":"2024-11-15T22:02:28Z","title":"AmoebaLLM: Constructing Any-Shape Large Language Models for Efficient and Instant Deployment","version":1},"reference_index":55,"source":"pdf_text","source_observed_at":"2026-08-12T19:35:30.233705Z"},"links":{"cited_paper":"/paper/2203.17189","citing_paper":"/paper/2411.10606"},"observation_digest":"sha256:0d0ed37b1df8e4b4b1b1626b72bf423669c7a4e2731e9a1c97f4bae259623e39","observation_id":"f88bf95c-bf63-4ef7-9be9-96c82bc7ec7c","resolution":{"observed_at":"2026-08-12T19:35:30.233705Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2307.14995","last_updated":"2024-01-19T07:47:01Z","snapshot_observed_at":"2026-08-16T15:12:45.594763Z","submitted_at":"2023-07-27T16:45:33Z","title":"TransNormerLLM: A Faster and Better Large Language Model with Improved TransNormer","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2307.14995","snapshot_observed_at":"2026-08-12T19:35:30.238787Z","title":"Scaling transnormer to 175 billion parameters","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2411.10606","last_updated":"2024-11-15T22:02:28Z","snapshot_observed_at":"2026-08-14T13:50:36.265901Z","submitted_at":"2024-11-15T22:02:28Z","title":"AmoebaLLM: Constructing Any-Shape Large Language Models for Efficient and Instant Deployment","version":1},"reference_index":56,"source":"pdf_text","source_observed_at":"2026-08-12T19:35:30.238787Z"},"links":{"cited_paper":"/paper/2307.14995","citing_paper":"/paper/2411.10606"},"observation_digest":"sha256:8f828c991637a70dd37d51df1d83825e748124a7c9564f7f126e6f80fa9f4911","observation_id":"fec69dda-b3ea-4cd3-aa62-1ba83d646261","resolution":{"observed_at":"2026-08-12T19:35:30.238787Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2001.08361","last_updated":"2020-01-23T03:59:20Z","snapshot_observed_at":"2026-08-13T17:41:53.092611Z","submitted_at":"2020-01-23T03:59:20Z","title":"Scaling Laws for Neural Language Models","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2001.08361","snapshot_observed_at":"2026-08-12T19:35:30.243328Z","title":"Scaling laws for neural language models","venue":null,"work_id":null,"year":2001},"citing_paper":{"arxiv_id":"2411.10606","last_updated":"2024-11-15T22:02:28Z","snapshot_observed_at":"2026-08-14T13:50:36.265901Z","submitted_at":"2024-11-15T22:02:28Z","title":"AmoebaLLM: Constructing Any-Shape Large Language Models for Efficient and Instant Deployment","version":1},"reference_index":57,"source":"pdf_text","source_observed_at":"2026-08-12T19:35:30.243328Z"},"links":{"cited_paper":"/paper/2001.08361","citing_paper":"/paper/2411.10606"},"observation_digest":"sha256:9e52ac76b81fc16197138b0fd41a1262e9af481864e21d124d4a80af6534e748","observation_id":"89a4e57a-69a7-4f85-8f7b-eb14d374a32f","resolution":{"observed_at":"2026-08-12T19:35:30.243328Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-12T19:35:30.248083Z","title":"Pythia: A suite for analyzing large language models across training and scaling","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2411.10606","last_updated":"2024-11-15T22:02:28Z","snapshot_observed_at":"2026-08-14T13:50:36.265901Z","submitted_at":"2024-11-15T22:02:28Z","title":"AmoebaLLM: Constructing Any-Shape Large Language Models for Efficient and Instant Deployment","version":1},"reference_index":58,"source":"pdf_text","source_observed_at":"2026-08-12T19:35:30.248083Z"},"links":{"citing_paper":"/paper/2411.10606"},"observation_digest":"sha256:11eafb06b3d3414ce2310dc758e3f9cca2094865980941e17d7c92e138ee7003","observation_id":"4108ab37-9d2f-4907-b9e5-871be0a33a11","resolution":{"observed_at":"2026-08-12T19:35:30.248083Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2103.10360","last_updated":"2022-03-17T11:49:55Z","snapshot_observed_at":"2026-08-16T18:37:57.151597Z","submitted_at":"2021-03-18T16:30:26Z","title":"GLM: General Language Model Pretraining with Autoregressive Blank Infilling","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2103.10360","snapshot_observed_at":"2026-08-12T19:35:30.252384Z","title":"Glm: General language model pretraining with autoregressive blank infilling","venue":null,"work_id":null,"year":2021},"citing_paper":{"arxiv_id":"2411.10606","last_updated":"2024-11-15T22:02:28Z","snapshot_observed_at":"2026-08-14T13:50:36.265901Z","submitted_at":"2024-11-15T22:02:28Z","title":"AmoebaLLM: Constructing Any-Shape Large Language Models for Efficient and Instant Deployment","version":1},"reference_index":59,"source":"pdf_text","source_observed_at":"2026-08-12T19:35:30.252384Z"},"links":{"cited_paper":"/paper/2103.10360","citing_paper":"/paper/2411.10606"},"observation_digest":"sha256:6f7993d9f4d9e41c7e97b0ac5304ae3b9a526c4d358345e7377be8e28f02a7a4","observation_id":"14985555-59ac-44e1-8547-bda8b4f4c931","resolution":{"observed_at":"2026-08-12T19:35:30.252384Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2205.01068","last_updated":"2022-06-21T17:04:40Z","snapshot_observed_at":"2026-08-06T03:13:37.403059Z","submitted_at":"2022-05-02T17:49:50Z","title":"OPT: Open Pre-trained Transformer Language Models","version":4},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2205.01068","snapshot_observed_at":"2026-08-12T19:35:30.257228Z","title":"Opt: Open pre-trained transformer language models","venue":null,"work_id":null,"year":2022},"citing_paper":{"arxiv_id":"2411.10606","last_updated":"2024-11-15T22:02:28Z","snapshot_observed_at":"2026-08-14T13:50:36.265901Z","submitted_at":"2024-11-15T22:02:28Z","title":"AmoebaLLM: Constructing Any-Shape Large Language Models for Efficient and Instant Deployment","version":1},"reference_index":60,"source":"pdf_text","source_observed_at":"2026-08-12T19:35:30.257228Z"},"links":{"cited_paper":"/paper/2205.01068","citing_paper":"/paper/2411.10606"},"observation_digest":"sha256:945a7b744889588f83b840056502c3bb2ac528ce5a5f123327072d2dd07cfcee","observation_id":"6248c978-7fbe-4e07-9a93-4dca453955fb","resolution":{"observed_at":"2026-08-12T19:35:30.257228Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2211.05100","last_updated":"2023-06-27T09:57:58Z","snapshot_observed_at":"2026-08-04T18:56:03.233715Z","submitted_at":"2022-11-09T18:48:09Z","title":"BLOOM: A 176B-Parameter Open-Access Multilingual Language Model","version":4},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2211.05100","snapshot_observed_at":"2026-08-12T19:35:30.261920Z","title":"Bloom: A 176b-parameter open-access multilingual language model","venue":null,"work_id":null,"year":2022},"citing_paper":{"arxiv_id":"2411.10606","last_updated":"2024-11-15T22:02:28Z","snapshot_observed_at":"2026-08-14T13:50:36.265901Z","submitted_at":"2024-11-15T22:02:28Z","title":"AmoebaLLM: Constructing Any-Shape Large Language Models for Efficient and Instant Deployment","version":1},"reference_index":61,"source":"pdf_text","source_observed_at":"2026-08-12T19:35:30.261920Z"},"links":{"cited_paper":"/paper/2211.05100","citing_paper":"/paper/2411.10606"},"observation_digest":"sha256:2225789aeac5099c817570bee603047b90d0390413881254deaa038f912099cf","observation_id":"0bf18923-75a3-43ca-81b0-1690dc5eaaf0","resolution":{"observed_at":"2026-08-12T19:35:30.261920Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-12T19:35:30.266571Z","title":"Specializing smaller language models towards multi-step reasoning","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2411.10606","last_updated":"2024-11-15T22:02:28Z","snapshot_observed_at":"2026-08-14T13:50:36.265901Z","submitted_at":"2024-11-15T22:02:28Z","title":"AmoebaLLM: Constructing Any-Shape Large Language Models for Efficient and Instant Deployment","version":1},"reference_index":62,"source":"pdf_text","source_observed_at":"2026-08-12T19:35:30.266571Z"},"links":{"citing_paper":"/paper/2411.10606"},"observation_digest":"sha256:1153daa7f50d56052771104535e7551c39692454379f92f1843261528fb3ac95","observation_id":"b33d280d-b0b7-4812-bf90-27c9eba21c60","resolution":{"observed_at":"2026-08-12T19:35:30.266571Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2305.02301","last_updated":"2023-07-05T16:59:31Z","snapshot_observed_at":"2026-08-17T13:50:40.586963Z","submitted_at":"2023-05-03T17:50:56Z","title":"Distilling Step-by-Step! Outperforming Larger Language Models with Less Training Data and Smaller Model Sizes","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2305.02301","snapshot_observed_at":"2026-08-12T19:35:30.271219Z","title":"Distilling step-by-step! outperform- ing larger language models with less training data and smaller model sizes","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2411.10606","last_updated":"2024-11-15T22:02:28Z","snapshot_observed_at":"2026-08-14T13:50:36.265901Z","submitted_at":"2024-11-15T22:02:28Z","title":"AmoebaLLM: Constructing Any-Shape Large Language Models for Efficient and Instant Deployment","version":1},"reference_index":63,"source":"pdf_text","source_observed_at":"2026-08-12T19:35:30.271219Z"},"links":{"cited_paper":"/paper/2305.02301","citing_paper":"/paper/2411.10606"},"observation_digest":"sha256:9f1e966226e75154930412f81a24027f0ff6f5d30017f725c91db3f59fc3cc33","observation_id":"2b37154b-71db-43cc-a101-e69071724284","resolution":{"observed_at":"2026-08-12T19:35:30.271219Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-12T19:35:30.276293Z","title":"Optq: Accurate quantization for generative pre-trained transformers","venue":null,"work_id":null,"year":2022},"citing_paper":{"arxiv_id":"2411.10606","last_updated":"2024-11-15T22:02:28Z","snapshot_observed_at":"2026-08-14T13:50:36.265901Z","submitted_at":"2024-11-15T22:02:28Z","title":"AmoebaLLM: Constructing Any-Shape Large Language Models for Efficient and Instant Deployment","version":1},"reference_index":64,"source":"pdf_text","source_observed_at":"2026-08-12T19:35:30.276293Z"},"links":{"citing_paper":"/paper/2411.10606"},"observation_digest":"sha256:7c4267d118ba2c9ad1213f2855bb0170220e50ce3e017a20bb9315d3f44289fd","observation_id":"00f0aae3-c47b-4802-98ec-a2ff828c0117","resolution":{"observed_at":"2026-08-12T19:35:30.276293Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-12T19:35:31.465309Z","title":"LLM.int8(): 8-bit matrix multiplication for transformers at scale","venue":null,"work_id":"8e9b1439-51cc-4726-8aaf-cc86fbfcff7d","year":2022},"citing_paper":{"arxiv_id":"2411.10606","last_updated":"2024-11-15T22:02:28Z","snapshot_observed_at":"2026-08-14T13:50:36.265901Z","submitted_at":"2024-11-15T22:02:28Z","title":"AmoebaLLM: Constructing Any-Shape Large Language Models for Efficient and Instant Deployment","version":1},"reference_index":65,"source":"pdf_text","source_observed_at":"2026-08-12T19:35:30.280919Z"},"links":{"citing_paper":"/paper/2411.10606"},"observation_digest":"sha256:ee101fa60e0d505149870c698a8afd3cc5a50bdab648b2091a1f3a323b3a7fce","observation_id":"02385843-c7da-4565-a2ed-347b0e363b5a","resolution":{"observed_at":"2026-08-12T19:35:31.470009Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-19T06:32:44.657259+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-19T06:32:44.657259+00:00","source":"crossref"},{"observed_at":"2026-08-19T06:32:39.956319+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-12T19:35:31.449948Z","title":"Smoothquant: Accurate and efficient post-training quantization for large language models","venue":null,"work_id":"13c38fb7-ffee-40ff-b7b7-db9aef81823e","year":2023},"citing_paper":{"arxiv_id":"2411.10606","last_updated":"2024-11-15T22:02:28Z","snapshot_observed_at":"2026-08-14T13:50:36.265901Z","submitted_at":"2024-11-15T22:02:28Z","title":"AmoebaLLM: Constructing Any-Shape Large Language Models for Efficient and Instant Deployment","version":1},"reference_index":66,"source":"pdf_text","source_observed_at":"2026-08-12T19:35:30.285078Z"},"links":{"citing_paper":"/paper/2411.10606"},"observation_digest":"sha256:3073e64eb8258b45b4e08c0568c71c08a791542a94de374fd18c56d42a7693a3","observation_id":"7767e87b-1cd3-4e64-bcd6-8c7d2ebda0bd","resolution":{"observed_at":"2026-08-12T19:35:31.454494Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-19T06:32:44.657259+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-19T06:32:44.657259+00:00","source":"crossref"},{"observed_at":"2026-08-19T06:32:39.956319+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-12T19:35:31.434638Z","title":"GPTQ: Accurate post-training compression for generative pretrained transformers","venue":null,"work_id":"e5e19ae6-06a4-4b3a-a4a1-949874639ae0","year":2023},"citing_paper":{"arxiv_id":"2411.10606","last_updated":"2024-11-15T22:02:28Z","snapshot_observed_at":"2026-08-14T13:50:36.265901Z","submitted_at":"2024-11-15T22:02:28Z","title":"AmoebaLLM: Constructing Any-Shape Large Language Models for Efficient and Instant Deployment","version":1},"reference_index":67,"source":"pdf_text","source_observed_at":"2026-08-12T19:35:30.289140Z"},"links":{"citing_paper":"/paper/2411.10606"},"observation_digest":"sha256:94801ea78cffd20c6b64c9dede09d65d73156817d6c7b53fbceaa0231cdec877","observation_id":"9b7f6f70-38e4-4c71-87fb-f8902d9004d2","resolution":{"observed_at":"2026-08-12T19:35:31.439474Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-19T06:32:44.657259+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-19T06:32:44.657259+00:00","source":"crossref"},{"observed_at":"2026-08-19T06:32:39.956319+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-12T19:35:30.292984Z","title":"Spqr: A sparse-quantized representation for near-lossless llm weight compression, 2023","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2411.10606","last_updated":"2024-11-15T22:02:28Z","snapshot_observed_at":"2026-08-14T13:50:36.265901Z","submitted_at":"2024-11-15T22:02:28Z","title":"AmoebaLLM: Constructing Any-Shape Large Language Models for Efficient and Instant Deployment","version":1},"reference_index":68,"source":"pdf_text","source_observed_at":"2026-08-12T19:35:30.292984Z"},"links":{"citing_paper":"/paper/2411.10606"},"observation_digest":"sha256:a3a9562850fb3eed8d2d2eb863f5675b6b5454365429c6bca14bf7d28c4e95e4","observation_id":"497a3f44-4ab2-4102-8b9b-231eb065c366","resolution":{"observed_at":"2026-08-12T19:35:30.292984Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2306.00978","last_updated":"2026-04-25T06:58:16Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2023-06-01T17:59:10Z","title":"AWQ: Activation-aware Weight Quantization for LLM Compression and Acceleration","version":6},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2306.00978","snapshot_observed_at":"2026-08-12T19:35:30.296985Z","title":"Awq: Activation-aware weight quantization for llm compression and acceleration","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2411.10606","last_updated":"2024-11-15T22:02:28Z","snapshot_observed_at":"2026-08-14T13:50:36.265901Z","submitted_at":"2024-11-15T22:02:28Z","title":"AmoebaLLM: Constructing Any-Shape Large Language Models for Efficient and Instant Deployment","version":1},"reference_index":69,"source":"pdf_text","source_observed_at":"2026-08-12T19:35:30.296985Z"},"links":{"cited_paper":"/paper/2306.00978","citing_paper":"/paper/2411.10606"},"observation_digest":"sha256:e61d621aef3dd9828cf890151f2be3e4a1faf4cacff82cec345764f0986392e7","observation_id":"c448eb37-af1f-47f7-874c-095e0719d7e3","resolution":{"observed_at":"2026-08-12T19:35:30.296985Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2307.08691","last_updated":"2023-07-17T17:50:36Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2023-07-17T17:50:36Z","title":"FlashAttention-2: Faster Attention with Better Parallelism and Work Partitioning","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2307.08691","snapshot_observed_at":"2026-08-12T19:35:30.301267Z","title":"Flashattention-2: Faster attention with better parallelism and work partitioning","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2411.10606","last_updated":"2024-11-15T22:02:28Z","snapshot_observed_at":"2026-08-14T13:50:36.265901Z","submitted_at":"2024-11-15T22:02:28Z","title":"AmoebaLLM: Constructing Any-Shape Large Language Models for Efficient and Instant Deployment","version":1},"reference_index":70,"source":"pdf_text","source_observed_at":"2026-08-12T19:35:30.301267Z"},"links":{"cited_paper":"/paper/2307.08691","citing_paper":"/paper/2411.10606"},"observation_digest":"sha256:2f035f0dcb0cd0ef8fee3e901a254f6feb474186fee370958396f1ddd59f099b","observation_id":"f8db1388-8254-44c6-9180-6c42739eeef2","resolution":{"observed_at":"2026-08-12T19:35:30.301267Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-12T19:35:30.305471Z","title":"Efficient memory management for large language model serving with pagedattention","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2411.10606","last_updated":"2024-11-15T22:02:28Z","snapshot_observed_at":"2026-08-14T13:50:36.265901Z","submitted_at":"2024-11-15T22:02:28Z","title":"AmoebaLLM: Constructing Any-Shape Large Language Models for Efficient and Instant Deployment","version":1},"reference_index":71,"source":"pdf_text","source_observed_at":"2026-08-12T19:35:30.305471Z"},"links":{"citing_paper":"/paper/2411.10606"},"observation_digest":"sha256:bc16e74c2d97cfb3096341378fc437c99e4b80ce5087a4d2fa0fc191600dec32","observation_id":"dbfbca65-b0e2-4c9b-849e-7eef92a2f4d1","resolution":{"observed_at":"2026-08-12T19:35:30.305471Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-12T19:35:31.396084Z","title":"Llm-pruner: On the structural pruning of large language models, 2023","venue":null,"work_id":"0155a16a-9e26-4391-8ba8-48d6d266382b","year":2023},"citing_paper":{"arxiv_id":"2411.10606","last_updated":"2024-11-15T22:02:28Z","snapshot_observed_at":"2026-08-14T13:50:36.265901Z","submitted_at":"2024-11-15T22:02:28Z","title":"AmoebaLLM: Constructing Any-Shape Large Language Models for Efficient and Instant Deployment","version":1},"reference_index":72,"source":"pdf_text","source_observed_at":"2026-08-12T19:35:30.309524Z"},"links":{"citing_paper":"/paper/2411.10606"},"observation_digest":"sha256:a01a8596abf605944122053816da0bb7b2e484b74e431ebee6a3d95e1750b603","observation_id":"57c18ccb-a4a9-463f-aead-949073402a78","resolution":{"observed_at":"2026-08-12T19:35:31.400730Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-19T06:32:44.657259+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-19T06:32:44.657259+00:00","source":"crossref"},{"observed_at":"2026-08-19T06:32:39.956319+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2102.04010","last_updated":"2021-04-18T10:18:00Z","snapshot_observed_at":"2026-08-16T18:47:07.825085Z","submitted_at":"2021-02-08T05:55:47Z","title":"Learning N:M Fine-grained Structured Sparse Neural Networks From Scratch","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2102.04010","snapshot_observed_at":"2026-08-12T19:35:30.313605Z","title":"Learning n: m fine-grained structured sparse neural networks from scratch","venue":null,"work_id":null,"year":2021},"citing_paper":{"arxiv_id":"2411.10606","last_updated":"2024-11-15T22:02:28Z","snapshot_observed_at":"2026-08-14T13:50:36.265901Z","submitted_at":"2024-11-15T22:02:28Z","title":"AmoebaLLM: Constructing Any-Shape Large Language Models for Efficient and Instant Deployment","version":1},"reference_index":73,"source":"pdf_text","source_observed_at":"2026-08-12T19:35:30.313605Z"},"links":{"cited_paper":"/paper/2102.04010","citing_paper":"/paper/2411.10606"},"observation_digest":"sha256:5a91783a8300eb1ee6c5661929b44cfe4fb54d63bdbab21511887fafeee72701","observation_id":"e6c46938-4012-4947-b6b3-73e52a18c3dd","resolution":{"observed_at":"2026-08-12T19:35:30.313605Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"1812.08928","last_updated":"2018-12-21T03:36:48Z","snapshot_observed_at":"2026-08-14T17:40:05.995607Z","submitted_at":"2018-12-21T03:36:48Z","title":"Slimmable Neural Networks","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"1812.08928","snapshot_observed_at":"2026-08-12T19:35:30.318165Z","title":"Slimmable neural networks","venue":null,"work_id":null,"year":2018},"citing_paper":{"arxiv_id":"2411.10606","last_updated":"2024-11-15T22:02:28Z","snapshot_observed_at":"2026-08-14T13:50:36.265901Z","submitted_at":"2024-11-15T22:02:28Z","title":"AmoebaLLM: Constructing Any-Shape Large Language Models for Efficient and Instant Deployment","version":1},"reference_index":74,"source":"pdf_text","source_observed_at":"2026-08-12T19:35:30.318165Z"},"links":{"cited_paper":"/paper/1812.08928","citing_paper":"/paper/2411.10606"},"observation_digest":"sha256:e4ed6b426d10ce7bd3faf4c6bd0e06b02cfaf13ffc5068630c02c9af9f3097b1","observation_id":"56f3729f-7d3a-4baa-b28a-a2f3d84f0c51","resolution":{"observed_at":"2026-08-12T19:35:30.318165Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-12T19:35:31.380980Z","title":"Universally slimmable networks and improved training techniques","venue":null,"work_id":"7154a955-5902-4de1-94b9-d4ce13528c1e","year":2019},"citing_paper":{"arxiv_id":"2411.10606","last_updated":"2024-11-15T22:02:28Z","snapshot_observed_at":"2026-08-14T13:50:36.265901Z","submitted_at":"2024-11-15T22:02:28Z","title":"AmoebaLLM: Constructing Any-Shape Large Language Models for Efficient and Instant Deployment","version":1},"reference_index":75,"source":"pdf_text","source_observed_at":"2026-08-12T19:35:30.322462Z"},"links":{"citing_paper":"/paper/2411.10606"},"observation_digest":"sha256:f43bec537a4e9ce053f4a285d4825adac27c86907c3e16d2ddc01433503b2138","observation_id":"774472ea-f6cd-4826-b023-fc7a80e3487f","resolution":{"observed_at":"2026-08-12T19:35:31.385443Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-19T06:32:44.657259+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-19T06:32:44.657259+00:00","source":"crossref"},{"observed_at":"2026-08-19T06:32:39.956319+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"1903.11728","last_updated":"2019-06-01T03:19:54Z","snapshot_observed_at":"2026-08-16T08:11:30.907868Z","submitted_at":"2019-03-27T23:17:28Z","title":"AutoSlim: Towards One-Shot Architecture Search for Channel Numbers","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"1903.11728","snapshot_observed_at":"2026-08-12T19:35:30.326810Z","title":"Autoslim: Towards one-shot architecture search for channel numbers","venue":null,"work_id":null,"year":1903},"citing_paper":{"arxiv_id":"2411.10606","last_updated":"2024-11-15T22:02:28Z","snapshot_observed_at":"2026-08-14T13:50:36.265901Z","submitted_at":"2024-11-15T22:02:28Z","title":"AmoebaLLM: Constructing Any-Shape Large Language Models for Efficient and Instant Deployment","version":1},"reference_index":76,"source":"pdf_text","source_observed_at":"2026-08-12T19:35:30.326810Z"},"links":{"cited_paper":"/paper/1903.11728","citing_paper":"/paper/2411.10606"},"observation_digest":"sha256:5f014781185628a0887eda6699edd046c4997b89aaa38e38d1d54c21ff0bc606","observation_id":"306ff83c-6aae-470e-bf53-c0e8dc670e27","resolution":{"observed_at":"2026-08-12T19:35:30.326810Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-12T19:35:31.365161Z","title":"Adabits: Neural network quantization with adaptive bit-widths","venue":null,"work_id":"80891bf4-b4fd-47ff-850e-ef6412d338ff","year":2020},"citing_paper":{"arxiv_id":"2411.10606","last_updated":"2024-11-15T22:02:28Z","snapshot_observed_at":"2026-08-14T13:50:36.265901Z","submitted_at":"2024-11-15T22:02:28Z","title":"AmoebaLLM: Constructing Any-Shape Large Language Models for Efficient and Instant Deployment","version":1},"reference_index":77,"source":"pdf_text","source_observed_at":"2026-08-12T19:35:30.331319Z"},"links":{"citing_paper":"/paper/2411.10606"},"observation_digest":"sha256:1ad098d431088bb341ed082c1da82c610cf5fd2c00f9a5e35af0002d51342d4a","observation_id":"ec61eaac-c47a-4463-99d2-fae350495337","resolution":{"observed_at":"2026-08-12T19:35:31.370401Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-19T06:32:44.657259+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-19T06:32:44.657259+00:00","source":"crossref"},{"observed_at":"2026-08-19T06:32:39.956319+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2002.02815","last_updated":"2020-02-07T14:43:44Z","snapshot_observed_at":"2026-08-16T01:15:34.952300Z","submitted_at":"2020-02-07T14:43:44Z","title":"Switchable Precision Neural Networks","version":1},"cited_work":{"arxiv_id":"2002.02815","doi":null,"metadata_source":"pith","pith_arxiv_id":"2002.02815","snapshot_observed_at":"2026-08-12T19:35:30.490956Z","title":"Switchable Precision Neural Networks","venue":"cs.CV","work_id":"a9e19ccd-d092-4dab-8103-4ac3da5c5a97","year":2020},"citing_paper":{"arxiv_id":"2411.10606","last_updated":"2024-11-15T22:02:28Z","snapshot_observed_at":"2026-08-14T13:50:36.265901Z","submitted_at":"2024-11-15T22:02:28Z","title":"AmoebaLLM: Constructing Any-Shape Large Language Models for Efficient and Instant Deployment","version":1},"reference_index":78,"source":"pdf_text","source_observed_at":"2026-08-12T19:35:30.335398Z"},"links":{"cited_paper":"/paper/2002.02815","citing_paper":"/paper/2411.10606"},"observation_digest":"sha256:641f1aa9c8f326016a4bb6b9be437e8b78af8db0064b53c9773f20c46b9faa83","observation_id":"1de285e0-fa7b-4d95-ad15-87b8462be3c1","resolution":{"observed_at":"2026-08-12T19:35:30.497574Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-19T06:32:44.657259+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-19T06:32:44.657259+00:00","source":"crossref"},{"observed_at":"2026-08-19T06:32:39.956319+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-12T19:35:31.348177Z","title":"Any-precision deep neural networks","venue":null,"work_id":"8916607e-0631-4675-807a-8f6aa14161b8","year":2021},"citing_paper":{"arxiv_id":"2411.10606","last_updated":"2024-11-15T22:02:28Z","snapshot_observed_at":"2026-08-14T13:50:36.265901Z","submitted_at":"2024-11-15T22:02:28Z","title":"AmoebaLLM: Constructing Any-Shape Large Language Models for Efficient and Instant Deployment","version":1},"reference_index":79,"source":"pdf_text","source_observed_at":"2026-08-12T19:35:30.340215Z"},"links":{"citing_paper":"/paper/2411.10606"},"observation_digest":"sha256:28d2e00d63d7d494d4cc9a5e8dca7e0e27f0ab5bc99fbacf8a3092546427ff9e","observation_id":"fc99e223-59ff-4770-8c17-212bfc2e0caa","resolution":{"observed_at":"2026-08-12T19:35:31.353646Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-19T06:32:44.657259+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-19T06:32:44.657259+00:00","source":"crossref"},{"observed_at":"2026-08-19T06:32:39.956319+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2402.10517","last_updated":"2024-06-21T05:20:56Z","snapshot_observed_at":"2026-08-16T14:18:04.510991Z","submitted_at":"2024-02-16T09:06:06Z","title":"Any-Precision LLM: Low-Cost Deployment of Multiple, Different-Sized LLMs","version":4},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2402.10517","snapshot_observed_at":"2026-08-12T19:35:30.344764Z","title":"Any-precision llm: Low-cost deployment of multiple, different-sized llms","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2411.10606","last_updated":"2024-11-15T22:02:28Z","snapshot_observed_at":"2026-08-14T13:50:36.265901Z","submitted_at":"2024-11-15T22:02:28Z","title":"AmoebaLLM: Constructing Any-Shape Large Language Models for Efficient and Instant Deployment","version":1},"reference_index":80,"source":"pdf_text","source_observed_at":"2026-08-12T19:35:30.344764Z"},"links":{"cited_paper":"/paper/2402.10517","citing_paper":"/paper/2411.10606"},"observation_digest":"sha256:017c8972e5c762a5fc59884fe8942a685d5e470b39141f00abc5972d33111412","observation_id":"07fa6e38-2884-47e1-9a07-4cc8ab5e4d49","resolution":{"observed_at":"2026-08-12T19:35:30.344764Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2406.10260","last_updated":"2024-08-28T17:26:03Z","snapshot_observed_at":"2026-08-16T13:44:21.253075Z","submitted_at":"2024-06-11T01:16:10Z","title":"Flextron: Many-in-One Flexible Large Language Model","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2406.10260","snapshot_observed_at":"2026-08-12T19:35:30.349664Z","title":"Flextron: Many-in-one flexible large language model","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2411.10606","last_updated":"2024-11-15T22:02:28Z","snapshot_observed_at":"2026-08-14T13:50:36.265901Z","submitted_at":"2024-11-15T22:02:28Z","title":"AmoebaLLM: Constructing Any-Shape Large Language Models for Efficient and Instant Deployment","version":1},"reference_index":81,"source":"pdf_text","source_observed_at":"2026-08-12T19:35:30.349664Z"},"links":{"cited_paper":"/paper/2406.10260","citing_paper":"/paper/2411.10606"},"observation_digest":"sha256:c44dc37488abe622a33ac8009d6007ca6953bc0c5b7e9516e95779b5f01059d0","observation_id":"10100ce0-99e4-4c7e-8823-e7e87bf0f52e","resolution":{"observed_at":"2026-08-12T19:35:30.349664Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-12T19:35:31.331296Z","title":"Guidelines: • The answer NA means that the abstract and introduction do not include the claims made in the paper","venue":null,"work_id":"1b8ce035-5a56-40b3-a89a-08c0a2c7fd28","year":null},"citing_paper":{"arxiv_id":"2411.10606","last_updated":"2024-11-15T22:02:28Z","snapshot_observed_at":"2026-08-14T13:50:36.265901Z","submitted_at":"2024-11-15T22:02:28Z","title":"AmoebaLLM: Constructing Any-Shape Large Language Models for Efficient and Instant Deployment","version":1},"reference_index":82,"source":"pdf_text","source_observed_at":"2026-08-12T19:35:30.354735Z"},"links":{"citing_paper":"/paper/2411.10606"},"observation_digest":"sha256:5139c42e7c32083c1e50f7b42b3effb86c5e4b01d3be73baae14084dce7c04eb","observation_id":"e446d2b3-ef50-410f-bd39-cf73787386e1","resolution":{"observed_at":"2026-08-12T19:35:31.336842Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-19T06:32:44.657259+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-19T06:32:44.657259+00:00","source":"crossref"},{"observed_at":"2026-08-19T06:32:39.956319+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-12T19:35:31.314676Z","title":"Limitations","venue":null,"work_id":"24e2b8c2-03fd-4a06-9cf8-416a8d11c26e","year":null},"citing_paper":{"arxiv_id":"2411.10606","last_updated":"2024-11-15T22:02:28Z","snapshot_observed_at":"2026-08-14T13:50:36.265901Z","submitted_at":"2024-11-15T22:02:28Z","title":"AmoebaLLM: Constructing Any-Shape Large Language Models for Efficient and Instant Deployment","version":1},"reference_index":83,"source":"pdf_text","source_observed_at":"2026-08-12T19:35:30.359545Z"},"links":{"citing_paper":"/paper/2411.10606"},"observation_digest":"sha256:e1892fd4abca9e58e610b917d0989713616c829aa2baba4cbd5ae01a3156165b","observation_id":"3c338bb7-ca03-4d79-b2b2-4685d57f63c6","resolution":{"observed_at":"2026-08-12T19:35:31.319618Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-19T06:32:44.657259+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-19T06:32:44.657259+00:00","source":"crossref"},{"observed_at":"2026-08-19T06:32:39.956319+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-12T19:35:31.299126Z","title":"Guidelines: • The answer NA means that the paper does not include theoretical results","venue":null,"work_id":"d3d7ec8f-8071-46ec-bb6d-52f9225ffdc5","year":null},"citing_paper":{"arxiv_id":"2411.10606","last_updated":"2024-11-15T22:02:28Z","snapshot_observed_at":"2026-08-14T13:50:36.265901Z","submitted_at":"2024-11-15T22:02:28Z","title":"AmoebaLLM: Constructing Any-Shape Large Language Models for Efficient and Instant Deployment","version":1},"reference_index":84,"source":"pdf_text","source_observed_at":"2026-08-12T19:35:30.364256Z"},"links":{"citing_paper":"/paper/2411.10606"},"observation_digest":"sha256:0ea53915fc2fbe82cfeb072e0fa3c0ba93f0d71a2ab0c68904424288039ca28a","observation_id":"d12b038f-019b-408d-bb0f-53af9e068fe5","resolution":{"observed_at":"2026-08-12T19:35:31.303621Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-19T06:32:44.657259+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-19T06:32:44.657259+00:00","source":"crossref"},{"observed_at":"2026-08-19T06:32:39.956319+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-12T19:35:31.283286Z","title":"Guidelines: • The answer NA means that the paper does not include experiments","venue":null,"work_id":"51c6c1ca-9988-4649-ab99-d621dcd184c7","year":null},"citing_paper":{"arxiv_id":"2411.10606","last_updated":"2024-11-15T22:02:28Z","snapshot_observed_at":"2026-08-14T13:50:36.265901Z","submitted_at":"2024-11-15T22:02:28Z","title":"AmoebaLLM: Constructing Any-Shape Large Language Models for Efficient and Instant Deployment","version":1},"reference_index":85,"source":"pdf_text","source_observed_at":"2026-08-12T19:35:30.368477Z"},"links":{"citing_paper":"/paper/2411.10606"},"observation_digest":"sha256:2bed3daf578152ddd06fb04c91e6ecd5f146bfb33063ea17e5e8462cbf68e9a0","observation_id":"67efbfa6-dfea-420d-bd68-9bca4378fce2","resolution":{"observed_at":"2026-08-12T19:35:31.288060Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-19T06:32:44.657259+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-19T06:32:44.657259+00:00","source":"crossref"},{"observed_at":"2026-08-19T06:32:39.956319+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-12T19:35:31.267931Z","title":"Guidelines: • The answer NA means that paper does not include experiments requiring code","venue":null,"work_id":"2f19fcc7-098f-4ee4-bd95-cc2b8e677028","year":null},"citing_paper":{"arxiv_id":"2411.10606","last_updated":"2024-11-15T22:02:28Z","snapshot_observed_at":"2026-08-14T13:50:36.265901Z","submitted_at":"2024-11-15T22:02:28Z","title":"AmoebaLLM: Constructing Any-Shape Large Language Models for Efficient and Instant Deployment","version":1},"reference_index":86,"source":"pdf_text","source_observed_at":"2026-08-12T19:35:30.373662Z"},"links":{"citing_paper":"/paper/2411.10606"},"observation_digest":"sha256:d35f2557b79d50515353708e2464064595eeae61060b2fb729230cf754dadcb8","observation_id":"bb66ef43-b795-4d21-b61c-b70ee5fb40e7","resolution":{"observed_at":"2026-08-12T19:35:31.272817Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-19T06:32:44.657259+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-19T06:32:44.657259+00:00","source":"crossref"},{"observed_at":"2026-08-19T06:32:39.956319+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-12T19:35:31.252526Z","title":"5.1 of our paper and also provided sufficient references","venue":null,"work_id":"34227827-98be-4c05-aaed-831e1c4fd9ab","year":null},"citing_paper":{"arxiv_id":"2411.10606","last_updated":"2024-11-15T22:02:28Z","snapshot_observed_at":"2026-08-14T13:50:36.265901Z","submitted_at":"2024-11-15T22:02:28Z","title":"AmoebaLLM: Constructing Any-Shape Large Language Models for Efficient and Instant Deployment","version":1},"reference_index":87,"source":"pdf_text","source_observed_at":"2026-08-12T19:35:30.378057Z"},"links":{"citing_paper":"/paper/2411.10606"},"observation_digest":"sha256:39245f362aaf856b8cb24c61b1f0083667d7a9c9083fc9382fef63145a7a91b0","observation_id":"80b9e18c-2625-447a-a961-25721f6d8a66","resolution":{"observed_at":"2026-08-12T19:35:31.257119Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-19T06:32:44.657259+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-19T06:32:44.657259+00:00","source":"crossref"},{"observed_at":"2026-08-19T06:32:39.956319+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-12T19:35:31.237642Z","title":"Guidelines: • The answer NA means that the paper does not include experiments","venue":null,"work_id":"179ed2b8-36ed-42f4-bd8e-8d35efc94a09","year":null},"citing_paper":{"arxiv_id":"2411.10606","last_updated":"2024-11-15T22:02:28Z","snapshot_observed_at":"2026-08-14T13:50:36.265901Z","submitted_at":"2024-11-15T22:02:28Z","title":"AmoebaLLM: Constructing Any-Shape Large Language Models for Efficient and Instant Deployment","version":1},"reference_index":88,"source":"pdf_text","source_observed_at":"2026-08-12T19:35:30.382864Z"},"links":{"citing_paper":"/paper/2411.10606"},"observation_digest":"sha256:a8fb4b5e3c582a089c9584ea28071735ee617e7d000eb39593f479f6fe6ddc22","observation_id":"ee4d9df4-a653-48fe-87d2-a4f5ff39be85","resolution":{"observed_at":"2026-08-12T19:35:31.242240Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-19T06:32:44.657259+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-19T06:32:44.657259+00:00","source":"crossref"},{"observed_at":"2026-08-19T06:32:39.956319+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-12T19:35:31.221845Z","title":"5.1 of our paper","venue":null,"work_id":"a74d73cc-976a-47ce-b436-d54ba779a5a2","year":null},"citing_paper":{"arxiv_id":"2411.10606","last_updated":"2024-11-15T22:02:28Z","snapshot_observed_at":"2026-08-14T13:50:36.265901Z","submitted_at":"2024-11-15T22:02:28Z","title":"AmoebaLLM: Constructing Any-Shape Large Language Models for Efficient and Instant Deployment","version":1},"reference_index":89,"source":"pdf_text","source_observed_at":"2026-08-12T19:35:30.387221Z"},"links":{"citing_paper":"/paper/2411.10606"},"observation_digest":"sha256:1c133563daf52ced9c82113c16dc330a55f1596a71b56ee21045f2ad6174cfda","observation_id":"7f4c3ed7-2993-43e5-8c27-2281d4b193b6","resolution":{"observed_at":"2026-08-12T19:35:31.226830Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-19T06:32:44.657259+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-19T06:32:44.657259+00:00","source":"crossref"},{"observed_at":"2026-08-19T06:32:39.956319+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-12T19:35:31.205026Z","title":"Guidelines: • The answer NA means that the authors have not reviewed the NeurIPS Code of Ethics","venue":null,"work_id":"cfebfdee-ea51-4aeb-a87c-d73b6ac29c12","year":null},"citing_paper":{"arxiv_id":"2411.10606","last_updated":"2024-11-15T22:02:28Z","snapshot_observed_at":"2026-08-14T13:50:36.265901Z","submitted_at":"2024-11-15T22:02:28Z","title":"AmoebaLLM: Constructing Any-Shape Large Language Models for Efficient and Instant Deployment","version":1},"reference_index":90,"source":"pdf_text","source_observed_at":"2026-08-12T19:35:30.391665Z"},"links":{"citing_paper":"/paper/2411.10606"},"observation_digest":"sha256:55efa2470c13b66c5b532deff3cfe8c56eb883eef7f68a3939e420ca9d8a0b38","observation_id":"d6d1ab37-1ce8-46b3-9fb7-811e0d605ea4","resolution":{"observed_at":"2026-08-12T19:35:31.210370Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-19T06:32:44.657259+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-19T06:32:44.657259+00:00","source":"crossref"},{"observed_at":"2026-08-19T06:32:39.956319+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-12T19:35:31.187513Z","title":"Guidelines: • The answer NA means that there is no societal impact of the work performed","venue":null,"work_id":"7efb0cbe-336e-492d-9f9d-160430c4a976","year":null},"citing_paper":{"arxiv_id":"2411.10606","last_updated":"2024-11-15T22:02:28Z","snapshot_observed_at":"2026-08-14T13:50:36.265901Z","submitted_at":"2024-11-15T22:02:28Z","title":"AmoebaLLM: Constructing Any-Shape Large Language Models for Efficient and Instant Deployment","version":1},"reference_index":91,"source":"pdf_text","source_observed_at":"2026-08-12T19:35:30.395935Z"},"links":{"citing_paper":"/paper/2411.10606"},"observation_digest":"sha256:f27b0b3058abda1ea4e364102ebbbeaf328f6c4c7289c86fa949a1fce68dc5d8","observation_id":"db2c3b4d-a307-4365-a513-5752fa96d014","resolution":{"observed_at":"2026-08-12T19:35:31.192850Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-19T06:32:44.657259+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-19T06:32:44.657259+00:00","source":"crossref"},{"observed_at":"2026-08-19T06:32:39.956319+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-12T19:35:31.170182Z","title":"Guidelines: • The answer NA means that the paper poses no such risks","venue":null,"work_id":"6af8b099-55bd-4e2d-bfaa-59569d13ca24","year":null},"citing_paper":{"arxiv_id":"2411.10606","last_updated":"2024-11-15T22:02:28Z","snapshot_observed_at":"2026-08-14T13:50:36.265901Z","submitted_at":"2024-11-15T22:02:28Z","title":"AmoebaLLM: Constructing Any-Shape Large Language Models for Efficient and Instant Deployment","version":1},"reference_index":92,"source":"pdf_text","source_observed_at":"2026-08-12T19:35:30.400170Z"},"links":{"citing_paper":"/paper/2411.10606"},"observation_digest":"sha256:64ac1a26d9bd0316866e6229832398624e31de0bf701eaee3867eba6ee274a3e","observation_id":"1053ed95-0f29-4744-bb99-43298f9855a5","resolution":{"observed_at":"2026-08-12T19:35:31.175198Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-19T06:32:44.657259+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-19T06:32:44.657259+00:00","source":"crossref"},{"observed_at":"2026-08-19T06:32:39.956319+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-12T19:35:31.154421Z","title":"Guidelines: • The answer NA means that the paper does not use existing assets","venue":null,"work_id":"3bcae2c8-c07b-4dc2-b608-e7c710d44bab","year":null},"citing_paper":{"arxiv_id":"2411.10606","last_updated":"2024-11-15T22:02:28Z","snapshot_observed_at":"2026-08-14T13:50:36.265901Z","submitted_at":"2024-11-15T22:02:28Z","title":"AmoebaLLM: Constructing Any-Shape Large Language Models for Efficient and Instant Deployment","version":1},"reference_index":93,"source":"pdf_text","source_observed_at":"2026-08-12T19:35:30.404513Z"},"links":{"citing_paper":"/paper/2411.10606"},"observation_digest":"sha256:804683614a2778a79ab0ce5e68c2b1d10b54899049b7d799e389da6e947943a2","observation_id":"3c702b83-c509-48f1-a92a-b8fe33c1a627","resolution":{"observed_at":"2026-08-12T19:35:31.159213Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-19T06:32:44.657259+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-19T06:32:44.657259+00:00","source":"crossref"},{"observed_at":"2026-08-19T06:32:39.956319+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-12T19:35:31.138667Z","title":"Guidelines: • The answer NA means that the paper does not release new assets","venue":null,"work_id":"ccf4f7e6-041c-410b-bac5-db7a442f04e2","year":null},"citing_paper":{"arxiv_id":"2411.10606","last_updated":"2024-11-15T22:02:28Z","snapshot_observed_at":"2026-08-14T13:50:36.265901Z","submitted_at":"2024-11-15T22:02:28Z","title":"AmoebaLLM: Constructing Any-Shape Large Language Models for Efficient and Instant Deployment","version":1},"reference_index":94,"source":"pdf_text","source_observed_at":"2026-08-12T19:35:30.408527Z"},"links":{"citing_paper":"/paper/2411.10606"},"observation_digest":"sha256:a3ef6e1495f46a94ec1634ee9567b0befccf9239a759fa6cb3cc466797f99709","observation_id":"c82ef08d-fa6a-4304-886c-de4e33e462af","resolution":{"observed_at":"2026-08-12T19:35:31.143558Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-19T06:32:44.657259+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-19T06:32:44.657259+00:00","source":"crossref"},{"observed_at":"2026-08-19T06:32:39.956319+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-12T19:35:31.122510Z","title":"Guidelines: • The answer NA means that the paper does not involve crowdsourcing nor research with human subjects","venue":null,"work_id":"44584124-2c39-466d-9948-933fee0081ac","year":null},"citing_paper":{"arxiv_id":"2411.10606","last_updated":"2024-11-15T22:02:28Z","snapshot_observed_at":"2026-08-14T13:50:36.265901Z","submitted_at":"2024-11-15T22:02:28Z","title":"AmoebaLLM: Constructing Any-Shape Large Language Models for Efficient and Instant Deployment","version":1},"reference_index":95,"source":"pdf_text","source_observed_at":"2026-08-12T19:35:30.412847Z"},"links":{"citing_paper":"/paper/2411.10606"},"observation_digest":"sha256:08e0c307480a03c4bd9af8820b85bc83638b17951cb7a435539d41fe1bc72a5b","observation_id":"3729f352-9fba-40bb-89d5-9291098ffdaa","resolution":{"observed_at":"2026-08-12T19:35:31.127838Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-19T06:32:44.657259+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-19T06:32:44.657259+00:00","source":"crossref"},{"observed_at":"2026-08-19T06:32:39.956319+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-12T19:35:31.105251Z","title":"Guidelines: • The answer NA means that the paper does not involve crowdsourcing nor research with human subjects","venue":null,"work_id":"fa75893a-fb1f-47d4-9a23-7323f26bad63","year":null},"citing_paper":{"arxiv_id":"2411.10606","last_updated":"2024-11-15T22:02:28Z","snapshot_observed_at":"2026-08-14T13:50:36.265901Z","submitted_at":"2024-11-15T22:02:28Z","title":"AmoebaLLM: Constructing Any-Shape Large Language Models for Efficient and Instant Deployment","version":1},"reference_index":96,"source":"pdf_text","source_observed_at":"2026-08-12T19:35:30.417149Z"},"links":{"citing_paper":"/paper/2411.10606"},"observation_digest":"sha256:fd83cb0351034947ce829fae0a4746cdaaacc2a2dc4f43f0dafc664457d73f69","observation_id":"9c54f489-23f4-41e9-b7e4-311f76b1a11e","resolution":{"observed_at":"2026-08-12T19:35:31.110658Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-19T06:32:44.657259+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-19T06:32:44.657259+00:00","source":"crossref"},{"observed_at":"2026-08-19T06:32:39.956319+00:00","source":"retraction_watch"}],"state":"measured"}}],"paper":{"arxiv_id":"2411.10606","last_updated":"2024-11-15T22:02:28Z","latest_version":1,"primary_category":"cs.LG","snapshot_observed_at":"2026-08-14T13:50:36.265901Z","submitted_at":"2024-11-15T22:02:28Z","title":"AmoebaLLM: Constructing Any-Shape Large Language Models for Efficient and Instant Deployment"},"reference_resolution":{"displayed":96,"state_counts":{"malformed_identifier":0,"metadata_mismatch":0,"parse_uncertain":1,"unresolved":65,"verified_exact":1,"verified_fuzzy":29},"total_outbound_references":96},"refusal":"A citation records a reference. It does not transfer a finding from one paper to another.","schema":"pith.paper-citation-record.v1","standing_sources":[{"observed_at":"2026-08-19T06:32:44.657259+00:00","source":"crossref"},{"observed_at":"2026-08-19T06:32:39.956319+00:00","source":"retraction_watch"}],"thesis":"As of 19 August 2026, this Paper Citation Record lists 96 of 96 outbound references and 0 inbound Pith citation observations for arXiv:2411.10606."}