{"as_of":"2026-08-19T20:27:00Z","caps":{"database_statements":6,"inbound":100,"outbound":100},"context_digest":"sha256:e683c88008cefb016e49368d32392d075a058a8b3e2d158127ba45117748faf4","coverage":[{"denominator":37,"lane":"reference_resolution","note":"Typed states for the displayed outbound observations.","records_observed":37,"source":"paper_references, paper_reference_links","source_observed_at":"2026-08-12T14:07:52.471138Z","state":"measured"},{"denominator":39,"lane":"standing_notices","note":"One-hop event checks from named stored sources.","records_observed":39,"source":"scholarly_work_events, retraction_status_cache","source_observed_at":"2026-08-19T06:32:44.657259+00:00","state":"measured"},{"denominator":2,"lane":"inbound_itemization","note":"Pith citing papers itemized under the disclosed page cap.","records_observed":2,"source":"paper_references, paper_reference_links","source_observed_at":"2026-08-16T11:07:59.365233Z","state":"measured"},{"denominator":1,"lane":"external_citation_measurements","note":"A source-named dated measurement, never combined with another source.","records_observed":0,"source":"pith","source_observed_at":"2026-08-10T20:34:35.675801Z","state":"measured"}],"external_citation_measurements":[],"inbound":[{"citation":{"cited_paper":{"arxiv_id":"2411.15664","last_updated":"2024-11-23T22:19:37Z","snapshot_observed_at":"2026-08-17T14:21:31.576613Z","submitted_at":"2024-11-23T22:19:37Z","title":"Enabling Efficient Serverless Inference Serving for LLM (Large Language Model) in the Cloud","version":1},"cited_work":{"arxiv_id":"2411.15664","doi":null,"metadata_source":"pith","pith_arxiv_id":"2411.15664","snapshot_observed_at":"2026-08-10T20:34:35.675801Z","title":"Enabling Efficient Serverless Inference Serving for LLM (Large Language Model) in the Cloud","venue":"cs.DC","work_id":"85663b51-fd5e-4061-90d6-dbdbfab74213","year":2024},"citing_paper":{"arxiv_id":"2501.08262","last_updated":"2025-01-14T17:21:16Z","snapshot_observed_at":"2026-08-14T19:33:53.600679Z","submitted_at":"2025-01-14T17:21:16Z","title":"Addressing the sustainable AI trilemma: a case study on LLM agents and RAG","version":1},"reference_index":2024,"source":"pdf_text","source_observed_at":"2026-08-10T20:34:34.676288Z"},"links":{"cited_paper":"/paper/2411.15664","citing_paper":"/paper/2501.08262"},"observation_digest":"sha256:4b3e9c6e5883d979452cba41ea49005e8c81d4057ee3523b81653553a83f5aff","observation_id":"36daebe2-8f27-449a-814d-24b40e8aae63","resolution":{"observed_at":"2026-08-10T20:34:35.680096Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-19T06:32:44.657259+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-19T06:32:44.657259+00:00","source":"crossref"},{"observed_at":"2026-08-19T06:32:39.956319+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2411.15664","last_updated":"2024-11-23T22:19:37Z","snapshot_observed_at":"2026-08-17T14:21:31.576613Z","submitted_at":"2024-11-23T22:19:37Z","title":"Enabling Efficient Serverless Inference Serving for LLM (Large Language Model) in the Cloud","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2411.15664","snapshot_observed_at":"2026-08-16T11:07:59.365233Z","title":"Enabling efficient serverless inference serving for llm (large language model) in the cloud,","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2504.16420","last_updated":"2025-04-23T05:02:51Z","snapshot_observed_at":"2026-08-17T14:22:18.404244Z","submitted_at":"2025-04-23T05:02:51Z","title":"A Survey of Foundation Model-Powered Recommender Systems: From Feature-Based, Generative to Agentic Paradigms","version":1},"reference_index":246,"source":"pdf_text","source_observed_at":"2026-08-16T11:07:59.365233Z"},"links":{"cited_paper":"/paper/2411.15664","citing_paper":"/paper/2504.16420"},"observation_digest":"sha256:0f36438bc3edd00f47e53631112f0f9b483a7b8dd064af2aee59e9f685994e5c","observation_id":"d8e55695-8058-4494-be69-13e34c0f8129","resolution":{"observed_at":"2026-08-16T11:07:59.365233Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}}],"links":{"evidence":"/evidence","html":"/paper/2411.15664/citation-record","integrity":"/paper/2411.15664/integrity","json":"/paper/2411.15664/citation-record.json","paper":"/paper/2411.15664"},"outbound":[{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-12T14:07:53.202718Z","title":"slideshare","venue":null,"work_id":"d06de8ff-6561-4e7c-a2a7-7ae51b04ed1a","year":2017},"citing_paper":{"arxiv_id":"2411.15664","last_updated":"2024-11-23T22:19:37Z","snapshot_observed_at":"2026-08-17T14:21:31.576613Z","submitted_at":"2024-11-23T22:19:37Z","title":"Enabling Efficient Serverless Inference Serving for LLM (Large Language Model) in the Cloud","version":1},"reference_index":1,"source":"pdf_text","source_observed_at":"2026-08-12T14:07:52.342326Z"},"links":{"citing_paper":"/paper/2411.15664"},"observation_digest":"sha256:aa2293456b989b34aaa53dc037de89862c2c746199cc5e5bbb7a9624a159f84b","observation_id":"364f3b2d-7051-41f6-93c8-dedeace87310","resolution":{"observed_at":"2026-08-12T14:07:53.205858Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-19T06:32:44.657259+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-19T06:32:44.657259+00:00","source":"crossref"},{"observed_at":"2026-08-19T06:32:39.956319+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-12T14:07:53.192069Z","title":null,"venue":null,"work_id":"b4812003-e656-4c3d-b95b-5e679edab27b","year":null},"citing_paper":{"arxiv_id":"2411.15664","last_updated":"2024-11-23T22:19:37Z","snapshot_observed_at":"2026-08-17T14:21:31.576613Z","submitted_at":"2024-11-23T22:19:37Z","title":"Enabling Efficient Serverless Inference Serving for LLM (Large Language Model) in the Cloud","version":1},"reference_index":2,"source":"pdf_text","source_observed_at":"2026-08-12T14:07:52.346076Z"},"links":{"citing_paper":"/paper/2411.15664"},"observation_digest":"sha256:b43d65d37e47d83472644deb10df226927422510af53f9263fefa250963b2d61","observation_id":"5df39438-f0dd-40b9-b554-337378de84b6","resolution":{"observed_at":"2026-08-12T14:07:53.196280Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-19T06:32:44.657259+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-19T06:32:44.657259+00:00","source":"crossref"},{"observed_at":"2026-08-19T06:32:39.956319+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-12T14:07:53.181923Z","title":null,"venue":null,"work_id":"06d3cb0e-d614-4b23-a7be-03f5cba4ed2d","year":null},"citing_paper":{"arxiv_id":"2411.15664","last_updated":"2024-11-23T22:19:37Z","snapshot_observed_at":"2026-08-17T14:21:31.576613Z","submitted_at":"2024-11-23T22:19:37Z","title":"Enabling Efficient Serverless Inference Serving for LLM (Large Language Model) in the Cloud","version":1},"reference_index":3,"source":"pdf_text","source_observed_at":"2026-08-12T14:07:52.349499Z"},"links":{"citing_paper":"/paper/2411.15664"},"observation_digest":"sha256:bab6c27372c9547e2c7a116d14a2f8ce4e9599a89d54f70e7473c459efdce1cb","observation_id":"1023118d-9954-4aca-9d91-080b1315d27e","resolution":{"observed_at":"2026-08-12T14:07:53.185421Z","resolver_source":"raw_fallback","status":"parse_uncertain"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-19T06:32:44.657259+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-19T06:32:44.657259+00:00","source":"crossref"},{"observed_at":"2026-08-19T06:32:39.956319+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-12T14:07:53.171066Z","title":"com/ jeremydaly/ lambda-warmer","venue":null,"work_id":"40d2fecf-b76f-4407-ad45-fa5c043797d8","year":null},"citing_paper":{"arxiv_id":"2411.15664","last_updated":"2024-11-23T22:19:37Z","snapshot_observed_at":"2026-08-17T14:21:31.576613Z","submitted_at":"2024-11-23T22:19:37Z","title":"Enabling Efficient Serverless Inference Serving for LLM (Large Language Model) in the Cloud","version":1},"reference_index":4,"source":"pdf_text","source_observed_at":"2026-08-12T14:07:52.353297Z"},"links":{"citing_paper":"/paper/2411.15664"},"observation_digest":"sha256:0e8c5a997c566fe56c08f2f03fe4ebe60fc89039a7d7b18b94af3c8bfb66d8a5","observation_id":"feb893f8-64be-4fa3-b183-850deaffcb9b","resolution":{"observed_at":"2026-08-12T14:07:53.175555Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-19T06:32:44.657259+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-19T06:32:44.657259+00:00","source":"crossref"},{"observed_at":"2026-08-19T06:32:39.956319+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-12T14:07:53.161090Z","title":"com/ google-cloud/ 3-solutions-to-mitigate-the-cold-starts-on-cloud-run-8c60f0ae7894","venue":null,"work_id":"f391a95b-bcfa-4adb-87c4-073064e0f270","year":null},"citing_paper":{"arxiv_id":"2411.15664","last_updated":"2024-11-23T22:19:37Z","snapshot_observed_at":"2026-08-17T14:21:31.576613Z","submitted_at":"2024-11-23T22:19:37Z","title":"Enabling Efficient Serverless Inference Serving for LLM (Large Language Model) in the Cloud","version":1},"reference_index":5,"source":"pdf_text","source_observed_at":"2026-08-12T14:07:52.356999Z"},"links":{"citing_paper":"/paper/2411.15664"},"observation_digest":"sha256:334073f4f4d3e610ed7b48eedfc1517b680dbb6d9dbac0a0b5461573d8b17d09","observation_id":"d3177d5b-5d1d-4900-a5f4-aa86aa6436e9","resolution":{"observed_at":"2026-08-12T14:07:53.164567Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-19T06:32:44.657259+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-19T06:32:44.657259+00:00","source":"crossref"},{"observed_at":"2026-08-19T06:32:39.956319+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-12T14:07:53.151642Z","title":null,"venue":null,"work_id":"baf5967d-b9b2-4034-b2f4-fce4c57433f6","year":null},"citing_paper":{"arxiv_id":"2411.15664","last_updated":"2024-11-23T22:19:37Z","snapshot_observed_at":"2026-08-17T14:21:31.576613Z","submitted_at":"2024-11-23T22:19:37Z","title":"Enabling Efficient Serverless Inference Serving for LLM (Large Language Model) in the Cloud","version":1},"reference_index":6,"source":"pdf_text","source_observed_at":"2026-08-12T14:07:52.360480Z"},"links":{"citing_paper":"/paper/2411.15664"},"observation_digest":"sha256:972fcb63feaac3eeba027a25c4ebd6fc0d48fcc8769f9179635e069ee1ada490","observation_id":"2f795c3d-195c-42f5-8547-405954d828c8","resolution":{"observed_at":"2026-08-12T14:07:53.154936Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-19T06:32:44.657259+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-19T06:32:44.657259+00:00","source":"crossref"},{"observed_at":"2026-08-19T06:32:39.956319+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-12T14:07:53.143081Z","title":null,"venue":null,"work_id":"f40c6d88-2783-4d67-8a2d-239fc40aca91","year":null},"citing_paper":{"arxiv_id":"2411.15664","last_updated":"2024-11-23T22:19:37Z","snapshot_observed_at":"2026-08-17T14:21:31.576613Z","submitted_at":"2024-11-23T22:19:37Z","title":"Enabling Efficient Serverless Inference Serving for LLM (Large Language Model) in the Cloud","version":1},"reference_index":7,"source":"pdf_text","source_observed_at":"2026-08-12T14:07:52.363413Z"},"links":{"citing_paper":"/paper/2411.15664"},"observation_digest":"sha256:aa7efb531ee4244bec30c4a9568a0ffe63f31063aca8607fef1c8c169b775f45","observation_id":"3d03005d-8e8d-46a6-8ae3-291f4f5dbaba","resolution":{"observed_at":"2026-08-12T14:07:53.145956Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-19T06:32:44.657259+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-19T06:32:44.657259+00:00","source":"crossref"},{"observed_at":"2026-08-19T06:32:39.956319+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-12T14:07:53.133732Z","title":"microsoft","venue":null,"work_id":"7f7c28d2-4115-4c9f-96d7-8675494c0b75","year":null},"citing_paper":{"arxiv_id":"2411.15664","last_updated":"2024-11-23T22:19:37Z","snapshot_observed_at":"2026-08-17T14:21:31.576613Z","submitted_at":"2024-11-23T22:19:37Z","title":"Enabling Efficient Serverless Inference Serving for LLM (Large Language Model) in the Cloud","version":1},"reference_index":8,"source":"pdf_text","source_observed_at":"2026-08-12T14:07:52.366346Z"},"links":{"citing_paper":"/paper/2411.15664"},"observation_digest":"sha256:5ccb256b0dd9bda2c29384e5b6c624f0e424e2358f035c098b7705bec1a0c73f","observation_id":"202f8e13-e3c3-4812-bcdc-c0c01875a679","resolution":{"observed_at":"2026-08-12T14:07:53.137112Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-19T06:32:44.657259+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-19T06:32:44.657259+00:00","source":"crossref"},{"observed_at":"2026-08-19T06:32:39.956319+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-12T14:07:53.123556Z","title":null,"venue":null,"work_id":"909adf15-1637-4748-ad59-16b4f8c42d4c","year":null},"citing_paper":{"arxiv_id":"2411.15664","last_updated":"2024-11-23T22:19:37Z","snapshot_observed_at":"2026-08-17T14:21:31.576613Z","submitted_at":"2024-11-23T22:19:37Z","title":"Enabling Efficient Serverless Inference Serving for LLM (Large Language Model) in the Cloud","version":1},"reference_index":9,"source":"pdf_text","source_observed_at":"2026-08-12T14:07:52.368997Z"},"links":{"citing_paper":"/paper/2411.15664"},"observation_digest":"sha256:624b6a0571881d6b8be83ef5d0afdfd2f84586715ac9edb2f437aa242f029b51","observation_id":"de4aaac9-7e90-4eba-a5dc-089d7fca0403","resolution":{"observed_at":"2026-08-12T14:07:53.127094Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-19T06:32:44.657259+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-19T06:32:44.657259+00:00","source":"crossref"},{"observed_at":"2026-08-19T06:32:39.956319+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-12T14:07:53.114066Z","title":"https: // www","venue":null,"work_id":"8f7c4cf4-a072-4cfe-9b4c-ceb148f18746","year":null},"citing_paper":{"arxiv_id":"2411.15664","last_updated":"2024-11-23T22:19:37Z","snapshot_observed_at":"2026-08-17T14:21:31.576613Z","submitted_at":"2024-11-23T22:19:37Z","title":"Enabling Efficient Serverless Inference Serving for LLM (Large Language Model) in the Cloud","version":1},"reference_index":10,"source":"pdf_text","source_observed_at":"2026-08-12T14:07:52.371989Z"},"links":{"citing_paper":"/paper/2411.15664"},"observation_digest":"sha256:64e8cf39763b2906a1837af07c1e258dc03dd926ee2ee654d6ff64def583613a","observation_id":"2fcba9bd-9c35-43a4-a845-12c0b7fa35f5","resolution":{"observed_at":"2026-08-12T14:07:53.117528Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-19T06:32:44.657259+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-19T06:32:44.657259+00:00","source":"crossref"},{"observed_at":"2026-08-19T06:32:39.956319+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-12T14:07:53.104296Z","title":"SedAI ( https: // www","venue":null,"work_id":"9c7c06ed-41e4-436b-b8c9-e77e6130a8b9","year":null},"citing_paper":{"arxiv_id":"2411.15664","last_updated":"2024-11-23T22:19:37Z","snapshot_observed_at":"2026-08-17T14:21:31.576613Z","submitted_at":"2024-11-23T22:19:37Z","title":"Enabling Efficient Serverless Inference Serving for LLM (Large Language Model) in the Cloud","version":1},"reference_index":11,"source":"pdf_text","source_observed_at":"2026-08-12T14:07:52.375096Z"},"links":{"citing_paper":"/paper/2411.15664"},"observation_digest":"sha256:7dc022c0b961b86cfb73d5f370e87b40efc64051adf51327f004dfed36d4846c","observation_id":"82eec83a-4d1c-40dd-b542-37ad77f34f47","resolution":{"observed_at":"2026-08-12T14:07:53.107679Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-19T06:32:44.657259+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-19T06:32:44.657259+00:00","source":"crossref"},{"observed_at":"2026-08-19T06:32:39.956319+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-12T14:07:53.094314Z","title":"Snowflake ( https: // www","venue":null,"work_id":"15b525ab-50b9-4a74-b18c-5df9536a331e","year":null},"citing_paper":{"arxiv_id":"2411.15664","last_updated":"2024-11-23T22:19:37Z","snapshot_observed_at":"2026-08-17T14:21:31.576613Z","submitted_at":"2024-11-23T22:19:37Z","title":"Enabling Efficient Serverless Inference Serving for LLM (Large Language Model) in the Cloud","version":1},"reference_index":12,"source":"pdf_text","source_observed_at":"2026-08-12T14:07:52.379176Z"},"links":{"citing_paper":"/paper/2411.15664"},"observation_digest":"sha256:fe6e1ebc6c8911e9535126cb4931410151b491648f790c13fad9f22e451788df","observation_id":"a4c88b13-02c7-44e8-a963-1a4f7d05878c","resolution":{"observed_at":"2026-08-12T14:07:53.097763Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-19T06:32:44.657259+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-19T06:32:44.657259+00:00","source":"crossref"},{"observed_at":"2026-08-19T06:32:39.956319+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-12T14:07:53.083961Z","title":"( https: // en","venue":null,"work_id":"183b65c4-b001-481f-b5c7-45c2ce11c37a","year":null},"citing_paper":{"arxiv_id":"2411.15664","last_updated":"2024-11-23T22:19:37Z","snapshot_observed_at":"2026-08-17T14:21:31.576613Z","submitted_at":"2024-11-23T22:19:37Z","title":"Enabling Efficient Serverless Inference Serving for LLM (Large Language Model) in the Cloud","version":1},"reference_index":13,"source":"pdf_text","source_observed_at":"2026-08-12T14:07:52.382791Z"},"links":{"citing_paper":"/paper/2411.15664"},"observation_digest":"sha256:ed49bcc172ebc2c3442b4b5ff1cede5aa65f3b90102f5cf50e45286bffa1acfa","observation_id":"4324c903-1289-4ab5-a36b-8b9f2f5b0cde","resolution":{"observed_at":"2026-08-12T14:07:53.088216Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-19T06:32:44.657259+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-19T06:32:44.657259+00:00","source":"crossref"},{"observed_at":"2026-08-19T06:32:39.956319+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2305.05920","last_updated":"2024-09-25T05:57:51Z","snapshot_observed_at":"2026-08-16T21:53:57.298060Z","submitted_at":"2023-05-10T06:17:50Z","title":"Fast Distributed Inference Serving for Large Language Models","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2305.05920","snapshot_observed_at":"2026-08-12T14:07:52.386091Z","title":null,"venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2411.15664","last_updated":"2024-11-23T22:19:37Z","snapshot_observed_at":"2026-08-17T14:21:31.576613Z","submitted_at":"2024-11-23T22:19:37Z","title":"Enabling Efficient Serverless Inference Serving for LLM (Large Language Model) in the Cloud","version":1},"reference_index":14,"source":"pdf_text","source_observed_at":"2026-08-12T14:07:52.386091Z"},"links":{"cited_paper":"/paper/2305.05920","citing_paper":"/paper/2411.15664"},"observation_digest":"sha256:e3a629774e69c9532c3593c61496d40ae850b50e3be831cebc36c65358fe28ba","observation_id":"45d9e220-c714-4936-8a6c-48e75dbbadc6","resolution":{"observed_at":"2026-08-12T14:07:52.386091Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-12T14:07:53.074773Z","title":null,"venue":null,"work_id":"decc58aa-b22a-4217-bd06-e412fa908837","year":null},"citing_paper":{"arxiv_id":"2411.15664","last_updated":"2024-11-23T22:19:37Z","snapshot_observed_at":"2026-08-17T14:21:31.576613Z","submitted_at":"2024-11-23T22:19:37Z","title":"Enabling Efficient Serverless Inference Serving for LLM (Large Language Model) in the Cloud","version":1},"reference_index":15,"source":"pdf_text","source_observed_at":"2026-08-12T14:07:52.390158Z"},"links":{"citing_paper":"/paper/2411.15664"},"observation_digest":"sha256:58854c8094078cea20c117f3f2152800fe3bd311e52fce816fb8ca10439fc5a0","observation_id":"de6462a3-8c62-4370-b450-504444aa0743","resolution":{"observed_at":"2026-08-12T14:07:53.077833Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-19T06:32:44.657259+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-19T06:32:44.657259+00:00","source":"crossref"},{"observed_at":"2026-08-19T06:32:39.956319+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-12T14:07:52.398677Z","title":"Catalyzer: Sub-millisecond startup for serverless computing with initialization-less booting","venue":null,"work_id":null,"year":2020},"citing_paper":{"arxiv_id":"2411.15664","last_updated":"2024-11-23T22:19:37Z","snapshot_observed_at":"2026-08-17T14:21:31.576613Z","submitted_at":"2024-11-23T22:19:37Z","title":"Enabling Efficient Serverless Inference Serving for LLM (Large Language Model) in the Cloud","version":1},"reference_index":16,"source":"pdf_text","source_observed_at":"2026-08-12T14:07:52.398677Z"},"links":{"citing_paper":"/paper/2411.15664"},"observation_digest":"sha256:f43d39e0b6aa1a39321b2dc00ea9cb592335531c88cadc833031c36a1ab351ae","observation_id":"68f846e0-6afe-439f-9d34-8d3fd4386f23","resolution":{"observed_at":"2026-08-12T14:07:52.398677Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-12T14:07:53.056065Z","title":null,"venue":null,"work_id":"82e5e216-739d-4fd7-93ed-e470860f459f","year":2018},"citing_paper":{"arxiv_id":"2411.15664","last_updated":"2024-11-23T22:19:37Z","snapshot_observed_at":"2026-08-17T14:21:31.576613Z","submitted_at":"2024-11-23T22:19:37Z","title":"Enabling Efficient Serverless Inference Serving for LLM (Large Language Model) in the Cloud","version":1},"reference_index":17,"source":"pdf_text","source_observed_at":"2026-08-12T14:07:52.402455Z"},"links":{"citing_paper":"/paper/2411.15664"},"observation_digest":"sha256:0701db634fdaa763a92697a7814c11dcf3967b296013d3c01215a050bcb6d3ae","observation_id":"4e46f8db-aba4-4517-846a-6e2e7015b7ef","resolution":{"observed_at":"2026-08-12T14:07:53.059267Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-19T06:32:44.657259+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-19T06:32:44.657259+00:00","source":"crossref"},{"observed_at":"2026-08-19T06:32:39.956319+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-12T14:07:52.406118Z","title":"Centralized core-granular scheduling for serverless functions","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2411.15664","last_updated":"2024-11-23T22:19:37Z","snapshot_observed_at":"2026-08-17T14:21:31.576613Z","submitted_at":"2024-11-23T22:19:37Z","title":"Enabling Efficient Serverless Inference Serving for LLM (Large Language Model) in the Cloud","version":1},"reference_index":18,"source":"pdf_text","source_observed_at":"2026-08-12T14:07:52.406118Z"},"links":{"citing_paper":"/paper/2411.15664"},"observation_digest":"sha256:c8d2c842b2e10b40cf214783b6710496aa57ade05acb9dc4c7730e4e4493f213","observation_id":"12f96f92-095c-409e-95ce-7c044a2ed2ec","resolution":{"observed_at":"2026-08-12T14:07:52.406118Z","resolver_source":null,"status":"malformed_identifier"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-12T14:07:53.065583Z","title":null,"venue":null,"work_id":"b01e7b51-acb0-406a-a835-507d97257b69","year":2019},"citing_paper":{"arxiv_id":"2411.15664","last_updated":"2024-11-23T22:19:37Z","snapshot_observed_at":"2026-08-17T14:21:31.576613Z","submitted_at":"2024-11-23T22:19:37Z","title":"Enabling Efficient Serverless Inference Serving for LLM (Large Language Model) in the Cloud","version":1},"reference_index":19,"source":"pdf_text","source_observed_at":"2026-08-12T14:07:52.394148Z"},"links":{"citing_paper":"/paper/2411.15664"},"observation_digest":"sha256:aeafb7317869cd83d1c124c250e2079012b166723cacfad3a98c03ab7fb3cded","observation_id":"0f092b08-7059-4c5c-ae16-0930f2ba1aff","resolution":{"observed_at":"2026-08-12T14:07:53.069167Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-19T06:32:44.657259+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-19T06:32:44.657259+00:00","source":"crossref"},{"observed_at":"2026-08-19T06:32:39.956319+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-12T14:07:53.046095Z","title":"Faaslight: General application-level cold-start latency optimization for function-as-a-service in serverless comput- ing","venue":null,"work_id":"d6955314-6f2b-4e0d-ab0a-07725747a928","year":2023},"citing_paper":{"arxiv_id":"2411.15664","last_updated":"2024-11-23T22:19:37Z","snapshot_observed_at":"2026-08-17T14:21:31.576613Z","submitted_at":"2024-11-23T22:19:37Z","title":"Enabling Efficient Serverless Inference Serving for LLM (Large Language Model) in the Cloud","version":1},"reference_index":20,"source":"pdf_text","source_observed_at":"2026-08-12T14:07:52.409303Z"},"links":{"citing_paper":"/paper/2411.15664"},"observation_digest":"sha256:fee3edea0019c946e0f56a2dcf865ea06f33ed7e52cf58a678281513ecd32aa8","observation_id":"db06b321-02d0-434d-8d5e-5a639575bc15","resolution":{"observed_at":"2026-08-12T14:07:53.049715Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-19T06:32:44.657259+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-19T06:32:44.657259+00:00","source":"crossref"},{"observed_at":"2026-08-19T06:32:39.956319+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-12T14:07:53.035802Z","title":"Rapid task provision- ing with Serverless-Optimized containers","venue":null,"work_id":"55043c20-d7f9-4873-845a-8da39b54feef","year":2018},"citing_paper":{"arxiv_id":"2411.15664","last_updated":"2024-11-23T22:19:37Z","snapshot_observed_at":"2026-08-17T14:21:31.576613Z","submitted_at":"2024-11-23T22:19:37Z","title":"Enabling Efficient Serverless Inference Serving for LLM (Large Language Model) in the Cloud","version":1},"reference_index":21,"source":"pdf_text","source_observed_at":"2026-08-12T14:07:52.413546Z"},"links":{"citing_paper":"/paper/2411.15664"},"observation_digest":"sha256:204b0531a229466b9f16b61209f9737d2755772a9fe21541900a8d178d87b6a0","observation_id":"5a4c2f0e-8182-4dd1-8406-34b232294aaf","resolution":{"observed_at":"2026-08-12T14:07:53.039443Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-19T06:32:44.657259+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-19T06:32:44.657259+00:00","source":"crossref"},{"observed_at":"2026-08-19T06:32:39.956319+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-12T14:07:53.025598Z","title":"Ghobaei-Arani","venue":null,"work_id":"84a8829e-22b7-42d3-9967-3d84e6b0965d","year":2024},"citing_paper":{"arxiv_id":"2411.15664","last_updated":"2024-11-23T22:19:37Z","snapshot_observed_at":"2026-08-17T14:21:31.576613Z","submitted_at":"2024-11-23T22:19:37Z","title":"Enabling Efficient Serverless Inference Serving for LLM (Large Language Model) in the Cloud","version":1},"reference_index":22,"source":"pdf_text","source_observed_at":"2026-08-12T14:07:52.416575Z"},"links":{"citing_paper":"/paper/2411.15664"},"observation_digest":"sha256:e8ac35a0603763e6e91414bada8daecba76d6ddfe0a1bede75841ec0eb815ce2","observation_id":"b4ba1b40-1a9f-4c4a-b1e0-f48fe07ae907","resolution":{"observed_at":"2026-08-12T14:07:53.029020Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-19T06:32:44.657259+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-19T06:32:44.657259+00:00","source":"crossref"},{"observed_at":"2026-08-19T06:32:39.956319+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-12T14:07:53.016253Z","title":"Cold start latency in serverless com- puting: A systematic review, taxonomy, and future directions","venue":null,"work_id":"81968d78-5070-4721-b6f3-66a6670c9a5b","year":null},"citing_paper":{"arxiv_id":"2411.15664","last_updated":"2024-11-23T22:19:37Z","snapshot_observed_at":"2026-08-17T14:21:31.576613Z","submitted_at":"2024-11-23T22:19:37Z","title":"Enabling Efficient Serverless Inference Serving for LLM (Large Language Model) in the Cloud","version":1},"reference_index":23,"source":"pdf_text","source_observed_at":"2026-08-12T14:07:52.419557Z"},"links":{"citing_paper":"/paper/2411.15664"},"observation_digest":"sha256:91410cd738de47a313e55bce2a36e633b4b743011bc655588d265b8131c3c24f","observation_id":"79557d58-7223-4ac3-8d6a-aa67466ddae0","resolution":{"observed_at":"2026-08-12T14:07:53.019552Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-19T06:32:44.657259+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-19T06:32:44.657259+00:00","source":"crossref"},{"observed_at":"2026-08-19T06:32:39.956319+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-12T14:07:53.005922Z","title":"What is serverless computing? IBM ( https: // www","venue":null,"work_id":"6d31a2d5-0567-4aad-ba56-3f06d6f9eab3","year":null},"citing_paper":{"arxiv_id":"2411.15664","last_updated":"2024-11-23T22:19:37Z","snapshot_observed_at":"2026-08-17T14:21:31.576613Z","submitted_at":"2024-11-23T22:19:37Z","title":"Enabling Efficient Serverless Inference Serving for LLM (Large Language Model) in the Cloud","version":1},"reference_index":24,"source":"pdf_text","source_observed_at":"2026-08-12T14:07:52.422766Z"},"links":{"citing_paper":"/paper/2411.15664"},"observation_digest":"sha256:49cfbfa170fa78ba7495235e9dcf322d938466495430646d50598cb5c3dea0b8","observation_id":"b542313f-fb5d-4c01-84dd-f9165069a116","resolution":{"observed_at":"2026-08-12T14:07:53.009677Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-19T06:32:44.657259+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-19T06:32:44.657259+00:00","source":"crossref"},{"observed_at":"2026-08-19T06:32:39.956319+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"1903.12221","last_updated":"2019-03-28T18:55:30Z","snapshot_observed_at":"2026-08-14T16:55:07.204261Z","submitted_at":"2019-03-28T18:55:30Z","title":"Mitigating Cold Starts in Serverless Platforms: A Pool-Based Approach","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"1903.12221","snapshot_observed_at":"2026-08-12T14:07:52.425584Z","title":"Mitigat- ing cold starts in serverless platforms: A pool-based approach https://arxiv.org/abs/ 1903.12221, 2019","venue":null,"work_id":null,"year":1903},"citing_paper":{"arxiv_id":"2411.15664","last_updated":"2024-11-23T22:19:37Z","snapshot_observed_at":"2026-08-17T14:21:31.576613Z","submitted_at":"2024-11-23T22:19:37Z","title":"Enabling Efficient Serverless Inference Serving for LLM (Large Language Model) in the Cloud","version":1},"reference_index":25,"source":"pdf_text","source_observed_at":"2026-08-12T14:07:52.425584Z"},"links":{"cited_paper":"/paper/1903.12221","citing_paper":"/paper/2411.15664"},"observation_digest":"sha256:4071b1be19600385b28f0b295c566b4ed3e146fc9085230ae1ea1a408351a11e","observation_id":"33d95291-1641-4918-a0e1-733e3aaf5e5f","resolution":{"observed_at":"2026-08-12T14:07:52.425584Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-12T14:07:52.994997Z","title":null,"venue":null,"work_id":"d651aa6d-9a27-4e8a-9cbc-aac003b843cd","year":2017},"citing_paper":{"arxiv_id":"2411.15664","last_updated":"2024-11-23T22:19:37Z","snapshot_observed_at":"2026-08-17T14:21:31.576613Z","submitted_at":"2024-11-23T22:19:37Z","title":"Enabling Efficient Serverless Inference Serving for LLM (Large Language Model) in the Cloud","version":1},"reference_index":26,"source":"pdf_text","source_observed_at":"2026-08-12T14:07:52.429465Z"},"links":{"citing_paper":"/paper/2411.15664"},"observation_digest":"sha256:38becb816bd45168000f8aaf1a96d3a738db7de8326efa4dab0a4a5365df0369","observation_id":"11398b45-2b56-40d3-9991-a9412e7f1d4f","resolution":{"observed_at":"2026-08-12T14:07:52.998521Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-19T06:32:44.657259+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-19T06:32:44.657259+00:00","source":"crossref"},{"observed_at":"2026-08-19T06:32:39.956319+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-12T14:07:52.983881Z","title":"Persson and W","venue":null,"work_id":"032c0c82-39b9-431e-b6ba-2523713213a4","year":null},"citing_paper":{"arxiv_id":"2411.15664","last_updated":"2024-11-23T22:19:37Z","snapshot_observed_at":"2026-08-17T14:21:31.576613Z","submitted_at":"2024-11-23T22:19:37Z","title":"Enabling Efficient Serverless Inference Serving for LLM (Large Language Model) in the Cloud","version":1},"reference_index":27,"source":"pdf_text","source_observed_at":"2026-08-12T14:07:52.433660Z"},"links":{"citing_paper":"/paper/2411.15664"},"observation_digest":"sha256:886233761018d8903bb440bc4beb3b8297c7b6bb3c3f4040c9bb02454abd8776","observation_id":"0e5868b2-1ac2-456a-bfa0-1064a03ebf44","resolution":{"observed_at":"2026-08-12T14:07:52.987551Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-19T06:32:44.657259+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-19T06:32:44.657259+00:00","source":"crossref"},{"observed_at":"2026-08-19T06:32:39.956319+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-12T14:07:52.437985Z","title":null,"venue":null,"work_id":null,"year":2021},"citing_paper":{"arxiv_id":"2411.15664","last_updated":"2024-11-23T22:19:37Z","snapshot_observed_at":"2026-08-17T14:21:31.576613Z","submitted_at":"2024-11-23T22:19:37Z","title":"Enabling Efficient Serverless Inference Serving for LLM (Large Language Model) in the Cloud","version":1},"reference_index":28,"source":"pdf_text","source_observed_at":"2026-08-12T14:07:52.437985Z"},"links":{"citing_paper":"/paper/2411.15664"},"observation_digest":"sha256:5e27e13622fc26e41d3e014fd1b8bcf13b4cb809b51728a4d9ad8f7a866c0fae","observation_id":"b4ad41ec-044c-48f0-96e5-3288dc7dee95","resolution":{"observed_at":"2026-08-12T14:07:52.437985Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-12T14:07:52.971603Z","title":"Pietzuch https: //api.semanticscholar.org/CorpusID: 11 51997872","venue":null,"work_id":"e1acfd42-61f7-44fe-93d3-642e295fc4e8","year":2018},"citing_paper":{"arxiv_id":"2411.15664","last_updated":"2024-11-23T22:19:37Z","snapshot_observed_at":"2026-08-17T14:21:31.576613Z","submitted_at":"2024-11-23T22:19:37Z","title":"Enabling Efficient Serverless Inference Serving for LLM (Large Language Model) in the Cloud","version":1},"reference_index":29,"source":"pdf_text","source_observed_at":"2026-08-12T14:07:52.441604Z"},"links":{"citing_paper":"/paper/2411.15664"},"observation_digest":"sha256:952be3e028ce59f00c8990dd851c890f0747bc1cd9a293539b28c8147592fd82","observation_id":"14f82162-c99a-4d65-b6f0-640b1c60b349","resolution":{"observed_at":"2026-08-12T14:07:52.975524Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-19T06:32:44.657259+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-19T06:32:44.657259+00:00","source":"crossref"},{"observed_at":"2026-08-19T06:32:39.956319+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-12T14:07:52.959966Z","title":null,"venue":null,"work_id":"a21c7f2c-acc1-44e2-a926-abefafcbe5b8","year":2020},"citing_paper":{"arxiv_id":"2411.15664","last_updated":"2024-11-23T22:19:37Z","snapshot_observed_at":"2026-08-17T14:21:31.576613Z","submitted_at":"2024-11-23T22:19:37Z","title":"Enabling Efficient Serverless Inference Serving for LLM (Large Language Model) in the Cloud","version":1},"reference_index":30,"source":"pdf_text","source_observed_at":"2026-08-12T14:07:52.445502Z"},"links":{"citing_paper":"/paper/2411.15664"},"observation_digest":"sha256:a2c330d7fac30a01e7d8db638fd2ccc9839507dd2a9ef1ae0d3210f06d17b05a","observation_id":"f63de849-60e9-4273-bc54-25fbc1fda6b5","resolution":{"observed_at":"2026-08-12T14:07:52.964010Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-19T06:32:44.657259+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-19T06:32:44.657259+00:00","source":"crossref"},{"observed_at":"2026-08-19T06:32:39.956319+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2206.12275","last_updated":"2022-12-16T13:22:26Z","snapshot_observed_at":"2026-08-16T16:50:36.791780Z","submitted_at":"2022-06-24T13:25:55Z","title":"Rise of the Planet of Serverless Computing: A Systematic Review","version":5},"cited_work":{"arxiv_id":"2206.12275","doi":null,"metadata_source":"pith","pith_arxiv_id":"2206.12275","snapshot_observed_at":"2026-08-12T14:07:52.668814Z","title":"Rise of the Planet of Serverless Computing: A Systematic Review","venue":"cs.SE","work_id":"6ddcb9d4-5d9e-4660-8b09-b6fb96148146","year":2022},"citing_paper":{"arxiv_id":"2411.15664","last_updated":"2024-11-23T22:19:37Z","snapshot_observed_at":"2026-08-17T14:21:31.576613Z","submitted_at":"2024-11-23T22:19:37Z","title":"Enabling Efficient Serverless Inference Serving for LLM (Large Language Model) in the Cloud","version":1},"reference_index":31,"source":"pdf_text","source_observed_at":"2026-08-12T14:07:52.448829Z"},"links":{"cited_paper":"/paper/2206.12275","citing_paper":"/paper/2411.15664"},"observation_digest":"sha256:715eccd2f1ada0bcf5cad1fb078edfd89f34dc2f77ed503e3c12c0d09b1da3d3","observation_id":"eb9f7d2c-ddeb-4acc-95b0-dd7eff3418a4","resolution":{"observed_at":"2026-08-12T14:07:52.674384Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-19T06:32:44.657259+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-19T06:32:44.657259+00:00","source":"crossref"},{"observed_at":"2026-08-19T06:32:39.956319+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-12T14:07:52.948797Z","title":null,"venue":null,"work_id":"9a3f6de9-7295-48fd-80e8-e89b8a8164a3","year":null},"citing_paper":{"arxiv_id":"2411.15664","last_updated":"2024-11-23T22:19:37Z","snapshot_observed_at":"2026-08-17T14:21:31.576613Z","submitted_at":"2024-11-23T22:19:37Z","title":"Enabling Efficient Serverless Inference Serving for LLM (Large Language Model) in the Cloud","version":1},"reference_index":32,"source":"pdf_text","source_observed_at":"2026-08-12T14:07:52.452351Z"},"links":{"citing_paper":"/paper/2411.15664"},"observation_digest":"sha256:d090d867c82a9c706e39f5f24ec1173a9298facec1a9cadafaa42645038d1320","observation_id":"727ca945-9b4f-450d-a8cd-e42ab8a30ace","resolution":{"observed_at":"2026-08-12T14:07:52.952132Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-19T06:32:44.657259+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-19T06:32:44.657259+00:00","source":"crossref"},{"observed_at":"2026-08-19T06:32:39.956319+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2303.06865","last_updated":"2023-06-12T07:48:53Z","snapshot_observed_at":"2026-08-17T14:20:32.390487Z","submitted_at":"2023-03-13T05:19:28Z","title":"FlexGen: High-Throughput Generative Inference of Large Language Models with a Single GPU","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2303.06865","snapshot_observed_at":"2026-08-12T14:07:52.463915Z","title":null,"venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2411.15664","last_updated":"2024-11-23T22:19:37Z","snapshot_observed_at":"2026-08-17T14:21:31.576613Z","submitted_at":"2024-11-23T22:19:37Z","title":"Enabling Efficient Serverless Inference Serving for LLM (Large Language Model) in the Cloud","version":1},"reference_index":33,"source":"pdf_text","source_observed_at":"2026-08-12T14:07:52.463915Z"},"links":{"cited_paper":"/paper/2303.06865","citing_paper":"/paper/2411.15664"},"observation_digest":"sha256:1d7c39d957ffd6a72561f6d64a8bad13d76538a01b6c44d92837b2084545a672","observation_id":"3aeb15bc-302b-4493-a87b-30bbfd068181","resolution":{"observed_at":"2026-08-12T14:07:52.463915Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-12T14:07:52.460464Z","title":null,"venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2411.15664","last_updated":"2024-11-23T22:19:37Z","snapshot_observed_at":"2026-08-17T14:21:31.576613Z","submitted_at":"2024-11-23T22:19:37Z","title":"Enabling Efficient Serverless Inference Serving for LLM (Large Language Model) in the Cloud","version":1},"reference_index":34,"source":"pdf_text","source_observed_at":"2026-08-12T14:07:52.460464Z"},"links":{"citing_paper":"/paper/2411.15664"},"observation_digest":"sha256:599a4533b46173284fb0ec85ebe32decf8aeb2976536c58fcbb97f63cc15c601","observation_id":"0daaa7ee-668f-4558-bc16-d7a6f5ed40da","resolution":{"observed_at":"2026-08-12T14:07:52.460464Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-12T14:07:52.938026Z","title":"Taming serverless cold start of cloud model inference with edge comput- ing","venue":null,"work_id":"f81e1691-ebf3-4e5c-bdaa-7b4e2790f252","year":2024},"citing_paper":{"arxiv_id":"2411.15664","last_updated":"2024-11-23T22:19:37Z","snapshot_observed_at":"2026-08-17T14:21:31.576613Z","submitted_at":"2024-11-23T22:19:37Z","title":"Enabling Efficient Serverless Inference Serving for LLM (Large Language Model) in the Cloud","version":1},"reference_index":35,"source":"pdf_text","source_observed_at":"2026-08-12T14:07:52.471138Z"},"links":{"citing_paper":"/paper/2411.15664"},"observation_digest":"sha256:75007d1860920218791119a1f01ee224ac169af389ce3467a4e6f13a80b14a62","observation_id":"ade5bb26-b9e9-4a20-97a5-b1e61e5ec008","resolution":{"observed_at":"2026-08-12T14:07:52.942131Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-19T06:32:44.657259+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-19T06:32:44.657259+00:00","source":"crossref"},{"observed_at":"2026-08-19T06:32:39.956319+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-12T14:07:52.467684Z","title":null,"venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2411.15664","last_updated":"2024-11-23T22:19:37Z","snapshot_observed_at":"2026-08-17T14:21:31.576613Z","submitted_at":"2024-11-23T22:19:37Z","title":"Enabling Efficient Serverless Inference Serving for LLM (Large Language Model) in the Cloud","version":1},"reference_index":36,"source":"pdf_text","source_observed_at":"2026-08-12T14:07:52.467684Z"},"links":{"citing_paper":"/paper/2411.15664"},"observation_digest":"sha256:6bd9c18759f2473635cdd543d7f0cf84b30a17b32186fbefbeff02f27c1a5676","observation_id":"11983ed9-bda9-45a0-bd46-fcb59feb7f7c","resolution":{"observed_at":"2026-08-12T14:07:52.467684Z","resolver_source":null,"status":"malformed_identifier"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2401.14351","last_updated":"2024-07-25T08:08:11Z","snapshot_observed_at":"2026-08-18T17:31:57.587960Z","submitted_at":"2024-01-25T17:55:07Z","title":"ServerlessLLM: Low-Latency Serverless Inference for Large Language Models","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2401.14351","snapshot_observed_at":"2026-08-12T14:07:52.456185Z","title":null,"venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2411.15664","last_updated":"2024-11-23T22:19:37Z","snapshot_observed_at":"2026-08-17T14:21:31.576613Z","submitted_at":"2024-11-23T22:19:37Z","title":"Enabling Efficient Serverless Inference Serving for LLM (Large Language Model) in the Cloud","version":1},"reference_index":2024,"source":"pdf_text","source_observed_at":"2026-08-12T14:07:52.456185Z"},"links":{"cited_paper":"/paper/2401.14351","citing_paper":"/paper/2411.15664"},"observation_digest":"sha256:7696f4a0f4b9e3e41882b1721f9ed6a45b98ef7f82bb43749d773e0382c4285b","observation_id":"ae7ea405-8a5d-4248-bc6f-6f03e4a0382b","resolution":{"observed_at":"2026-08-12T14:07:52.456185Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}}],"paper":{"arxiv_id":"2411.15664","last_updated":"2024-11-23T22:19:37Z","latest_version":1,"primary_category":"cs.DC","snapshot_observed_at":"2026-08-17T14:21:31.576613Z","submitted_at":"2024-11-23T22:19:37Z","title":"Enabling Efficient Serverless Inference Serving for LLM (Large Language Model) in the Cloud"},"reference_resolution":{"displayed":37,"state_counts":{"malformed_identifier":2,"metadata_mismatch":0,"parse_uncertain":1,"unresolved":17,"verified_exact":1,"verified_fuzzy":16},"total_outbound_references":37},"refusal":"A citation records a reference. It does not transfer a finding from one paper to another.","schema":"pith.paper-citation-record.v1","standing_sources":[{"observed_at":"2026-08-19T06:32:44.657259+00:00","source":"crossref"},{"observed_at":"2026-08-19T06:32:39.956319+00:00","source":"retraction_watch"}],"thesis":"As of 19 August 2026, this Paper Citation Record lists 37 of 37 outbound references and 2 inbound Pith citation observations for arXiv:2411.15664."}