[{"data":1,"prerenderedAt":-1},["ShallowReactive",2],{"glossary-autoscaling::en":3,"gloss-cluster-autoscaling::en":26,"gloss-next-autoscaling::en":9},{"slug":4,"category":5,"name":6,"definition":7,"meta_desc":8,"faq":9,"schema_markup":9,"related":10},"autoscaling","cloud","Autoscaling","Autoscaling automatically adds or removes compute capacity in response to real-time demand, so your app keeps up during traffic spikes and stops paying for idle machines when things go quiet. Horizontal autoscaling spins up more identical instances behind a load balancer; vertical autoscaling resizes an instance to a bigger machine. Rules trigger on metrics like CPU, memory, request rate, or queue depth. Why it matters: for a SaaS with uneven load — a Monday-morning rush, a launch on Product Hunt — autoscaling is the difference between a graceful ramp and a crashed site or a wastefully over-provisioned fleet. Practical note: scaling is not instant. New instances take seconds to minutes to boot and warm caches (see cold start), so set thresholds to scale up early, not at 100% CPU. Set a sane maximum to cap runaway bills from a bug or bot flood, and load-test your scale-up path before you actually need it.","Autoscaling adds and removes compute as demand moves, so you survive traffic spikes and stop paying for idle machines — but cold starts and scale-down lag are real.",null,[11,14,17,20,23],{"slug":12,"name":13},"cold-start","Cold Start",{"slug":15,"name":16},"load-balancer","Load Balancer",{"slug":18,"name":19},"serverless","Serverless",{"slug":21,"name":22},"spot-instances","Spot Instances",{"slug":24,"name":25},"uptime","Uptime",[27,31,34,38,41,44,47,50,53,56,59,62],{"slug":28,"category":5,"name":29,"updated_at":30},"availability-zone","Availability Zone (AZ)","2026-08-24T02:46:37+00:00",{"slug":32,"category":5,"name":33,"updated_at":30},"block-storage","Block Storage",{"slug":35,"category":5,"name":36,"updated_at":37},"disaster-recovery","Disaster Recovery","2026-08-24T02:46:38+00:00",{"slug":39,"category":5,"name":40,"updated_at":37},"edge-ai","Edge AI",{"slug":42,"category":5,"name":43,"updated_at":30},"egress-fees","Egress Fees (Data Transfer Out)",{"slug":45,"category":5,"name":46,"updated_at":30},"finops","FinOps (Cloud Financial Operations)",{"slug":48,"category":5,"name":49,"updated_at":37},"immutable-infrastructure","Immutable Infrastructure",{"slug":51,"category":5,"name":52,"updated_at":37},"infrastructure-drift","Infrastructure Drift",{"slug":54,"category":5,"name":55,"updated_at":30},"managed-kubernetes","Managed Kubernetes",{"slug":57,"category":5,"name":58,"updated_at":30},"multi-region","Multi-Region",{"slug":60,"category":5,"name":61,"updated_at":37},"noisy-neighbor","Noisy Neighbor",{"slug":63,"category":5,"name":64,"updated_at":30},"platform-as-a-service","Platform as a Service (PaaS)"]