exit: Put an upper limit on how often we can oops

author Jann Horn <jannh@google.com>

Thu, 2 Feb 2023 04:42:48 +0000 (20:42 -0800)

committer Greg Kroah-Hartman <gregkh@linuxfoundation.org>

Mon, 6 Feb 2023 06:52:49 +0000 (07:52 +0100)
author Jann Horn <jannh@google.com>
Thu, 2 Feb 2023 04:42:48 +0000 (20:42 -0800)
committer Greg Kroah-Hartman <gregkh@linuxfoundation.org>
Mon, 6 Feb 2023 06:52:49 +0000 (07:52 +0100)
diff --git a/Documentation/admin-guide/sysctl/kernel.rst b/Documentation/admin-guide/sysctl/kernel.rst

index 9715685be6e3b7c531fd489eeba4672a113a772f..4bdf845c79aa3d78b1b1fde3ce2e0ef2cc71b06f 100644 (file)
--- a/Documentation/admin-guide/sysctl/kernel.rst
+++ b/Documentation/admin-guide/sysctl/kernel.rst
@@ -557,6 +557,14 @@ numa_balancing_scan_size_mb is how many megabytes worth of pages are
  scanned for a given scan.
  
  
+oops_limit
+==========
+
+Number of kernel oopses after which the kernel should panic when
+``panic_on_oops`` is not set. Setting this to 0 or 1 has the same effect
+as setting ``panic_on_oops=1``.
+
+
  osrelease, ostype & version:
  ============================
  
diff --git a/kernel/exit.c b/kernel/exit.c

index 6512d82b4d9b0a1aff1db55bdc1df885762e45b1..4236970aa43849fa0c79e2ba5832fe034b68478c 100644 (file)
--- a/kernel/exit.c
+++ b/kernel/exit.c
@@ -69,6 +69,33 @@
  #include <asm/pgtable.h>
  #include <asm/mmu_context.h>
  
+/*
+ * The default value should be high enough to not crash a system that randomly
+ * crashes its kernel from time to time, but low enough to at least not permit
+ * overflowing 32-bit refcounts or the ldsem writer count.
+ */
+static unsigned int oops_limit = 10000;
+
+#ifdef CONFIG_SYSCTL
+static struct ctl_table kern_exit_table[] = {
+       {
+               .procname       = "oops_limit",
+               .data           = &oops_limit,
+               .maxlen         = sizeof(oops_limit),
+               .mode           = 0644,
+               .proc_handler   = proc_douintvec,
+       },
+       { }
+};
+
+static __init int kernel_exit_sysctls_init(void)
+{
+       register_sysctl_init("kernel", kern_exit_table);
+       return 0;
+}
+late_initcall(kernel_exit_sysctls_init);
+#endif
+
  static void __unhash_process(struct task_struct *p, bool group_dead)
  {
         nr_threads--;
@@ -866,10 +893,26 @@ EXPORT_SYMBOL_GPL(do_exit);
  
  void __noreturn make_task_dead(int signr)
  {
+       static atomic_t oops_count = ATOMIC_INIT(0);
+
         /*
          * Take the task off the cpu after something catastrophic has
          * happened.
          */
+
+       /*
+        * Every time the system oopses, if the oops happens while a reference
+        * to an object was held, the reference leaks.
+        * If the oops doesn't also leak memory, repeated oopsing can cause
+        * reference counters to wrap around (if they're not using refcount_t).
+        * This means that repeated oopsing can make unexploitable-looking bugs
+        * exploitable through repeated oopsing.
+        * To make sure this can't happen, place an upper bound on how often the
+        * kernel may oops without panic().
+        */
+       if (atomic_inc_return(&oops_count) >= READ_ONCE(oops_limit))
+               panic("Oopsed too often (kernel.oops_limit is %d)", oops_limit);
+
         do_exit(signr);
  }
author	Jann Horn <jannh@google.com>
	Thu, 2 Feb 2023 04:42:48 +0000 (20:42 -0800)
committer	Greg Kroah-Hartman <gregkh@linuxfoundation.org>
	Mon, 6 Feb 2023 06:52:49 +0000 (07:52 +0100)
Documentation/admin-guide/sysctl/kernel.rst		patch \| blob \| history
kernel/exit.c		patch \| blob \| history