exit: Put an upper limit on how often we can oops

author Jann Horn <jannh@google.com>

Thu, 17 Nov 2022 23:43:22 +0000 (15:43 -0800)

committer Greg Kroah-Hartman <gregkh@linuxfoundation.org>

Tue, 24 Jan 2023 06:24:41 +0000 (07:24 +0100)
author Jann Horn <jannh@google.com>
Thu, 17 Nov 2022 23:43:22 +0000 (15:43 -0800)
committer Greg Kroah-Hartman <gregkh@linuxfoundation.org>
Tue, 24 Jan 2023 06:24:41 +0000 (07:24 +0100)
diff --git a/Documentation/admin-guide/sysctl/kernel.rst b/Documentation/admin-guide/sysctl/kernel.rst

index c2c64c1b706ff6be4859cf1be24bbb6cc0364a22..6075fa604fc9d9fd0b79afc3f9acb9cd87668048 100644 (file)
--- a/Documentation/admin-guide/sysctl/kernel.rst
+++ b/Documentation/admin-guide/sysctl/kernel.rst
@@ -667,6 +667,14 @@ This is the default behavior.
  an oops event is detected.
  
  
+oops_limit
+==========
+
+Number of kernel oopses after which the kernel should panic when
+``panic_on_oops`` is not set. Setting this to 0 or 1 has the same effect
+as setting ``panic_on_oops=1``.
+
+
  osrelease, ostype & version
  ===========================
  
diff --git a/kernel/exit.c b/kernel/exit.c

index 35e0a31a0315c1f4f3aa86c17903d118b16fa1d9..2ab3ead621189af9115279e685339a1bac8814a7 100644 (file)
--- a/kernel/exit.c
+++ b/kernel/exit.c
@@ -72,6 +72,33 @@
  #include <asm/unistd.h>
  #include <asm/mmu_context.h>
  
+/*
+ * The default value should be high enough to not crash a system that randomly
+ * crashes its kernel from time to time, but low enough to at least not permit
+ * overflowing 32-bit refcounts or the ldsem writer count.
+ */
+static unsigned int oops_limit = 10000;
+
+#ifdef CONFIG_SYSCTL
+static struct ctl_table kern_exit_table[] = {
+       {
+               .procname       = "oops_limit",
+               .data           = &oops_limit,
+               .maxlen         = sizeof(oops_limit),
+               .mode           = 0644,
+               .proc_handler   = proc_douintvec,
+       },
+       { }
+};
+
+static __init int kernel_exit_sysctls_init(void)
+{
+       register_sysctl_init("kernel", kern_exit_table);
+       return 0;
+}
+late_initcall(kernel_exit_sysctls_init);
+#endif
+
  static void __unhash_process(struct task_struct *p, bool group_dead)
  {
         nr_threads--;
@@ -874,6 +901,8 @@ void __noreturn do_exit(long code)
  
  void __noreturn make_task_dead(int signr)
  {
+       static atomic_t oops_count = ATOMIC_INIT(0);
+
         /*
          * Take the task off the cpu after something catastrophic has
          * happened.
@@ -897,6 +926,19 @@ void __noreturn make_task_dead(int signr)
                 preempt_count_set(PREEMPT_ENABLED);
         }
  
+       /*
+        * Every time the system oopses, if the oops happens while a reference
+        * to an object was held, the reference leaks.
+        * If the oops doesn't also leak memory, repeated oopsing can cause
+        * reference counters to wrap around (if they're not using refcount_t).
+        * This means that repeated oopsing can make unexploitable-looking bugs
+        * exploitable through repeated oopsing.
+        * To make sure this can't happen, place an upper bound on how often the
+        * kernel may oops without panic().
+        */
+       if (atomic_inc_return(&oops_count) >= READ_ONCE(oops_limit))
+               panic("Oopsed too often (kernel.oops_limit is %d)", oops_limit);
+
         /*
          * We're taking recursive faults here in make_task_dead. Safest is to just
          * leave this task alone and wait for reboot.
author	Jann Horn <jannh@google.com>
	Thu, 17 Nov 2022 23:43:22 +0000 (15:43 -0800)
committer	Greg Kroah-Hartman <gregkh@linuxfoundation.org>
	Tue, 24 Jan 2023 06:24:41 +0000 (07:24 +0100)
Documentation/admin-guide/sysctl/kernel.rst		patch \| blob \| history
kernel/exit.c		patch \| blob \| history