blob: ee1b2da9e5e7d168de7de4053eaaed737dd4c3e9 [file] [log] [blame]
Linus Torvalds1da177e2005-04-16 15:20:36 -07001/*
2 * linux/arch/x86-64/traps.c
3 *
4 * Copyright (C) 1991, 1992 Linus Torvalds
5 * Copyright (C) 2000, 2001, 2002 Andi Kleen, SuSE Labs
6 *
7 * Pentium III FXSR, SSE support
8 * Gareth Hughes <gareth@valinux.com>, May 2000
9 *
10 * $Id: traps.c,v 1.36 2002/03/24 11:09:10 ak Exp $
11 */
12
13/*
14 * 'Traps.c' handles hardware traps and faults after we have saved some
15 * state in 'entry.S'.
16 */
17#include <linux/config.h>
18#include <linux/sched.h>
19#include <linux/kernel.h>
20#include <linux/string.h>
21#include <linux/errno.h>
22#include <linux/ptrace.h>
23#include <linux/timer.h>
24#include <linux/mm.h>
25#include <linux/init.h>
26#include <linux/delay.h>
27#include <linux/spinlock.h>
28#include <linux/interrupt.h>
29#include <linux/module.h>
30#include <linux/moduleparam.h>
akpm@osdl.org35faa712005-04-16 15:24:54 -070031#include <linux/nmi.h>
Prasanna S Panchamukhi0f2fbdc2005-09-06 15:19:28 -070032#include <linux/kprobes.h>
Linus Torvalds1da177e2005-04-16 15:20:36 -070033
34#include <asm/system.h>
35#include <asm/uaccess.h>
36#include <asm/io.h>
37#include <asm/atomic.h>
38#include <asm/debugreg.h>
39#include <asm/desc.h>
40#include <asm/i387.h>
41#include <asm/kdebug.h>
42#include <asm/processor.h>
43
44#include <asm/smp.h>
45#include <asm/pgalloc.h>
46#include <asm/pda.h>
47#include <asm/proto.h>
48#include <asm/nmi.h>
49
Linus Torvalds1da177e2005-04-16 15:20:36 -070050extern struct gate_struct idt_table[256];
51
52asmlinkage void divide_error(void);
53asmlinkage void debug(void);
54asmlinkage void nmi(void);
55asmlinkage void int3(void);
56asmlinkage void overflow(void);
57asmlinkage void bounds(void);
58asmlinkage void invalid_op(void);
59asmlinkage void device_not_available(void);
60asmlinkage void double_fault(void);
61asmlinkage void coprocessor_segment_overrun(void);
62asmlinkage void invalid_TSS(void);
63asmlinkage void segment_not_present(void);
64asmlinkage void stack_segment(void);
65asmlinkage void general_protection(void);
66asmlinkage void page_fault(void);
67asmlinkage void coprocessor_error(void);
68asmlinkage void simd_coprocessor_error(void);
69asmlinkage void reserved(void);
70asmlinkage void alignment_check(void);
71asmlinkage void machine_check(void);
72asmlinkage void spurious_interrupt_bug(void);
Linus Torvalds1da177e2005-04-16 15:20:36 -070073
74struct notifier_block *die_chain;
75static DEFINE_SPINLOCK(die_notifier_lock);
76
77int register_die_notifier(struct notifier_block *nb)
78{
79 int err = 0;
80 unsigned long flags;
81 spin_lock_irqsave(&die_notifier_lock, flags);
82 err = notifier_chain_register(&die_chain, nb);
83 spin_unlock_irqrestore(&die_notifier_lock, flags);
84 return err;
85}
86
87static inline void conditional_sti(struct pt_regs *regs)
88{
89 if (regs->eflags & X86_EFLAGS_IF)
90 local_irq_enable();
91}
92
93static int kstack_depth_to_print = 10;
94
95#ifdef CONFIG_KALLSYMS
96#include <linux/kallsyms.h>
97int printk_address(unsigned long address)
98{
99 unsigned long offset = 0, symsize;
100 const char *symname;
101 char *modname;
102 char *delim = ":";
103 char namebuf[128];
104
105 symname = kallsyms_lookup(address, &symsize, &offset, &modname, namebuf);
106 if (!symname)
107 return printk("[<%016lx>]", address);
108 if (!modname)
109 modname = delim = "";
110 return printk("<%016lx>{%s%s%s%s%+ld}",
111 address,delim,modname,delim,symname,offset);
112}
113#else
114int printk_address(unsigned long address)
115{
116 return printk("[<%016lx>]", address);
117}
118#endif
119
Andi Kleen0a658002005-04-16 15:25:17 -0700120static unsigned long *in_exception_stack(unsigned cpu, unsigned long stack,
121 unsigned *usedp, const char **idp)
122{
Jan Beulichb556b352006-01-11 22:43:00 +0100123 static char ids[][8] = {
Andi Kleen0a658002005-04-16 15:25:17 -0700124 [DEBUG_STACK - 1] = "#DB",
125 [NMI_STACK - 1] = "NMI",
126 [DOUBLEFAULT_STACK - 1] = "#DF",
127 [STACKFAULT_STACK - 1] = "#SS",
128 [MCE_STACK - 1] = "#MC",
Jan Beulichb556b352006-01-11 22:43:00 +0100129#if DEBUG_STKSZ > EXCEPTION_STKSZ
130 [N_EXCEPTION_STACKS ... N_EXCEPTION_STACKS + DEBUG_STKSZ / EXCEPTION_STKSZ - 2] = "#DB[?]"
131#endif
Andi Kleen0a658002005-04-16 15:25:17 -0700132 };
133 unsigned k;
Linus Torvalds1da177e2005-04-16 15:20:36 -0700134
Andi Kleen0a658002005-04-16 15:25:17 -0700135 for (k = 0; k < N_EXCEPTION_STACKS; k++) {
136 unsigned long end;
137
Jan Beulichb556b352006-01-11 22:43:00 +0100138 switch (k + 1) {
139#if DEBUG_STKSZ > EXCEPTION_STKSZ
140 case DEBUG_STACK:
Ravikiran G Thirumalaidf79efd2006-01-11 22:45:39 +0100141 end = cpu_pda(cpu)->debugstack + DEBUG_STKSZ;
Jan Beulichb556b352006-01-11 22:43:00 +0100142 break;
143#endif
144 default:
145 end = per_cpu(init_tss, cpu).ist[k];
146 break;
147 }
Andi Kleen0a658002005-04-16 15:25:17 -0700148 if (stack >= end)
149 continue;
150 if (stack >= end - EXCEPTION_STKSZ) {
151 if (*usedp & (1U << k))
152 break;
153 *usedp |= 1U << k;
154 *idp = ids[k];
155 return (unsigned long *)end;
156 }
Jan Beulichb556b352006-01-11 22:43:00 +0100157#if DEBUG_STKSZ > EXCEPTION_STKSZ
158 if (k == DEBUG_STACK - 1 && stack >= end - DEBUG_STKSZ) {
159 unsigned j = N_EXCEPTION_STACKS - 1;
160
161 do {
162 ++j;
163 end -= EXCEPTION_STKSZ;
164 ids[j][4] = '1' + (j - N_EXCEPTION_STACKS);
165 } while (stack < end - EXCEPTION_STKSZ);
166 if (*usedp & (1U << j))
167 break;
168 *usedp |= 1U << j;
169 *idp = ids[j];
170 return (unsigned long *)end;
171 }
172#endif
Linus Torvalds1da177e2005-04-16 15:20:36 -0700173 }
174 return NULL;
Andi Kleen0a658002005-04-16 15:25:17 -0700175}
Linus Torvalds1da177e2005-04-16 15:20:36 -0700176
177/*
178 * x86-64 can have upto three kernel stacks:
179 * process stack
180 * interrupt stack
Andi Kleen0a658002005-04-16 15:25:17 -0700181 * severe exception (double fault, nmi, stack fault, debug, mce) hardware stack
Linus Torvalds1da177e2005-04-16 15:20:36 -0700182 */
183
184void show_trace(unsigned long *stack)
185{
Andi Kleen0a658002005-04-16 15:25:17 -0700186 const unsigned cpu = safe_smp_processor_id();
Ravikiran G Thirumalaidf79efd2006-01-11 22:45:39 +0100187 unsigned long *irqstack_end = (unsigned long *)cpu_pda(cpu)->irqstackptr;
Linus Torvalds1da177e2005-04-16 15:20:36 -0700188 int i;
Andi Kleen0a658002005-04-16 15:25:17 -0700189 unsigned used = 0;
Linus Torvalds1da177e2005-04-16 15:20:36 -0700190
191 printk("\nCall Trace:");
Andi Kleen0a658002005-04-16 15:25:17 -0700192
193#define HANDLE_STACK(cond) \
194 do while (cond) { \
Jan Beulich1b2f6302006-01-11 22:46:45 +0100195 unsigned long addr = *stack++; \
Andi Kleen0a658002005-04-16 15:25:17 -0700196 if (kernel_text_address(addr)) { \
Jan Beulich1b2f6302006-01-11 22:46:45 +0100197 if (i > 50) { \
198 printk("\n "); \
199 i = 0; \
200 } \
201 else \
202 i += printk(" "); \
Andi Kleen0a658002005-04-16 15:25:17 -0700203 /* \
204 * If the address is either in the text segment of the \
205 * kernel, or in the region which contains vmalloc'ed \
206 * memory, it *may* be the address of a calling \
207 * routine; if so, print it so that someone tracing \
208 * down the cause of the crash will be able to figure \
209 * out the call path that was taken. \
210 */ \
211 i += printk_address(addr); \
Andi Kleen0a658002005-04-16 15:25:17 -0700212 } \
213 } while (0)
214
Jan Beulich1b2f6302006-01-11 22:46:45 +0100215 for(i = 11; ; ) {
Andi Kleen0a658002005-04-16 15:25:17 -0700216 const char *id;
217 unsigned long *estack_end;
218 estack_end = in_exception_stack(cpu, (unsigned long)stack,
219 &used, &id);
220
221 if (estack_end) {
Jan Beulich1b2f6302006-01-11 22:46:45 +0100222 i += printk(" <%s>", id);
Andi Kleen0a658002005-04-16 15:25:17 -0700223 HANDLE_STACK (stack < estack_end);
Jan Beulich1b2f6302006-01-11 22:46:45 +0100224 i += printk(" <EOE>");
Andi Kleen0a658002005-04-16 15:25:17 -0700225 stack = (unsigned long *) estack_end[-2];
226 continue;
227 }
228 if (irqstack_end) {
229 unsigned long *irqstack;
230 irqstack = irqstack_end -
231 (IRQSTACKSIZE - 64) / sizeof(*irqstack);
232
233 if (stack >= irqstack && stack < irqstack_end) {
Jan Beulich1b2f6302006-01-11 22:46:45 +0100234 i += printk(" <IRQ>");
Andi Kleen0a658002005-04-16 15:25:17 -0700235 HANDLE_STACK (stack < irqstack_end);
236 stack = (unsigned long *) (irqstack_end[-1]);
237 irqstack_end = NULL;
Jan Beulich1b2f6302006-01-11 22:46:45 +0100238 i += printk(" <EOI>");
Andi Kleen0a658002005-04-16 15:25:17 -0700239 continue;
Linus Torvalds1da177e2005-04-16 15:20:36 -0700240 }
241 }
Andi Kleen0a658002005-04-16 15:25:17 -0700242 break;
Linus Torvalds1da177e2005-04-16 15:20:36 -0700243 }
Andi Kleen0a658002005-04-16 15:25:17 -0700244
245 HANDLE_STACK (((long) stack & (THREAD_SIZE-1)) != 0);
246#undef HANDLE_STACK
Linus Torvalds1da177e2005-04-16 15:20:36 -0700247 printk("\n");
248}
249
250void show_stack(struct task_struct *tsk, unsigned long * rsp)
251{
252 unsigned long *stack;
253 int i;
254 const int cpu = safe_smp_processor_id();
Ravikiran G Thirumalaidf79efd2006-01-11 22:45:39 +0100255 unsigned long *irqstack_end = (unsigned long *) (cpu_pda(cpu)->irqstackptr);
256 unsigned long *irqstack = (unsigned long *) (cpu_pda(cpu)->irqstackptr - IRQSTACKSIZE);
Linus Torvalds1da177e2005-04-16 15:20:36 -0700257
258 // debugging aid: "show_stack(NULL, NULL);" prints the
259 // back trace for this cpu.
260
261 if (rsp == NULL) {
262 if (tsk)
263 rsp = (unsigned long *)tsk->thread.rsp;
264 else
265 rsp = (unsigned long *)&rsp;
266 }
267
268 stack = rsp;
269 for(i=0; i < kstack_depth_to_print; i++) {
270 if (stack >= irqstack && stack <= irqstack_end) {
271 if (stack == irqstack_end) {
272 stack = (unsigned long *) (irqstack_end[-1]);
273 printk(" <EOI> ");
274 }
275 } else {
276 if (((long) stack & (THREAD_SIZE-1)) == 0)
277 break;
278 }
279 if (i && ((i % 4) == 0))
280 printk("\n ");
281 printk("%016lx ", *stack++);
akpm@osdl.org35faa712005-04-16 15:24:54 -0700282 touch_nmi_watchdog();
Linus Torvalds1da177e2005-04-16 15:20:36 -0700283 }
284 show_trace((unsigned long *)rsp);
285}
286
287/*
288 * The architecture-independent dump_stack generator
289 */
290void dump_stack(void)
291{
292 unsigned long dummy;
293 show_trace(&dummy);
294}
295
296EXPORT_SYMBOL(dump_stack);
297
298void show_registers(struct pt_regs *regs)
299{
300 int i;
Vincent Hanquez76381fe2005-06-23 00:08:46 -0700301 int in_kernel = !user_mode(regs);
Linus Torvalds1da177e2005-04-16 15:20:36 -0700302 unsigned long rsp;
303 const int cpu = safe_smp_processor_id();
Ravikiran G Thirumalaidf79efd2006-01-11 22:45:39 +0100304 struct task_struct *cur = cpu_pda(cpu)->pcurrent;
Linus Torvalds1da177e2005-04-16 15:20:36 -0700305
306 rsp = regs->rsp;
307
308 printk("CPU %d ", cpu);
309 __show_regs(regs);
310 printk("Process %s (pid: %d, threadinfo %p, task %p)\n",
Al Viroe4f17c42006-01-12 01:05:38 -0800311 cur->comm, cur->pid, task_thread_info(cur), cur);
Linus Torvalds1da177e2005-04-16 15:20:36 -0700312
313 /*
314 * When in-kernel, we also print out the stack and code at the
315 * time of the fault..
316 */
317 if (in_kernel) {
318
319 printk("Stack: ");
320 show_stack(NULL, (unsigned long*)rsp);
321
322 printk("\nCode: ");
323 if(regs->rip < PAGE_OFFSET)
324 goto bad;
325
326 for(i=0;i<20;i++)
327 {
328 unsigned char c;
329 if(__get_user(c, &((unsigned char*)regs->rip)[i])) {
330bad:
331 printk(" Bad RIP value.");
332 break;
333 }
334 printk("%02x ", c);
335 }
336 }
337 printk("\n");
338}
339
340void handle_BUG(struct pt_regs *regs)
341{
342 struct bug_frame f;
Jan Beulich5f1d1892006-01-11 22:46:48 +0100343 long len;
344 const char *prefix = "";
Linus Torvalds1da177e2005-04-16 15:20:36 -0700345
Vincent Hanquez76381fe2005-06-23 00:08:46 -0700346 if (user_mode(regs))
Linus Torvalds1da177e2005-04-16 15:20:36 -0700347 return;
Stephen Hemminger77a75332006-01-11 22:46:30 +0100348 if (__copy_from_user(&f, (const void __user *) regs->rip,
Linus Torvalds1da177e2005-04-16 15:20:36 -0700349 sizeof(struct bug_frame)))
350 return;
Jan Beulich049cdef2005-09-12 18:49:25 +0200351 if (f.filename >= 0 ||
Linus Torvalds1da177e2005-04-16 15:20:36 -0700352 f.ud2[0] != 0x0f || f.ud2[1] != 0x0b)
353 return;
Jan Beulich5f1d1892006-01-11 22:46:48 +0100354 len = __strnlen_user((char *)(long)f.filename, PATH_MAX) - 1;
355 if (len < 0 || len >= PATH_MAX)
Jan Beulich049cdef2005-09-12 18:49:25 +0200356 f.filename = (int)(long)"unmapped filename";
Jan Beulich5f1d1892006-01-11 22:46:48 +0100357 else if (len > 50) {
358 f.filename += len - 50;
359 prefix = "...";
360 }
Linus Torvalds1da177e2005-04-16 15:20:36 -0700361 printk("----------- [cut here ] --------- [please bite here ] ---------\n");
Jan Beulich5f1d1892006-01-11 22:46:48 +0100362 printk(KERN_ALERT "Kernel BUG at %s%.50s:%d\n", prefix, (char *)(long)f.filename, f.line);
Linus Torvalds1da177e2005-04-16 15:20:36 -0700363}
364
Alexander Nyberg4f60fdf2005-05-25 12:31:28 -0700365#ifdef CONFIG_BUG
Linus Torvalds1da177e2005-04-16 15:20:36 -0700366void out_of_line_bug(void)
367{
368 BUG();
369}
Alexander Nyberg4f60fdf2005-05-25 12:31:28 -0700370#endif
Linus Torvalds1da177e2005-04-16 15:20:36 -0700371
372static DEFINE_SPINLOCK(die_lock);
373static int die_owner = -1;
374
Andi Kleeneddb6fb2006-02-03 21:50:41 +0100375unsigned __kprobes long oops_begin(void)
Linus Torvalds1da177e2005-04-16 15:20:36 -0700376{
Jan Beulich12091402005-09-12 18:49:24 +0200377 int cpu = safe_smp_processor_id();
378 unsigned long flags;
379
380 /* racy, but better than risking deadlock. */
381 local_irq_save(flags);
Linus Torvalds1da177e2005-04-16 15:20:36 -0700382 if (!spin_trylock(&die_lock)) {
383 if (cpu == die_owner)
384 /* nested oops. should stop eventually */;
385 else
Jan Beulich12091402005-09-12 18:49:24 +0200386 spin_lock(&die_lock);
Linus Torvalds1da177e2005-04-16 15:20:36 -0700387 }
Jan Beulich12091402005-09-12 18:49:24 +0200388 die_owner = cpu;
Linus Torvalds1da177e2005-04-16 15:20:36 -0700389 console_verbose();
Jan Beulich12091402005-09-12 18:49:24 +0200390 bust_spinlocks(1);
391 return flags;
Linus Torvalds1da177e2005-04-16 15:20:36 -0700392}
393
Andi Kleeneddb6fb2006-02-03 21:50:41 +0100394void __kprobes oops_end(unsigned long flags)
Linus Torvalds1da177e2005-04-16 15:20:36 -0700395{
396 die_owner = -1;
Jan Beulich12091402005-09-12 18:49:24 +0200397 bust_spinlocks(0);
398 spin_unlock_irqrestore(&die_lock, flags);
Linus Torvalds1da177e2005-04-16 15:20:36 -0700399 if (panic_on_oops)
Jan Beulich12091402005-09-12 18:49:24 +0200400 panic("Oops");
401}
Linus Torvalds1da177e2005-04-16 15:20:36 -0700402
Andi Kleeneddb6fb2006-02-03 21:50:41 +0100403void __kprobes __die(const char * str, struct pt_regs * regs, long err)
Linus Torvalds1da177e2005-04-16 15:20:36 -0700404{
405 static int die_counter;
406 printk(KERN_EMERG "%s: %04lx [%u] ", str, err & 0xffff,++die_counter);
407#ifdef CONFIG_PREEMPT
408 printk("PREEMPT ");
409#endif
410#ifdef CONFIG_SMP
411 printk("SMP ");
412#endif
413#ifdef CONFIG_DEBUG_PAGEALLOC
414 printk("DEBUG_PAGEALLOC");
415#endif
416 printk("\n");
Jan Beulich6e3f3612006-01-11 22:42:14 +0100417 notify_die(DIE_OOPS, str, regs, err, current->thread.trap_no, SIGSEGV);
Linus Torvalds1da177e2005-04-16 15:20:36 -0700418 show_registers(regs);
419 /* Executive summary in case the oops scrolled away */
420 printk(KERN_ALERT "RIP ");
421 printk_address(regs->rip);
422 printk(" RSP <%016lx>\n", regs->rsp);
423}
424
425void die(const char * str, struct pt_regs * regs, long err)
426{
Jan Beulich12091402005-09-12 18:49:24 +0200427 unsigned long flags = oops_begin();
428
Linus Torvalds1da177e2005-04-16 15:20:36 -0700429 handle_BUG(regs);
430 __die(str, regs, err);
Jan Beulich12091402005-09-12 18:49:24 +0200431 oops_end(flags);
Linus Torvalds1da177e2005-04-16 15:20:36 -0700432 do_exit(SIGSEGV);
433}
Linus Torvalds1da177e2005-04-16 15:20:36 -0700434
Andi Kleeneddb6fb2006-02-03 21:50:41 +0100435void __kprobes die_nmi(char *str, struct pt_regs *regs)
Linus Torvalds1da177e2005-04-16 15:20:36 -0700436{
Jan Beulich12091402005-09-12 18:49:24 +0200437 unsigned long flags = oops_begin();
438
Linus Torvalds1da177e2005-04-16 15:20:36 -0700439 /*
440 * We are in trouble anyway, lets at least try
441 * to get a message out.
442 */
443 printk(str, safe_smp_processor_id());
444 show_registers(regs);
445 if (panic_on_timeout || panic_on_oops)
446 panic("nmi watchdog");
447 printk("console shuts up ...\n");
Jan Beulich12091402005-09-12 18:49:24 +0200448 oops_end(flags);
Linus Torvalds1da177e2005-04-16 15:20:36 -0700449 do_exit(SIGSEGV);
450}
451
Prasanna S Panchamukhi0f2fbdc2005-09-06 15:19:28 -0700452static void __kprobes do_trap(int trapnr, int signr, char *str,
453 struct pt_regs * regs, long error_code,
454 siginfo_t *info)
Linus Torvalds1da177e2005-04-16 15:20:36 -0700455{
Jan Beulich6e3f3612006-01-11 22:42:14 +0100456 struct task_struct *tsk = current;
457
Linus Torvalds1da177e2005-04-16 15:20:36 -0700458 conditional_sti(regs);
459
Jan Beulich6e3f3612006-01-11 22:42:14 +0100460 tsk->thread.error_code = error_code;
461 tsk->thread.trap_no = trapnr;
Linus Torvalds1da177e2005-04-16 15:20:36 -0700462
Jan Beulich6e3f3612006-01-11 22:42:14 +0100463 if (user_mode(regs)) {
Linus Torvalds1da177e2005-04-16 15:20:36 -0700464 if (exception_trace && unhandled_signal(tsk, signr))
465 printk(KERN_INFO
466 "%s[%d] trap %s rip:%lx rsp:%lx error:%lx\n",
467 tsk->comm, tsk->pid, str,
468 regs->rip,regs->rsp,error_code);
469
Linus Torvalds1da177e2005-04-16 15:20:36 -0700470 if (info)
471 force_sig_info(signr, info, tsk);
472 else
473 force_sig(signr, tsk);
474 return;
475 }
476
477
478 /* kernel trap */
479 {
480 const struct exception_table_entry *fixup;
481 fixup = search_exception_tables(regs->rip);
482 if (fixup) {
483 regs->rip = fixup->fixup;
484 } else
485 die(str, regs, error_code);
486 return;
487 }
488}
489
490#define DO_ERROR(trapnr, signr, str, name) \
491asmlinkage void do_##name(struct pt_regs * regs, long error_code) \
492{ \
493 if (notify_die(DIE_TRAP, str, regs, error_code, trapnr, signr) \
494 == NOTIFY_STOP) \
495 return; \
496 do_trap(trapnr, signr, str, regs, error_code, NULL); \
497}
498
499#define DO_ERROR_INFO(trapnr, signr, str, name, sicode, siaddr) \
500asmlinkage void do_##name(struct pt_regs * regs, long error_code) \
501{ \
502 siginfo_t info; \
503 info.si_signo = signr; \
504 info.si_errno = 0; \
505 info.si_code = sicode; \
506 info.si_addr = (void __user *)siaddr; \
507 if (notify_die(DIE_TRAP, str, regs, error_code, trapnr, signr) \
508 == NOTIFY_STOP) \
509 return; \
510 do_trap(trapnr, signr, str, regs, error_code, &info); \
511}
512
513DO_ERROR_INFO( 0, SIGFPE, "divide error", divide_error, FPE_INTDIV, regs->rip)
514DO_ERROR( 4, SIGSEGV, "overflow", overflow)
515DO_ERROR( 5, SIGSEGV, "bounds", bounds)
Chuck Ebbert100c0e32006-01-11 22:46:00 +0100516DO_ERROR_INFO( 6, SIGILL, "invalid opcode", invalid_op, ILL_ILLOPN, regs->rip)
Linus Torvalds1da177e2005-04-16 15:20:36 -0700517DO_ERROR( 7, SIGSEGV, "device not available", device_not_available)
518DO_ERROR( 9, SIGFPE, "coprocessor segment overrun", coprocessor_segment_overrun)
519DO_ERROR(10, SIGSEGV, "invalid TSS", invalid_TSS)
520DO_ERROR(11, SIGBUS, "segment not present", segment_not_present)
521DO_ERROR_INFO(17, SIGBUS, "alignment check", alignment_check, BUS_ADRALN, 0)
522DO_ERROR(18, SIGSEGV, "reserved", reserved)
Andi Kleen6fefb0d2005-04-16 15:25:03 -0700523DO_ERROR(12, SIGBUS, "stack segment", stack_segment)
Jan Beulicheca37c12006-01-11 22:42:17 +0100524
525asmlinkage void do_double_fault(struct pt_regs * regs, long error_code)
526{
527 static const char str[] = "double fault";
528 struct task_struct *tsk = current;
529
530 /* Return not checked because double check cannot be ignored */
531 notify_die(DIE_TRAP, str, regs, error_code, 8, SIGSEGV);
532
533 tsk->thread.error_code = error_code;
534 tsk->thread.trap_no = 8;
535
536 /* This is always a kernel trap and never fixable (and thus must
537 never return). */
538 for (;;)
539 die(str, regs, error_code);
540}
Linus Torvalds1da177e2005-04-16 15:20:36 -0700541
Prasanna S Panchamukhi0f2fbdc2005-09-06 15:19:28 -0700542asmlinkage void __kprobes do_general_protection(struct pt_regs * regs,
543 long error_code)
Linus Torvalds1da177e2005-04-16 15:20:36 -0700544{
Jan Beulich6e3f3612006-01-11 22:42:14 +0100545 struct task_struct *tsk = current;
546
Linus Torvalds1da177e2005-04-16 15:20:36 -0700547 conditional_sti(regs);
548
Jan Beulich6e3f3612006-01-11 22:42:14 +0100549 tsk->thread.error_code = error_code;
550 tsk->thread.trap_no = 13;
Linus Torvalds1da177e2005-04-16 15:20:36 -0700551
Jan Beulich6e3f3612006-01-11 22:42:14 +0100552 if (user_mode(regs)) {
Linus Torvalds1da177e2005-04-16 15:20:36 -0700553 if (exception_trace && unhandled_signal(tsk, SIGSEGV))
554 printk(KERN_INFO
555 "%s[%d] general protection rip:%lx rsp:%lx error:%lx\n",
556 tsk->comm, tsk->pid,
557 regs->rip,regs->rsp,error_code);
558
Linus Torvalds1da177e2005-04-16 15:20:36 -0700559 force_sig(SIGSEGV, tsk);
560 return;
561 }
562
563 /* kernel gp */
564 {
565 const struct exception_table_entry *fixup;
566 fixup = search_exception_tables(regs->rip);
567 if (fixup) {
568 regs->rip = fixup->fixup;
569 return;
570 }
571 if (notify_die(DIE_GPF, "general protection fault", regs,
572 error_code, 13, SIGSEGV) == NOTIFY_STOP)
573 return;
574 die("general protection fault", regs, error_code);
575 }
576}
577
Andi Kleeneddb6fb2006-02-03 21:50:41 +0100578static __kprobes void
579mem_parity_error(unsigned char reason, struct pt_regs * regs)
Linus Torvalds1da177e2005-04-16 15:20:36 -0700580{
581 printk("Uhhuh. NMI received. Dazed and confused, but trying to continue\n");
582 printk("You probably have a hardware problem with your RAM chips\n");
583
584 /* Clear and disable the memory parity error line. */
585 reason = (reason & 0xf) | 4;
586 outb(reason, 0x61);
587}
588
Andi Kleeneddb6fb2006-02-03 21:50:41 +0100589static __kprobes void
590io_check_error(unsigned char reason, struct pt_regs * regs)
Linus Torvalds1da177e2005-04-16 15:20:36 -0700591{
592 printk("NMI: IOCK error (debug interrupt?)\n");
593 show_registers(regs);
594
595 /* Re-enable the IOCK line, wait for a few seconds */
596 reason = (reason & 0xf) | 8;
597 outb(reason, 0x61);
598 mdelay(2000);
599 reason &= ~8;
600 outb(reason, 0x61);
601}
602
Andi Kleeneddb6fb2006-02-03 21:50:41 +0100603static __kprobes void
604unknown_nmi_error(unsigned char reason, struct pt_regs * regs)
Linus Torvalds1da177e2005-04-16 15:20:36 -0700605{ printk("Uhhuh. NMI received for unknown reason %02x.\n", reason);
606 printk("Dazed and confused, but trying to continue\n");
607 printk("Do you have a strange power saving mode enabled?\n");
608}
609
Andi Kleen6fefb0d2005-04-16 15:25:03 -0700610/* Runs on IST stack. This code must keep interrupts off all the time.
611 Nested NMIs are prevented by the CPU. */
Andi Kleeneddb6fb2006-02-03 21:50:41 +0100612asmlinkage __kprobes void default_do_nmi(struct pt_regs *regs)
Linus Torvalds1da177e2005-04-16 15:20:36 -0700613{
614 unsigned char reason = 0;
Ashok Raj76e4f662005-06-25 14:55:00 -0700615 int cpu;
616
617 cpu = smp_processor_id();
Linus Torvalds1da177e2005-04-16 15:20:36 -0700618
619 /* Only the BSP gets external NMIs from the system. */
Ashok Raj76e4f662005-06-25 14:55:00 -0700620 if (!cpu)
Linus Torvalds1da177e2005-04-16 15:20:36 -0700621 reason = get_nmi_reason();
622
623 if (!(reason & 0xc0)) {
Jan Beulich6e3f3612006-01-11 22:42:14 +0100624 if (notify_die(DIE_NMI_IPI, "nmi_ipi", regs, reason, 2, SIGINT)
Linus Torvalds1da177e2005-04-16 15:20:36 -0700625 == NOTIFY_STOP)
626 return;
627#ifdef CONFIG_X86_LOCAL_APIC
628 /*
629 * Ok, so this is none of the documented NMI sources,
630 * so it must be the NMI watchdog.
631 */
632 if (nmi_watchdog > 0) {
633 nmi_watchdog_tick(regs,reason);
634 return;
635 }
636#endif
637 unknown_nmi_error(reason, regs);
638 return;
639 }
Jan Beulich6e3f3612006-01-11 22:42:14 +0100640 if (notify_die(DIE_NMI, "nmi", regs, reason, 2, SIGINT) == NOTIFY_STOP)
Linus Torvalds1da177e2005-04-16 15:20:36 -0700641 return;
642
643 /* AK: following checks seem to be broken on modern chipsets. FIXME */
644
645 if (reason & 0x80)
646 mem_parity_error(reason, regs);
647 if (reason & 0x40)
648 io_check_error(reason, regs);
649}
650
Jan Beulichb556b352006-01-11 22:43:00 +0100651/* runs on IST stack. */
Prasanna S Panchamukhi0f2fbdc2005-09-06 15:19:28 -0700652asmlinkage void __kprobes do_int3(struct pt_regs * regs, long error_code)
Linus Torvalds1da177e2005-04-16 15:20:36 -0700653{
654 if (notify_die(DIE_INT3, "int3", regs, error_code, 3, SIGTRAP) == NOTIFY_STOP) {
655 return;
656 }
657 do_trap(3, SIGTRAP, "int3", regs, error_code, NULL);
658 return;
659}
660
Andi Kleen6fefb0d2005-04-16 15:25:03 -0700661/* Help handler running on IST stack to switch back to user stack
662 for scheduling or signal handling. The actual stack switch is done in
663 entry.S */
Andi Kleeneddb6fb2006-02-03 21:50:41 +0100664asmlinkage __kprobes struct pt_regs *sync_regs(struct pt_regs *eregs)
Linus Torvalds1da177e2005-04-16 15:20:36 -0700665{
Andi Kleen6fefb0d2005-04-16 15:25:03 -0700666 struct pt_regs *regs = eregs;
667 /* Did already sync */
668 if (eregs == (struct pt_regs *)eregs->rsp)
669 ;
670 /* Exception from user space */
Vincent Hanquez76381fe2005-06-23 00:08:46 -0700671 else if (user_mode(eregs))
Al Virobb049232006-01-12 01:05:38 -0800672 regs = task_pt_regs(current);
Andi Kleen6fefb0d2005-04-16 15:25:03 -0700673 /* Exception from kernel and interrupts are enabled. Move to
674 kernel process stack. */
675 else if (eregs->eflags & X86_EFLAGS_IF)
676 regs = (struct pt_regs *)(eregs->rsp -= sizeof(struct pt_regs));
677 if (eregs != regs)
678 *regs = *eregs;
679 return regs;
680}
681
682/* runs on IST stack. */
Prasanna S Panchamukhi0f2fbdc2005-09-06 15:19:28 -0700683asmlinkage void __kprobes do_debug(struct pt_regs * regs,
684 unsigned long error_code)
Andi Kleen6fefb0d2005-04-16 15:25:03 -0700685{
Linus Torvalds1da177e2005-04-16 15:20:36 -0700686 unsigned long condition;
687 struct task_struct *tsk = current;
688 siginfo_t info;
689
Vincent Hanqueze9129e52005-06-23 00:08:46 -0700690 get_debugreg(condition, 6);
Linus Torvalds1da177e2005-04-16 15:20:36 -0700691
692 if (notify_die(DIE_DEBUG, "debug", regs, condition, error_code,
Andi Kleendaeeafe2005-04-16 15:25:13 -0700693 SIGTRAP) == NOTIFY_STOP)
Andi Kleen6fefb0d2005-04-16 15:25:03 -0700694 return;
Andi Kleendaeeafe2005-04-16 15:25:13 -0700695
Linus Torvalds1da177e2005-04-16 15:20:36 -0700696 conditional_sti(regs);
697
698 /* Mask out spurious debug traps due to lazy DR7 setting */
699 if (condition & (DR_TRAP0|DR_TRAP1|DR_TRAP2|DR_TRAP3)) {
700 if (!tsk->thread.debugreg7) {
701 goto clear_dr7;
702 }
703 }
704
705 tsk->thread.debugreg6 = condition;
706
707 /* Mask out spurious TF errors due to lazy TF clearing */
Andi Kleendaeeafe2005-04-16 15:25:13 -0700708 if (condition & DR_STEP) {
Linus Torvalds1da177e2005-04-16 15:20:36 -0700709 /*
710 * The TF error should be masked out only if the current
711 * process is not traced and if the TRAP flag has been set
712 * previously by a tracing process (condition detected by
713 * the PT_DTRACE flag); remember that the i386 TRAP flag
714 * can be modified by the process itself in user mode,
715 * allowing programs to debug themselves without the ptrace()
716 * interface.
717 */
Vincent Hanquez76381fe2005-06-23 00:08:46 -0700718 if (!user_mode(regs))
Linus Torvalds1da177e2005-04-16 15:20:36 -0700719 goto clear_TF_reenable;
Andi Kleenbe61bff2005-04-16 15:24:57 -0700720 /*
721 * Was the TF flag set by a debugger? If so, clear it now,
722 * so that register information is correct.
723 */
724 if (tsk->ptrace & PT_DTRACE) {
725 regs->eflags &= ~TF_MASK;
726 tsk->ptrace &= ~PT_DTRACE;
727 }
Linus Torvalds1da177e2005-04-16 15:20:36 -0700728 }
729
730 /* Ok, finally something we can handle */
731 tsk->thread.trap_no = 1;
732 tsk->thread.error_code = error_code;
733 info.si_signo = SIGTRAP;
734 info.si_errno = 0;
735 info.si_code = TRAP_BRKPT;
John Blackwood01b8faa2006-01-11 22:44:15 +0100736 info.si_addr = user_mode(regs) ? (void __user *)regs->rip : NULL;
737 force_sig_info(SIGTRAP, &info, tsk);
Linus Torvalds1da177e2005-04-16 15:20:36 -0700738
Linus Torvalds1da177e2005-04-16 15:20:36 -0700739clear_dr7:
Vincent Hanqueze9129e52005-06-23 00:08:46 -0700740 set_debugreg(0UL, 7);
Andi Kleen6fefb0d2005-04-16 15:25:03 -0700741 return;
Linus Torvalds1da177e2005-04-16 15:20:36 -0700742
743clear_TF_reenable:
744 set_tsk_thread_flag(tsk, TIF_SINGLESTEP);
Linus Torvalds1da177e2005-04-16 15:20:36 -0700745 regs->eflags &= ~TF_MASK;
Linus Torvalds1da177e2005-04-16 15:20:36 -0700746}
747
Jan Beulich6e3f3612006-01-11 22:42:14 +0100748static int kernel_math_error(struct pt_regs *regs, const char *str, int trapnr)
Linus Torvalds1da177e2005-04-16 15:20:36 -0700749{
750 const struct exception_table_entry *fixup;
751 fixup = search_exception_tables(regs->rip);
752 if (fixup) {
753 regs->rip = fixup->fixup;
754 return 1;
755 }
Jan Beulich6e3f3612006-01-11 22:42:14 +0100756 notify_die(DIE_GPF, str, regs, 0, trapnr, SIGFPE);
Andi Kleen3a848f62005-04-16 15:25:06 -0700757 /* Illegal floating point operation in the kernel */
Jan Beulich6e3f3612006-01-11 22:42:14 +0100758 current->thread.trap_no = trapnr;
Linus Torvalds1da177e2005-04-16 15:20:36 -0700759 die(str, regs, 0);
Linus Torvalds1da177e2005-04-16 15:20:36 -0700760 return 0;
761}
762
763/*
764 * Note that we play around with the 'TS' bit in an attempt to get
765 * the correct behaviour even in the presence of the asynchronous
766 * IRQ13 behaviour
767 */
768asmlinkage void do_coprocessor_error(struct pt_regs *regs)
769{
770 void __user *rip = (void __user *)(regs->rip);
771 struct task_struct * task;
772 siginfo_t info;
773 unsigned short cwd, swd;
774
775 conditional_sti(regs);
Vincent Hanquez76381fe2005-06-23 00:08:46 -0700776 if (!user_mode(regs) &&
Jan Beulich6e3f3612006-01-11 22:42:14 +0100777 kernel_math_error(regs, "kernel x87 math error", 16))
Linus Torvalds1da177e2005-04-16 15:20:36 -0700778 return;
779
780 /*
781 * Save the info for the exception handler and clear the error.
782 */
783 task = current;
784 save_init_fpu(task);
785 task->thread.trap_no = 16;
786 task->thread.error_code = 0;
787 info.si_signo = SIGFPE;
788 info.si_errno = 0;
789 info.si_code = __SI_FAULT;
790 info.si_addr = rip;
791 /*
792 * (~cwd & swd) will mask out exceptions that are not set to unmasked
793 * status. 0x3f is the exception bits in these regs, 0x200 is the
794 * C1 reg you need in case of a stack fault, 0x040 is the stack
795 * fault bit. We should only be taking one exception at a time,
796 * so if this combination doesn't produce any single exception,
797 * then we have a bad program that isn't synchronizing its FPU usage
798 * and it will suffer the consequences since we won't be able to
799 * fully reproduce the context of the exception
800 */
801 cwd = get_fpu_cwd(task);
802 swd = get_fpu_swd(task);
Chuck Ebbertff347b22005-09-12 18:49:25 +0200803 switch (swd & ~cwd & 0x3f) {
Linus Torvalds1da177e2005-04-16 15:20:36 -0700804 case 0x000:
805 default:
806 break;
807 case 0x001: /* Invalid Op */
Chuck Ebbertff347b22005-09-12 18:49:25 +0200808 /*
809 * swd & 0x240 == 0x040: Stack Underflow
810 * swd & 0x240 == 0x240: Stack Overflow
811 * User must clear the SF bit (0x40) if set
812 */
Linus Torvalds1da177e2005-04-16 15:20:36 -0700813 info.si_code = FPE_FLTINV;
814 break;
815 case 0x002: /* Denormalize */
816 case 0x010: /* Underflow */
817 info.si_code = FPE_FLTUND;
818 break;
819 case 0x004: /* Zero Divide */
820 info.si_code = FPE_FLTDIV;
821 break;
822 case 0x008: /* Overflow */
823 info.si_code = FPE_FLTOVF;
824 break;
825 case 0x020: /* Precision */
826 info.si_code = FPE_FLTRES;
827 break;
828 }
829 force_sig_info(SIGFPE, &info, task);
830}
831
832asmlinkage void bad_intr(void)
833{
834 printk("bad interrupt");
835}
836
837asmlinkage void do_simd_coprocessor_error(struct pt_regs *regs)
838{
839 void __user *rip = (void __user *)(regs->rip);
840 struct task_struct * task;
841 siginfo_t info;
842 unsigned short mxcsr;
843
844 conditional_sti(regs);
Vincent Hanquez76381fe2005-06-23 00:08:46 -0700845 if (!user_mode(regs) &&
Jan Beulich6e3f3612006-01-11 22:42:14 +0100846 kernel_math_error(regs, "kernel simd math error", 19))
Linus Torvalds1da177e2005-04-16 15:20:36 -0700847 return;
848
849 /*
850 * Save the info for the exception handler and clear the error.
851 */
852 task = current;
853 save_init_fpu(task);
854 task->thread.trap_no = 19;
855 task->thread.error_code = 0;
856 info.si_signo = SIGFPE;
857 info.si_errno = 0;
858 info.si_code = __SI_FAULT;
859 info.si_addr = rip;
860 /*
861 * The SIMD FPU exceptions are handled a little differently, as there
862 * is only a single status/control register. Thus, to determine which
863 * unmasked exception was caught we must mask the exception mask bits
864 * at 0x1f80, and then use these to mask the exception bits at 0x3f.
865 */
866 mxcsr = get_fpu_mxcsr(task);
867 switch (~((mxcsr & 0x1f80) >> 7) & (mxcsr & 0x3f)) {
868 case 0x000:
869 default:
870 break;
871 case 0x001: /* Invalid Op */
872 info.si_code = FPE_FLTINV;
873 break;
874 case 0x002: /* Denormalize */
875 case 0x010: /* Underflow */
876 info.si_code = FPE_FLTUND;
877 break;
878 case 0x004: /* Zero Divide */
879 info.si_code = FPE_FLTDIV;
880 break;
881 case 0x008: /* Overflow */
882 info.si_code = FPE_FLTOVF;
883 break;
884 case 0x020: /* Precision */
885 info.si_code = FPE_FLTRES;
886 break;
887 }
888 force_sig_info(SIGFPE, &info, task);
889}
890
891asmlinkage void do_spurious_interrupt_bug(struct pt_regs * regs)
892{
893}
894
895asmlinkage void __attribute__((weak)) smp_thermal_interrupt(void)
896{
897}
898
Jacob Shin89b831e2005-11-05 17:25:53 +0100899asmlinkage void __attribute__((weak)) mce_threshold_interrupt(void)
900{
901}
902
Linus Torvalds1da177e2005-04-16 15:20:36 -0700903/*
904 * 'math_state_restore()' saves the current math information in the
905 * old math state array, and gets the new ones from the current task
906 *
907 * Careful.. There are problems with IBM-designed IRQ13 behaviour.
908 * Don't touch unless you *really* know how it works.
909 */
910asmlinkage void math_state_restore(void)
911{
912 struct task_struct *me = current;
913 clts(); /* Allow maths ops (or we recurse) */
914
915 if (!used_math())
916 init_fpu(me);
917 restore_fpu_checking(&me->thread.i387.fxsave);
Al Viroe4f17c42006-01-12 01:05:38 -0800918 task_thread_info(me)->status |= TS_USEDFPU;
Linus Torvalds1da177e2005-04-16 15:20:36 -0700919}
920
Linus Torvalds1da177e2005-04-16 15:20:36 -0700921void __init trap_init(void)
922{
923 set_intr_gate(0,&divide_error);
924 set_intr_gate_ist(1,&debug,DEBUG_STACK);
925 set_intr_gate_ist(2,&nmi,NMI_STACK);
Jan Beulichb556b352006-01-11 22:43:00 +0100926 set_system_gate_ist(3,&int3,DEBUG_STACK); /* int3 can be called from all */
Jan Beulich0a521582006-01-11 22:42:08 +0100927 set_system_gate(4,&overflow); /* int4 can be called from all */
928 set_intr_gate(5,&bounds);
Linus Torvalds1da177e2005-04-16 15:20:36 -0700929 set_intr_gate(6,&invalid_op);
930 set_intr_gate(7,&device_not_available);
931 set_intr_gate_ist(8,&double_fault, DOUBLEFAULT_STACK);
932 set_intr_gate(9,&coprocessor_segment_overrun);
933 set_intr_gate(10,&invalid_TSS);
934 set_intr_gate(11,&segment_not_present);
935 set_intr_gate_ist(12,&stack_segment,STACKFAULT_STACK);
936 set_intr_gate(13,&general_protection);
937 set_intr_gate(14,&page_fault);
938 set_intr_gate(15,&spurious_interrupt_bug);
939 set_intr_gate(16,&coprocessor_error);
940 set_intr_gate(17,&alignment_check);
941#ifdef CONFIG_X86_MCE
942 set_intr_gate_ist(18,&machine_check, MCE_STACK);
943#endif
944 set_intr_gate(19,&simd_coprocessor_error);
945
946#ifdef CONFIG_IA32_EMULATION
947 set_system_gate(IA32_SYSCALL_VECTOR, ia32_syscall);
948#endif
949
Linus Torvalds1da177e2005-04-16 15:20:36 -0700950 /*
951 * Should be a barrier for any external CPU state.
952 */
953 cpu_init();
954}
955
956
957/* Actual parsing is done early in setup.c. */
958static int __init oops_dummy(char *s)
959{
960 panic_on_oops = 1;
961 return -1;
962}
963__setup("oops=", oops_dummy);
964
965static int __init kstack_setup(char *s)
966{
967 kstack_depth_to_print = simple_strtoul(s,NULL,0);
968 return 0;
969}
970__setup("kstack=", kstack_setup);
971