9977e886cb
Improve the save and restore behavior of FPU register contents to use the vector extension within the kernel. The kernel does not use floating-point or vector registers and, therefore, saving and restoring the FPU register contents are performed for handling signals or switching processes only. To prepare for using vector instructions and vector registers within the kernel, enhance the save behavior and implement a lazy restore at return to user space from a system call or interrupt. To implement the lazy restore, the save_fpu_regs() sets a CPU information flag, CIF_FPU, to indicate that the FPU registers must be restored. Saving and setting CIF_FPU is performed in an atomic fashion to be interrupt-safe. When the kernel wants to use the vector extension or wants to change the FPU register state for a task during signal handling, the save_fpu_regs() must be called first. The CIF_FPU flag is also set at process switch. At return to user space, the FPU state is restored. In particular, the FPU state includes the floating-point or vector register contents, as well as, vector-enablement and floating-point control. The FPU state restore and clearing CIF_FPU is also performed in an atomic fashion. For KVM, the restore of the FPU register state is performed when restoring the general-purpose guest registers before the SIE instructions is started. Because the path towards the SIE instruction is interruptible, the CIF_FPU flag must be checked again right before going into SIE. If set, the guest registers must be reloaded again by re-entering the outer SIE loop. This is the same behavior as if the SIE critical section is interrupted. Signed-off-by: Hendrik Brueckner <brueckner@linux.vnet.ibm.com> Signed-off-by: Martin Schwidefsky <schwidefsky@de.ibm.com>
78 lines
1.9 KiB
C
78 lines
1.9 KiB
C
/*
|
|
* Copyright IBM Corp. 1999, 2009
|
|
*
|
|
* Author(s): Martin Schwidefsky <schwidefsky@de.ibm.com>
|
|
*/
|
|
|
|
#ifndef __ASM_CTL_REG_H
|
|
#define __ASM_CTL_REG_H
|
|
|
|
#include <linux/bug.h>
|
|
|
|
#define __ctl_load(array, low, high) { \
|
|
typedef struct { char _[sizeof(array)]; } addrtype; \
|
|
\
|
|
BUILD_BUG_ON(sizeof(addrtype) != (high - low + 1) * sizeof(long));\
|
|
asm volatile( \
|
|
" lctlg %1,%2,%0\n" \
|
|
: : "Q" (*(addrtype *)(&array)), "i" (low), "i" (high));\
|
|
}
|
|
|
|
#define __ctl_store(array, low, high) { \
|
|
typedef struct { char _[sizeof(array)]; } addrtype; \
|
|
\
|
|
BUILD_BUG_ON(sizeof(addrtype) != (high - low + 1) * sizeof(long));\
|
|
asm volatile( \
|
|
" stctg %1,%2,%0\n" \
|
|
: "=Q" (*(addrtype *)(&array)) \
|
|
: "i" (low), "i" (high)); \
|
|
}
|
|
|
|
static inline void __ctl_set_bit(unsigned int cr, unsigned int bit)
|
|
{
|
|
unsigned long reg;
|
|
|
|
__ctl_store(reg, cr, cr);
|
|
reg |= 1UL << bit;
|
|
__ctl_load(reg, cr, cr);
|
|
}
|
|
|
|
static inline void __ctl_clear_bit(unsigned int cr, unsigned int bit)
|
|
{
|
|
unsigned long reg;
|
|
|
|
__ctl_store(reg, cr, cr);
|
|
reg &= ~(1UL << bit);
|
|
__ctl_load(reg, cr, cr);
|
|
}
|
|
|
|
void __ctl_set_vx(void);
|
|
|
|
void smp_ctl_set_bit(int cr, int bit);
|
|
void smp_ctl_clear_bit(int cr, int bit);
|
|
|
|
union ctlreg0 {
|
|
unsigned long val;
|
|
struct {
|
|
unsigned long : 32;
|
|
unsigned long : 3;
|
|
unsigned long lap : 1; /* Low-address-protection control */
|
|
unsigned long : 4;
|
|
unsigned long edat : 1; /* Enhanced-DAT-enablement control */
|
|
unsigned long : 4;
|
|
unsigned long afp : 1; /* AFP-register control */
|
|
unsigned long vx : 1; /* Vector enablement control */
|
|
unsigned long : 17;
|
|
};
|
|
};
|
|
|
|
#ifdef CONFIG_SMP
|
|
# define ctl_set_bit(cr, bit) smp_ctl_set_bit(cr, bit)
|
|
# define ctl_clear_bit(cr, bit) smp_ctl_clear_bit(cr, bit)
|
|
#else
|
|
# define ctl_set_bit(cr, bit) __ctl_set_bit(cr, bit)
|
|
# define ctl_clear_bit(cr, bit) __ctl_clear_bit(cr, bit)
|
|
#endif
|
|
|
|
#endif /* __ASM_CTL_REG_H */
|