Blame SOURCES/gcc48-rh1469697-14.patch

56d343
commit 21397732bbcef3347c0d5ff8a0ee5163e803e2fb
56d343
Author: Jeff Law <law@redhat.com>
56d343
Date:   Mon Oct 2 12:30:26 2017 -0600
56d343
56d343
    Dependencies for aarch64 work
56d343
56d343
diff --git a/gcc/config/aarch64/aarch64-protos.h b/gcc/config/aarch64/aarch64-protos.h
56d343
index 07ff7031b35..91dd5b7fc02 100644
56d343
--- a/gcc/config/aarch64/aarch64-protos.h
56d343
+++ b/gcc/config/aarch64/aarch64-protos.h
56d343
@@ -181,6 +181,7 @@ unsigned aarch64_dbx_register_number (unsigned);
56d343
 unsigned aarch64_trampoline_size (void);
56d343
 void aarch64_asm_output_labelref (FILE *, const char *);
56d343
 void aarch64_elf_asm_named_section (const char *, unsigned, tree);
56d343
+const char * aarch64_output_probe_stack_range (rtx, rtx);
56d343
 void aarch64_expand_epilogue (bool);
56d343
 void aarch64_expand_mov_immediate (rtx, rtx);
56d343
 void aarch64_expand_prologue (void);
56d343
diff --git a/gcc/config/aarch64/aarch64.c b/gcc/config/aarch64/aarch64.c
56d343
index 5afc167d569..cadf193cfcf 100644
56d343
--- a/gcc/config/aarch64/aarch64.c
56d343
+++ b/gcc/config/aarch64/aarch64.c
56d343
@@ -969,6 +969,199 @@ aarch64_function_ok_for_sibcall (tree decl, tree exp ATTRIBUTE_UNUSED)
56d343
   return true;
56d343
 }
56d343
 
56d343
+static int
56d343
+aarch64_internal_mov_immediate (rtx dest, rtx imm, bool generate,
56d343
+				enum machine_mode mode)
56d343
+{
56d343
+  int i;
56d343
+  unsigned HOST_WIDE_INT val, val2, mask;
56d343
+  int one_match, zero_match;
56d343
+  int num_insns;
56d343
+
56d343
+  val = INTVAL (imm);
56d343
+
56d343
+  if (aarch64_move_imm (val, mode))
56d343
+    {
56d343
+      if (generate)
56d343
+	emit_insn (gen_rtx_SET (VOIDmode, dest, imm));
56d343
+      return 1;
56d343
+    }
56d343
+
56d343
+  /* Check to see if the low 32 bits are either 0xffffXXXX or 0xXXXXffff
56d343
+     (with XXXX non-zero). In that case check to see if the move can be done in
56d343
+     a smaller mode.  */
56d343
+  val2 = val & 0xffffffff;
56d343
+  if (mode == DImode
56d343
+      && aarch64_move_imm (val2, SImode)
56d343
+      && (((val >> 32) & 0xffff) == 0 || (val >> 48) == 0))
56d343
+    {
56d343
+      if (generate)
56d343
+	emit_insn (gen_rtx_SET (VOIDmode, dest, GEN_INT (val2)));
56d343
+
56d343
+      /* Check if we have to emit a second instruction by checking to see
56d343
+         if any of the upper 32 bits of the original DI mode value is set.  */
56d343
+      if (val == val2)
56d343
+	return 1;
56d343
+
56d343
+      i = (val >> 48) ? 48 : 32;
56d343
+
56d343
+      if (generate)
56d343
+	 emit_insn (gen_insv_immdi (dest, GEN_INT (i),
56d343
+				    GEN_INT ((val >> i) & 0xffff)));
56d343
+
56d343
+      return 2;
56d343
+    }
56d343
+
56d343
+  if ((val >> 32) == 0 || mode == SImode)
56d343
+    {
56d343
+      if (generate)
56d343
+	{
56d343
+	  emit_insn (gen_rtx_SET (VOIDmode, dest, GEN_INT (val & 0xffff)));
56d343
+	  if (mode == SImode)
56d343
+	    emit_insn (gen_insv_immsi (dest, GEN_INT (16),
56d343
+				       GEN_INT ((val >> 16) & 0xffff)));
56d343
+	  else
56d343
+	    emit_insn (gen_insv_immdi (dest, GEN_INT (16),
56d343
+				       GEN_INT ((val >> 16) & 0xffff)));
56d343
+	}
56d343
+      return 2;
56d343
+    }
56d343
+
56d343
+  /* Remaining cases are all for DImode.  */
56d343
+
56d343
+  mask = 0xffff;
56d343
+  zero_match = ((val & mask) == 0) + ((val & (mask << 16)) == 0) +
56d343
+    ((val & (mask << 32)) == 0) + ((val & (mask << 48)) == 0);
56d343
+  one_match = ((~val & mask) == 0) + ((~val & (mask << 16)) == 0) +
56d343
+    ((~val & (mask << 32)) == 0) + ((~val & (mask << 48)) == 0);
56d343
+
56d343
+  if (zero_match != 2 && one_match != 2)
56d343
+    {
56d343
+      /* Try emitting a bitmask immediate with a movk replacing 16 bits.
56d343
+	 For a 64-bit bitmask try whether changing 16 bits to all ones or
56d343
+	 zeroes creates a valid bitmask.  To check any repeated bitmask,
56d343
+	 try using 16 bits from the other 32-bit half of val.  */
56d343
+
56d343
+      for (i = 0; i < 64; i += 16, mask <<= 16)
56d343
+	{
56d343
+	  val2 = val & ~mask;
56d343
+	  if (val2 != val && aarch64_bitmask_imm (val2, mode))
56d343
+	    break;
56d343
+	  val2 = val | mask;
56d343
+	  if (val2 != val && aarch64_bitmask_imm (val2, mode))
56d343
+	    break;
56d343
+	  val2 = val2 & ~mask;
56d343
+	  val2 = val2 | (((val2 >> 32) | (val2 << 32)) & mask);
56d343
+	  if (val2 != val && aarch64_bitmask_imm (val2, mode))
56d343
+	    break;
56d343
+	}
56d343
+      if (i != 64)
56d343
+	{
56d343
+	  if (generate)
56d343
+	    {
56d343
+	      emit_insn (gen_rtx_SET (VOIDmode, dest, GEN_INT (val2)));
56d343
+	      emit_insn (gen_insv_immdi (dest, GEN_INT (i),
56d343
+					 GEN_INT ((val >> i) & 0xffff)));
56d343
+	    }
56d343
+	  return 2;
56d343
+	}
56d343
+    }
56d343
+
56d343
+  /* Generate 2-4 instructions, skipping 16 bits of all zeroes or ones which
56d343
+     are emitted by the initial mov.  If one_match > zero_match, skip set bits,
56d343
+     otherwise skip zero bits.  */
56d343
+
56d343
+  num_insns = 1;
56d343
+  mask = 0xffff;
56d343
+  val2 = one_match > zero_match ? ~val : val;
56d343
+  i = (val2 & mask) != 0 ? 0 : (val2 & (mask << 16)) != 0 ? 16 : 32;
56d343
+
56d343
+  if (generate)
56d343
+    emit_insn (gen_rtx_SET (VOIDmode, dest, GEN_INT (one_match > zero_match
56d343
+					   ? (val | ~(mask << i))
56d343
+					   : (val & (mask << i)))));
56d343
+  for (i += 16; i < 64; i += 16)
56d343
+    {
56d343
+      if ((val2 & (mask << i)) == 0)
56d343
+	continue;
56d343
+      if (generate)
56d343
+	emit_insn (gen_insv_immdi (dest, GEN_INT (i),
56d343
+				   GEN_INT ((val >> i) & 0xffff)));
56d343
+      num_insns ++;
56d343
+    }
56d343
+
56d343
+  return num_insns;
56d343
+}
56d343
+
56d343
+/* Add DELTA to REGNUM in mode MODE.  SCRATCHREG can be used to hold a
56d343
+   temporary value if necessary.  FRAME_RELATED_P should be true if
56d343
+   the RTX_FRAME_RELATED flag should be set and CFA adjustments added
56d343
+   to the generated instructions.  If SCRATCHREG is known to hold
56d343
+   abs (delta), EMIT_MOVE_IMM can be set to false to avoid emitting the
56d343
+   immediate again.
56d343
+
56d343
+   Since this function may be used to adjust the stack pointer, we must
56d343
+   ensure that it cannot cause transient stack deallocation (for example
56d343
+   by first incrementing SP and then decrementing when adjusting by a
56d343
+   large immediate).  */
56d343
+
56d343
+static void
56d343
+aarch64_add_constant_internal (enum machine_mode mode, int regnum,
56d343
+			       int scratchreg, HOST_WIDE_INT delta,
56d343
+			       bool frame_related_p, bool emit_move_imm)
56d343
+{
56d343
+  HOST_WIDE_INT mdelta = abs_hwi (delta);
56d343
+  rtx this_rtx = gen_rtx_REG (mode, regnum);
56d343
+  rtx insn;
56d343
+
56d343
+  if (!mdelta)
56d343
+    return;
56d343
+
56d343
+  /* Single instruction adjustment.  */
56d343
+  if (aarch64_uimm12_shift (mdelta))
56d343
+    {
56d343
+      insn = emit_insn (gen_add2_insn (this_rtx, GEN_INT (delta)));
56d343
+      RTX_FRAME_RELATED_P (insn) = frame_related_p;
56d343
+      return;
56d343
+    }
56d343
+
56d343
+  /* Emit 2 additions/subtractions if the adjustment is less than 24 bits.
56d343
+     Only do this if mdelta is not a 16-bit move as adjusting using a move
56d343
+     is better.  */
56d343
+  if (mdelta < 0x1000000 && !aarch64_move_imm (mdelta, mode))
56d343
+    {
56d343
+      HOST_WIDE_INT low_off = mdelta & 0xfff;
56d343
+
56d343
+      low_off = delta < 0 ? -low_off : low_off;
56d343
+      insn = emit_insn (gen_add2_insn (this_rtx, GEN_INT (low_off)));
56d343
+      RTX_FRAME_RELATED_P (insn) = frame_related_p;
56d343
+      insn = emit_insn (gen_add2_insn (this_rtx, GEN_INT (delta - low_off)));
56d343
+      RTX_FRAME_RELATED_P (insn) = frame_related_p;
56d343
+      return;
56d343
+    }
56d343
+
56d343
+  /* Emit a move immediate if required and an addition/subtraction.  */
56d343
+  rtx scratch_rtx = gen_rtx_REG (mode, scratchreg);
56d343
+  if (emit_move_imm)
56d343
+    aarch64_internal_mov_immediate (scratch_rtx, GEN_INT (mdelta), true, mode);
56d343
+  insn = emit_insn (delta < 0 ? gen_sub2_insn (this_rtx, scratch_rtx)
56d343
+			      : gen_add2_insn (this_rtx, scratch_rtx));
56d343
+  if (frame_related_p)
56d343
+    {
56d343
+      RTX_FRAME_RELATED_P (insn) = frame_related_p;
56d343
+      rtx adj = plus_constant (mode, this_rtx, delta);
56d343
+      add_reg_note (insn , REG_CFA_ADJUST_CFA,
56d343
+		    gen_rtx_SET (VOIDmode, this_rtx, adj));
56d343
+    }
56d343
+}
56d343
+
56d343
+static inline void
56d343
+aarch64_sub_sp (int scratchreg, HOST_WIDE_INT delta, bool frame_related_p)
56d343
+{
56d343
+  aarch64_add_constant_internal (Pmode, SP_REGNUM, scratchreg, -delta,
56d343
+				 frame_related_p, true);
56d343
+}
56d343
+
56d343
 /* Implement TARGET_PASS_BY_REFERENCE.  */
56d343
 
56d343
 static bool
56d343
@@ -1476,6 +1669,47 @@ aarch64_libgcc_cmp_return_mode (void)
56d343
   return SImode;
56d343
 }
56d343
 
56d343
+#define PROBE_INTERVAL (1 << STACK_CHECK_PROBE_INTERVAL_EXP)
56d343
+
56d343
+/* We use the 12-bit shifted immediate arithmetic instructions so values
56d343
+   must be multiple of (1 << 12), i.e. 4096.  */
56d343
+#define ARITH_FACTOR 4096
56d343
+
56d343
+/* Probe a range of stack addresses from REG1 to REG2 inclusive.  These are
56d343
+   absolute addresses.  */
56d343
+
56d343
+const char *
56d343
+aarch64_output_probe_stack_range (rtx reg1, rtx reg2)
56d343
+{
56d343
+  static int labelno = 0;
56d343
+  char loop_lab[32];
56d343
+  rtx xops[2];
56d343
+
56d343
+  ASM_GENERATE_INTERNAL_LABEL (loop_lab, "LPSRL", labelno++);
56d343
+
56d343
+  /* Loop.  */
56d343
+  ASM_OUTPUT_INTERNAL_LABEL (asm_out_file, loop_lab);
56d343
+
56d343
+  /* TEST_ADDR = TEST_ADDR + PROBE_INTERVAL.  */
56d343
+  xops[0] = reg1;
56d343
+  xops[1] = GEN_INT (PROBE_INTERVAL);
56d343
+  output_asm_insn ("sub\t%0, %0, %1", xops);
56d343
+
56d343
+  /* Probe at TEST_ADDR.  */
56d343
+  output_asm_insn ("str\txzr, [%0]", xops);
56d343
+
56d343
+  /* Test if TEST_ADDR == LAST_ADDR.  */
56d343
+  xops[1] = reg2;
56d343
+  output_asm_insn ("cmp\t%0, %1", xops);
56d343
+
56d343
+  /* Branch.  */
56d343
+  fputs ("\tb.ne\t", asm_out_file);
56d343
+  assemble_name_raw (asm_out_file, loop_lab);
56d343
+  fputc ('\n', asm_out_file);
56d343
+
56d343
+  return "";
56d343
+}
56d343
+
56d343
 static bool
56d343
 aarch64_frame_pointer_required (void)
56d343
 {
56d343
diff --git a/gcc/config/aarch64/aarch64.md b/gcc/config/aarch64/aarch64.md
56d343
index 91299901bbf..17082486ac8 100644
56d343
--- a/gcc/config/aarch64/aarch64.md
56d343
+++ b/gcc/config/aarch64/aarch64.md
56d343
@@ -88,6 +88,7 @@
56d343
     UNSPEC_ST4
56d343
     UNSPEC_TLS
56d343
     UNSPEC_TLSDESC
56d343
+    UNSPECV_PROBE_STACK_RANGE   ; Represent stack range probing.
56d343
     UNSPEC_VSTRUCTDUMMY
56d343
 ])
56d343
 
56d343
@@ -3399,6 +3400,18 @@
56d343
   [(set_attr "length" "0")]
56d343
 )
56d343
 
56d343
+(define_insn "probe_stack_range"
56d343
+  [(set (match_operand:DI 0 "register_operand" "=r")
56d343
+	(unspec_volatile:DI [(match_operand:DI 1 "register_operand" "0")
56d343
+			     (match_operand:DI 2 "register_operand" "r")]
56d343
+			      UNSPECV_PROBE_STACK_RANGE))]
56d343
+  ""
56d343
+{
56d343
+  return aarch64_output_probe_stack_range (operands[0], operands[2]);
56d343
+}
56d343
+  [(set_attr "length" "32")]
56d343
+)
56d343
+
56d343
 ;; Named pattern for expanding thread pointer reference.
56d343
 (define_expand "get_thread_pointerdi"
56d343
   [(match_operand:DI 0 "register_operand" "=r")]