Blame sysdeps/sparc/sparc64/multiarch/memset-niagara4.S

Packit 6c4009
/* Set a block of memory to some byte value.  For SUN4V Niagara-4.
Packit 6c4009
   Copyright (C) 2012-2018 Free Software Foundation, Inc.
Packit 6c4009
   This file is part of the GNU C Library.
Packit 6c4009
   Contributed by David S. Miller (davem@davemloft.net)
Packit 6c4009
Packit 6c4009
   The GNU C Library is free software; you can redistribute it and/or
Packit 6c4009
   modify it under the terms of the GNU Lesser General Public
Packit 6c4009
   License as published by the Free Software Foundation; either
Packit 6c4009
   version 2.1 of the License, or (at your option) any later version.
Packit 6c4009
Packit 6c4009
   The GNU C Library is distributed in the hope that it will be useful,
Packit 6c4009
   but WITHOUT ANY WARRANTY; without even the implied warranty of
Packit 6c4009
   MERCHANTABILITY or FITNESS FOR A PARTICULAR PURPOSE.  See the GNU
Packit 6c4009
   Lesser General Public License for more details.
Packit 6c4009
Packit 6c4009
   You should have received a copy of the GNU Lesser General Public
Packit 6c4009
   License along with the GNU C Library; if not, see
Packit 6c4009
   <http://www.gnu.org/licenses/>.  */
Packit 6c4009
Packit 6c4009
#include <sysdep.h>
Packit 6c4009
Packit 6c4009
#define ASI_BLK_INIT_QUAD_LDD_P	0xe2
Packit 6c4009
Packit 6c4009
#if IS_IN (libc)
Packit 6c4009
Packit 6c4009
	.register	%g2, #scratch
Packit 6c4009
	.register	%g3, #scratch
Packit 6c4009
Packit 6c4009
	.text
Packit 6c4009
	.align		32
Packit 6c4009
Packit 6c4009
ENTRY(__memset_niagara4)
Packit 6c4009
	andcc		%o1, 0xff, %o4
Packit 6c4009
	be,pt		%icc, 1f
Packit 6c4009
	 mov		%o2, %o1
Packit 6c4009
	sllx		%o4, 8, %g1
Packit 6c4009
	or		%g1, %o4, %o2
Packit 6c4009
	sllx		%o2, 16, %g1
Packit 6c4009
	or		%g1, %o2, %o2
Packit 6c4009
	sllx		%o2, 32, %g1
Packit 6c4009
	ba,pt		%icc, 1f
Packit 6c4009
	 or		%g1, %o2, %o4
Packit 6c4009
END(__memset_niagara4)
Packit 6c4009
Packit 6c4009
	.align		32
Packit 6c4009
ENTRY(__bzero_niagara4)
Packit 6c4009
	clr		%o4
Packit 6c4009
1:	cmp		%o1, 16
Packit 6c4009
	ble		%icc, .Ltiny
Packit 6c4009
	 mov		%o0, %o3
Packit 6c4009
	sub		%g0, %o0, %g1
Packit 6c4009
	and		%g1, 0x7, %g1
Packit 6c4009
	brz,pt		%g1, .Laligned8
Packit 6c4009
	 sub		%o1, %g1, %o1
Packit 6c4009
1:	stb		%o4, [%o0 + 0x00]
Packit 6c4009
	subcc		%g1, 1, %g1
Packit 6c4009
	bne,pt		%icc, 1b
Packit 6c4009
	 add		%o0, 1, %o0
Packit 6c4009
.Laligned8:
Packit 6c4009
	cmp		%o1, 64 + (64 - 8)
Packit 6c4009
	ble		.Lmedium
Packit 6c4009
	 sub		%g0, %o0, %g1
Packit 6c4009
	andcc		%g1, (64 - 1), %g1
Packit 6c4009
	brz,pn		%g1, .Laligned64
Packit 6c4009
	 sub		%o1, %g1, %o1
Packit 6c4009
1:	stx		%o4, [%o0 + 0x00]
Packit 6c4009
	subcc		%g1, 8, %g1
Packit 6c4009
	bne,pt		%icc, 1b
Packit 6c4009
	 add		%o0, 0x8, %o0
Packit 6c4009
.Laligned64:
Packit 6c4009
	andn		%o1, 64 - 1, %g1
Packit 6c4009
	sub		%o1, %g1, %o1
Packit 6c4009
	brnz,pn		%o4, .Lnon_bzero_loop
Packit 6c4009
	 mov		0x20, %g2
Packit 6c4009
1:	stxa		%o4, [%o0 + %g0] ASI_BLK_INIT_QUAD_LDD_P
Packit 6c4009
	subcc		%g1, 0x40, %g1
Packit 6c4009
	stxa		%o4, [%o0 + %g2] ASI_BLK_INIT_QUAD_LDD_P
Packit 6c4009
	bne,pt		%icc, 1b
Packit 6c4009
	 add		%o0, 0x40, %o0
Packit 6c4009
.Lpostloop:
Packit 6c4009
	cmp		%o1, 8
Packit 6c4009
	bl,pn		%icc, .Ltiny
Packit 6c4009
	 membar		#StoreStore|#StoreLoad
Packit 6c4009
.Lmedium:
Packit 6c4009
	andn		%o1, 0x7, %g1
Packit 6c4009
	sub		%o1, %g1, %o1
Packit 6c4009
1:	stx		%o4, [%o0 + 0x00]
Packit 6c4009
	subcc		%g1, 0x8, %g1
Packit 6c4009
	bne,pt		%icc, 1b
Packit 6c4009
	 add		%o0, 0x08, %o0
Packit 6c4009
	andcc		%o1, 0x4, %g1
Packit 6c4009
	be,pt		%icc, .Ltiny
Packit 6c4009
	 sub		%o1, %g1, %o1
Packit 6c4009
	stw		%o4, [%o0 + 0x00]
Packit 6c4009
	add		%o0, 0x4, %o0
Packit 6c4009
.Ltiny:
Packit 6c4009
	cmp		%o1, 0
Packit 6c4009
	be,pn		%icc, .Lexit
Packit 6c4009
1:	 subcc		%o1, 1, %o1
Packit 6c4009
	stb		%o4, [%o0 + 0x00]
Packit 6c4009
	bne,pt		%icc, 1b
Packit 6c4009
	 add		%o0, 1, %o0
Packit 6c4009
.Lexit:
Packit 6c4009
	retl
Packit 6c4009
	 mov		%o3, %o0
Packit 6c4009
.Lnon_bzero_loop:
Packit 6c4009
	mov		0x08, %g3
Packit 6c4009
	mov		0x28, %o5
Packit 6c4009
1:	stxa		%o4, [%o0 + %g0] ASI_BLK_INIT_QUAD_LDD_P
Packit 6c4009
	subcc		%g1, 0x40, %g1
Packit 6c4009
	stxa		%o4, [%o0 + %g2] ASI_BLK_INIT_QUAD_LDD_P
Packit 6c4009
	stxa		%o4, [%o0 + %g3] ASI_BLK_INIT_QUAD_LDD_P
Packit 6c4009
	stxa		%o4, [%o0 + %o5] ASI_BLK_INIT_QUAD_LDD_P
Packit 6c4009
	add		%o0, 0x10, %o0
Packit 6c4009
	stxa		%o4, [%o0 + %g0] ASI_BLK_INIT_QUAD_LDD_P
Packit 6c4009
	stxa		%o4, [%o0 + %g2] ASI_BLK_INIT_QUAD_LDD_P
Packit 6c4009
	stxa		%o4, [%o0 + %g3] ASI_BLK_INIT_QUAD_LDD_P
Packit 6c4009
	stxa		%o4, [%o0 + %o5] ASI_BLK_INIT_QUAD_LDD_P
Packit 6c4009
	bne,pt		%icc, 1b
Packit 6c4009
	 add		%o0, 0x30, %o0
Packit 6c4009
	ba,a,pt		%icc, .Lpostloop
Packit 6c4009
END(__bzero_niagara4)
Packit 6c4009
Packit 6c4009
#endif