diff options
author | Noah Goldstein <goldstein.w.n@gmail.com> | 2022-06-29 16:07:05 -0700 |
---|---|---|
committer | Noah Goldstein <goldstein.w.n@gmail.com> | 2022-06-29 19:47:52 -0700 |
commit | 4a3f29e7e475dd4e7cce2a24c187e6fb7b5b0a05 (patch) | |
tree | d3fcb94ad6ec06cc3d2ebd5d302a24c84ca8af65 /sysdeps/x86_64/multiarch/memset-erms.S | |
parent | 2a1099020cdc1e4c9c928156aa85c8cf9d540291 (diff) | |
download | glibc-4a3f29e7e475dd4e7cce2a24c187e6fb7b5b0a05.tar.gz glibc-4a3f29e7e475dd4e7cce2a24c187e6fb7b5b0a05.tar.xz glibc-4a3f29e7e475dd4e7cce2a24c187e6fb7b5b0a05.zip |
x86: Move and slightly improve memset_erms
Implementation wise: 1. Remove the VZEROUPPER as memset_{impl}_unaligned_erms does not use the L(stosb) label that was previously defined. 2. Don't give the hotpath (fallthrough) to zero size. Code positioning wise: Move memset_{chk}_erms to its own file. Leaving it in between the memset_{impl}_unaligned both adds unnecessary complexity to the file and wastes space in a relatively hot cache section.
Diffstat (limited to 'sysdeps/x86_64/multiarch/memset-erms.S')
-rw-r--r-- | sysdeps/x86_64/multiarch/memset-erms.S | 44 |
1 files changed, 44 insertions, 0 deletions
diff --git a/sysdeps/x86_64/multiarch/memset-erms.S b/sysdeps/x86_64/multiarch/memset-erms.S new file mode 100644 index 0000000000..e83cccc731 --- /dev/null +++ b/sysdeps/x86_64/multiarch/memset-erms.S @@ -0,0 +1,44 @@ +/* memset implement with rep stosb + Copyright (C) 2022 Free Software Foundation, Inc. + This file is part of the GNU C Library. + + The GNU C Library is free software; you can redistribute it and/or + modify it under the terms of the GNU Lesser General Public + License as published by the Free Software Foundation; either + version 2.1 of the License, or (at your option) any later version. + + The GNU C Library is distributed in the hope that it will be useful, + but WITHOUT ANY WARRANTY; without even the implied warranty of + MERCHANTABILITY or FITNESS FOR A PARTICULAR PURPOSE. See the GNU + Lesser General Public License for more details. + + You should have received a copy of the GNU Lesser General Public + License along with the GNU C Library; if not, see + <https://www.gnu.org/licenses/>. */ + + +#include <sysdep.h> + +#if defined USE_MULTIARCH && IS_IN (libc) + .text +ENTRY (__memset_chk_erms) + cmp %RDX_LP, %RCX_LP + jb HIDDEN_JUMPTARGET (__chk_fail) +END (__memset_chk_erms) + +/* Only used to measure performance of REP STOSB. */ +ENTRY (__memset_erms) + /* Skip zero length. */ + test %RDX_LP, %RDX_LP + jz L(stosb_return_zero) + mov %RDX_LP, %RCX_LP + movzbl %sil, %eax + mov %RDI_LP, %RDX_LP + rep stosb + mov %RDX_LP, %RAX_LP + ret +L(stosb_return_zero): + movq %rdi, %rax + ret +END (__memset_erms) +#endif |