mirror of
https://github.com/SCST-project/scst.git
synced 2026-09-19 14:34:18 +00:00
svn+ssh://yanb123@svn.code.sf.net/p/scst/svn/trunk
........
r5536 | vlnb | 2014-05-22 06:06:46 +0300 (Thu, 22 May 2014) | 3 lines
Version changed to 3.1.0-pre1
........
r5537 | vlnb | 2014-05-22 06:18:27 +0300 (Thu, 22 May 2014) | 3 lines
Web updates
........
r5538 | bvassche | 2014-05-22 10:16:04 +0300 (Thu, 22 May 2014) | 1 line
nightly build: Update kernel versions
........
r5539 | vlnb | 2014-05-23 05:20:35 +0300 (Fri, 23 May 2014) | 9 lines
vdisk_nullio: Add "read_zero" attribute
Add an attribute called "read_zero" to vdisk_nullio devices that
controls whether or not READs from a vdisk_nullio device return
zeroed data buffers.
Signed-off-by: Bart Van Assche <bvanassche@acm.org>
........
r5543 | bvassche | 2014-05-23 10:33:53 +0300 (Fri, 23 May 2014) | 1 line
RHEL 7 build fixes
........
r5545 | bvassche | 2014-05-23 11:36:36 +0300 (Fri, 23 May 2014) | 1 line
scripts/rebuild-rhel-kernel-rpm: Add RHEL 7 RC support
........
r5547 | vlnb | 2014-05-24 06:10:34 +0300 (Sat, 24 May 2014) | 3 lines
Optimize read_zero functionality
........
r5555 | bvassche | 2014-05-27 14:59:11 +0300 (Tue, 27 May 2014) | 5 lines
qla2x00t: Documentation / source code comment / log messages spelling fix
Change a few occurrences of "conformation" into "confirmation". See also the
QLogic 2500 Series Firmware Interface Specification.
........
r5557 | vlnb | 2014-05-30 03:42:34 +0300 (Fri, 30 May 2014) | 5 lines
Small code reorganization.
No functionality changed
........
r5558 | vlnb | 2014-05-30 06:00:07 +0300 (Fri, 30 May 2014) | 3 lines
Logging fixes
........
r5560 | bvassche | 2014-06-02 18:31:50 +0300 (Mon, 02 Jun 2014) | 1 line
Makefile: Only report which RPMs have been built if "make rpm" is run as a non-privileged user
........
r5561 | bvassche | 2014-06-03 09:04:47 +0300 (Tue, 03 Jun 2014) | 1 line
nightly build: Update kernel versions
........
r5562 | vlnb | 2014-06-04 04:54:21 +0300 (Wed, 04 Jun 2014) | 3 lines
Decrease max WRITE SAME length for better latencies
........
r5563 | vlnb | 2014-06-04 05:16:51 +0300 (Wed, 04 Jun 2014) | 3 lines
Enforce limit on max unmap LBAs
........
r5566 | bvassche | 2014-06-04 18:14:22 +0300 (Wed, 04 Jun 2014) | 1 line
ib_srpt: Fix an error message
........
r5567 | bvassche | 2014-06-04 18:17:59 +0300 (Wed, 04 Jun 2014) | 1 line
ib_srpt: Avoid triggering a SCSI command timeout after login
........
r5568 | bvassche | 2014-06-05 09:34:19 +0300 (Thu, 05 Jun 2014) | 1 line
scst_vdisk: Build fix for kernel versions <= 2.6.32
........
r5569 | bvassche | 2014-06-05 09:46:57 +0300 (Thu, 05 Jun 2014) | 1 line
scst_vdisk: Fix a kernel version < 2.6.38 compiler warning
........
r5570 | vlnb | 2014-06-06 06:20:26 +0300 (Fri, 06 Jun 2014) | 8 lines
scst_lib: Fix a compiler warning triggered by the WRITE SAME implementation
Avoid for release builds that the compiler reports that the variable
'ws_sg_cnt' is not used.
Signed-off-by: Bart Van Assche <bvanassche@acm.org>
........
r5571 | vlnb | 2014-06-06 06:22:14 +0300 (Fri, 06 Jun 2014) | 7 lines
nullio_exec_read(): Fix kunmap() argument
The argument of kunmap() is of type struct page *. Detected by smatch.
Signed-off-by: Bart Van Assche <bvanassche@acm.org>
........
r5572 | vlnb | 2014-06-06 06:24:03 +0300 (Fri, 06 Jun 2014) | 11 lines
scst: Leave out FSF mail address
This avoids that the following checkpatch complaint is triggered:
Do not include the paragraph about writing to the Free Software Foundation's
mailing address from the sample GPL notice. The FSF has changed addresses in
the past, and may do so again. Linux already includes a copy of the GPL.
Signed-off-by: Bart Van Assche <bvanassche@acm.org>
........
r5573 | vlnb | 2014-06-06 06:26:55 +0300 (Fri, 06 Jun 2014) | 10 lines
scst: Make lockdep_assert_held() easier to use
The lockdep_assert_held() macro is a convenient debugging tool.
However, it is inconvenient to surround each invocation of that
macro by an #ifdef/#endif pair. Hence make it easier to use this
macro with older kernel versions.
Signed-off-by: Bart Van Assche <bvanassche@acm.org>
........
r5574 | vlnb | 2014-06-07 00:59:24 +0300 (Sat, 07 Jun 2014) | 3 lines
Use limits.discard_zeroes_data to set LBPRZ
........
r5575 | bvassche | 2014-06-07 13:46:49 +0300 (Sat, 07 Jun 2014) | 1 line
nightly build: Update kernel versions
........
r5577 | bvassche | 2014-06-10 17:16:14 +0300 (Tue, 10 Jun 2014) | 1 line
ib_srpt: Make the test for IB_EVENT_GID_CHANGE support more robust
........
r5578 | bvassche | 2014-06-10 17:49:59 +0300 (Tue, 10 Jun 2014) | 1 line
ib_srpt: Make IB_EVENT_GID_CHANGE test independent of the OFED detection code
........
r5579 | bvassche | 2014-06-11 13:02:15 +0300 (Wed, 11 Jun 2014) | 1 line
ib_srpt: RHEL 5 build fix
........
r5581 | bvassche | 2014-06-11 18:27:06 +0300 (Wed, 11 Jun 2014) | 1 line
regression tests: Sync with a recent sysfs change
........
r5582 | bvassche | 2014-06-11 18:27:48 +0300 (Wed, 11 Jun 2014) | 1 line
regression tests: Sort hash keys before comparing
........
r5583 | bvassche | 2014-06-11 18:41:01 +0300 (Wed, 11 Jun 2014) | 1 line
nightly build: Update kernel versions
........
r5584 | vlnb | 2014-06-11 22:33:18 +0300 (Wed, 11 Jun 2014) | 8 lines
scst: RHEL 5 build fix
Avoid that building the scst kernel module fails on RHEL 5 due to
a missing kvasprintf() implementation.
Signed-off-by: Bart Van Assche <bvanassche@acm.org>
........
r5585 | vlnb | 2014-06-11 22:38:10 +0300 (Wed, 11 Jun 2014) | 11 lines
scst: Remove unused variables
Avoid that building scst with W=1 triggers compiler warnings about
variables that are set but not used. See also the documentation of
the gcc compiler flag -Wunused-but-set-variable.
This patch does not change any functionality.
Signed-off-by: Bart Van Assche <bvanassche@acm.org>
........
r5586 | vlnb | 2014-06-11 22:39:51 +0300 (Wed, 11 Jun 2014) | 9 lines
scst_lib: Introduce additional temporary variables
Make the code slightly easier to read by introducing temporary
variables for the expressions 'tgt_dev->sess' and 'sess->tgt->tgtt'.
This patch does not change any functionality.
Signed-off-by: Bart Van Assche <bvanassche@acm.org>
........
r5587 | vlnb | 2014-06-11 23:57:03 +0300 (Wed, 11 Jun 2014) | 10 lines
scst: Add support for 64-bit LUNs
The datatype of scsi_device.lun will be changed from u32 into u64
in the near future. Update SCST accordingly. These changes have
been implemented such that these are compatible with 32-bit and
64-bit LUNs.
Signed-off-by: Bart Van Assche <bvanassche@acm.org>
........
r5588 | vlnb | 2014-06-12 00:00:16 +0300 (Thu, 12 Jun 2014) | 9 lines
scst_local: Support LUN numbers >= 16384
Add support for 32-bit LUN numbers. As soon as the patches that add
64-bit LUN support are upstream this patch will also make 64-bit
LUN support available in scst_local.
Signed-off-by: Bart Van Assche <bvanassche@acm.org>
........
r5589 | vlnb | 2014-06-12 00:42:08 +0300 (Thu, 12 Jun 2014) | 8 lines
scst: Clean up __scst_resume_activity()
Move all management commands from scst_delayed_mgmt_cmd_list to the
active command list during resume instead of only the first one.
Signed-off-by: Bart Van Assche <bvanassche@acm.org>
........
r5590 | vlnb | 2014-06-12 01:07:00 +0300 (Thu, 12 Jun 2014) | 9 lines
scst: Introduce scst_lookup_tgt_dev()
This patch does not change any functionality.
Signed-off-by: Bart Van Assche <bvanassche@acm.org>
with some improvements
........
r5592 | bvassche | 2014-06-12 11:38:45 +0300 (Thu, 12 Jun 2014) | 7 lines
scst.h: Move definition of swap()
Make sure that the definition of swap() is guarded by
"#if LINUX_VERSION_CODE < KERNEL_VERSION(2, 6, 29)" only instead
of "#if LINUX_VERSION_CODE < KERNEL_VERSION(2, 6, 28)" and
"#if LINUX_VERSION_CODE < KERNEL_VERSION(2, 6, 29)".
........
r5593 | bvassche | 2014-06-12 12:15:50 +0300 (Thu, 12 Jun 2014) | 1 line
nightly build: Update kernel versions
........
r5594 | bvassche | 2014-06-12 14:33:00 +0300 (Thu, 12 Jun 2014) | 1 line
ib_srpt: Set MOFED include path correctly if MOFED has been installed with --add-kernel-support
........
r5595 | bvassche | 2014-06-12 16:38:38 +0300 (Thu, 12 Jun 2014) | 1 line
ib_srpt: Make non-OFED build work again
........
r5596 | vlnb | 2014-06-13 07:52:18 +0300 (Fri, 13 Jun 2014) | 16 lines
scst: Switch from the cpu_*() to the cpumask_*() API
The cpus_*() functions were deprecated via patch "cpumask:
introduce new API, without changing anything" (November 2008,
commit ID 2d3854a37e8b). Hence switch from the cpus_*() API to
the cpumask_*() API.
This patch has the intended side effect of not adding the "[key]"
property to cpumask sysfs attributes that contain the default
cpumask. The current code namely reads uninitialized bits on
systems where nr_cpu_ids < NR_CPUS because cpus_equal() compares
more bits than those that were set by cpumask_copy().
Signed-off-by: Bart Van Assche <bvanassche@acm.org>
........
r5597 | vlnb | 2014-06-13 08:03:17 +0300 (Fri, 13 Jun 2014) | 3 lines
Forgotten versions updated
........
r5598 | bvassche | 2014-06-13 09:55:23 +0300 (Fri, 13 Jun 2014) | 1 line
ib_srpt: Make one_target_per_port the default mode
........
r5600 | vlnb | 2014-06-14 01:24:06 +0300 (Sat, 14 Jun 2014) | 9 lines
scst: Avoid that W=1 triggers complaints about unused variables
Avoid that building scst with W=1 triggers compiler warnings about
variables that are set but not used. See also the documentation of
the gcc compiler flag -Wunused-but-set-variable.
Signed-off-by: Bart Van Assche <bvanassche@acm.org>
........
r5601 | vlnb | 2014-06-14 01:31:42 +0300 (Sat, 14 Jun 2014) | 8 lines
scst_local: Add close_session() callback function
This is useful for triggering the session reassignment code via
the scst_local driver.
Signed-off-by: Bart Van Assche <bvanassche@acm.org>
........
r5602 | vlnb | 2014-06-14 02:57:26 +0300 (Sat, 14 Jun 2014) | 8 lines
scst_pr_read_reservation(): Initialize returned buffer
Avoid that this function returns an uninitialized buffer to the
initiator if buffer_size < 8. Detected by Coverity.
Signed-off-by: Bart Van Assche <bvanassche@acm.org>
........
r5603 | vlnb | 2014-06-14 02:58:28 +0300 (Sat, 14 Jun 2014) | 5 lines
scst: Help Coverity recognize that vmalloc(0) returns NULL
Signed-off-by: Bart Van Assche <bvanassche@acm.org>
........
r5605 | bvassche | 2014-06-14 20:10:58 +0300 (Sat, 14 Jun 2014) | 1 line
fcst: Remove an unused variable
........
r5606 | bvassche | 2014-06-14 20:17:56 +0300 (Sat, 14 Jun 2014) | 5 lines
fcst: Move exch_done() calls into ft_cmd_done()
This patch ensures that exch_done() gets called if an fcst
callback returns SCST_TGT_RES_FATAL_ERROR.
........
r5607 | bvassche | 2014-06-14 20:18:34 +0300 (Sat, 14 Jun 2014) | 10 lines
fcst: Handle frame send failures properly
Retry sending XFER_RDY, data and response frames if the network
driver reports that sending failed (-ENOMEM) instead of reporting
a kernel warning (WARN_ON(1)). If sending XFER_RDY or data frames
failed for another reason, report this to the initiator as a
write error (ASC = 03; ASCQ = 00 which stands for PERIPHERAL
DEVICE WRITE FAULT). If sending a response frame failed with
another error code than -ENOMEM, do not send a response.
........
r5608 | vlnb | 2014-06-17 03:50:46 +0300 (Tue, 17 Jun 2014) | 10 lines
scst: Make access control group removal behavior configurable
SCST rejects removal of an access control group with one or more
sessions with error code -EBUSY. Make it easy to change this
behavior into forcibly closing sessions when an access control
group is removed.
Signed-off-by: Bart Van Assche <bvanassche@acm.org>
........
r5609 | bvassche | 2014-06-17 09:37:08 +0300 (Tue, 17 Jun 2014) | 1 line
nightly build: Update kernel versions
........
r5610 | vlnb | 2014-06-19 06:51:48 +0300 (Thu, 19 Jun 2014) | 3 lines
Update for 3.15 kernels
........
r5611 | bvassche | 2014-06-19 10:09:53 +0300 (Thu, 19 Jun 2014) | 1 line
nightly build: Add kernel 3.15 build infrastructure
........
r5612 | bvassche | 2014-06-19 15:48:25 +0300 (Thu, 19 Jun 2014) | 1 line
kernel module installation: Skip "depmod" when building an RPM
........
r5613 | vlnb | 2014-06-20 07:00:41 +0300 (Fri, 20 Jun 2014) | 8 lines
scst: Convert a loop to keep smatch happy
Avoid that smatch reports the following warning:
scst_init_session() info: loop could be replaced with if statement.
Signed-off-by: Bart Van Assche <bvanassche@acm.org>
........
r5614 | vlnb | 2014-06-20 07:02:00 +0300 (Fri, 20 Jun 2014) | 13 lines
iscsi-scst: Suppress a compiler warning
Avoid that the following compiler warning is reported when compiling
iscsi-scst:
chap.c: In function 'chap_rand':
chap.c:348:5: warning: ignoring return value of 'read', declared with attribute warn_unused_result [-Wunused-result]
(void)read(fd, &r, sizeof(r));
^
Signed-off-by: Bart Van Assche <bvanassche@acm.org>
........
r5615 | vlnb | 2014-06-20 07:03:40 +0300 (Fri, 20 Jun 2014) | 5 lines
scst, iscsi-scst: Fix RHEL 5 compilation warnings
Signed-off-by: Bart Van Assche <bvanassche@acm.org>
........
r5616 | vlnb | 2014-06-20 07:05:11 +0300 (Fri, 20 Jun 2014) | 10 lines
scst: Exclude certain locking code from static analysis
Loops with locking statements and also lock and unlock
statements guarded by an if-statement trigger false positive
warnings when analyzing the SCST code with smatch and/or sparse.
Hence exclude such code from static analysis.
Signed-off-by: Bart Van Assche <bvanassche@acm.org>
........
r5617 | vlnb | 2014-06-20 07:09:11 +0300 (Fri, 20 Jun 2014) | 10 lines
scst: Avoid that sparse complains about unreachable code
Remove the code after BUG() statements to avoid that smatch
complains about unreachable code. Hide the spin_unlock() statements
before BUG() statements for static analysis tools to avoid that
sparse complains about locking imbalances.
Signed-off-by: Bart Van Assche <bvanassche@acm.org>
........
r5618 | vlnb | 2014-06-20 07:10:40 +0300 (Fri, 20 Jun 2014) | 12 lines
Change BUG_ON(1) into BUG()
With CONFIG_BUG=y both BUG() and BUG_ON(1) halt the system. However,
with CONFIG_BUG=n BUG() halts the system but BUG_ON(1) not. To avoid
such subtleties, change BUG_ON(1) into BUG().
See also patch Josh Triplett, "bug: Make BUG() always stop the machine",
7 April 2014 (commit ID a4b5d580e07875f9be29f62a57c67fbbdbb40ba2).
Signed-off-by: Bart Van Assche <bvanassche@acm.org>
........
r5619 | bvassche | 2014-06-20 08:56:36 +0300 (Fri, 20 Jun 2014) | 1 line
nightly build: Add kernel version 3.15.1
........
r5620 | vlnb | 2014-06-24 07:45:08 +0300 (Tue, 24 Jun 2014) | 9 lines
scst_vdisk: Split vdisk_exec_inquiry()
Make vdisk_exec_inquiry() easier to read by moving the code
for the implementation of each VPD page into a separate function.
This patch does not change any functionality.
Signed-off-by: Bart Van Assche <bvanassche@acm.org>
........
r5621 | bvassche | 2014-06-24 16:32:18 +0300 (Tue, 24 Jun 2014) | 1 line
ib_srpt: Complain if another ib_srpt.ko kernel module already exists
........
r5622 | bvassche | 2014-06-24 16:33:23 +0300 (Tue, 24 Jun 2014) | 2 lines
ib_srpt: Set SCSI residual fields in SRP_CMD reply
........
r5624 | bvassche | 2014-06-25 14:50:40 +0300 (Wed, 25 Jun 2014) | 1 line
nightly build: Use http instead of ftp for downloading kernel source code
........
r5625 | vlnb | 2014-06-26 00:38:19 +0300 (Thu, 26 Jun 2014) | 5 lines
scst_debug.h: Make EXTRACHECKS_*_ON() statements visible to Coverity
Signed-off-by: Bart Van Assche <bvanassche@acm.org>
........
r5626 | vlnb | 2014-06-27 02:26:25 +0300 (Fri, 27 Jun 2014) | 10 lines
scst_vdisk: Three more put_unaligned_*() conversions
Convert three more *(__be16 *)p = cpu_to_be16(v) statements into
put_unaligned_be16(v, p) since the latter is easier to read. Also
convert one "cmd->dev" into "dev" expression. This patch does not
change any functionality.
Signed-off-by: Bart Van Assche <bvanassche@acm.org>
........
r5627 | bvassche | 2014-06-27 13:32:02 +0300 (Fri, 27 Jun 2014) | 1 line
nightly build: Update kernel versions
........
r5628 | bvassche | 2014-06-28 22:56:36 +0300 (Sat, 28 Jun 2014) | 1 line
ib_srpt: Remove existing ib_srpt.ko kernel modules before installation
........
r5629 | bvassche | 2014-06-28 22:58:44 +0300 (Sat, 28 Jun 2014) | 6 lines
scst_vdisk: Fix 32-bit build
Avoid 64-bit modulo computations since these result in undefined symbol
errors on 32-bit systems (__moddi3 / __umoddi3). Support sizes >= 2**32
bytes on 32-bit systems.
........
r5630 | bvassche | 2014-06-28 23:00:22 +0300 (Sat, 28 Jun 2014) | 1 line
scst.spec.in: Follow-up for r5628
........
r5631 | bvassche | 2014-06-28 23:15:45 +0300 (Sat, 28 Jun 2014) | 1 line
scst_local: Avoid that session deletion triggers a kernel warning
........
r5647 | bvassche | 2014-06-30 10:18:53 +0300 (Mon, 30 Jun 2014) | 1 line
scst: Build fix for Linux kernel versions 2.6.33 and 2.6.34
........
r5648 | bvassche | 2014-06-30 10:28:18 +0300 (Mon, 30 Jun 2014) | 1 line
scst: Build fix for kernel versions <= 2.6.31
........
r5649 | bvassche | 2014-06-30 11:40:23 +0300 (Mon, 30 Jun 2014) | 6 lines
scst_vdisk: Fix a checkpatch warning
Address the following checkpatch warning:
char * array declaration might be better as static const
........
r5650 | bvassche | 2014-06-30 11:52:06 +0300 (Mon, 30 Jun 2014) | 1 line
nightly build: Correct a kernel version
........
r5651 | bvassche | 2014-06-30 12:18:41 +0300 (Mon, 30 Jun 2014) | 1 line
nightly build: Correct a kernel version
........
r5654 | bvassche | 2014-07-01 09:38:13 +0300 (Tue, 01 Jul 2014) | 6 lines
scst_vdisk: Fix a checkpatch warning
Avoid that checkpatch reports the following warning:
WARNING: static const char * array should probably be static const char * const
........
r5657 | bvassche | 2014-07-01 19:46:12 +0300 (Tue, 01 Jul 2014) | 1 line
nightly build: Update kernel versions
........
r5658 | bvassche | 2014-07-03 11:36:48 +0300 (Thu, 03 Jul 2014) | 1 line
scripts/kernel-functions: Handle 3.x.0 kernel versions correctly
........
r5659 | bvassche | 2014-07-03 11:42:08 +0300 (Thu, 03 Jul 2014) | 1 line
scripts/generate-patched-kernel: Clean up
........
r5661 | bvassche | 2014-07-04 08:39:28 +0300 (Fri, 04 Jul 2014) | 1 line
Make scripts/kernel-functions again compatible with 2.6.x kernels
........
r5662 | bvassche | 2014-07-06 11:02:28 +0300 (Sun, 06 Jul 2014) | 1 line
scripts/run-regression-tests: Add command-line option -4 (disable IPv6)
........
git-svn-id: http://svn.code.sf.net/p/scst/svn/branches/iser@5666 d57e44dd-8a1f-0410-8b47-8ef2f437770f
1867 lines
46 KiB
C
1867 lines
46 KiB
C
/*
|
|
* Network threads.
|
|
*
|
|
* Copyright (C) 2004 - 2005 FUJITA Tomonori <tomof@acm.org>
|
|
* Copyright (C) 2007 - 2014 Vladislav Bolkhovitin
|
|
* Copyright (C) 2007 - 2014 Fusion-io, Inc.
|
|
*
|
|
* This program is free software; you can redistribute it and/or
|
|
* modify it under the terms of the GNU General Public License
|
|
* as published by the Free Software Foundation.
|
|
*
|
|
* This program is distributed in the hope that it will be useful,
|
|
* but WITHOUT ANY WARRANTY; without even the implied warranty of
|
|
* MERCHANTABILITY or FITNESS FOR A PARTICULAR PURPOSE. See the
|
|
* GNU General Public License for more details.
|
|
*/
|
|
|
|
#include <linux/sched.h>
|
|
#include <linux/file.h>
|
|
#include <linux/kthread.h>
|
|
#include <linux/delay.h>
|
|
#include <net/tcp_states.h>
|
|
|
|
#include "iscsi.h"
|
|
#include "digest.h"
|
|
#include "iscsit_transport.h"
|
|
|
|
/* Read data states */
|
|
enum rx_state {
|
|
RX_INIT_BHS, /* Must be zero for better "switch" optimization. */
|
|
RX_BHS,
|
|
RX_CMD_START,
|
|
RX_DATA,
|
|
RX_END,
|
|
|
|
RX_CMD_CONTINUE,
|
|
RX_INIT_HDIGEST,
|
|
RX_CHECK_HDIGEST,
|
|
RX_INIT_DDIGEST,
|
|
RX_CHECK_DDIGEST,
|
|
RX_AHS,
|
|
RX_PADDING,
|
|
};
|
|
|
|
enum tx_state {
|
|
TX_INIT = 0, /* Must be zero for better "switch" optimization. */
|
|
TX_BHS_DATA,
|
|
TX_INIT_PADDING,
|
|
TX_PADDING,
|
|
TX_INIT_DDIGEST,
|
|
TX_DDIGEST,
|
|
TX_END,
|
|
};
|
|
|
|
#if defined(CONFIG_TCP_ZERO_COPY_TRANSFER_COMPLETION_NOTIFICATION)
|
|
static void iscsi_check_closewait(struct iscsi_conn *conn)
|
|
{
|
|
struct iscsi_cmnd *cmnd;
|
|
|
|
TRACE_ENTRY();
|
|
|
|
TRACE_CONN_CLOSE_DBG("conn %p, sk_state %d", conn,
|
|
conn->sock->sk->sk_state);
|
|
|
|
if (conn->sock->sk->sk_state != TCP_CLOSE) {
|
|
TRACE_CONN_CLOSE_DBG("conn %p, skipping", conn);
|
|
goto out;
|
|
}
|
|
|
|
/*
|
|
* No data are going to be sent, so all queued buffers can be freed
|
|
* now. In many cases TCP does that only in close(), but we can't rely
|
|
* on user space on calling it.
|
|
*/
|
|
|
|
again:
|
|
spin_lock_bh(&conn->cmd_list_lock);
|
|
list_for_each_entry(cmnd, &conn->cmd_list, cmd_list_entry) {
|
|
struct iscsi_cmnd *rsp;
|
|
int restart = 0;
|
|
|
|
TRACE_CONN_CLOSE_DBG("cmd %p, scst_state %x, "
|
|
"r2t_len_to_receive %d, ref_cnt %d, parent_req %p, "
|
|
"net_ref_cnt %d, sg %p", cmnd, cmnd->scst_state,
|
|
cmnd->r2t_len_to_receive, atomic_read(&cmnd->ref_cnt),
|
|
cmnd->parent_req, atomic_read(&cmnd->net_ref_cnt),
|
|
cmnd->sg);
|
|
|
|
sBUG_ON(cmnd->parent_req != NULL);
|
|
|
|
if (cmnd->sg != NULL) {
|
|
int i;
|
|
|
|
if (cmnd_get_check(cmnd))
|
|
continue;
|
|
|
|
for (i = 0; i < cmnd->sg_cnt; i++) {
|
|
struct page *page = sg_page(&cmnd->sg[i]);
|
|
TRACE_CONN_CLOSE_DBG("page %p, net_priv %p, "
|
|
"_count %d", page, page->net_priv,
|
|
atomic_read(&page->_count));
|
|
|
|
if (page->net_priv != NULL) {
|
|
if (restart == 0) {
|
|
spin_unlock_bh(&conn->cmd_list_lock);
|
|
restart = 1;
|
|
}
|
|
while (page->net_priv != NULL)
|
|
iscsi_put_page_callback(page);
|
|
}
|
|
}
|
|
cmnd_put(cmnd);
|
|
|
|
if (restart)
|
|
goto again;
|
|
}
|
|
|
|
list_for_each_entry(rsp, &cmnd->rsp_cmd_list,
|
|
rsp_cmd_list_entry) {
|
|
TRACE_CONN_CLOSE_DBG(" rsp %p, ref_cnt %d, "
|
|
"net_ref_cnt %d, sg %p",
|
|
rsp, atomic_read(&rsp->ref_cnt),
|
|
atomic_read(&rsp->net_ref_cnt), rsp->sg);
|
|
|
|
if ((rsp->sg != cmnd->sg) && (rsp->sg != NULL)) {
|
|
int i;
|
|
|
|
if (cmnd_get_check(rsp))
|
|
continue;
|
|
|
|
for (i = 0; i < rsp->sg_cnt; i++) {
|
|
struct page *page =
|
|
sg_page(&rsp->sg[i]);
|
|
TRACE_CONN_CLOSE_DBG(
|
|
" page %p, net_priv %p, "
|
|
"_count %d",
|
|
page, page->net_priv,
|
|
atomic_read(&page->_count));
|
|
|
|
if (page->net_priv != NULL) {
|
|
if (restart == 0) {
|
|
spin_unlock_bh(&conn->cmd_list_lock);
|
|
restart = 1;
|
|
}
|
|
while (page->net_priv != NULL)
|
|
iscsi_put_page_callback(page);
|
|
}
|
|
}
|
|
cmnd_put(rsp);
|
|
|
|
if (restart)
|
|
goto again;
|
|
}
|
|
}
|
|
}
|
|
spin_unlock_bh(&conn->cmd_list_lock);
|
|
|
|
out:
|
|
TRACE_EXIT();
|
|
return;
|
|
}
|
|
#else
|
|
static inline void iscsi_check_closewait(struct iscsi_conn *conn) {};
|
|
#endif
|
|
|
|
static void free_pending_commands(struct iscsi_conn *conn)
|
|
{
|
|
struct iscsi_session *session = conn->session;
|
|
struct list_head *pending_list = &session->pending_list;
|
|
int req_freed;
|
|
struct iscsi_cmnd *cmnd;
|
|
|
|
spin_lock(&session->sn_lock);
|
|
do {
|
|
req_freed = 0;
|
|
list_for_each_entry(cmnd, pending_list, pending_list_entry) {
|
|
TRACE_CONN_CLOSE_DBG("Pending cmd %p"
|
|
"(conn %p, cmd_sn %u, exp_cmd_sn %u)",
|
|
cmnd, conn, cmnd->pdu.bhs.sn,
|
|
session->exp_cmd_sn);
|
|
if ((cmnd->conn == conn) &&
|
|
(session->exp_cmd_sn == cmnd->pdu.bhs.sn)) {
|
|
TRACE_MGMT_DBG("Freeing pending cmd %p "
|
|
"(cmd_sn %u, exp_cmd_sn %u)",
|
|
cmnd, cmnd->pdu.bhs.sn,
|
|
session->exp_cmd_sn);
|
|
|
|
list_del(&cmnd->pending_list_entry);
|
|
cmnd->pending = 0;
|
|
|
|
session->exp_cmd_sn++;
|
|
|
|
spin_unlock(&session->sn_lock);
|
|
|
|
req_cmnd_release_force(cmnd);
|
|
|
|
req_freed = 1;
|
|
spin_lock(&session->sn_lock);
|
|
break;
|
|
}
|
|
}
|
|
} while (req_freed);
|
|
spin_unlock(&session->sn_lock);
|
|
|
|
return;
|
|
}
|
|
|
|
static void free_orphaned_pending_commands(struct iscsi_conn *conn)
|
|
{
|
|
struct iscsi_session *session = conn->session;
|
|
struct list_head *pending_list = &session->pending_list;
|
|
int req_freed;
|
|
struct iscsi_cmnd *cmnd;
|
|
|
|
spin_lock(&session->sn_lock);
|
|
do {
|
|
req_freed = 0;
|
|
list_for_each_entry(cmnd, pending_list, pending_list_entry) {
|
|
TRACE_CONN_CLOSE_DBG("Pending cmd %p"
|
|
"(conn %p, cmd_sn %u, exp_cmd_sn %u)",
|
|
cmnd, conn, cmnd->pdu.bhs.sn,
|
|
session->exp_cmd_sn);
|
|
if (cmnd->conn == conn) {
|
|
TRACE_MGMT_DBG("Freeing orphaned pending "
|
|
"cmnd %p (cmd_sn %u, exp_cmd_sn %u)",
|
|
cmnd, cmnd->pdu.bhs.sn,
|
|
session->exp_cmd_sn);
|
|
|
|
list_del(&cmnd->pending_list_entry);
|
|
cmnd->pending = 0;
|
|
|
|
if (session->exp_cmd_sn == cmnd->pdu.bhs.sn)
|
|
session->exp_cmd_sn++;
|
|
|
|
spin_unlock(&session->sn_lock);
|
|
|
|
req_cmnd_release_force(cmnd);
|
|
|
|
req_freed = 1;
|
|
spin_lock(&session->sn_lock);
|
|
break;
|
|
}
|
|
}
|
|
} while (req_freed);
|
|
spin_unlock(&session->sn_lock);
|
|
|
|
return;
|
|
}
|
|
|
|
#ifdef CONFIG_SCST_DEBUG
|
|
static void trace_conn_close(struct iscsi_conn *conn)
|
|
{
|
|
struct iscsi_cmnd *cmnd;
|
|
#if defined(CONFIG_TCP_ZERO_COPY_TRANSFER_COMPLETION_NOTIFICATION)
|
|
struct iscsi_cmnd *rsp;
|
|
#endif
|
|
|
|
#if 0
|
|
if (time_after(jiffies, start_waiting + 10*HZ))
|
|
trace_flag |= TRACE_CONN_OC_DBG;
|
|
#endif
|
|
|
|
spin_lock_bh(&conn->cmd_list_lock);
|
|
list_for_each_entry(cmnd, &conn->cmd_list,
|
|
cmd_list_entry) {
|
|
TRACE_CONN_CLOSE_DBG(
|
|
"cmd %p, scst_cmd %p, scst_state %x, scst_cmd state "
|
|
"%d, r2t_len_to_receive %d, ref_cnt %d, sn %u, "
|
|
"parent_req %p, pending %d",
|
|
cmnd, cmnd->scst_cmd, cmnd->scst_state,
|
|
((cmnd->parent_req == NULL) && cmnd->scst_cmd) ?
|
|
cmnd->scst_cmd->state : -1,
|
|
cmnd->r2t_len_to_receive, atomic_read(&cmnd->ref_cnt),
|
|
cmnd->pdu.bhs.sn, cmnd->parent_req, cmnd->pending);
|
|
#if defined(CONFIG_TCP_ZERO_COPY_TRANSFER_COMPLETION_NOTIFICATION)
|
|
TRACE_CONN_CLOSE_DBG("net_ref_cnt %d, sg %p",
|
|
atomic_read(&cmnd->net_ref_cnt),
|
|
cmnd->sg);
|
|
if (cmnd->sg != NULL) {
|
|
int i;
|
|
for (i = 0; i < cmnd->sg_cnt; i++) {
|
|
struct page *page = sg_page(&cmnd->sg[i]);
|
|
TRACE_CONN_CLOSE_DBG("page %p, "
|
|
"net_priv %p, _count %d",
|
|
page, page->net_priv,
|
|
atomic_read(&page->_count));
|
|
}
|
|
}
|
|
|
|
sBUG_ON(cmnd->parent_req != NULL);
|
|
|
|
list_for_each_entry(rsp, &cmnd->rsp_cmd_list,
|
|
rsp_cmd_list_entry) {
|
|
TRACE_CONN_CLOSE_DBG(" rsp %p, "
|
|
"ref_cnt %d, net_ref_cnt %d, sg %p",
|
|
rsp, atomic_read(&rsp->ref_cnt),
|
|
atomic_read(&rsp->net_ref_cnt), rsp->sg);
|
|
if (rsp->sg != cmnd->sg && rsp->sg) {
|
|
int i;
|
|
for (i = 0; i < rsp->sg_cnt; i++) {
|
|
TRACE_CONN_CLOSE_DBG(" page %p, "
|
|
"net_priv %p, _count %d",
|
|
sg_page(&rsp->sg[i]),
|
|
sg_page(&rsp->sg[i])->net_priv,
|
|
atomic_read(&sg_page(&rsp->sg[i])->
|
|
_count));
|
|
}
|
|
}
|
|
}
|
|
#endif /* CONFIG_TCP_ZERO_COPY_TRANSFER_COMPLETION_NOTIFICATION */
|
|
}
|
|
spin_unlock_bh(&conn->cmd_list_lock);
|
|
return;
|
|
}
|
|
#else /* CONFIG_SCST_DEBUG */
|
|
static void trace_conn_close(struct iscsi_conn *conn) {}
|
|
#endif /* CONFIG_SCST_DEBUG */
|
|
|
|
void iscsi_task_mgmt_affected_cmds_done(struct scst_mgmt_cmd *scst_mcmd)
|
|
{
|
|
int fn = scst_mgmt_cmd_get_fn(scst_mcmd);
|
|
void *priv = scst_mgmt_cmd_get_tgt_priv(scst_mcmd);
|
|
|
|
TRACE_MGMT_DBG("scst_mcmd %p, fn %d, priv %p", scst_mcmd, fn, priv);
|
|
|
|
switch (fn) {
|
|
case SCST_NEXUS_LOSS_SESS:
|
|
{
|
|
struct iscsi_conn *conn = (struct iscsi_conn *)priv;
|
|
struct iscsi_session *sess = conn->session;
|
|
struct iscsi_conn *c;
|
|
|
|
if (sess->sess_reinst_successor != NULL)
|
|
scst_reassign_retained_sess_states(
|
|
sess->sess_reinst_successor->scst_sess,
|
|
sess->scst_sess);
|
|
|
|
mutex_lock(&sess->target->target_mutex);
|
|
|
|
/*
|
|
* We can't mark sess as shutting down earlier, because until
|
|
* now it might have pending commands. Otherwise, in case of
|
|
* reinstatement, it might lead to data corruption, because
|
|
* commands in being reinstated session can be executed
|
|
* after commands in the new session.
|
|
*/
|
|
sess->sess_shutting_down = 1;
|
|
list_for_each_entry(c, &sess->conn_list, conn_list_entry) {
|
|
if (!test_bit(ISCSI_CONN_SHUTTINGDOWN, &c->conn_aflags)) {
|
|
sess->sess_shutting_down = 0;
|
|
break;
|
|
}
|
|
}
|
|
|
|
if (conn->conn_reinst_successor != NULL) {
|
|
sBUG_ON(!test_bit(ISCSI_CONN_REINSTATING,
|
|
&conn->conn_reinst_successor->conn_aflags));
|
|
conn_reinst_finished(conn->conn_reinst_successor);
|
|
conn->conn_reinst_successor = NULL;
|
|
} else if (sess->sess_reinst_successor != NULL) {
|
|
sess_reinst_finished(sess->sess_reinst_successor);
|
|
sess->sess_reinst_successor = NULL;
|
|
}
|
|
mutex_unlock(&sess->target->target_mutex);
|
|
|
|
complete_all(&conn->ready_to_free);
|
|
break;
|
|
}
|
|
case SCST_ABORT_ALL_TASKS_SESS:
|
|
case SCST_ABORT_ALL_TASKS:
|
|
case SCST_NEXUS_LOSS:
|
|
sBUG();
|
|
break;
|
|
default:
|
|
/* Nothing to do */
|
|
break;
|
|
}
|
|
|
|
return;
|
|
}
|
|
|
|
/* No locks */
|
|
static void close_conn(struct iscsi_conn *conn)
|
|
{
|
|
struct iscsi_session *session = conn->session;
|
|
struct iscsi_target *target = conn->target;
|
|
typeof(jiffies) start_waiting = jiffies;
|
|
typeof(jiffies) shut_start_waiting = start_waiting;
|
|
bool pending_reported = 0, wait_expired = 0, shut_expired = 0;
|
|
uint32_t tid, cid;
|
|
uint64_t sid;
|
|
int rc;
|
|
int lun = 0;
|
|
|
|
#define CONN_PENDING_TIMEOUT ((typeof(jiffies))10*HZ)
|
|
#define CONN_WAIT_TIMEOUT ((typeof(jiffies))10*HZ)
|
|
#define CONN_REG_SHUT_TIMEOUT ((typeof(jiffies))125*HZ)
|
|
#define CONN_DEL_SHUT_TIMEOUT ((typeof(jiffies))10*HZ)
|
|
|
|
TRACE_ENTRY();
|
|
|
|
TRACE_MGMT_DBG("Closing connection %p (conn_ref_cnt=%d)", conn,
|
|
atomic_read(&conn->conn_ref_cnt));
|
|
|
|
iscsi_extracheck_is_rd_thread(conn);
|
|
|
|
sBUG_ON(!conn->closing);
|
|
|
|
if (conn->active_close) {
|
|
/* We want all our already send operations to complete */
|
|
conn->transport->iscsit_conn_close(conn, RCV_SHUTDOWN);
|
|
} else {
|
|
conn->transport->iscsit_conn_close(conn,
|
|
RCV_SHUTDOWN|SEND_SHUTDOWN);
|
|
}
|
|
|
|
mutex_lock(&session->target->target_mutex);
|
|
|
|
set_bit(ISCSI_CONN_SHUTTINGDOWN, &conn->conn_aflags);
|
|
|
|
mutex_unlock(&session->target->target_mutex);
|
|
|
|
rc = scst_rx_mgmt_fn_lun(session->scst_sess,
|
|
SCST_NEXUS_LOSS_SESS, &lun, sizeof(lun),
|
|
SCST_NON_ATOMIC, conn);
|
|
if (rc != 0)
|
|
PRINT_ERROR("SCST_NEXUS_LOSS_SESS failed %d", rc);
|
|
|
|
if (conn->read_state != RX_INIT_BHS) {
|
|
struct iscsi_cmnd *cmnd = conn->read_cmnd;
|
|
|
|
if (cmnd->scst_state == ISCSI_CMD_STATE_RX_CMD) {
|
|
TRACE_CONN_CLOSE_DBG("Going to wait for cmnd %p to "
|
|
"change state from RX_CMD", cmnd);
|
|
}
|
|
wait_event(conn->read_state_waitQ,
|
|
cmnd->scst_state != ISCSI_CMD_STATE_RX_CMD);
|
|
|
|
TRACE_CONN_CLOSE_DBG("Releasing conn->read_cmnd %p (conn %p)",
|
|
conn->read_cmnd, conn);
|
|
|
|
conn->read_cmnd = NULL;
|
|
conn->read_state = RX_INIT_BHS;
|
|
req_cmnd_release_force(cmnd);
|
|
}
|
|
|
|
conn_abort(conn);
|
|
|
|
/* ToDo: not the best way to wait */
|
|
while (atomic_read(&conn->conn_ref_cnt) != 0) {
|
|
if (conn->conn_tm_active)
|
|
iscsi_check_tm_data_wait_timeouts(conn, true);
|
|
|
|
mutex_lock(&target->target_mutex);
|
|
spin_lock(&session->sn_lock);
|
|
if (session->tm_rsp && session->tm_rsp->conn == conn) {
|
|
struct iscsi_cmnd *tm_rsp = session->tm_rsp;
|
|
TRACE_MGMT_DBG("Dropping delayed TM rsp %p", tm_rsp);
|
|
session->tm_rsp = NULL;
|
|
session->tm_active--;
|
|
WARN_ON(session->tm_active < 0);
|
|
spin_unlock(&session->sn_lock);
|
|
mutex_unlock(&target->target_mutex);
|
|
|
|
rsp_cmnd_release(tm_rsp);
|
|
} else {
|
|
spin_unlock(&session->sn_lock);
|
|
mutex_unlock(&target->target_mutex);
|
|
}
|
|
|
|
/* It's safe to check it without sn_lock */
|
|
if (!list_empty(&session->pending_list)) {
|
|
TRACE_CONN_CLOSE_DBG("Disposing pending commands on "
|
|
"connection %p (conn_ref_cnt=%d)", conn,
|
|
atomic_read(&conn->conn_ref_cnt));
|
|
|
|
free_pending_commands(conn);
|
|
|
|
if (time_after(jiffies,
|
|
start_waiting + CONN_PENDING_TIMEOUT)) {
|
|
if (!pending_reported) {
|
|
TRACE_CONN_CLOSE("%s",
|
|
"Pending wait time expired");
|
|
pending_reported = 1;
|
|
}
|
|
free_orphaned_pending_commands(conn);
|
|
}
|
|
}
|
|
|
|
conn->transport->iscsit_make_conn_wr_active(conn);
|
|
|
|
/* That's for active close only, actually */
|
|
if (time_after(jiffies, start_waiting + CONN_WAIT_TIMEOUT) &&
|
|
!wait_expired) {
|
|
TRACE_CONN_CLOSE("Wait time expired (conn %p)", conn);
|
|
conn->transport->iscsit_conn_close(conn, SEND_SHUTDOWN);
|
|
wait_expired = 1;
|
|
shut_start_waiting = jiffies;
|
|
}
|
|
|
|
if (wait_expired && !shut_expired &&
|
|
time_after(jiffies, shut_start_waiting +
|
|
conn->deleting ? CONN_DEL_SHUT_TIMEOUT :
|
|
CONN_REG_SHUT_TIMEOUT)) {
|
|
TRACE_CONN_CLOSE("Wait time after shutdown expired "
|
|
"(conn %p)", conn);
|
|
conn->transport->iscsit_conn_close(conn, 0);
|
|
shut_expired = 1;
|
|
}
|
|
|
|
if (conn->deleting)
|
|
msleep(200);
|
|
else
|
|
msleep(1000);
|
|
|
|
TRACE_CONN_CLOSE_DBG("conn %p, conn_ref_cnt %d left, "
|
|
"wr_state %d, exp_cmd_sn %u",
|
|
conn, atomic_read(&conn->conn_ref_cnt),
|
|
conn->wr_state, session->exp_cmd_sn);
|
|
|
|
trace_conn_close(conn);
|
|
|
|
if (!conn->session->sess_params.rdma_extensions) {
|
|
/* It might never be called for being closed conn */
|
|
__iscsi_write_space_ready(conn);
|
|
|
|
iscsi_check_closewait(conn);
|
|
}
|
|
}
|
|
|
|
if (!conn->session->sess_params.rdma_extensions) {
|
|
write_lock_bh(&conn->sock->sk->sk_callback_lock);
|
|
conn->sock->sk->sk_state_change = conn->old_state_change;
|
|
conn->sock->sk->sk_data_ready = conn->old_data_ready;
|
|
conn->sock->sk->sk_write_space = conn->old_write_space;
|
|
write_unlock_bh(&conn->sock->sk->sk_callback_lock);
|
|
}
|
|
|
|
while (1) {
|
|
bool t;
|
|
|
|
spin_lock_bh(&conn->conn_thr_pool->wr_lock);
|
|
t = (conn->wr_state == ISCSI_CONN_WR_STATE_IDLE);
|
|
spin_unlock_bh(&conn->conn_thr_pool->wr_lock);
|
|
|
|
if (t && (atomic_read(&conn->conn_ref_cnt) == 0))
|
|
break;
|
|
|
|
TRACE_CONN_CLOSE_DBG("Waiting for wr thread (conn %p), "
|
|
"wr_state %x", conn, conn->wr_state);
|
|
msleep(50);
|
|
}
|
|
|
|
wait_for_completion(&conn->ready_to_free);
|
|
|
|
tid = target->tid;
|
|
sid = session->sid;
|
|
cid = conn->cid;
|
|
|
|
mutex_lock(&target->target_mutex);
|
|
conn_free(conn);
|
|
mutex_unlock(&target->target_mutex);
|
|
|
|
/*
|
|
* We can't send E_CONN_CLOSE earlier, because otherwise we would have
|
|
* a race, when the user space tried to destroy session, which still
|
|
* has connections.
|
|
*
|
|
* !! All target, session and conn can be already dead here !!
|
|
*/
|
|
TRACE_CONN_CLOSE("Notifying user space about closing connection %p",
|
|
conn);
|
|
event_send(tid, sid, cid, 0, E_CONN_CLOSE, NULL, NULL);
|
|
|
|
TRACE_EXIT();
|
|
return;
|
|
}
|
|
|
|
static int close_conn_thr(void *arg)
|
|
{
|
|
struct iscsi_conn *conn = (struct iscsi_conn *)arg;
|
|
|
|
TRACE_ENTRY();
|
|
|
|
#ifdef CONFIG_SCST_EXTRACHECKS
|
|
/*
|
|
* To satisfy iscsi_extracheck_is_rd_thread() in functions called
|
|
* on the connection close. It is safe, because at this point conn
|
|
* can't be used by any other thread.
|
|
*/
|
|
conn->rd_task = current;
|
|
#endif
|
|
close_conn(conn);
|
|
|
|
TRACE_EXIT();
|
|
return 0;
|
|
}
|
|
|
|
/* No locks */
|
|
void start_close_conn(struct iscsi_conn *conn)
|
|
{
|
|
struct task_struct *t;
|
|
|
|
TRACE_ENTRY();
|
|
|
|
t = kthread_run(close_conn_thr, conn, "iscsi_conn_cleanup");
|
|
if (IS_ERR(t)) {
|
|
PRINT_ERROR("kthread_run() failed (%ld), closing conn %p "
|
|
"directly", PTR_ERR(t), conn);
|
|
close_conn(conn);
|
|
}
|
|
|
|
TRACE_EXIT();
|
|
return;
|
|
}
|
|
EXPORT_SYMBOL(start_close_conn);
|
|
|
|
static inline void iscsi_conn_init_read(struct iscsi_conn *conn,
|
|
void __user *data, size_t len)
|
|
{
|
|
conn->read_iov[0].iov_base = data;
|
|
conn->read_iov[0].iov_len = len;
|
|
conn->read_msg.msg_iov = conn->read_iov;
|
|
conn->read_msg.msg_iovlen = 1;
|
|
conn->read_size = len;
|
|
return;
|
|
}
|
|
|
|
static void iscsi_conn_prepare_read_ahs(struct iscsi_conn *conn,
|
|
struct iscsi_cmnd *cmnd)
|
|
{
|
|
int asize = (cmnd->pdu.ahssize + 3) & -4;
|
|
|
|
/* ToDo: __GFP_NOFAIL ?? */
|
|
cmnd->pdu.ahs = kmalloc(asize, __GFP_NOFAIL|GFP_KERNEL);
|
|
sBUG_ON(cmnd->pdu.ahs == NULL);
|
|
iscsi_conn_init_read(conn, (void __force __user *)cmnd->pdu.ahs, asize);
|
|
return;
|
|
}
|
|
|
|
struct iscsi_cmnd *iscsi_get_send_cmnd(struct iscsi_conn *conn)
|
|
{
|
|
struct iscsi_cmnd *cmnd = NULL;
|
|
|
|
spin_lock_bh(&conn->write_list_lock);
|
|
if (!list_empty(&conn->write_list)) {
|
|
cmnd = list_first_entry(&conn->write_list, struct iscsi_cmnd,
|
|
write_list_entry);
|
|
cmd_del_from_write_list(cmnd);
|
|
cmnd->write_processing_started = 1;
|
|
} else {
|
|
spin_unlock_bh(&conn->write_list_lock);
|
|
goto out;
|
|
}
|
|
spin_unlock_bh(&conn->write_list_lock);
|
|
|
|
if (unlikely(test_bit(ISCSI_CMD_ABORTED,
|
|
&cmnd->parent_req->prelim_compl_flags))) {
|
|
TRACE_MGMT_DBG("Going to send acmd %p (scst cmd %p, "
|
|
"state %d, parent_req %p)", cmnd, cmnd->scst_cmd,
|
|
cmnd->scst_state, cmnd->parent_req);
|
|
}
|
|
|
|
if (unlikely(cmnd_opcode(cmnd) == ISCSI_OP_SCSI_TASK_MGT_RSP)) {
|
|
#ifdef CONFIG_SCST_DEBUG
|
|
struct iscsi_task_mgt_hdr *req_hdr =
|
|
(struct iscsi_task_mgt_hdr *)&cmnd->parent_req->pdu.bhs;
|
|
struct iscsi_task_rsp_hdr *rsp_hdr =
|
|
(struct iscsi_task_rsp_hdr *)&cmnd->pdu.bhs;
|
|
TRACE_MGMT_DBG("Going to send TM response %p (status %d, "
|
|
"fn %d, parent_req %p)", cmnd, rsp_hdr->response,
|
|
req_hdr->function & ISCSI_FUNCTION_MASK,
|
|
cmnd->parent_req);
|
|
#endif
|
|
}
|
|
|
|
out:
|
|
return cmnd;
|
|
}
|
|
EXPORT_SYMBOL(iscsi_get_send_cmnd);
|
|
|
|
/* Returns number of bytes left to receive or <0 for error */
|
|
static int do_recv(struct iscsi_conn *conn)
|
|
{
|
|
int res;
|
|
mm_segment_t oldfs;
|
|
struct msghdr msg;
|
|
int first_len;
|
|
|
|
EXTRACHECKS_BUG_ON(conn->read_cmnd == NULL);
|
|
|
|
if (unlikely(conn->closing)) {
|
|
res = -EIO;
|
|
goto out;
|
|
}
|
|
|
|
/*
|
|
* We suppose that if sock_recvmsg() returned less data than requested,
|
|
* then next time it will return -EAGAIN, so there's no point to call
|
|
* it again.
|
|
*/
|
|
|
|
restart:
|
|
memset(&msg, 0, sizeof(msg));
|
|
msg.msg_iov = conn->read_msg.msg_iov;
|
|
msg.msg_iovlen = conn->read_msg.msg_iovlen;
|
|
first_len = msg.msg_iov->iov_len;
|
|
|
|
oldfs = get_fs();
|
|
set_fs(get_ds());
|
|
res = sock_recvmsg(conn->sock, &msg, conn->read_size,
|
|
MSG_DONTWAIT | MSG_NOSIGNAL);
|
|
set_fs(oldfs);
|
|
|
|
TRACE_DBG("msg_iovlen %zd, first_len %d, read_size %d, res %d",
|
|
msg.msg_iovlen, first_len, conn->read_size, res);
|
|
|
|
if (res > 0) {
|
|
/*
|
|
* To save some considerable effort and CPU power we
|
|
* suppose that TCP functions adjust
|
|
* conn->read_msg.msg_iov and conn->read_msg.msg_iovlen
|
|
* on amount of copied data. This BUG_ON is intended
|
|
* to catch if it is changed in the future.
|
|
*/
|
|
sBUG_ON((res >= first_len) &&
|
|
(conn->read_msg.msg_iov->iov_len != 0));
|
|
conn->read_size -= res;
|
|
if (conn->read_size != 0) {
|
|
if (res >= first_len) {
|
|
int done = 1 + ((res - first_len) >> PAGE_SHIFT);
|
|
TRACE_DBG("done %d", done);
|
|
conn->read_msg.msg_iov += done;
|
|
conn->read_msg.msg_iovlen -= done;
|
|
}
|
|
}
|
|
res = conn->read_size;
|
|
} else {
|
|
switch (res) {
|
|
case -EAGAIN:
|
|
TRACE_DBG("EAGAIN received for conn %p", conn);
|
|
res = conn->read_size;
|
|
break;
|
|
case -ERESTARTSYS:
|
|
TRACE_DBG("ERESTARTSYS received for conn %p", conn);
|
|
goto restart;
|
|
default:
|
|
if (!conn->closing) {
|
|
PRINT_ERROR("sock_recvmsg() failed: %d", res);
|
|
mark_conn_closed(conn);
|
|
}
|
|
if (res == 0)
|
|
res = -EIO;
|
|
break;
|
|
}
|
|
}
|
|
|
|
out:
|
|
TRACE_EXIT_RES(res);
|
|
return res;
|
|
}
|
|
|
|
static int iscsi_rx_check_ddigest(struct iscsi_conn *conn)
|
|
{
|
|
struct iscsi_cmnd *cmnd = conn->read_cmnd;
|
|
int res;
|
|
|
|
res = do_recv(conn);
|
|
if (res == 0) {
|
|
conn->read_state = RX_END;
|
|
|
|
if (cmnd->pdu.datasize <= 16*1024) {
|
|
/*
|
|
* It's cache hot, so let's compute it inline. The
|
|
* choice here about what will expose more latency:
|
|
* possible cache misses or the digest calculation.
|
|
*/
|
|
TRACE_DBG("cmnd %p, opcode %x: checking RX "
|
|
"ddigest inline", cmnd, cmnd_opcode(cmnd));
|
|
cmnd->ddigest_checked = 1;
|
|
res = digest_rx_data(cmnd);
|
|
if (unlikely(res != 0)) {
|
|
struct iscsi_cmnd *orig_req;
|
|
if (cmnd_opcode(cmnd) == ISCSI_OP_SCSI_DATA_OUT)
|
|
orig_req = cmnd->cmd_req;
|
|
else
|
|
orig_req = cmnd;
|
|
if (unlikely(orig_req->scst_cmd == NULL)) {
|
|
/* Just drop it */
|
|
iscsi_preliminary_complete(cmnd, orig_req, false);
|
|
} else {
|
|
set_scst_preliminary_status_rsp(orig_req, false,
|
|
SCST_LOAD_SENSE(iscsi_sense_crc_error));
|
|
/*
|
|
* Let's prelim complete cmnd too to
|
|
* handle the DATA OUT case
|
|
*/
|
|
iscsi_preliminary_complete(cmnd, orig_req, false);
|
|
}
|
|
res = 0;
|
|
}
|
|
} else if (cmnd_opcode(cmnd) == ISCSI_OP_SCSI_CMD) {
|
|
cmd_add_on_rx_ddigest_list(cmnd, cmnd);
|
|
cmnd_get(cmnd);
|
|
} else if (cmnd_opcode(cmnd) != ISCSI_OP_SCSI_DATA_OUT) {
|
|
/*
|
|
* We could get here only for Nop-Out. ISCSI RFC
|
|
* doesn't specify how to deal with digest errors in
|
|
* this case. Let's just drop the command.
|
|
*/
|
|
TRACE_DBG("cmnd %p, opcode %x: checking NOP RX "
|
|
"ddigest", cmnd, cmnd_opcode(cmnd));
|
|
res = digest_rx_data(cmnd);
|
|
if (unlikely(res != 0)) {
|
|
iscsi_preliminary_complete(cmnd, cmnd, false);
|
|
res = 0;
|
|
}
|
|
}
|
|
}
|
|
|
|
return res;
|
|
}
|
|
|
|
/* No locks, conn is rd processing */
|
|
static int process_read_io(struct iscsi_conn *conn, int *closed)
|
|
{
|
|
struct iscsi_cmnd *cmnd = conn->read_cmnd;
|
|
int res;
|
|
|
|
TRACE_ENTRY();
|
|
|
|
/* In case of error cmnd will be freed in close_conn() */
|
|
|
|
do {
|
|
switch (conn->read_state) {
|
|
case RX_INIT_BHS:
|
|
EXTRACHECKS_BUG_ON(conn->read_cmnd != NULL);
|
|
cmnd = conn->transport->iscsit_alloc_cmd(conn, NULL);
|
|
conn->read_cmnd = cmnd;
|
|
iscsi_conn_init_read(cmnd->conn,
|
|
(void __force __user *)&cmnd->pdu.bhs,
|
|
sizeof(cmnd->pdu.bhs));
|
|
conn->read_state = RX_BHS;
|
|
/* go through */
|
|
|
|
case RX_BHS:
|
|
res = do_recv(conn);
|
|
if (res == 0) {
|
|
/*
|
|
* This command not yet received on the aborted
|
|
* time, so shouldn't be affected by any abort.
|
|
*/
|
|
EXTRACHECKS_BUG_ON(cmnd->prelim_compl_flags != 0);
|
|
|
|
iscsi_cmnd_get_length(&cmnd->pdu);
|
|
|
|
if (cmnd->pdu.ahssize == 0) {
|
|
if ((conn->hdigest_type & DIGEST_NONE) == 0)
|
|
conn->read_state = RX_INIT_HDIGEST;
|
|
else
|
|
conn->read_state = RX_CMD_START;
|
|
} else {
|
|
iscsi_conn_prepare_read_ahs(conn, cmnd);
|
|
conn->read_state = RX_AHS;
|
|
}
|
|
}
|
|
break;
|
|
|
|
case RX_CMD_START:
|
|
res = cmnd_rx_start(cmnd);
|
|
if (res == 0) {
|
|
if (cmnd->pdu.datasize == 0)
|
|
conn->read_state = RX_END;
|
|
else
|
|
conn->read_state = RX_DATA;
|
|
} else if (res > 0)
|
|
conn->read_state = RX_CMD_CONTINUE;
|
|
else
|
|
sBUG_ON(!conn->closing);
|
|
break;
|
|
|
|
case RX_CMD_CONTINUE:
|
|
if (cmnd->scst_state == ISCSI_CMD_STATE_RX_CMD) {
|
|
TRACE_DBG("cmnd %p is still in RX_CMD state",
|
|
cmnd);
|
|
res = 1;
|
|
break;
|
|
}
|
|
res = cmnd_rx_continue(cmnd);
|
|
if (unlikely(res != 0))
|
|
sBUG_ON(!conn->closing);
|
|
else {
|
|
if (cmnd->pdu.datasize == 0)
|
|
conn->read_state = RX_END;
|
|
else
|
|
conn->read_state = RX_DATA;
|
|
}
|
|
break;
|
|
|
|
case RX_DATA:
|
|
res = do_recv(conn);
|
|
if (res == 0) {
|
|
int psz = ((cmnd->pdu.datasize + 3) & -4) - cmnd->pdu.datasize;
|
|
if (psz != 0) {
|
|
TRACE_DBG("padding %d bytes", psz);
|
|
iscsi_conn_init_read(conn,
|
|
(void __force __user *)&conn->rpadding, psz);
|
|
conn->read_state = RX_PADDING;
|
|
} else if ((conn->ddigest_type & DIGEST_NONE) != 0)
|
|
conn->read_state = RX_END;
|
|
else
|
|
conn->read_state = RX_INIT_DDIGEST;
|
|
}
|
|
break;
|
|
|
|
case RX_END:
|
|
if (unlikely(conn->read_size != 0)) {
|
|
PRINT_CRIT_ERROR("conn read_size !=0 on RX_END "
|
|
"(conn %p, op %x, read_size %d)", conn,
|
|
cmnd_opcode(cmnd), conn->read_size);
|
|
sBUG();
|
|
}
|
|
conn->read_cmnd = NULL;
|
|
conn->read_state = RX_INIT_BHS;
|
|
|
|
cmnd_rx_end(cmnd);
|
|
|
|
EXTRACHECKS_BUG_ON(conn->read_size != 0);
|
|
|
|
/*
|
|
* To maintain fairness. Res must be 0 here anyway, the
|
|
* assignment is only to remove compiler warning about
|
|
* uninitialized variable.
|
|
*/
|
|
res = 0;
|
|
goto out;
|
|
|
|
case RX_INIT_HDIGEST:
|
|
iscsi_conn_init_read(conn,
|
|
(void __force __user *)&cmnd->hdigest, sizeof(u32));
|
|
conn->read_state = RX_CHECK_HDIGEST;
|
|
/* go through */
|
|
|
|
case RX_CHECK_HDIGEST:
|
|
res = do_recv(conn);
|
|
if (res == 0) {
|
|
res = digest_rx_header(cmnd);
|
|
if (unlikely(res != 0)) {
|
|
PRINT_ERROR("rx header digest for "
|
|
"initiator %s failed (%d)",
|
|
conn->session->initiator_name,
|
|
res);
|
|
mark_conn_closed(conn);
|
|
} else
|
|
conn->read_state = RX_CMD_START;
|
|
}
|
|
break;
|
|
|
|
case RX_INIT_DDIGEST:
|
|
iscsi_conn_init_read(conn,
|
|
(void __force __user *)&cmnd->ddigest,
|
|
sizeof(u32));
|
|
conn->read_state = RX_CHECK_DDIGEST;
|
|
/* go through */
|
|
|
|
case RX_CHECK_DDIGEST:
|
|
res = iscsi_rx_check_ddigest(conn);
|
|
break;
|
|
|
|
case RX_AHS:
|
|
res = do_recv(conn);
|
|
if (res == 0) {
|
|
if ((conn->hdigest_type & DIGEST_NONE) == 0)
|
|
conn->read_state = RX_INIT_HDIGEST;
|
|
else
|
|
conn->read_state = RX_CMD_START;
|
|
}
|
|
break;
|
|
|
|
case RX_PADDING:
|
|
res = do_recv(conn);
|
|
if (res == 0) {
|
|
if ((conn->ddigest_type & DIGEST_NONE) == 0)
|
|
conn->read_state = RX_INIT_DDIGEST;
|
|
else
|
|
conn->read_state = RX_END;
|
|
}
|
|
break;
|
|
|
|
default:
|
|
PRINT_CRIT_ERROR("%d %x", conn->read_state, cmnd_opcode(cmnd));
|
|
res = -1; /* to keep compiler happy */
|
|
sBUG();
|
|
}
|
|
} while (res == 0);
|
|
|
|
if (unlikely(conn->closing)) {
|
|
start_close_conn(conn);
|
|
*closed = 1;
|
|
}
|
|
|
|
out:
|
|
TRACE_EXIT_RES(res);
|
|
return res;
|
|
}
|
|
|
|
/*
|
|
* Called under rd_lock and BHs disabled, but will drop it inside,
|
|
* then reacquire.
|
|
*/
|
|
static void scst_do_job_rd(struct iscsi_thread_pool *p)
|
|
__acquires(&rd_lock)
|
|
__releases(&rd_lock)
|
|
{
|
|
TRACE_ENTRY();
|
|
|
|
/*
|
|
* We delete/add to tail connections to maintain fairness between them.
|
|
*/
|
|
|
|
while (!list_empty(&p->rd_list)) {
|
|
int closed = 0, rc;
|
|
struct iscsi_conn *conn = list_first_entry(&p->rd_list,
|
|
typeof(*conn), rd_list_entry);
|
|
|
|
list_del(&conn->rd_list_entry);
|
|
|
|
sBUG_ON(conn->rd_state == ISCSI_CONN_RD_STATE_PROCESSING);
|
|
conn->rd_data_ready = 0;
|
|
conn->rd_state = ISCSI_CONN_RD_STATE_PROCESSING;
|
|
#ifdef CONFIG_SCST_EXTRACHECKS
|
|
conn->rd_task = current;
|
|
#endif
|
|
spin_unlock_bh(&p->rd_lock);
|
|
|
|
rc = process_read_io(conn, &closed);
|
|
|
|
spin_lock_bh(&p->rd_lock);
|
|
|
|
if (unlikely(closed))
|
|
continue;
|
|
|
|
if (unlikely(conn->conn_tm_active)) {
|
|
spin_unlock_bh(&p->rd_lock);
|
|
iscsi_check_tm_data_wait_timeouts(conn, false);
|
|
spin_lock_bh(&p->rd_lock);
|
|
}
|
|
|
|
#ifdef CONFIG_SCST_EXTRACHECKS
|
|
conn->rd_task = NULL;
|
|
#endif
|
|
if ((rc == 0) || conn->rd_data_ready) {
|
|
list_add_tail(&conn->rd_list_entry, &p->rd_list);
|
|
conn->rd_state = ISCSI_CONN_RD_STATE_IN_LIST;
|
|
} else
|
|
conn->rd_state = ISCSI_CONN_RD_STATE_IDLE;
|
|
}
|
|
|
|
TRACE_EXIT();
|
|
return;
|
|
}
|
|
|
|
static inline int test_rd_list(struct iscsi_thread_pool *p)
|
|
{
|
|
int res = !list_empty(&p->rd_list) ||
|
|
unlikely(kthread_should_stop());
|
|
return res;
|
|
}
|
|
|
|
int istrd(void *arg)
|
|
{
|
|
struct iscsi_thread_pool *p = arg;
|
|
int rc;
|
|
|
|
TRACE_ENTRY();
|
|
|
|
PRINT_INFO("Read thread for pool %p started, PID %d", p, current->pid);
|
|
|
|
current->flags |= PF_NOFREEZE;
|
|
#if defined(RHEL_MAJOR) && RHEL_MAJOR -0 <= 5
|
|
rc = set_cpus_allowed(current, p->cpu_mask);
|
|
#else
|
|
rc = set_cpus_allowed_ptr(current, &p->cpu_mask);
|
|
#endif
|
|
if (rc != 0)
|
|
PRINT_ERROR("Setting CPU affinity failed: %d", rc);
|
|
|
|
spin_lock_bh(&p->rd_lock);
|
|
while (!kthread_should_stop()) {
|
|
wait_event_locked(p->rd_waitQ, test_rd_list(p), lock_bh,
|
|
p->rd_lock);
|
|
scst_do_job_rd(p);
|
|
}
|
|
spin_unlock_bh(&p->rd_lock);
|
|
|
|
/*
|
|
* If kthread_should_stop() is true, we are guaranteed to be
|
|
* on the module unload, so rd_list must be empty.
|
|
*/
|
|
sBUG_ON(!list_empty(&p->rd_list));
|
|
|
|
PRINT_INFO("Read thread for PID %d for pool %p finished", current->pid, p);
|
|
|
|
TRACE_EXIT();
|
|
return 0;
|
|
}
|
|
|
|
#if defined(CONFIG_TCP_ZERO_COPY_TRANSFER_COMPLETION_NOTIFICATION)
|
|
static inline void __iscsi_get_page_callback(struct iscsi_cmnd *cmd)
|
|
{
|
|
int v;
|
|
|
|
TRACE_NET_PAGE("cmd %p, new net_ref_cnt %d",
|
|
cmd, atomic_read(&cmd->net_ref_cnt)+1);
|
|
|
|
v = atomic_inc_return(&cmd->net_ref_cnt);
|
|
if (v == 1) {
|
|
TRACE_NET_PAGE("getting cmd %p", cmd);
|
|
cmnd_get(cmd);
|
|
}
|
|
return;
|
|
}
|
|
|
|
void iscsi_get_page_callback(struct page *page)
|
|
{
|
|
struct iscsi_cmnd *cmd = (struct iscsi_cmnd *)page->net_priv;
|
|
|
|
TRACE_NET_PAGE("page %p, _count %d", page,
|
|
atomic_read(&page->_count));
|
|
|
|
__iscsi_get_page_callback(cmd);
|
|
return;
|
|
}
|
|
|
|
static inline void __iscsi_put_page_callback(struct iscsi_cmnd *cmd)
|
|
{
|
|
TRACE_NET_PAGE("cmd %p, new net_ref_cnt %d", cmd,
|
|
atomic_read(&cmd->net_ref_cnt)-1);
|
|
|
|
if (atomic_dec_and_test(&cmd->net_ref_cnt)) {
|
|
int i, sg_cnt = cmd->sg_cnt;
|
|
for (i = 0; i < sg_cnt; i++) {
|
|
struct page *page = sg_page(&cmd->sg[i]);
|
|
TRACE_NET_PAGE("Clearing page %p", page);
|
|
if (page->net_priv == cmd)
|
|
page->net_priv = NULL;
|
|
}
|
|
cmnd_put(cmd);
|
|
}
|
|
return;
|
|
}
|
|
|
|
void iscsi_put_page_callback(struct page *page)
|
|
{
|
|
struct iscsi_cmnd *cmd = (struct iscsi_cmnd *)page->net_priv;
|
|
|
|
TRACE_NET_PAGE("page %p, _count %d", page,
|
|
atomic_read(&page->_count));
|
|
|
|
__iscsi_put_page_callback(cmd);
|
|
return;
|
|
}
|
|
|
|
static void check_net_priv(struct iscsi_cmnd *cmd, struct page *page)
|
|
{
|
|
smp_rmb(); /* to sync with __iscsi_get_page_callback() */
|
|
if ((atomic_read(&cmd->net_ref_cnt) == 1) && (page->net_priv == cmd)) {
|
|
TRACE_DBG("sendpage() not called get_page(), zeroing net_priv "
|
|
"%p (page %p)", page->net_priv, page);
|
|
page->net_priv = NULL;
|
|
}
|
|
return;
|
|
}
|
|
#else
|
|
static inline void check_net_priv(struct iscsi_cmnd *cmd, struct page *page) {}
|
|
static inline void __iscsi_get_page_callback(struct iscsi_cmnd *cmd) {}
|
|
static inline void __iscsi_put_page_callback(struct iscsi_cmnd *cmd) {}
|
|
#endif
|
|
|
|
void req_add_to_write_timeout_list(struct iscsi_cmnd *req)
|
|
{
|
|
struct iscsi_conn *conn;
|
|
bool set_conn_tm_active = false;
|
|
|
|
TRACE_ENTRY();
|
|
|
|
if (req->on_write_timeout_list)
|
|
goto out;
|
|
|
|
conn = req->conn;
|
|
|
|
TRACE_DBG("Adding req %p to conn %p write_timeout_list",
|
|
req, conn);
|
|
|
|
spin_lock_bh(&conn->write_list_lock);
|
|
|
|
/* Recheck, since it can be changed behind us */
|
|
if (unlikely(req->on_write_timeout_list)) {
|
|
spin_unlock_bh(&conn->write_list_lock);
|
|
goto out;
|
|
}
|
|
|
|
req->on_write_timeout_list = 1;
|
|
req->write_start = jiffies;
|
|
|
|
if (unlikely(cmnd_opcode(req) == ISCSI_OP_NOP_OUT)) {
|
|
unsigned long req_tt = iscsi_get_timeout_time(req);
|
|
struct iscsi_cmnd *r;
|
|
bool inserted = false;
|
|
list_for_each_entry(r, &conn->write_timeout_list,
|
|
write_timeout_list_entry) {
|
|
unsigned long tt = iscsi_get_timeout_time(r);
|
|
if (time_after(tt, req_tt)) {
|
|
TRACE_DBG("Add NOP IN req %p (tt %ld) before "
|
|
"req %p (tt %ld)", req, req_tt, r, tt);
|
|
list_add_tail(&req->write_timeout_list_entry,
|
|
&r->write_timeout_list_entry);
|
|
inserted = true;
|
|
break;
|
|
} else
|
|
TRACE_DBG("Skipping op %x req %p (tt %ld)",
|
|
cmnd_opcode(r), r, tt);
|
|
}
|
|
if (!inserted) {
|
|
TRACE_DBG("Add NOP IN req %p in the tail", req);
|
|
list_add_tail(&req->write_timeout_list_entry,
|
|
&conn->write_timeout_list);
|
|
}
|
|
|
|
/* We suppose that nop_in_timeout must be <= data_rsp_timeout */
|
|
req_tt += ISCSI_ADD_SCHED_TIME;
|
|
if (timer_pending(&conn->rsp_timer) &&
|
|
time_after(conn->rsp_timer.expires, req_tt)) {
|
|
TRACE_DBG("Timer adjusted for sooner expired NOP IN "
|
|
"req %p", req);
|
|
mod_timer(&conn->rsp_timer, req_tt);
|
|
}
|
|
} else
|
|
list_add_tail(&req->write_timeout_list_entry,
|
|
&conn->write_timeout_list);
|
|
|
|
if (!timer_pending(&conn->rsp_timer)) {
|
|
unsigned long timeout_time;
|
|
if (unlikely(conn->conn_tm_active ||
|
|
test_bit(ISCSI_CMD_ABORTED,
|
|
&req->prelim_compl_flags))) {
|
|
set_conn_tm_active = true;
|
|
timeout_time = req->write_start +
|
|
ISCSI_TM_DATA_WAIT_TIMEOUT;
|
|
} else
|
|
timeout_time = iscsi_get_timeout_time(req);
|
|
|
|
timeout_time += ISCSI_ADD_SCHED_TIME;
|
|
|
|
TRACE_DBG("Starting timer on %ld (con %p, write_start %ld)",
|
|
timeout_time, conn, req->write_start);
|
|
|
|
conn->rsp_timer.expires = timeout_time;
|
|
add_timer(&conn->rsp_timer);
|
|
} else if (unlikely(test_bit(ISCSI_CMD_ABORTED,
|
|
&req->prelim_compl_flags))) {
|
|
unsigned long timeout_time = jiffies +
|
|
ISCSI_TM_DATA_WAIT_TIMEOUT + ISCSI_ADD_SCHED_TIME;
|
|
set_conn_tm_active = true;
|
|
if (time_after(conn->rsp_timer.expires, timeout_time)) {
|
|
TRACE_MGMT_DBG("Mod timer on %ld (conn %p)",
|
|
timeout_time, conn);
|
|
mod_timer(&conn->rsp_timer, timeout_time);
|
|
}
|
|
}
|
|
|
|
spin_unlock_bh(&conn->write_list_lock);
|
|
|
|
/*
|
|
* conn_tm_active can be already cleared by
|
|
* iscsi_check_tm_data_wait_timeouts(). write_list_lock is an inner
|
|
* lock for rd_lock.
|
|
*/
|
|
if (unlikely(set_conn_tm_active)) {
|
|
spin_lock_bh(&conn->conn_thr_pool->rd_lock);
|
|
TRACE_MGMT_DBG("Setting conn_tm_active for conn %p", conn);
|
|
conn->conn_tm_active = 1;
|
|
spin_unlock_bh(&conn->conn_thr_pool->rd_lock);
|
|
}
|
|
|
|
out:
|
|
TRACE_EXIT();
|
|
return;
|
|
}
|
|
|
|
static int write_data(struct iscsi_conn *conn)
|
|
{
|
|
mm_segment_t oldfs;
|
|
struct file *file;
|
|
struct iovec *iop;
|
|
struct socket *sock;
|
|
ssize_t (*sock_sendpage)(struct socket *, struct page *, int, size_t,
|
|
int);
|
|
ssize_t (*sendpage)(struct socket *, struct page *, int, size_t, int);
|
|
struct iscsi_cmnd *write_cmnd = conn->write_cmnd;
|
|
struct iscsi_cmnd *ref_cmd;
|
|
struct page *page;
|
|
struct scatterlist *sg;
|
|
int saved_size, size, sendsize;
|
|
int length, offset, idx;
|
|
int flags, res, count, sg_size;
|
|
bool do_put = false, ref_cmd_to_parent;
|
|
|
|
TRACE_ENTRY();
|
|
|
|
iscsi_extracheck_is_wr_thread(conn);
|
|
|
|
if (!write_cmnd->own_sg) {
|
|
ref_cmd = write_cmnd->parent_req;
|
|
ref_cmd_to_parent = true;
|
|
} else {
|
|
ref_cmd = write_cmnd;
|
|
ref_cmd_to_parent = false;
|
|
}
|
|
|
|
req_add_to_write_timeout_list(write_cmnd->parent_req);
|
|
|
|
file = conn->file;
|
|
size = conn->write_size;
|
|
saved_size = size;
|
|
iop = conn->write_iop;
|
|
count = conn->write_iop_used;
|
|
|
|
if (iop) {
|
|
while (1) {
|
|
loff_t off = 0;
|
|
int rest;
|
|
|
|
sBUG_ON(count > ARRAY_SIZE(conn->write_iov));
|
|
retry:
|
|
oldfs = get_fs();
|
|
set_fs(KERNEL_DS);
|
|
res = vfs_writev(file,
|
|
(struct iovec __force __user *)iop,
|
|
count, &off);
|
|
set_fs(oldfs);
|
|
TRACE_WRITE("sid %#Lx, cid %u, res %d, iov_len %zd",
|
|
(long long unsigned int)conn->session->sid,
|
|
conn->cid, res, iop->iov_len);
|
|
if (unlikely(res <= 0)) {
|
|
if (res == -EAGAIN) {
|
|
conn->write_iop = iop;
|
|
conn->write_iop_used = count;
|
|
goto out_iov;
|
|
} else if (res == -EINTR)
|
|
goto retry;
|
|
goto out_err;
|
|
}
|
|
|
|
rest = res;
|
|
size -= res;
|
|
while ((typeof(rest))iop->iov_len <= rest && rest) {
|
|
rest -= iop->iov_len;
|
|
iop++;
|
|
count--;
|
|
}
|
|
if (count == 0) {
|
|
conn->write_iop = NULL;
|
|
conn->write_iop_used = 0;
|
|
if (size)
|
|
break;
|
|
goto out_iov;
|
|
}
|
|
sBUG_ON(iop >
|
|
conn->write_iov + ARRAY_SIZE(conn->write_iov));
|
|
iop->iov_base += rest;
|
|
iop->iov_len -= rest;
|
|
}
|
|
}
|
|
|
|
sg = write_cmnd->sg;
|
|
if (unlikely(sg == NULL)) {
|
|
PRINT_INFO("WARNING: Data missed (cmd %p)!", write_cmnd);
|
|
res = 0;
|
|
goto out;
|
|
}
|
|
|
|
/* To protect from too early transfer completion race */
|
|
__iscsi_get_page_callback(ref_cmd);
|
|
do_put = true;
|
|
|
|
sock = conn->sock;
|
|
|
|
#if defined(CONFIG_TCP_ZERO_COPY_TRANSFER_COMPLETION_NOTIFICATION)
|
|
sock_sendpage = sock->ops->sendpage;
|
|
#else
|
|
if ((write_cmnd->parent_req->scst_cmd != NULL) &&
|
|
scst_cmd_get_dh_data_buff_alloced(write_cmnd->parent_req->scst_cmd))
|
|
sock_sendpage = sock_no_sendpage;
|
|
else
|
|
sock_sendpage = sock->ops->sendpage;
|
|
#endif
|
|
|
|
flags = MSG_DONTWAIT;
|
|
sg_size = size;
|
|
|
|
if (sg != write_cmnd->rsp_sg) {
|
|
offset = conn->write_offset + sg[0].offset;
|
|
idx = offset >> PAGE_SHIFT;
|
|
if (offset + sg[0].offset >= PAGE_SIZE)
|
|
offset += sg[0].offset;
|
|
offset &= ~PAGE_MASK;
|
|
length = min(size, (int)PAGE_SIZE - offset);
|
|
TRACE_WRITE("write_offset %d, sg_size %d, idx %d, offset %d, "
|
|
"length %d", conn->write_offset, sg_size, idx, offset,
|
|
length);
|
|
} else {
|
|
idx = 0;
|
|
offset = conn->write_offset;
|
|
while (offset >= sg[idx].length) {
|
|
offset -= sg[idx].length;
|
|
idx++;
|
|
}
|
|
length = sg[idx].length - offset;
|
|
offset += sg[idx].offset;
|
|
sock_sendpage = sock_no_sendpage;
|
|
TRACE_WRITE("rsp_sg: write_offset %d, sg_size %d, idx %d, "
|
|
"offset %d, length %d", conn->write_offset, sg_size,
|
|
idx, offset, length);
|
|
}
|
|
page = sg_page(&sg[idx]);
|
|
|
|
while (1) {
|
|
sendpage = sock_sendpage;
|
|
|
|
#if defined(CONFIG_TCP_ZERO_COPY_TRANSFER_COMPLETION_NOTIFICATION)
|
|
{
|
|
static DEFINE_SPINLOCK(net_priv_lock);
|
|
spin_lock(&net_priv_lock);
|
|
if (unlikely(page->net_priv != NULL)) {
|
|
if (page->net_priv != ref_cmd) {
|
|
/*
|
|
* This might happen if user space
|
|
* supplies to scst_user the same
|
|
* pages in different commands or in
|
|
* case of zero-copy FILEIO, when
|
|
* several initiators request the same
|
|
* data simultaneously.
|
|
*/
|
|
TRACE_DBG("net_priv isn't NULL and != "
|
|
"ref_cmd (write_cmnd %p, ref_cmd "
|
|
"%p, sg %p, idx %d, page %p, "
|
|
"net_priv %p)",
|
|
write_cmnd, ref_cmd, sg, idx,
|
|
page, page->net_priv);
|
|
sendpage = sock_no_sendpage;
|
|
}
|
|
} else
|
|
page->net_priv = ref_cmd;
|
|
spin_unlock(&net_priv_lock);
|
|
}
|
|
#endif
|
|
sendsize = min(size, length);
|
|
if (size <= sendsize) {
|
|
retry2:
|
|
res = sendpage(sock, page, offset, size, flags);
|
|
TRACE_WRITE("Final %s sid %#Lx, cid %u, res %d (page "
|
|
"index %lu, offset %u, size %u, cmd %p, "
|
|
"page %p)", (sendpage != sock_no_sendpage) ?
|
|
"sendpage" : "sock_no_sendpage",
|
|
(long long unsigned int)conn->session->sid,
|
|
conn->cid, res, page->index,
|
|
offset, size, write_cmnd, page);
|
|
if (unlikely(res <= 0)) {
|
|
if (res == -EINTR)
|
|
goto retry2;
|
|
else
|
|
goto out_res;
|
|
}
|
|
|
|
check_net_priv(ref_cmd, page);
|
|
if (res == size) {
|
|
conn->write_size = 0;
|
|
res = saved_size;
|
|
goto out_put;
|
|
}
|
|
|
|
offset += res;
|
|
size -= res;
|
|
goto retry2;
|
|
}
|
|
|
|
retry1:
|
|
res = sendpage(sock, page, offset, sendsize, flags | MSG_MORE);
|
|
TRACE_WRITE("%s sid %#Lx, cid %u, res %d (page index %lu, "
|
|
"offset %u, sendsize %u, size %u, cmd %p, page %p)",
|
|
(sendpage != sock_no_sendpage) ? "sendpage" :
|
|
"sock_no_sendpage",
|
|
(unsigned long long)conn->session->sid, conn->cid,
|
|
res, page->index, offset, sendsize, size,
|
|
write_cmnd, page);
|
|
if (unlikely(res <= 0)) {
|
|
if (res == -EINTR)
|
|
goto retry1;
|
|
else
|
|
goto out_res;
|
|
}
|
|
|
|
check_net_priv(ref_cmd, page);
|
|
|
|
size -= res;
|
|
|
|
if (res == sendsize) {
|
|
idx++;
|
|
EXTRACHECKS_BUG_ON(idx >= ref_cmd->sg_cnt);
|
|
page = sg_page(&sg[idx]);
|
|
length = sg[idx].length;
|
|
offset = sg[idx].offset;
|
|
} else {
|
|
offset += res;
|
|
sendsize -= res;
|
|
goto retry1;
|
|
}
|
|
}
|
|
|
|
out_off:
|
|
conn->write_offset += sg_size - size;
|
|
|
|
out_iov:
|
|
conn->write_size = size;
|
|
if ((saved_size == size) && res == -EAGAIN)
|
|
goto out_put;
|
|
|
|
res = saved_size - size;
|
|
|
|
out_put:
|
|
if (do_put)
|
|
__iscsi_put_page_callback(ref_cmd);
|
|
|
|
out:
|
|
TRACE_EXIT_RES(res);
|
|
return res;
|
|
|
|
out_res:
|
|
check_net_priv(ref_cmd, page);
|
|
if (res == -EAGAIN)
|
|
goto out_off;
|
|
/* else go through */
|
|
|
|
out_err:
|
|
#ifndef CONFIG_SCST_DEBUG
|
|
if (!conn->closing) {
|
|
#else
|
|
{
|
|
#endif
|
|
PRINT_ERROR("error %d at sid:cid %#Lx:%u, cmnd %p", res,
|
|
(long long unsigned int)conn->session->sid,
|
|
conn->cid, conn->write_cmnd);
|
|
}
|
|
if (ref_cmd_to_parent &&
|
|
((ref_cmd->scst_cmd != NULL) || (ref_cmd->scst_aen != NULL))) {
|
|
if (ref_cmd->scst_state == ISCSI_CMD_STATE_AEN)
|
|
scst_set_aen_delivery_status(ref_cmd->scst_aen,
|
|
SCST_AEN_RES_FAILED);
|
|
else
|
|
scst_set_delivery_status(ref_cmd->scst_cmd,
|
|
SCST_CMD_DELIVERY_FAILED);
|
|
}
|
|
goto out_put;
|
|
}
|
|
|
|
static int exit_tx(struct iscsi_conn *conn, int res)
|
|
{
|
|
iscsi_extracheck_is_wr_thread(conn);
|
|
|
|
switch (res) {
|
|
case -EAGAIN:
|
|
case -ERESTARTSYS:
|
|
break;
|
|
default:
|
|
#ifndef CONFIG_SCST_DEBUG
|
|
if (!conn->closing) {
|
|
#else
|
|
{
|
|
#endif
|
|
PRINT_ERROR("Sending data failed: initiator %s, "
|
|
"write_size %d, write_state %d, res %d",
|
|
conn->session->initiator_name,
|
|
conn->write_size,
|
|
conn->write_state, res);
|
|
}
|
|
conn->write_state = TX_END;
|
|
conn->write_size = 0;
|
|
mark_conn_closed(conn);
|
|
break;
|
|
}
|
|
return res;
|
|
}
|
|
|
|
static int tx_ddigest(struct iscsi_cmnd *cmnd, int state)
|
|
{
|
|
int res, rest = cmnd->conn->write_size;
|
|
struct msghdr msg = {.msg_flags = MSG_NOSIGNAL | MSG_DONTWAIT};
|
|
struct kvec iov;
|
|
|
|
iscsi_extracheck_is_wr_thread(cmnd->conn);
|
|
|
|
TRACE_DBG("Sending data digest %x (cmd %p)", cmnd->ddigest, cmnd);
|
|
|
|
iov.iov_base = (char *)(&cmnd->ddigest) + (sizeof(u32) - rest);
|
|
iov.iov_len = rest;
|
|
|
|
res = kernel_sendmsg(cmnd->conn->sock, &msg, &iov, 1, rest);
|
|
if (res > 0) {
|
|
cmnd->conn->write_size -= res;
|
|
if (!cmnd->conn->write_size)
|
|
cmnd->conn->write_state = state;
|
|
} else
|
|
res = exit_tx(cmnd->conn, res);
|
|
|
|
return res;
|
|
}
|
|
|
|
static void init_tx_hdigest(struct iscsi_cmnd *cmnd)
|
|
{
|
|
struct iscsi_conn *conn = cmnd->conn;
|
|
struct iovec *iop;
|
|
|
|
iscsi_extracheck_is_wr_thread(conn);
|
|
|
|
digest_tx_header(cmnd);
|
|
|
|
sBUG_ON(conn->write_iop_used >= ARRAY_SIZE(conn->write_iov));
|
|
|
|
iop = &conn->write_iop[conn->write_iop_used];
|
|
conn->write_iop_used++;
|
|
iop->iov_base = (void __force __user *)&(cmnd->hdigest);
|
|
iop->iov_len = sizeof(u32);
|
|
conn->write_size += sizeof(u32);
|
|
|
|
return;
|
|
}
|
|
|
|
static int tx_padding(struct iscsi_cmnd *cmnd, int state)
|
|
{
|
|
int res, rest = cmnd->conn->write_size;
|
|
struct msghdr msg = {.msg_flags = MSG_NOSIGNAL | MSG_DONTWAIT};
|
|
struct kvec iov;
|
|
static const uint32_t padding;
|
|
|
|
iscsi_extracheck_is_wr_thread(cmnd->conn);
|
|
|
|
TRACE_DBG("Sending %d padding bytes (cmd %p)", rest, cmnd);
|
|
|
|
iov.iov_base = (char *)(&padding) + (sizeof(uint32_t) - rest);
|
|
iov.iov_len = rest;
|
|
|
|
res = kernel_sendmsg(cmnd->conn->sock, &msg, &iov, 1, rest);
|
|
if (res > 0) {
|
|
cmnd->conn->write_size -= res;
|
|
if (!cmnd->conn->write_size)
|
|
cmnd->conn->write_state = state;
|
|
} else
|
|
res = exit_tx(cmnd->conn, res);
|
|
|
|
return res;
|
|
}
|
|
|
|
static int iscsi_do_send(struct iscsi_conn *conn, int state)
|
|
{
|
|
int res;
|
|
|
|
iscsi_extracheck_is_wr_thread(conn);
|
|
|
|
res = write_data(conn);
|
|
if (res > 0) {
|
|
if (!conn->write_size)
|
|
conn->write_state = state;
|
|
} else
|
|
res = exit_tx(conn, res);
|
|
|
|
return res;
|
|
}
|
|
|
|
/*
|
|
* No locks, conn is wr processing.
|
|
*
|
|
* IMPORTANT! Connection conn must be protected by additional conn_get()
|
|
* upon entrance in this function, because otherwise it could be destroyed
|
|
* inside as a result of cmnd release.
|
|
*/
|
|
int iscsi_send(struct iscsi_conn *conn)
|
|
{
|
|
struct iscsi_cmnd *cmnd = conn->write_cmnd;
|
|
int ddigest, res = 0;
|
|
|
|
TRACE_ENTRY();
|
|
|
|
TRACE_DBG("conn %p, write_cmnd %p", conn, cmnd);
|
|
|
|
iscsi_extracheck_is_wr_thread(conn);
|
|
|
|
ddigest = conn->ddigest_type != DIGEST_NONE ? 1 : 0;
|
|
|
|
switch (conn->write_state) {
|
|
case TX_INIT:
|
|
sBUG_ON(cmnd != NULL);
|
|
cmnd = conn->write_cmnd = iscsi_get_send_cmnd(conn);
|
|
if (!cmnd)
|
|
goto out;
|
|
cmnd_tx_start(cmnd);
|
|
if (!(conn->hdigest_type & DIGEST_NONE))
|
|
init_tx_hdigest(cmnd);
|
|
conn->write_state = TX_BHS_DATA;
|
|
case TX_BHS_DATA:
|
|
res = iscsi_do_send(conn, cmnd->pdu.datasize ?
|
|
TX_INIT_PADDING : TX_END);
|
|
if (res <= 0 || conn->write_state != TX_INIT_PADDING)
|
|
break;
|
|
case TX_INIT_PADDING:
|
|
cmnd->conn->write_size = ((cmnd->pdu.datasize + 3) & -4) -
|
|
cmnd->pdu.datasize;
|
|
if (cmnd->conn->write_size != 0)
|
|
conn->write_state = TX_PADDING;
|
|
else if (ddigest)
|
|
conn->write_state = TX_INIT_DDIGEST;
|
|
else
|
|
conn->write_state = TX_END;
|
|
break;
|
|
case TX_PADDING:
|
|
res = tx_padding(cmnd, ddigest ? TX_INIT_DDIGEST : TX_END);
|
|
if (res <= 0 || conn->write_state != TX_INIT_DDIGEST)
|
|
break;
|
|
case TX_INIT_DDIGEST:
|
|
cmnd->conn->write_size = sizeof(u32);
|
|
conn->write_state = TX_DDIGEST;
|
|
case TX_DDIGEST:
|
|
res = tx_ddigest(cmnd, TX_END);
|
|
break;
|
|
default:
|
|
PRINT_CRIT_ERROR("%d %d %x", res, conn->write_state,
|
|
cmnd_opcode(cmnd));
|
|
sBUG();
|
|
}
|
|
|
|
if (res == 0)
|
|
goto out;
|
|
|
|
if (conn->write_state != TX_END)
|
|
goto out;
|
|
|
|
if (unlikely(conn->write_size)) {
|
|
PRINT_CRIT_ERROR("%d %x %u", res, cmnd_opcode(cmnd),
|
|
conn->write_size);
|
|
sBUG();
|
|
}
|
|
cmnd_tx_end(cmnd);
|
|
|
|
rsp_cmnd_release(cmnd);
|
|
|
|
conn->write_cmnd = NULL;
|
|
conn->write_state = TX_INIT;
|
|
|
|
out:
|
|
TRACE_EXIT_RES(res);
|
|
return res;
|
|
}
|
|
|
|
/*
|
|
* Called under wr_lock and BHs disabled, but will drop it inside,
|
|
* then reacquire.
|
|
*/
|
|
static void scst_do_job_wr(struct iscsi_thread_pool *p)
|
|
__acquires(&wr_lock)
|
|
__releases(&wr_lock)
|
|
{
|
|
TRACE_ENTRY();
|
|
|
|
/*
|
|
* We delete/add to tail connections to maintain fairness between them.
|
|
*/
|
|
|
|
while (!list_empty(&p->wr_list)) {
|
|
int rc;
|
|
struct iscsi_conn *conn = list_first_entry(&p->wr_list,
|
|
typeof(*conn), wr_list_entry);
|
|
|
|
TRACE_DBG("conn %p, wr_state %x, wr_space_ready %d, "
|
|
"write ready %d", conn, conn->wr_state,
|
|
conn->wr_space_ready, test_write_ready(conn));
|
|
|
|
list_del(&conn->wr_list_entry);
|
|
|
|
sBUG_ON(conn->wr_state == ISCSI_CONN_WR_STATE_PROCESSING);
|
|
|
|
conn->wr_state = ISCSI_CONN_WR_STATE_PROCESSING;
|
|
conn->wr_space_ready = 0;
|
|
#ifdef CONFIG_SCST_EXTRACHECKS
|
|
conn->wr_task = current;
|
|
#endif
|
|
spin_unlock_bh(&p->wr_lock);
|
|
|
|
conn_get(conn);
|
|
|
|
rc = iscsi_send(conn);
|
|
|
|
spin_lock_bh(&p->wr_lock);
|
|
#ifdef CONFIG_SCST_EXTRACHECKS
|
|
conn->wr_task = NULL;
|
|
#endif
|
|
if ((rc == -EAGAIN) && !conn->wr_space_ready) {
|
|
TRACE_DBG("EAGAIN, setting WR_STATE_SPACE_WAIT "
|
|
"(conn %p)", conn);
|
|
conn->wr_state = ISCSI_CONN_WR_STATE_SPACE_WAIT;
|
|
} else if (test_write_ready(conn)) {
|
|
list_add_tail(&conn->wr_list_entry, &p->wr_list);
|
|
conn->wr_state = ISCSI_CONN_WR_STATE_IN_LIST;
|
|
} else
|
|
conn->wr_state = ISCSI_CONN_WR_STATE_IDLE;
|
|
|
|
conn_put(conn);
|
|
}
|
|
|
|
TRACE_EXIT();
|
|
return;
|
|
}
|
|
|
|
static inline int test_wr_list(struct iscsi_thread_pool *p)
|
|
{
|
|
int res = !list_empty(&p->wr_list) ||
|
|
unlikely(kthread_should_stop());
|
|
return res;
|
|
}
|
|
|
|
int istwr(void *arg)
|
|
{
|
|
struct iscsi_thread_pool *p = arg;
|
|
int rc;
|
|
|
|
TRACE_ENTRY();
|
|
|
|
PRINT_INFO("Write thread for pool %p started, PID %d", p, current->pid);
|
|
|
|
current->flags |= PF_NOFREEZE;
|
|
#if defined(RHEL_MAJOR) && RHEL_MAJOR -0 <= 5
|
|
rc = set_cpus_allowed(current, p->cpu_mask);
|
|
#else
|
|
rc = set_cpus_allowed_ptr(current, &p->cpu_mask);
|
|
#endif
|
|
if (rc != 0)
|
|
PRINT_ERROR("Setting CPU affinity failed: %d", rc);
|
|
|
|
spin_lock_bh(&p->wr_lock);
|
|
while (!kthread_should_stop()) {
|
|
wait_event_locked(p->wr_waitQ, test_wr_list(p), lock_bh,
|
|
p->wr_lock);
|
|
scst_do_job_wr(p);
|
|
}
|
|
spin_unlock_bh(&p->wr_lock);
|
|
|
|
/*
|
|
* If kthread_should_stop() is true, we are guaranteed to be
|
|
* on the module unload, so wr_list must be empty.
|
|
*/
|
|
sBUG_ON(!list_empty(&p->wr_list));
|
|
|
|
PRINT_INFO("Write thread PID %d for pool %p finished", current->pid, p);
|
|
|
|
TRACE_EXIT();
|
|
return 0;
|
|
}
|