Changes in 4.9.212 xfs: Sanity check flags of Q_XQUOTARM call powerpc/archrandom: fix arch_get_random_seed_int() mt7601u: fix bbp version check in mt7601u_wait_bbp_ready drm/sti: do not remove the drm_bridge that was never added drm/virtio: fix bounds check in virtio_gpu_cmd_get_capset() ALSA: hda: fix unused variable warning IB/rxe: replace kvfree with vfree ALSA: usb-audio: update quirk for B&W PX to remove microphone staging: comedi: ni_mio_common: protect register write overflow pwm: lpss: Release runtime-pm reference from the driver's remove callback mlxsw: reg: QEEC: Add minimum shaper fields pcrypt: use format specifier in kobject_add exportfs: fix 'passing zero to ERR_PTR()' warning drm/dp_mst: Skip validating ports during destruction, just ref net: phy: Fix not to call phy_resume() if PHY is not attached pinctrl: sh-pfc: r8a7740: Add missing REF125CK pin to gether_gmii group pinctrl: sh-pfc: r8a7740: Add missing LCD0 marks to lcd0_data24_1 group pinctrl: sh-pfc: r8a7791: Remove bogus ctrl marks from qspi_data4_b group pinctrl: sh-pfc: r8a7791: Remove bogus marks from vin1_b_data18 group pinctrl: sh-pfc: sh73a0: Add missing TO pin to tpu4_to3 group pinctrl: sh-pfc: r8a7794: Remove bogus IPSR9 field pinctrl: sh-pfc: sh7734: Add missing IPSR11 field pinctrl: sh-pfc: sh7269: Add missing PCIOR0 field pinctrl: sh-pfc: sh7734: Remove bogus IPSR10 value Input: nomadik-ske-keypad - fix a loop timeout test clk: highbank: fix refcount leak in hb_clk_init() clk: qoriq: fix refcount leak in clockgen_init() clk: socfpga: fix refcount leak clk: samsung: exynos4: fix refcount leak in exynos4_get_xom() clk: imx6q: fix refcount leak in imx6q_clocks_init() clk: imx6sx: fix refcount leak in imx6sx_clocks_init() clk: imx7d: fix refcount leak in imx7d_clocks_init() clk: vf610: fix refcount leak in vf610_clocks_init() clk: armada-370: fix refcount leak in a370_clk_init() clk: kirkwood: fix refcount leak in kirkwood_clk_init() clk: armada-xp: fix refcount leak in axp_clk_init() clk: dove: fix refcount leak in dove_clk_init() IB/usnic: Fix out of bounds index check in query pkey RDMA/ocrdma: Fix out of bounds index check in query pkey RDMA/qedr: Fix out of bounds index check in query pkey arm64: dts: apq8016-sbc: Increase load on l11 for SDCARD drm/etnaviv: NULL vs IS_ERR() buf in etnaviv_core_dump() media: s5p-jpeg: Correct step and max values for V4L2_CID_JPEG_RESTART_INTERVAL crypto: tgr192 - fix unaligned memory access ASoC: imx-sgtl5000: put of nodes if finding codec fails IB/iser: Pass the correct number of entries for dma mapped SGL rtc: cmos: ignore bogus century byte clk: sunxi-ng: sun8i-a23: Enable PLL-MIPI LDOs when ungating it iwlwifi: mvm: fix A-MPDU reference assignment tty: ipwireless: Fix potential NULL pointer dereference crypto: crypto4xx - Fix wrong ppc4xx_trng_probe()/ppc4xx_trng_remove() arguments ARM: dts: lpc32xx: add required clocks property to keypad device node ARM: dts: lpc32xx: reparent keypad controller to SIC1 ARM: dts: lpc32xx: fix ARM PrimeCell LCD controller variant ARM: dts: lpc32xx: fix ARM PrimeCell LCD controller clocks property ARM: dts: lpc32xx: phy3250: fix SD card regulator voltage iwlwifi: mvm: fix RSS config command staging: most: cdev: add missing check for cdev_add failure rtc: ds1672: fix unintended sign extension thermal: mediatek: fix register index error net: phy: fixed_phy: Fix fixed_phy not checking GPIO rtc: 88pm860x: fix unintended sign extension rtc: 88pm80x: fix unintended sign extension rtc: pm8xxx: fix unintended sign extension fbdev: chipsfb: remove set but not used variable 'size' iw_cxgb4: use tos when importing the endpoint iw_cxgb4: use tos when finding ipv6 routes pinctrl: sh-pfc: emev2: Add missing pinmux functions pinctrl: sh-pfc: r8a7791: Fix scifb2_data_c pin group pinctrl: sh-pfc: r8a7792: Fix vin1_data18_b pin group pinctrl: sh-pfc: sh73a0: Fix fsic_spdif pin groups usb: phy: twl6030-usb: fix possible use-after-free on remove block: don't use bio->bi_vcnt to figure out segment number keys: Timestamp new keys vfio_pci: Enable memory accesses before calling pci_map_rom dmaengine: mv_xor: Use correct device for DMA API cdc-wdm: pass return value of recover_from_urb_loss regulator: pv88060: Fix array out-of-bounds access regulator: pv88080: Fix array out-of-bounds access regulator: pv88090: Fix array out-of-bounds access net: dsa: qca8k: Enable delay for RGMII_ID mode drm/nouveau/bios/ramcfg: fix missing parentheses when calculating RON drm/nouveau/pmu: don't print reply values if exec is false ASoC: qcom: Fix of-node refcount unbalance in apq8016_sbc_parse_of() fs/nfs: Fix nfs_parse_devname to not modify it's argument NFS: Fix a soft lockup in the delegation recovery code clocksource/drivers/sun5i: Fail gracefully when clock rate is unavailable clocksource/drivers/exynos_mct: Fix error path in timer resources initialization mmc: sdhci-brcmstb: handle mmc_of_parse() errors during probe ARM: 8847/1: pm: fix HYP/SVC mode mismatch when MCPM is used ARM: 8848/1: virt: Align GIC version check with arm64 counterpart regulator: wm831x-dcdc: Fix list of wm831x_dcdc_ilim from mA to uA nios2: ksyms: Add missing symbol exports scsi: megaraid_sas: reduce module load time drivers/rapidio/rio_cm.c: fix potential oops in riocm_ch_listen() xen, cpu_hotplug: Prevent an out of bounds access net: sh_eth: fix a missing check of of_get_phy_mode media: ivtv: update *pos correctly in ivtv_read_pos() media: cx18: update *pos correctly in cx18_read_pos() media: wl128x: Fix an error code in fm_download_firmware() media: cx23885: check allocation return regulator: tps65086: Fix tps65086_ldoa1_ranges for selector 0xB jfs: fix bogus variable self-initialization tipc: tipc clang warning m68k: mac: Fix VIA timer counter accesses ARM: OMAP2+: Fix potentially uninitialized return value for _setup_reset() media: davinci-isif: avoid uninitialized variable use media: tw5864: Fix possible NULL pointer dereference in tw5864_handle_frame spi: tegra114: clear packed bit for unpacked mode spi: tegra114: fix for unpacked mode transfers soc/fsl/qe: Fix an error code in qe_pin_request() spi: bcm2835aux: fix driver to not allow 65535 (=-1) cs-gpios ehea: Fix a copy-paste err in ehea_init_port_res scsi: qla2xxx: Unregister chrdev if module initialization fails ARM: pxa: ssp: Fix "WARNING: invalid free of devm_ allocated data" hwmon: (w83627hf) Use request_muxed_region for Super-IO accesses tipc: set sysctl_tipc_rmem and named_timeout right range powerpc: vdso: Make vdso32 installation conditional in vdso_install ARM: dts: ls1021: Fix SGMII PCS link remaining down after PHY disconnect media: ov2659: fix unbalanced mutex_lock/unlock 6lowpan: Off by one handling ->nexthdr dmaengine: axi-dmac: Don't check the number of frames for alignment ALSA: usb-audio: Handle the error from snd_usb_mixer_apply_create_quirk() packet: in recvmsg msg_name return at least sizeof sockaddr_ll ASoC: fix valid stream condition usb: gadget: fsl: fix link error against usb-gadget module IB/mlx5: Add missing XRC options to QP optional params mask iommu/vt-d: Make kernel parameter igfx_off work with vIOMMU net: ena: fix swapped parameters when calling ena_com_indirect_table_fill_entry net: ena: fix: Free napi resources when ena_up() fails net: ena: fix incorrect test of supported hash function net: ena: fix ena_com_fill_hash_function() implementation dmaengine: tegra210-adma: restore channel status l2tp: Fix possible NULL pointer dereference media: omap_vout: potential buffer overflow in vidioc_dqbuf() media: davinci/vpbe: array underflow in vpbe_enum_outputs() platform/x86: alienware-wmi: printing the wrong error code netfilter: ebtables: CONFIG_COMPAT: reject trailing data after last rule pwm: meson: Don't disable PWM when setting duty repeatedly ARM: riscpc: fix lack of keyboard interrupts after irq conversion kdb: do a sanity check on the cpu in kdb_per_cpu() backlight: lm3630a: Return 0 on success in update_status functions thermal: cpu_cooling: Actually trace CPU load in thermal_power_cpu_get_power dmaengine: tegra210-adma: Fix crash during probe spi: spi-fsl-spi: call spi_finalize_current_message() at the end crypto: ccp - fix AES CFB error exposed by new test vectors serial: stm32: fix transmit_chars when tx is stopped misc: sgi-xp: Properly initialize buf in xpc_get_rsvd_page_pa iommu: Use right function to get group for device signal/cifs: Fix cifs_put_tcp_session to call send_sig instead of force_sig inet: frags: call inet_frags_fini() after unregister_pernet_subsys() media: vivid: fix incorrect assignment operation when setting video mode powerpc/cacheinfo: add cacheinfo_teardown, cacheinfo_rebuild drm/msm/mdp5: Fix mdp5_cfg_init error return net: netem: fix backlog accounting for corrupted GSO frames net/af_iucv: always register net_device notifier ASoC: ti: davinci-mcasp: Fix slot mask settings when using multiple AXRs rtc: pcf8563: Clear event flags and disable interrupts before requesting irq drm/msm/a3xx: remove TPL1 regs from snapshot perf/ioctl: Add check for the sample_period value dmaengine: hsu: Revert "set HSU_CH_MTSR to memory width" clk: qcom: Fix -Wunused-const-variable iommu/amd: Make iommu_disable safer mfd: intel-lpss: Release IDA resources rxrpc: Fix uninitialized error code in rxrpc_send_data_packet() devres: allow const resource arguments RDMA/hns: Fixs hw access invalid dma memory error net: pasemi: fix an use-after-free in pasemi_mac_phy_init() scsi: libfc: fix null pointer dereference on a null lport libertas_tf: Use correct channel range in lbtf_geo_init qed: reduce maximum stack frame size usb: host: xhci-hub: fix extra endianness conversion mic: avoid statically declaring a 'struct device'. x86/kgbd: Use NMI_VECTOR not APIC_DM_NMI ALSA: aoa: onyx: always initialize register read value net/mlx5: Fix mlx5_ifc_query_lag_out_bits cifs: fix rmmod regression in cifs.ko caused by force_sig changes crypto: caam - free resources in case caam_rng registration failed ext4: set error return correctly when ext4_htree_store_dirent fails ASoC: es8328: Fix copy-paste error in es8328_right_line_controls ASoC: cs4349: Use PM ops 'cs4349_runtime_pm' ASoC: wm8737: Fix copy-paste error in wm8737_snd_controls signal: Allow cifs and drbd to receive their terminating signals ASoC: sun4i-i2s: RX and TX counter registers are swapped dmaengine: dw: platform: Switch to acpi_dma_controller_register() mac80211: minstrel_ht: fix per-group max throughput rate initialization mips: avoid explicit UB in assignment of mips_io_port_base ahci: Do not export local variable ahci_em_messages Partially revert "kfifo: fix kfifo_alloc() and kfifo_init()" hwmon: (lm75) Fix write operations for negative temperatures power: supply: Init device wakeup after device_add() x86, perf: Fix the dependency of the x86 insn decoder selftest staging: greybus: light: fix a couple double frees bcma: fix incorrect update of BCMA_CORE_PCI_MDIO_DATA iio: dac: ad5380: fix incorrect assignment to val ath9k: dynack: fix possible deadlock in ath_dynack_node_{de}init net: sonic: return NETDEV_TX_OK if failed to map buffer Btrfs: fix hang when loading existing inode cache off disk hwmon: (shtc1) fix shtc1 and shtw1 id mask net: sonic: replace dev_kfree_skb in sonic_send_packet net/rds: Fix 'ib_evt_handler_call' element in 'rds_ib_stat_names' iommu/amd: Wait for completion of IOTLB flush in attach_device net: hisilicon: Fix signedness bug in hix5hd2_dev_probe() net: broadcom/bcmsysport: Fix signedness in bcm_sysport_probe() net: stmmac: dwmac-meson8b: Fix signedness bug in probe of: mdio: Fix a signedness bug in of_phy_get_and_connect() net: ethernet: stmmac: Fix signedness bug in ipq806x_gmac_of_parse() nvme: retain split access workaround for capability reads net: stmmac: gmac4+: Not all Unicast addresses may be available mac80211: accept deauth frames in IBSS mode llc: fix another potential sk_buff leak in llc_ui_sendmsg() llc: fix sk_buff refcounting in llc_conn_state_process() net: stmmac: fix length of PTP clock's name string act_mirred: Fix mirred_init_module error handling drm/msm/dsi: Implement reset correctly dmaengine: imx-sdma: fix size check for sdma script_number net: netem: fix error path for corrupted GSO frames net: netem: correct the parent's backlog when corrupted packet was dropped net: qca_spi: Move reset_count to struct qcaspi afs: Fix large file support media: ov6650: Fix incorrect use of JPEG colorspace media: ov6650: Fix some format attributes not under control media: ov6650: Fix .get_fmt() V4L2_SUBDEV_FORMAT_TRY support MIPS: Loongson: Fix return value of loongson_hwmon_init net: neigh: use long type to store jiffies delta packet: fix data-race in fanout_flow_is_huge() dmaengine: ti: edma: fix missed failure handling drm/radeon: fix bad DMA from INTERRUPT_CNTL2 arm64: dts: juno: Fix UART frequency IB/iser: Fix dma_nents type definition m68k: Call timer_interrupt() with interrupts disabled net: ethtool: Add back transceiver type net: phy: Keep reporting transceiver type can, slip: Protect tty->disc_data in write_wakeup and close with RCU firestream: fix memory leaks net: cxgb3_main: Add CAP_NET_ADMIN check to CHELSIO_GET_MEM net, ip6_tunnel: fix namespaces move net, ip_tunnel: fix namespaces move net_sched: fix datalen for ematch tcp_bbr: improve arithmetic division in bbr_update_bw() net: usb: lan78xx: Add .ndo_features_check gtp: make sure only SOCK_DGRAM UDP sockets are accepted hwmon: (adt7475) Make volt2reg return same reg as reg2volt input hwmon: (core) Simplify sysfs attribute name allocation hwmon: Deal with errors from the thermal subsystem hwmon: (core) Fix double-free in __hwmon_device_register() hwmon: (core) Do not use device managed functions for memory allocations Input: keyspan-remote - fix control-message timeouts ARM: 8950/1: ftrace/recordmcount: filter relocation types mmc: tegra: fix SDR50 tuning override mmc: sdhci: fix minimum clock rate for v3 controller Input: sur40 - fix interface sanity checks Input: gtco - fix endpoint sanity check Input: aiptek - fix endpoint sanity check Input: pegasus_notetaker - fix endpoint sanity check Input: sun4i-ts - add a check for devm_thermal_zone_of_sensor_register hwmon: (nct7802) Fix voltage limits to wrong registers scsi: RDMA/isert: Fix a recently introduced regression related to logout tracing: xen: Ordered comparison of function pointers do_last(): fetch directory ->i_mode and ->i_uid before it's too late Documentation: Document arm64 kpti control arm64: kpti: Whitelist Cortex-A CPUs that don't implement the CSV3 field coresight: etb10: Do not call smp_processor_id from preemptible coresight: tmc-etf: Do not call smp_processor_id from preemptible libertas: Fix two buffer overflows at parsing bss descriptor bcache: silence static checker warning scsi: iscsi: Avoid potential deadlock in iscsi_if_rx func md: Avoid namespace collision with bitmap API bitmap: Add bitmap_alloc(), bitmap_zalloc() and bitmap_free() netfilter: ipset: use bitmap infrastructure completely net/x25: fix nonblocking connect Linux 4.9.212 Signed-off-by: Greg Kroah-Hartman <gregkh@google.com> Change-Id: I2e83a05c5f119a7467a4d6984045d45d0c06b764
600 lines
14 KiB
C
600 lines
14 KiB
C
/*
|
|
* IPv6 fragment reassembly
|
|
* Linux INET6 implementation
|
|
*
|
|
* Authors:
|
|
* Pedro Roque <roque@di.fc.ul.pt>
|
|
*
|
|
* Based on: net/ipv4/ip_fragment.c
|
|
*
|
|
* This program is free software; you can redistribute it and/or
|
|
* modify it under the terms of the GNU General Public License
|
|
* as published by the Free Software Foundation; either version
|
|
* 2 of the License, or (at your option) any later version.
|
|
*/
|
|
|
|
/*
|
|
* Fixes:
|
|
* Andi Kleen Make it work with multiple hosts.
|
|
* More RFC compliance.
|
|
*
|
|
* Horst von Brand Add missing #include <linux/string.h>
|
|
* Alexey Kuznetsov SMP races, threading, cleanup.
|
|
* Patrick McHardy LRU queue of frag heads for evictor.
|
|
* Mitsuru KANDA @USAGI Register inet6_protocol{}.
|
|
* David Stevens and
|
|
* YOSHIFUJI,H. @USAGI Always remove fragment header to
|
|
* calculate ICV correctly.
|
|
*/
|
|
|
|
#define pr_fmt(fmt) "IPv6: " fmt
|
|
|
|
#include <linux/errno.h>
|
|
#include <linux/types.h>
|
|
#include <linux/string.h>
|
|
#include <linux/socket.h>
|
|
#include <linux/sockios.h>
|
|
#include <linux/jiffies.h>
|
|
#include <linux/net.h>
|
|
#include <linux/list.h>
|
|
#include <linux/netdevice.h>
|
|
#include <linux/in6.h>
|
|
#include <linux/ipv6.h>
|
|
#include <linux/icmpv6.h>
|
|
#include <linux/random.h>
|
|
#include <linux/jhash.h>
|
|
#include <linux/skbuff.h>
|
|
#include <linux/slab.h>
|
|
#include <linux/export.h>
|
|
|
|
#include <net/sock.h>
|
|
#include <net/snmp.h>
|
|
|
|
#include <net/ipv6.h>
|
|
#include <net/ip6_route.h>
|
|
#include <net/protocol.h>
|
|
#include <net/transp_v6.h>
|
|
#include <net/rawv6.h>
|
|
#include <net/ndisc.h>
|
|
#include <net/addrconf.h>
|
|
#include <net/ipv6_frag.h>
|
|
#include <net/inet_ecn.h>
|
|
|
|
static const char ip6_frag_cache_name[] = "ip6-frags";
|
|
|
|
static u8 ip6_frag_ecn(const struct ipv6hdr *ipv6h)
|
|
{
|
|
return 1 << (ipv6_get_dsfield(ipv6h) & INET_ECN_MASK);
|
|
}
|
|
|
|
static struct inet_frags ip6_frags;
|
|
|
|
static int ip6_frag_reasm(struct frag_queue *fq, struct sk_buff *skb,
|
|
struct sk_buff *prev_tail, struct net_device *dev);
|
|
|
|
static void ip6_frag_expire(unsigned long data)
|
|
{
|
|
struct frag_queue *fq;
|
|
struct net *net;
|
|
|
|
fq = container_of((struct inet_frag_queue *)data, struct frag_queue, q);
|
|
net = container_of(fq->q.net, struct net, ipv6.frags);
|
|
|
|
ip6frag_expire_frag_queue(net, fq);
|
|
}
|
|
|
|
static struct frag_queue *
|
|
fq_find(struct net *net, __be32 id, const struct ipv6hdr *hdr, int iif)
|
|
{
|
|
struct frag_v6_compare_key key = {
|
|
.id = id,
|
|
.saddr = hdr->saddr,
|
|
.daddr = hdr->daddr,
|
|
.user = IP6_DEFRAG_LOCAL_DELIVER,
|
|
.iif = iif,
|
|
};
|
|
struct inet_frag_queue *q;
|
|
|
|
if (!(ipv6_addr_type(&hdr->daddr) & (IPV6_ADDR_MULTICAST |
|
|
IPV6_ADDR_LINKLOCAL)))
|
|
key.iif = 0;
|
|
|
|
q = inet_frag_find(&net->ipv6.frags, &key);
|
|
if (!q)
|
|
return NULL;
|
|
|
|
return container_of(q, struct frag_queue, q);
|
|
}
|
|
|
|
static int ip6_frag_queue(struct frag_queue *fq, struct sk_buff *skb,
|
|
struct frag_hdr *fhdr, int nhoff,
|
|
u32 *prob_offset)
|
|
{
|
|
struct net *net = dev_net(skb_dst(skb)->dev);
|
|
int offset, end, fragsize;
|
|
struct sk_buff *prev_tail;
|
|
struct net_device *dev;
|
|
int err = -ENOENT;
|
|
u8 ecn;
|
|
|
|
if (fq->q.flags & INET_FRAG_COMPLETE)
|
|
goto err;
|
|
|
|
err = -EINVAL;
|
|
offset = ntohs(fhdr->frag_off) & ~0x7;
|
|
end = offset + (ntohs(ipv6_hdr(skb)->payload_len) -
|
|
((u8 *)(fhdr + 1) - (u8 *)(ipv6_hdr(skb) + 1)));
|
|
|
|
if ((unsigned int)end > IPV6_MAXPLEN) {
|
|
*prob_offset = (u8 *)&fhdr->frag_off - skb_network_header(skb);
|
|
/* note that if prob_offset is set, the skb is freed elsewhere,
|
|
* we do not free it here.
|
|
*/
|
|
return -1;
|
|
}
|
|
|
|
ecn = ip6_frag_ecn(ipv6_hdr(skb));
|
|
|
|
if (skb->ip_summed == CHECKSUM_COMPLETE) {
|
|
const unsigned char *nh = skb_network_header(skb);
|
|
skb->csum = csum_sub(skb->csum,
|
|
csum_partial(nh, (u8 *)(fhdr + 1) - nh,
|
|
0));
|
|
}
|
|
|
|
/* Is this the final fragment? */
|
|
if (!(fhdr->frag_off & htons(IP6_MF))) {
|
|
/* If we already have some bits beyond end
|
|
* or have different end, the segment is corrupted.
|
|
*/
|
|
if (end < fq->q.len ||
|
|
((fq->q.flags & INET_FRAG_LAST_IN) && end != fq->q.len))
|
|
goto discard_fq;
|
|
fq->q.flags |= INET_FRAG_LAST_IN;
|
|
fq->q.len = end;
|
|
} else {
|
|
/* Check if the fragment is rounded to 8 bytes.
|
|
* Required by the RFC.
|
|
*/
|
|
if (end & 0x7) {
|
|
/* RFC2460 says always send parameter problem in
|
|
* this case. -DaveM
|
|
*/
|
|
*prob_offset = offsetof(struct ipv6hdr, payload_len);
|
|
return -1;
|
|
}
|
|
if (end > fq->q.len) {
|
|
/* Some bits beyond end -> corruption. */
|
|
if (fq->q.flags & INET_FRAG_LAST_IN)
|
|
goto discard_fq;
|
|
fq->q.len = end;
|
|
}
|
|
}
|
|
|
|
if (end == offset)
|
|
goto discard_fq;
|
|
|
|
err = -ENOMEM;
|
|
/* Point into the IP datagram 'data' part. */
|
|
if (!pskb_pull(skb, (u8 *) (fhdr + 1) - skb->data))
|
|
goto discard_fq;
|
|
|
|
err = pskb_trim_rcsum(skb, end - offset);
|
|
if (err)
|
|
goto discard_fq;
|
|
|
|
/* Note : skb->rbnode and skb->dev share the same location. */
|
|
dev = skb->dev;
|
|
/* Makes sure compiler wont do silly aliasing games */
|
|
barrier();
|
|
|
|
prev_tail = fq->q.fragments_tail;
|
|
err = inet_frag_queue_insert(&fq->q, skb, offset, end);
|
|
if (err)
|
|
goto insert_error;
|
|
|
|
if (dev)
|
|
fq->iif = dev->ifindex;
|
|
|
|
fq->q.stamp = skb->tstamp;
|
|
fq->q.meat += skb->len;
|
|
fq->ecn |= ecn;
|
|
add_frag_mem_limit(fq->q.net, skb->truesize);
|
|
|
|
fragsize = -skb_network_offset(skb) + skb->len;
|
|
if (fragsize > fq->q.max_size)
|
|
fq->q.max_size = fragsize;
|
|
|
|
/* The first fragment.
|
|
* nhoffset is obtained from the first fragment, of course.
|
|
*/
|
|
if (offset == 0) {
|
|
fq->nhoffset = nhoff;
|
|
fq->q.flags |= INET_FRAG_FIRST_IN;
|
|
}
|
|
|
|
if (fq->q.flags == (INET_FRAG_FIRST_IN | INET_FRAG_LAST_IN) &&
|
|
fq->q.meat == fq->q.len) {
|
|
unsigned long orefdst = skb->_skb_refdst;
|
|
|
|
skb->_skb_refdst = 0UL;
|
|
err = ip6_frag_reasm(fq, skb, prev_tail, dev);
|
|
skb->_skb_refdst = orefdst;
|
|
return err;
|
|
}
|
|
|
|
skb_dst_drop(skb);
|
|
return -EINPROGRESS;
|
|
|
|
insert_error:
|
|
if (err == IPFRAG_DUP) {
|
|
kfree_skb(skb);
|
|
return -EINVAL;
|
|
}
|
|
err = -EINVAL;
|
|
__IP6_INC_STATS(net, ip6_dst_idev(skb_dst(skb)),
|
|
IPSTATS_MIB_REASM_OVERLAPS);
|
|
discard_fq:
|
|
inet_frag_kill(&fq->q);
|
|
__IP6_INC_STATS(net, ip6_dst_idev(skb_dst(skb)),
|
|
IPSTATS_MIB_REASMFAILS);
|
|
err:
|
|
kfree_skb(skb);
|
|
return err;
|
|
}
|
|
|
|
/*
|
|
* Check if this packet is complete.
|
|
*
|
|
* It is called with locked fq, and caller must check that
|
|
* queue is eligible for reassembly i.e. it is not COMPLETE,
|
|
* the last and the first frames arrived and all the bits are here.
|
|
*/
|
|
static int ip6_frag_reasm(struct frag_queue *fq, struct sk_buff *skb,
|
|
struct sk_buff *prev_tail, struct net_device *dev)
|
|
{
|
|
struct net *net = container_of(fq->q.net, struct net, ipv6.frags);
|
|
unsigned int nhoff;
|
|
void *reasm_data;
|
|
int payload_len;
|
|
u8 ecn;
|
|
|
|
inet_frag_kill(&fq->q);
|
|
|
|
ecn = ip_frag_ecn_table[fq->ecn];
|
|
if (unlikely(ecn == 0xff))
|
|
goto out_fail;
|
|
|
|
reasm_data = inet_frag_reasm_prepare(&fq->q, skb, prev_tail);
|
|
if (!reasm_data)
|
|
goto out_oom;
|
|
|
|
payload_len = ((skb->data - skb_network_header(skb)) -
|
|
sizeof(struct ipv6hdr) + fq->q.len -
|
|
sizeof(struct frag_hdr));
|
|
if (payload_len > IPV6_MAXPLEN)
|
|
goto out_oversize;
|
|
|
|
/* We have to remove fragment header from datagram and to relocate
|
|
* header in order to calculate ICV correctly. */
|
|
nhoff = fq->nhoffset;
|
|
skb_network_header(skb)[nhoff] = skb_transport_header(skb)[0];
|
|
memmove(skb->head + sizeof(struct frag_hdr), skb->head,
|
|
(skb->data - skb->head) - sizeof(struct frag_hdr));
|
|
if (skb_mac_header_was_set(skb))
|
|
skb->mac_header += sizeof(struct frag_hdr);
|
|
skb->network_header += sizeof(struct frag_hdr);
|
|
|
|
skb_reset_transport_header(skb);
|
|
|
|
inet_frag_reasm_finish(&fq->q, skb, reasm_data);
|
|
|
|
skb->dev = dev;
|
|
ipv6_hdr(skb)->payload_len = htons(payload_len);
|
|
ipv6_change_dsfield(ipv6_hdr(skb), 0xff, ecn);
|
|
IP6CB(skb)->nhoff = nhoff;
|
|
IP6CB(skb)->flags |= IP6SKB_FRAGMENTED;
|
|
IP6CB(skb)->frag_max_size = fq->q.max_size;
|
|
|
|
/* Yes, and fold redundant checksum back. 8) */
|
|
skb_postpush_rcsum(skb, skb_network_header(skb),
|
|
skb_network_header_len(skb));
|
|
|
|
rcu_read_lock();
|
|
__IP6_INC_STATS(net, __in6_dev_get(dev), IPSTATS_MIB_REASMOKS);
|
|
rcu_read_unlock();
|
|
fq->q.fragments = NULL;
|
|
fq->q.rb_fragments = RB_ROOT;
|
|
fq->q.fragments_tail = NULL;
|
|
fq->q.last_run_head = NULL;
|
|
return 1;
|
|
|
|
out_oversize:
|
|
net_dbg_ratelimited("ip6_frag_reasm: payload len = %d\n", payload_len);
|
|
goto out_fail;
|
|
out_oom:
|
|
net_dbg_ratelimited("ip6_frag_reasm: no memory for reassembly\n");
|
|
out_fail:
|
|
rcu_read_lock();
|
|
__IP6_INC_STATS(net, __in6_dev_get(dev), IPSTATS_MIB_REASMFAILS);
|
|
rcu_read_unlock();
|
|
inet_frag_kill(&fq->q);
|
|
return -1;
|
|
}
|
|
|
|
static int ipv6_frag_rcv(struct sk_buff *skb)
|
|
{
|
|
struct frag_hdr *fhdr;
|
|
struct frag_queue *fq;
|
|
const struct ipv6hdr *hdr = ipv6_hdr(skb);
|
|
struct net *net = dev_net(skb_dst(skb)->dev);
|
|
int iif;
|
|
|
|
if (IP6CB(skb)->flags & IP6SKB_FRAGMENTED)
|
|
goto fail_hdr;
|
|
|
|
__IP6_INC_STATS(net, ip6_dst_idev(skb_dst(skb)), IPSTATS_MIB_REASMREQDS);
|
|
|
|
/* Jumbo payload inhibits frag. header */
|
|
if (hdr->payload_len == 0)
|
|
goto fail_hdr;
|
|
|
|
if (!pskb_may_pull(skb, (skb_transport_offset(skb) +
|
|
sizeof(struct frag_hdr))))
|
|
goto fail_hdr;
|
|
|
|
hdr = ipv6_hdr(skb);
|
|
fhdr = (struct frag_hdr *)skb_transport_header(skb);
|
|
|
|
if (!(fhdr->frag_off & htons(0xFFF9))) {
|
|
/* It is not a fragmented frame */
|
|
skb->transport_header += sizeof(struct frag_hdr);
|
|
__IP6_INC_STATS(net,
|
|
ip6_dst_idev(skb_dst(skb)), IPSTATS_MIB_REASMOKS);
|
|
|
|
IP6CB(skb)->nhoff = (u8 *)fhdr - skb_network_header(skb);
|
|
IP6CB(skb)->flags |= IP6SKB_FRAGMENTED;
|
|
return 1;
|
|
}
|
|
|
|
iif = skb->dev ? skb->dev->ifindex : 0;
|
|
fq = fq_find(net, fhdr->identification, hdr, iif);
|
|
if (fq) {
|
|
u32 prob_offset = 0;
|
|
int ret;
|
|
|
|
spin_lock(&fq->q.lock);
|
|
|
|
fq->iif = iif;
|
|
ret = ip6_frag_queue(fq, skb, fhdr, IP6CB(skb)->nhoff,
|
|
&prob_offset);
|
|
|
|
spin_unlock(&fq->q.lock);
|
|
inet_frag_put(&fq->q);
|
|
if (prob_offset) {
|
|
__IP6_INC_STATS(net, ip6_dst_idev(skb_dst(skb)),
|
|
IPSTATS_MIB_INHDRERRORS);
|
|
/* icmpv6_param_prob() calls kfree_skb(skb) */
|
|
icmpv6_param_prob(skb, ICMPV6_HDR_FIELD, prob_offset);
|
|
}
|
|
return ret;
|
|
}
|
|
|
|
__IP6_INC_STATS(net, ip6_dst_idev(skb_dst(skb)), IPSTATS_MIB_REASMFAILS);
|
|
kfree_skb(skb);
|
|
return -1;
|
|
|
|
fail_hdr:
|
|
__IP6_INC_STATS(net, ip6_dst_idev(skb_dst(skb)),
|
|
IPSTATS_MIB_INHDRERRORS);
|
|
icmpv6_param_prob(skb, ICMPV6_HDR_FIELD, skb_network_header_len(skb));
|
|
return -1;
|
|
}
|
|
|
|
static const struct inet6_protocol frag_protocol = {
|
|
.handler = ipv6_frag_rcv,
|
|
.flags = INET6_PROTO_NOPOLICY,
|
|
};
|
|
|
|
#ifdef CONFIG_SYSCTL
|
|
|
|
static struct ctl_table ip6_frags_ns_ctl_table[] = {
|
|
{
|
|
.procname = "ip6frag_high_thresh",
|
|
.data = &init_net.ipv6.frags.high_thresh,
|
|
.maxlen = sizeof(unsigned long),
|
|
.mode = 0644,
|
|
.proc_handler = proc_doulongvec_minmax,
|
|
.extra1 = &init_net.ipv6.frags.low_thresh
|
|
},
|
|
{
|
|
.procname = "ip6frag_low_thresh",
|
|
.data = &init_net.ipv6.frags.low_thresh,
|
|
.maxlen = sizeof(unsigned long),
|
|
.mode = 0644,
|
|
.proc_handler = proc_doulongvec_minmax,
|
|
.extra2 = &init_net.ipv6.frags.high_thresh
|
|
},
|
|
{
|
|
.procname = "ip6frag_time",
|
|
.data = &init_net.ipv6.frags.timeout,
|
|
.maxlen = sizeof(int),
|
|
.mode = 0644,
|
|
.proc_handler = proc_dointvec_jiffies,
|
|
},
|
|
{ }
|
|
};
|
|
|
|
/* secret interval has been deprecated */
|
|
static int ip6_frags_secret_interval_unused;
|
|
static struct ctl_table ip6_frags_ctl_table[] = {
|
|
{
|
|
.procname = "ip6frag_secret_interval",
|
|
.data = &ip6_frags_secret_interval_unused,
|
|
.maxlen = sizeof(int),
|
|
.mode = 0644,
|
|
.proc_handler = proc_dointvec_jiffies,
|
|
},
|
|
{ }
|
|
};
|
|
|
|
static int __net_init ip6_frags_ns_sysctl_register(struct net *net)
|
|
{
|
|
struct ctl_table *table;
|
|
struct ctl_table_header *hdr;
|
|
|
|
table = ip6_frags_ns_ctl_table;
|
|
if (!net_eq(net, &init_net)) {
|
|
table = kmemdup(table, sizeof(ip6_frags_ns_ctl_table), GFP_KERNEL);
|
|
if (!table)
|
|
goto err_alloc;
|
|
|
|
table[0].data = &net->ipv6.frags.high_thresh;
|
|
table[0].extra1 = &net->ipv6.frags.low_thresh;
|
|
table[0].extra2 = &init_net.ipv6.frags.high_thresh;
|
|
table[1].data = &net->ipv6.frags.low_thresh;
|
|
table[1].extra2 = &net->ipv6.frags.high_thresh;
|
|
table[2].data = &net->ipv6.frags.timeout;
|
|
}
|
|
|
|
hdr = register_net_sysctl(net, "net/ipv6", table);
|
|
if (!hdr)
|
|
goto err_reg;
|
|
|
|
net->ipv6.sysctl.frags_hdr = hdr;
|
|
return 0;
|
|
|
|
err_reg:
|
|
if (!net_eq(net, &init_net))
|
|
kfree(table);
|
|
err_alloc:
|
|
return -ENOMEM;
|
|
}
|
|
|
|
static void __net_exit ip6_frags_ns_sysctl_unregister(struct net *net)
|
|
{
|
|
struct ctl_table *table;
|
|
|
|
table = net->ipv6.sysctl.frags_hdr->ctl_table_arg;
|
|
unregister_net_sysctl_table(net->ipv6.sysctl.frags_hdr);
|
|
if (!net_eq(net, &init_net))
|
|
kfree(table);
|
|
}
|
|
|
|
static struct ctl_table_header *ip6_ctl_header;
|
|
|
|
static int ip6_frags_sysctl_register(void)
|
|
{
|
|
ip6_ctl_header = register_net_sysctl(&init_net, "net/ipv6",
|
|
ip6_frags_ctl_table);
|
|
return ip6_ctl_header == NULL ? -ENOMEM : 0;
|
|
}
|
|
|
|
static void ip6_frags_sysctl_unregister(void)
|
|
{
|
|
unregister_net_sysctl_table(ip6_ctl_header);
|
|
}
|
|
#else
|
|
static int ip6_frags_ns_sysctl_register(struct net *net)
|
|
{
|
|
return 0;
|
|
}
|
|
|
|
static void ip6_frags_ns_sysctl_unregister(struct net *net)
|
|
{
|
|
}
|
|
|
|
static int ip6_frags_sysctl_register(void)
|
|
{
|
|
return 0;
|
|
}
|
|
|
|
static void ip6_frags_sysctl_unregister(void)
|
|
{
|
|
}
|
|
#endif
|
|
|
|
static int __net_init ipv6_frags_init_net(struct net *net)
|
|
{
|
|
int res;
|
|
|
|
net->ipv6.frags.high_thresh = IPV6_FRAG_HIGH_THRESH;
|
|
net->ipv6.frags.low_thresh = IPV6_FRAG_LOW_THRESH;
|
|
net->ipv6.frags.timeout = IPV6_FRAG_TIMEOUT;
|
|
net->ipv6.frags.f = &ip6_frags;
|
|
|
|
res = inet_frags_init_net(&net->ipv6.frags);
|
|
if (res < 0)
|
|
return res;
|
|
|
|
res = ip6_frags_ns_sysctl_register(net);
|
|
if (res < 0)
|
|
inet_frags_exit_net(&net->ipv6.frags);
|
|
return res;
|
|
}
|
|
|
|
static void __net_exit ipv6_frags_exit_net(struct net *net)
|
|
{
|
|
ip6_frags_ns_sysctl_unregister(net);
|
|
inet_frags_exit_net(&net->ipv6.frags);
|
|
}
|
|
|
|
static struct pernet_operations ip6_frags_ops = {
|
|
.init = ipv6_frags_init_net,
|
|
.exit = ipv6_frags_exit_net,
|
|
};
|
|
|
|
static const struct rhashtable_params ip6_rhash_params = {
|
|
.head_offset = offsetof(struct inet_frag_queue, node),
|
|
.hashfn = ip6frag_key_hashfn,
|
|
.obj_hashfn = ip6frag_obj_hashfn,
|
|
.obj_cmpfn = ip6frag_obj_cmpfn,
|
|
.automatic_shrinking = true,
|
|
};
|
|
|
|
int __init ipv6_frag_init(void)
|
|
{
|
|
int ret;
|
|
|
|
ip6_frags.constructor = ip6frag_init;
|
|
ip6_frags.destructor = NULL;
|
|
ip6_frags.qsize = sizeof(struct frag_queue);
|
|
ip6_frags.frag_expire = ip6_frag_expire;
|
|
ip6_frags.frags_cache_name = ip6_frag_cache_name;
|
|
ip6_frags.rhash_params = ip6_rhash_params;
|
|
ret = inet_frags_init(&ip6_frags);
|
|
if (ret)
|
|
goto out;
|
|
|
|
ret = inet6_add_protocol(&frag_protocol, IPPROTO_FRAGMENT);
|
|
if (ret)
|
|
goto err_protocol;
|
|
|
|
ret = ip6_frags_sysctl_register();
|
|
if (ret)
|
|
goto err_sysctl;
|
|
|
|
ret = register_pernet_subsys(&ip6_frags_ops);
|
|
if (ret)
|
|
goto err_pernet;
|
|
|
|
out:
|
|
return ret;
|
|
|
|
err_pernet:
|
|
ip6_frags_sysctl_unregister();
|
|
err_sysctl:
|
|
inet6_del_protocol(&frag_protocol, IPPROTO_FRAGMENT);
|
|
err_protocol:
|
|
inet_frags_fini(&ip6_frags);
|
|
goto out;
|
|
}
|
|
|
|
void ipv6_frag_exit(void)
|
|
{
|
|
ip6_frags_sysctl_unregister();
|
|
unregister_pernet_subsys(&ip6_frags_ops);
|
|
inet6_del_protocol(&frag_protocol, IPPROTO_FRAGMENT);
|
|
inet_frags_fini(&ip6_frags);
|
|
}
|